搜索
首页php教程php手册PHP采集程序中常用的函数

PHP采集程序中常用的函数

Jun 13, 2016 am 10:59 AM
php例子关键字函数查询实际的程序获得采集

函数描述及例子 PHP采集程序中常用的函数 查询关键字 PHP采集程序中常用的函数

//获得当前的脚本网址 
function get_php_url(){ 
        if(!empty($_SERVER["REQUEST_URI"])){ 
                $scriptName = $_SERVER["REQUEST_URI"]; 
                $nowurl = $scriptName; 
        }else{ 
                $scriptName = $_SERVER["PHP_SELF"]; 
                if(empty($_SERVER["QUERY_STRING"])) $nowurl = $scriptName; 
                else $nowurl = $scriptName."?".$_SERVER["QUERY_STRING"]; 
        } 
        return $nowurl; 
} 
//把全角数字转为半角数字 
function GetAlabNum($fnum){ 
        $nums = array("0","1","2","3","4","5","6","7","8","9"); 
        $fnums = "0123456789"; 
        for($i=0;$i<=9;$i++) $fnum = str_replace($nums[$i],$fnums[$i],$fnum); 
        $fnum = ereg_replace("[^0-9\.]|^0{1,}","",$fnum); 
        if($fnum=="") $fnum=0; 
        return $fnum; 
} 
//去除HTML标记 
function Text2Html($txt){ 
        $txt = str_replace("  "," ",$txt); 
        $txt = str_replace("<","<",$txt); 
        $txt = str_replace(">",">",$txt); 
        $txt = preg_replace("/[\r\n]{1,}/isU","
\r\n",$txt); 
        return $txt; 
}
//清除HTML标记 
function ClearHtml($str){ 
        $str = str_replace(&#39;<&#39;,&#39;<&#39;,$str); 
        $str = str_replace(&#39;>&#39;,&#39;>&#39;,$str); 
        return $str; 
} 
//相对路径转化成绝对路径 
function relative_to_absolute($content, $feed_url) { 
    preg_match(&#39;/(http|https|ftp):\/\//&#39;, $feed_url, $protocol); 
    $server_url = preg_replace("/(http|https|ftp|news):\/\//", "", $feed_url); 
    $server_url = preg_replace("/\/.*/", "", $server_url);
    if ($server_url == &#39;&#39;) { 
        return $content; 
    }
    if (isset($protocol[0])) { 
        $new_content = preg_replace(&#39;/href="\//&#39;, &#39;href="&#39;.$protocol[0].$server_url.&#39;/&#39;, $content); 
        $new_content = preg_replace(&#39;/src="\//&#39;, &#39;src="&#39;.$protocol[0].$server_url.&#39;/&#39;, $new_content); 
    } else { 
        $new_content = $content; 
    } 
    return $new_content; 
} 
//取得所有链接 
function get_all_url($code){ 
        preg_match_all(&#39;/"\&#39; ]+)["|\&#39;]?\s*[^>]*>([^>]+)<\/a>/i&#39;,$code,$arr); 
        return array(&#39;name&#39;=>$arr[2],&#39;url&#39;=>$arr[1]); 
}
//获取指定标记中的内容 
function get_tag_data($str, $start, $end){ 
        if ( $start == &#39;&#39; || $end == &#39;&#39; ){ 
               return; 
        } 
        $str = explode($start, $str); 
        $str = explode($end, $str[1]); 
        return $str[0]; 
} 
//HTML表格的每行转为CSV格式数组 
function get_tr_array($table) { 
        $table = preg_replace("&#39;<td[^>]*?>&#39;si",&#39;"&#39;,$table); 
        $table = str_replace("",&#39;",&#39;,$table); 
        $table = str_replace("","{tr}",$table); 
        //去掉 HTML 标记 
        $table = preg_replace("&#39;<[\/\!]*?[^<>]*?>&#39;si","",$table); 
        //去掉空白字符 
        $table = preg_replace("&#39;([\r\n])[\s]+&#39;","",$table); 
        $table = str_replace(" ","",$table); 
        $table = str_replace(" ","",$table);
        $table = explode(",{tr}",$table); 
        array_pop($table); 
        return $table; 
}
//将HTML表格的每行每列转为数组,采集表格数据 
function get_td_array($table) { 
        $table = preg_replace("&#39;<table[^>]*?>&#39;si","",$table); 
        $table = preg_replace("&#39;<tr[^>]*?>&#39;si","",$table); 
        $table = preg_replace("&#39;<td[^>]*?>&#39;si","",$table); 
        $table = str_replace("","{tr}",$table); 
        $table = str_replace("","{td}",$table); 
        //去掉 HTML 标记 
        $table = preg_replace("&#39;<[\/\!]*?[^<>]*?>&#39;si","",$table); 
        //去掉空白字符 
        $table = preg_replace("&#39;([\r\n])[\s]+&#39;","",$table); 
        $table = str_replace(" ","",$table); 
        $table = str_replace(" ","",$table); 
        
        $table = explode(&#39;{tr}&#39;, $table); 
        array_pop($table); 
        foreach ($table as $key=>$tr) { 
                $td = explode(&#39;{td}&#39;, $tr); 
                array_pop($td); 
            $td_array[] = $td; 
        } 
        return $td_array; 
}
//返回字符串中的所有单词 $distinct=true 去除重复 
function split_en_str($str,$distinct=true) { 
        preg_match_all(&#39;/([a-zA-Z]+)/&#39;,$str,$match); 
        if ($distinct == true) { 
                $match[1] = array_unique($match[1]); 
        } 
        sort($match[1]); 
        return $match[1]; 
}
 
函数描述及例子
 
PHP采集程序中常用的函数

查询关键字
 
PHP采集程序中常用的函数
<!--?
//获得当前的脚本网址 
function get_php_url(){ 
        if(!empty($_SERVER["REQUEST_URI"])){ 
                $scriptName = $_SERVER["REQUEST_URI"]; 
                $nowurl = $scriptName; 
        }else{ 
                $scriptName = $_SERVER["PHP_SELF"]; 
                if(empty($_SERVER["QUERY_STRING"])) $nowurl = $scriptName; 
                else $nowurl = $scriptName."?".$_SERVER["QUERY_STRING"]; 
        } 
        return $nowurl; 
} 
//把全角数字转为半角数字 
function GetAlabNum($fnum){ 
        $nums = array("0","1","2","3","4","5","6","7","8","9"); 
        $fnums = "0123456789"; 
        for($i=0;$i<=9;$i++) $fnum = str_replace($nums[$i],$fnums[$i],$fnum); 
        $fnum = ereg_replace("[^0-9\.]|^0{1,}","",$fnum); 
        if($fnum=="") $fnum=0; 
        return $fnum; 
} 
//去除HTML标记 
function Text2Html($txt){ 
        $txt = str_replace("  "," ",$txt); 
        $txt = str_replace("<","<",$txt); 
        $txt = str_replace("-->",">",$txt); 
        $txt = preg_replace("/[\r\n]{1,}/isU","
\r\n",$txt); 
        return $txt; 
}
//清除HTML标记 
function ClearHtml($str){ 
        $str = str_replace(&#39;<&#39;,&#39;<&#39;,$str); 
        $str = str_replace(&#39;>&#39;,&#39;>&#39;,$str); 
        return $str; 
} 
//相对路径转化成绝对路径 
function relative_to_absolute($content, $feed_url) { 
    preg_match(&#39;/(http|https|ftp):\/\//&#39;, $feed_url, $protocol); 
    $server_url = preg_replace("/(http|https|ftp|news):\/\//", "", $feed_url); 
    $server_url = preg_replace("/\/.*/", "", $server_url);
    if ($server_url == &#39;&#39;) { 
        return $content; 
    }
    if (isset($protocol[0])) { 
        $new_content = preg_replace(&#39;/href="\//&#39;, &#39;href="&#39;.$protocol[0].$server_url.&#39;/&#39;, $content); 
        $new_content = preg_replace(&#39;/src="\//&#39;, &#39;src="&#39;.$protocol[0].$server_url.&#39;/&#39;, $new_content); 
    } else { 
        $new_content = $content; 
    } 
    return $new_content; 
} 
//取得所有链接 
function get_all_url($code){ 
        preg_match_all(&#39;/"\&#39; ]+)["|\&#39;]?\s*[^>]*>([^>]+)<\/a>/i&#39;,$code,$arr); 
        return array(&#39;name&#39;=>$arr[2],&#39;url&#39;=>$arr[1]); 
}
//获取指定标记中的内容 
function get_tag_data($str, $start, $end){ 
        if ( $start == &#39;&#39; || $end == &#39;&#39; ){ 
               return; 
        } 
        $str = explode($start, $str); 
        $str = explode($end, $str[1]); 
        return $str[0]; 
} 
//HTML表格的每行转为CSV格式数组 
function get_tr_array($table) { 
        $table = preg_replace("&#39;<td[^>]*?>&#39;si",&#39;"&#39;,$table); 
        $table = str_replace("",&#39;",&#39;,$table); 
        $table = str_replace("","{tr}",$table); 
        //去掉 HTML 标记 
        $table = preg_replace("&#39;<[\/\!]*?[^<>]*?>&#39;si","",$table); 
        //去掉空白字符 
        $table = preg_replace("&#39;([\r\n])[\s]+&#39;","",$table); 
        $table = str_replace(" ","",$table); 
        $table = str_replace(" ","",$table);
        $table = explode(",{tr}",$table); 
        array_pop($table); 
        return $table; 
}
//将HTML表格的每行每列转为数组,采集表格数据 
function get_td_array($table) { 
        $table = preg_replace("&#39;<table[^>]*?>&#39;si","",$table); 
        $table = preg_replace("&#39;<tr[^>]*?>&#39;si","",$table); 
        $table = preg_replace("&#39;<td[^>]*?>&#39;si","",$table); 
        $table = str_replace("","{tr}",$table); 
        $table = str_replace("","{td}",$table); 
        //去掉 HTML 标记 
        $table = preg_replace("&#39;<[\/\!]*?[^<>]*?>&#39;si","",$table); 
        //去掉空白字符 
        $table = preg_replace("&#39;([\r\n])[\s]+&#39;","",$table); 
        $table = str_replace(" ","",$table); 
        $table = str_replace(" ","",$table); 
        
        $table = explode(&#39;{tr}&#39;, $table); 
        array_pop($table); 
        foreach ($table as $key=>$tr) { 
                $td = explode(&#39;{td}&#39;, $tr); 
                array_pop($td); 
            $td_array[] = $td; 
        } 
        return $td_array; 
}
//返回字符串中的所有单词 $distinct=true 去除重复 
function split_en_str($str,$distinct=true) { 
        preg_match_all(&#39;/([a-zA-Z]+)/&#39;,$str,$match); 
        if ($distinct == true) { 
                $match[1] = array_unique($match[1]); 
        } 
        sort($match[1]); 
        return $match[1]; 
}
 
</td[^></tr[^></table[^></td[^></a\s+href=["|\&#39;]?([^></td[^></tr[^></table[^></td[^></a\s+href=["|\&#39;]?([^>

声明
本文内容由网友自发贡献,版权归原作者所有,本站不承担相应法律责任。如您发现有涉嫌抄袭侵权的内容,请联系admin@php.cn

热AI工具

Undresser.AI Undress

Undresser.AI Undress

人工智能驱动的应用程序,用于创建逼真的裸体照片

AI Clothes Remover

AI Clothes Remover

用于从照片中去除衣服的在线人工智能工具。

Undress AI Tool

Undress AI Tool

免费脱衣服图片

Clothoff.io

Clothoff.io

AI脱衣机

Video Face Swap

Video Face Swap

使用我们完全免费的人工智能换脸工具轻松在任何视频中换脸!

热工具

VSCode Windows 64位 下载

VSCode Windows 64位 下载

微软推出的免费、功能强大的一款IDE编辑器

MinGW - 适用于 Windows 的极简 GNU

MinGW - 适用于 Windows 的极简 GNU

这个项目正在迁移到osdn.net/projects/mingw的过程中,你可以继续在那里关注我们。MinGW:GNU编译器集合(GCC)的本地Windows移植版本,可自由分发的导入库和用于构建本地Windows应用程序的头文件;包括对MSVC运行时的扩展,以支持C99功能。MinGW的所有软件都可以在64位Windows平台上运行。

mPDF

mPDF

mPDF是一个PHP库,可以从UTF-8编码的HTML生成PDF文件。原作者Ian Back编写mPDF以从他的网站上“即时”输出PDF文件,并处理不同的语言。与原始脚本如HTML2FPDF相比,它的速度较慢,并且在使用Unicode字体时生成的文件较大,但支持CSS样式等,并进行了大量增强。支持几乎所有语言,包括RTL(阿拉伯语和希伯来语)和CJK(中日韩)。支持嵌套的块级元素(如P、DIV),

PhpStorm Mac 版本

PhpStorm Mac 版本

最新(2018.2.1 )专业的PHP集成开发工具

SublimeText3 英文版

SublimeText3 英文版

推荐:为Win版本,支持代码提示!