一段采集程序代碼
更新時(shí)間:2006年06月26日 00:00:00 作者:
<%@LANGUAGE="JScript" CODEPAGE="936"%>
<script language=VBScript runat="Server">
Function bytes2BSTR(vIn)
strReturn = ""
For i = 1 To LenB(vIn)
ThisCharCode = AscB(MidB(vIn,i,1))
If ThisCharCode < &H80 Then
strReturn = strReturn & Chr(ThisCharCode)
Else
NextCharCode = AscB(MidB(vIn,i+1,1))
strReturn = strReturn & Chr(CLng(ThisCharCode) * &H100 + CInt(NextCharCode))
i = i + 1
End If
Next
bytes2BSTR = strReturn
End Function
Function ajaxRead(theURL)
dim XmlHttp
set XmlHttp = CreateObject("Microsoft.XMLHTTP")
XmlHttp.Open "GET", theURL, false
XmlHttp.setRequestHeader "Content-Type","text/HTML"
XmlHttp.Send
dim htmlstr
htmlstr = bytes2BSTR(XmlHttp.responseBody)
ajaxRead = htmlstr
End Function
</script>
<%
var ADOConn;
function OpenDatabase(){
try{
ADOConn = new ActiveXObject("ADODB.Connection");
ADOConn.Open ("Provider=Microsoft.Jet.Oledb.4.0;Data Source="+Server.MapPath("getcaiku.mdb"));
}catch(e){
ADOConn.close;
Response.Write("數(shù)據(jù)庫(kù)連接出錯(cuò),請(qǐng)檢查連接字串。");
Response.End;
}
}
function CloseDatabase(){
ADOConn.close;
}
Response.Buffer = 1;
Server.ScriptTimeout = 99999;
//////////可修改以下參數(shù)////////////////
var beginid = 230;//開(kāi)始ID
var endid = 500;//結(jié)束ID
////////////////////////////////////////
var arr,tstr,tid,getdata;
var countid = 0;
Response.Write ("開(kāi)始采集:從"+beginid+"到"+endid+"<hr>");
Response.Flush;
OpenDatabase();
var re=new RegExp("<title>(.*?) - 彩酷</title>","ig");
for(var fi=beginid;fi<(endid+1);fi++){
tid = String(fi);
getdata = ajaxRead("http://mms.caiku.com/sendcring.aspx?uid=0&id="+tid);
if(arr = re.exec(getdata)!=null){
tstr = String(RegExp.$1);
if(tstr!=null&&tstr!="undefined"&&tstr!="")
tstr = tstr.replace("'","");
ADOConn.execute("INSERT INTO getdata(title,tid)VALUES('"+tstr+"',"+tid+")");
Response.Write (tid+":"+tstr+" ___>OK!<br>");
countid++;
Response.Flush
}
}
re.close;
CloseDatabase();
Response.Write ("<hr>采集完畢!共錄入數(shù)據(jù)"+countid+"條。");
%>
<script language=VBScript runat="Server">
Function bytes2BSTR(vIn)
strReturn = ""
For i = 1 To LenB(vIn)
ThisCharCode = AscB(MidB(vIn,i,1))
If ThisCharCode < &H80 Then
strReturn = strReturn & Chr(ThisCharCode)
Else
NextCharCode = AscB(MidB(vIn,i+1,1))
strReturn = strReturn & Chr(CLng(ThisCharCode) * &H100 + CInt(NextCharCode))
i = i + 1
End If
Next
bytes2BSTR = strReturn
End Function
Function ajaxRead(theURL)
dim XmlHttp
set XmlHttp = CreateObject("Microsoft.XMLHTTP")
XmlHttp.Open "GET", theURL, false
XmlHttp.setRequestHeader "Content-Type","text/HTML"
XmlHttp.Send
dim htmlstr
htmlstr = bytes2BSTR(XmlHttp.responseBody)
ajaxRead = htmlstr
End Function
</script>
<%
var ADOConn;
function OpenDatabase(){
try{
ADOConn = new ActiveXObject("ADODB.Connection");
ADOConn.Open ("Provider=Microsoft.Jet.Oledb.4.0;Data Source="+Server.MapPath("getcaiku.mdb"));
}catch(e){
ADOConn.close;
Response.Write("數(shù)據(jù)庫(kù)連接出錯(cuò),請(qǐng)檢查連接字串。");
Response.End;
}
}
function CloseDatabase(){
ADOConn.close;
}
Response.Buffer = 1;
Server.ScriptTimeout = 99999;
//////////可修改以下參數(shù)////////////////
var beginid = 230;//開(kāi)始ID
var endid = 500;//結(jié)束ID
////////////////////////////////////////
var arr,tstr,tid,getdata;
var countid = 0;
Response.Write ("開(kāi)始采集:從"+beginid+"到"+endid+"<hr>");
Response.Flush;
OpenDatabase();
var re=new RegExp("<title>(.*?) - 彩酷</title>","ig");
for(var fi=beginid;fi<(endid+1);fi++){
tid = String(fi);
getdata = ajaxRead("http://mms.caiku.com/sendcring.aspx?uid=0&id="+tid);
if(arr = re.exec(getdata)!=null){
tstr = String(RegExp.$1);
if(tstr!=null&&tstr!="undefined"&&tstr!="")
tstr = tstr.replace("'","");
ADOConn.execute("INSERT INTO getdata(title,tid)VALUES('"+tstr+"',"+tid+")");
Response.Write (tid+":"+tstr+" ___>OK!<br>");
countid++;
Response.Flush
}
}
re.close;
CloseDatabase();
Response.Write ("<hr>采集完畢!共錄入數(shù)據(jù)"+countid+"條。");
%>
相關(guān)文章
使用xmlHttp結(jié)合ASP實(shí)現(xiàn)網(wǎng)頁(yè)的異步調(diào)用
使用xmlHttp結(jié)合ASP實(shí)現(xiàn)網(wǎng)頁(yè)的異步調(diào)用...2006-06-06
一個(gè)帶采集遠(yuǎn)程文章內(nèi)容,保存圖片,生成文件等完整的采集功能
本文提供了一套完整的ASP采集功能函數(shù),包含提取地址的原字符,保存遠(yuǎn)程的文件到本地模擬登錄,獲取網(wǎng)頁(yè)源碼等功能函數(shù)2009-06-06
asp下利用XMLHTTP 從其他頁(yè)面獲取數(shù)據(jù)的代碼
asp下利用XMLHTTP 從其他頁(yè)面獲取數(shù)據(jù)的代碼...2007-11-11
利用MSXML2.XmlHttp和Adodb.Stream采集圖片
asp下經(jīng)常用來(lái)采集的兩個(gè)組件結(jié)合使用例子2008-05-05

