




版權(quán)說明:本文檔由用戶提供并上傳,收益歸屬內(nèi)容提供方,若內(nèi)容存在侵權(quán),請進(jìn)行舉報或認(rèn)領(lǐng)
文檔簡介
1、機(jī)器學(xué)習(xí)課內(nèi)實(shí)驗(yàn)報告(1) ID算法實(shí)現(xiàn)決策樹2015 - 2016學(xué)年 第 2 學(xué)期專業(yè):智能科學(xué)與技術(shù)班級:智能1301班學(xué)號:06133029姓名:張爭輝一、 實(shí)驗(yàn)?zāi)康模豪斫釯D3算法的基本原理,并且編程實(shí)現(xiàn)。二、 實(shí)驗(yàn)要求:使用C/C+/MATLAB實(shí)現(xiàn)ID3算法。輸入:若干行,每行 5 個字符串,表示Outlook Temperature Humidity Wind Play ball如上表。輸出:決策樹。實(shí)驗(yàn)結(jié)果如下:輸入: Sunny Hot High Weak No Sunny Hot High Strong No Overcast Hot High Weak Yes Rain
2、 Mild High Weak Yes Rain Cool Normal Weak Yes Rain Cool Normal Strong No Overcast Cool Normal Strong Yes Sunny Mild High Weak No Sunny Cool Normal Weak Yes Rain Mild Normal Weak Yes Sunny Mild Normal Strong Yes Overcast Mild High Strong Yes Overcast Hot Normal Weak Yes Rain Mild High Strong No輸出:Out
3、look Rain Wind Strong No Weak Yes Overcast Yes Sunny Humidity Normal Yes High No 三、 具體實(shí)現(xiàn):實(shí)現(xiàn)算法如下:#include <iostream>#include <fstream>#include <math.h>#include <string>using namespace std;#define ROW 14#define COL 5#define log2 0.69314718055typedef struct TNode char data15; ch
4、ar weight15; TNode * firstchild,*nextsibling;*tree;typedef struct LNode char OutLook15; char Temperature15; char Humidity15; char Wind15; char PlayTennis5; LNode *next;*link;typedef struct AttrNode char attributes15;/屬性 int attr_Num;/屬性的個數(shù) AttrNode *next;*Attributes;char * ExamplesROWCOL = /"Ov
5、erCast","Cool","High","Strong","No", /"Rain","Hot","Normal","Strong","Yes", "Sunny","Hot","High","Weak","No", "Sunny","Hot","High&
6、quot;,"Strong","No", "OverCast","Hot","High","Weak","Yes", "Rain","Mild","High","Weak","Yes", "Rain","Cool","Normal","Weak","Yes",
7、 "Rain","Cool","Normal","Strong","No", "OverCast","Cool","Normal","Strong","Yes", "Sunny","Mild","High","Weak","No", "Sunny","Cool"
8、;,"Normal","Weak","Yes", "Rain","Mild","Normal","Weak","Yes", "Sunny","Mild","Normal","Strong","Yes", "OverCast","Mild","Normal","Stron
9、g","Yes", "OverCast","Hot","Normal","Weak","Yes", "Rain","Mild","High","Strong","No" ;char * Attributes_kind4 = "OutLook","Temperature","Humidity","Wi
10、nd"int Attr_kind4 = 3,3,2,2;char * OutLook_kind3 = "Sunny","OverCast","Rain"char * Temperature_kind3 = "Hot","Mild","Cool"char * Humidity_kind2 = "High","Normal"char * Wind_kind2 = "Weak","Strong"
11、;/*int i_Exampple145 = 0,0,0,0,1, 0,0,0,1,1, 1,0,0,1,0, 2,1,0,0,0, 2,2,1,0,0, 2,2,1,1,1, 1,2,1,1,0, 0,1,0,0,1, 0,2,1,0,0, 2,1,1,0,0, 0,1,1,1,0, 1,1,1,1,0, 1,1,1,0,0, 2,1,0,0,1 ;*/void treelists(tree T);void InitAttr(Attributes &attr_link,char * Attributes_kind,int Attr_kind);void InitLink(link &
12、amp;L,char * ExamplesCOL);void ID3(tree &T,link L,link Target_Attr,Attributes attr);void PN_Num(link L,int &positve,int &negative);double Gain(int positive,int negative,char * atrribute,link L,Attributes attr_L);void main() link LL,p; Attributes attr_L,q; tree T; T = new TNode; T->fir
13、stchild = T->nextsibling = NULL; strcpy(T->weight,""); strcpy(T->data,""); attr_L = new AttrNode; attr_L->next = NULL; LL = new LNode; LL->next = NULL; /成功建立兩個鏈表 InitLink(LL,Examples); InitAttr(attr_L,Attributes_kind,Attr_kind); ID3(T,LL,NULL,attr_L); cout<<&
14、quot;決策樹以廣義表形式輸出如下:"<<endl; treelists(T);/以廣義表的形式輸出樹/cout<<Gain(9,5,"OutLook",LL,attr_L)<<endl; cout<<endl;/以廣義表的形式輸出樹void treelists(tree T) tree p; if(!T) return; cout<<""<<T->weight<<"" cout<<T->data; p = T-&g
15、t;firstchild; if (p) cout<<"(" while (p) treelists(p); p = p->nextsibling; if (p)cout<<',' cout<<")" void InitAttr(Attributes &attr_link,char * Attributes_kind,int Attr_kind) Attributes p; for (int i =0;i < 4;i+) p = new AttrNode; p->next =
16、NULL; strcpy(p->attributes,Attributes_kindi); p->attr_Num = Attr_kindi; p->next = attr_link->next; attr_link->next = p; void InitLink(link &LL,char * ExamplesCOL) link p; for (int i = 0;i < ROW;i+) p = new LNode; p->next = NULL; strcpy(p->OutLook,Examplesi0); strcpy(p->
17、;Temperature,Examplesi1); strcpy(p->Humidity,Examplesi2); strcpy(p->Wind,Examplesi3); strcpy(p->PlayTennis,Examplesi4); p->next = LL->next; LL->next = p; void PN_Num(link L,int &positve,int &negative) positve = 0; negative = 0; link p; p = L->next; while (p) if (strcmp(p
18、->PlayTennis,"No") = 0) negative+; else if(strcmp(p->PlayTennis,"Yes") = 0) positve+; p = p->next; /計算信息增益/link L: 樣本集合S/attr_L:屬性集合double Gain(int positive,int negative,char * atrribute,link L,Attributes attr_L) int atrr_kinds;/每個屬性中的值的個數(shù) Attributes p = attr_L->next;
19、 link q = L->next; int attr_th = 0;/第幾個屬性 while (p) if (strcmp(p->attributes,atrribute) = 0) atrr_kinds = p->attr_Num; break; p = p->next; attr_th+; double entropy,gain=0; double p1 = 1.0*positive/(positive + negative); double p2 = 1.0*negative/(positive + negative); entropy = -p1*log(p1
20、)/log2 - p2*log(p2)/log2;/集合熵 gain = entropy; /獲取每個屬性值在訓(xùn)練樣本中出現(xiàn)的個數(shù) /獲取每個屬性值所對應(yīng)的正例和反例的個數(shù) /聲明一個3*atrr_kinds的數(shù)組 int * kinds= new int * 3; for (int j =0;j < 3;j+) kindsj = new intatrr_kinds;/保存每個屬性值在訓(xùn)練樣本中出現(xiàn)的個數(shù) /初始化 for (int j = 0;j< 3;j+) for (int i =0;i < atrr_kinds;i+) kindsji = 0; while (q) i
21、f (strcmp("OutLook",atrribute) = 0) for (int i = 0;i < atrr_kinds;i+) if(strcmp(q->OutLook,OutLook_kindi) = 0) kinds0i+; if(strcmp(q->PlayTennis,"Yes") = 0) kinds1i+; else kinds2i+; else if (strcmp("Temperature",atrribute) = 0) for (int i = 0;i < atrr_kinds;
22、i+) if(strcmp(q->Temperature,Temperature_kindi) = 0) kinds0i+; if(strcmp(q->PlayTennis,"Yes") = 0) kinds1i+; else kinds2i+; else if (strcmp("Humidity",atrribute) = 0) for (int i = 0;i < atrr_kinds;i+) if(strcmp(q->Humidity,Humidity_kindi) = 0) kinds0i+; if(strcmp(q-&g
23、t;PlayTennis,"Yes") = 0) kinds1i+;/ else kinds2i+; else if (strcmp("Wind",atrribute) = 0) for (int i = 0;i < atrr_kinds;i+) if(strcmp(q->Wind,Wind_kindi) = 0) kinds0i+; if(strcmp(q->PlayTennis,"Yes") = 0) kinds1i+; else kinds2i+; q = q->next; /計算信息增益 double
24、* gain_kind = new doubleatrr_kinds; int positive_kind = 0,negative_kind = 0; for (int j = 0;j < atrr_kinds;j+) if (kinds0j != 0 && kinds1j != 0 && kinds2j != 0) p1 = 1.0*kinds1j/kinds0j; p2 = 1.0*kinds2j/kinds0j; gain_kindj = -p1*log(p1)/log2-p2*log(p2)/log2; gain = gain - (1.0*ki
25、nds0j/(positive + negative)*gain_kindj; else gain_kindj = 0; return gain;/在ID3算法中的訓(xùn)練樣本子集合與屬性子集合的鏈表需要進(jìn)行清空void FreeLink(link &Link) link p,q; p = Link->next; Link->next = NULL; while (p) q = p; p = p->next; free(q); void ID3(tree &T,link L,link Target_Attr,Attributes attr) Attributes
26、p,max,attr_child,p1; link q,link_child,q1; tree r,tree_p; int positive =0,negative =0; PN_Num(L,positive,negative); /初始化兩個子集合 attr_child = new AttrNode; attr_child->next = NULL; link_child = new LNode; link_child->next = NULL; if (positive = 0)/全是反例 strcpy(T->data,"No"); return; e
27、lse if( negative = 0)/全是正例 strcpy(T->data,"Yes"); return; p = attr->next; /屬性鏈表 double gain,g = 0; /*/ /* 建立屬性子集合與訓(xùn)練樣本子集合有兩個方案: 一:在原來鏈表的基礎(chǔ)上進(jìn)行刪除; 二:另外申請空間進(jìn)行存儲子集合; 采用第二種方法雖然浪費(fèi)了空間,但也省了很多事情,避免了變量之間的應(yīng)用混亂 */ /*/ if(p) while (p) gain = Gain(positive,negative,p->attributes,L,attr); cout&l
28、t;<p->attributes<<" "<<gain<<endl; if(gain > g) g = gain; max = p;/尋找信息增益最大的屬性 p = p->next; strcpy(T->data,max->attributes);/增加決策樹的節(jié)點(diǎn) cout<<"信息增益最大的屬性:max->attributes = "<<max->attributes<<endl<<endl; /下面開始建立決策樹 /創(chuàng)
29、建屬性子集合 p = attr->next; while (p) if (strcmp(p->attributes,max->attributes) != 0) p1 = new AttrNode; strcpy(p1->attributes,p->attributes); p1->attr_Num = p->attr_Num; p1->next = NULL; p1->next = attr_child->next; attr_child->next = p1; p = p->next; /需要區(qū)分出是哪一種屬性 /建立
30、每一層的第一個節(jié)點(diǎn) if (strcmp("OutLook",max->attributes) = 0) r = new TNode; r->firstchild = r->nextsibling = NULL; strcpy(r->weight,OutLook_kind0); T->firstchild = r; /獲取與屬性值相關(guān)的訓(xùn)練樣例Example(vi),建立一個新的訓(xùn)練樣本鏈表link_child q = L->next; while (q) if (strcmp(q->OutLook,OutLook_kind0) =
31、 0) q1 = new LNode; strcpy(q1->OutLook,q->OutLook); strcpy(q1->Humidity,q->Humidity); strcpy(q1->Temperature,q->Temperature); strcpy(q1->Wind,q->Wind); strcpy(q1->PlayTennis,q->PlayTennis); q1->next = NULL; q1->next = link_child->next; link_child->next = q1;
32、 q = q->next; else if (strcmp("Temperature",max->attributes) = 0) r = new TNode; r->firstchild = r->nextsibling = NULL; strcpy(r->weight,Temperature_kind0); T->firstchild = r; /獲取與屬性值相關(guān)的訓(xùn)練樣例Example(vi),建立一個新的訓(xùn)練樣本鏈表link_child q = L->next; while (q) if (strcmp(q->Temp
33、erature,Temperature_kind0) = 0) q1 = new LNode; strcpy(q1->OutLook,q->OutLook); strcpy(q1->Humidity,q->Humidity); strcpy(q1->Temperature,q->Temperature); strcpy(q1->Wind,q->Wind); strcpy(q1->PlayTennis,q->PlayTennis); q1->next = NULL; q1->next = link_child->nex
34、t; link_child->next = q1; q = q->next; else if (strcmp("Humidity",max->attributes) = 0) r = new TNode; r->firstchild = r->nextsibling = NULL; strcpy(r->weight,Humidity_kind0); T->firstchild = r; /獲取與屬性值相關(guān)的訓(xùn)練樣例Example(vi),建立一個新的訓(xùn)練樣本鏈表link_child q = L->next; while (q)
35、 if (strcmp(q->Humidity,Humidity_kind0) = 0) q1 = new LNode; strcpy(q1->OutLook,q->OutLook); strcpy(q1->Humidity,q->Humidity); strcpy(q1->Temperature,q->Temperature); strcpy(q1->Wind,q->Wind); strcpy(q1->PlayTennis,q->PlayTennis); q1->next = NULL; q1->next = li
36、nk_child->next; link_child->next = q1; q = q->next; else if (strcmp("Wind",max->attributes) = 0) r = new TNode; r->firstchild = r->nextsibling = NULL; strcpy(r->weight,Wind_kind0); T->firstchild = r; /獲取與屬性值相關(guān)的訓(xùn)練樣例Example(vi),建立一個新的訓(xùn)練樣本鏈表link_child q = L->next; w
37、hile (q) if (strcmp(q->Wind,Wind_kind0) = 0) q1 = new LNode; strcpy(q1->OutLook,q->OutLook); strcpy(q1->Humidity,q->Humidity); strcpy(q1->Temperature,q->Temperature); strcpy(q1->Wind,q->Wind); strcpy(q1->PlayTennis,q->PlayTennis); q1->next = NULL; q1->next = li
38、nk_child->next; link_child->next = q1; q = q->next; int p = 0,n = 0; PN_Num(link_child,p,n); if (p != 0 && n != 0) ID3(T->firstchild,link_child,Target_Attr,attr_child); FreeLink(link_child); else if(p = 0) strcpy(T->firstchild->data,"No"); FreeLink(link_child); /s
39、trcpy(T->firstchild->data,q1->PlayTennis);/-此處應(yīng)該需要修改-:) else if(n = 0) strcpy(T->firstchild->data,"Yes"); FreeLink(link_child); /建立每一層上的其他節(jié)點(diǎn) tree_p = T->firstchild; for (int i = 1;i < max->attr_Num;i+) /需要區(qū)分出是哪一種屬性 if (strcmp("OutLook",max->attributes)
40、= 0) r = new TNode; r->firstchild = r->nextsibling = NULL; strcpy(r->weight,OutLook_kindi); tree_p->nextsibling = r; /獲取與屬性值相關(guān)的訓(xùn)練樣例Example(vi),建立一個新的訓(xùn)練樣本鏈表link_child q = L->next; while (q) if (strcmp(q->OutLook,OutLook_kindi) = 0) q1 = new LNode; strcpy(q1->OutLook,q->OutLook
41、); strcpy(q1->Humidity,q->Humidity); strcpy(q1->Temperature,q->Temperature); strcpy(q1->Wind,q->Wind); strcpy(q1->PlayTennis,q->PlayTennis); q1->next = NULL; q1->next = link_child->next; link_child->next = q1; q = q->next; else if (strcmp("Temperature"
42、;,max->attributes) = 0) r = new TNode; r->firstchild = r->nextsibling = NULL; strcpy(r->weight,Temperature_kindi); tree_p->nextsibling = r; /獲取與屬性值相關(guān)的訓(xùn)練樣例Example(vi),建立一個新的訓(xùn)練樣本鏈表link_child q = L->next; while (q) if (strcmp(q->Temperature,Temperature_kindi) = 0) q1 = new LNode; s
43、trcpy(q1->OutLook,q->OutLook); strcpy(q1->Humidity,q->Humidity); strcpy(q1->Temperature,q->Temperature); strcpy(q1->Wind,q->Wind); strcpy(q1->PlayTennis,q->PlayTennis); q1->next = NULL; q1->next = link_child->next; link_child->next = q1; q = q->next; else
44、 if (strcmp("Humidity",max->attributes) = 0) r = new TNode; r->firstchild = r->nextsibling = NULL; strcpy(r->weight,Humidity_kindi); tree_p->nextsibling = r; /獲取與屬性值相關(guān)的訓(xùn)練樣例Example(vi),建立一個新的訓(xùn)練樣本鏈表link_child q = L->next; while (q) if (strcmp(q->Humidity,Humidity_kindi) = 0) q1 = new LNode; strcpy(q1->OutLook,q->OutLook); strcpy(q1->Humidity,q->Humidity); strcpy(q1-&
溫馨提示
- 1. 本站所有資源如無特殊說明,都需要本地電腦安裝OFFICE2007和PDF閱讀器。圖紙軟件為CAD,CAXA,PROE,UG,SolidWorks等.壓縮文件請下載最新的WinRAR軟件解壓。
- 2. 本站的文檔不包含任何第三方提供的附件圖紙等,如果需要附件,請聯(lián)系上傳者。文件的所有權(quán)益歸上傳用戶所有。
- 3. 本站RAR壓縮包中若帶圖紙,網(wǎng)頁內(nèi)容里面會有圖紙預(yù)覽,若沒有圖紙預(yù)覽就沒有圖紙。
- 4. 未經(jīng)權(quán)益所有人同意不得將文件中的內(nèi)容挪作商業(yè)或盈利用途。
- 5. 人人文庫網(wǎng)僅提供信息存儲空間,僅對用戶上傳內(nèi)容的表現(xiàn)方式做保護(hù)處理,對用戶上傳分享的文檔內(nèi)容本身不做任何修改或編輯,并不能對任何下載內(nèi)容負(fù)責(zé)。
- 6. 下載文件中如有侵權(quán)或不適當(dāng)內(nèi)容,請與我們聯(lián)系,我們立即糾正。
- 7. 本站不保證下載資源的準(zhǔn)確性、安全性和完整性, 同時也不承擔(dān)用戶因使用這些下載資源對自己和他人造成任何形式的傷害或損失。
最新文檔
- 廠長職位聘任與勞動合同
- 美食街運(yùn)營管理合作協(xié)議范本
- 鄉(xiāng)下人家教學(xué)課件下載
- 2024-2025學(xué)年山東省聊城市高一下學(xué)期期中考政治試題及答案
- 高中一年級生物《基因指導(dǎo)蛋白質(zhì)的合成(第1課時)》
- 建筑信息模型與人工智能融合技術(shù)考核試卷
- 化妝品中的酒精成分對皮膚屏障損害研究考核試卷
- 國際體育賽事規(guī)則與賽事轉(zhuǎn)播限制考核試卷
- 高中英文試題及答案
- 智能家居紡織品應(yīng)用分析考核試卷
- 2022 消化內(nèi)科專業(yè) 藥物臨床試驗(yàn)GCP管理制度操作規(guī)程設(shè)計規(guī)范應(yīng)急預(yù)案
- 三級安全教育試題(公司級、部門級、班組級)
- 整流器并聯(lián)運(yùn)行控制策略
- 農(nóng)業(yè)土壤檢測技術(shù)行業(yè)發(fā)展前景及投資風(fēng)險預(yù)測分析報告
- 廣東省深圳市羅湖區(qū)2023-2024學(xué)年二年級下學(xué)期期末考試數(shù)學(xué)試題
- 初級美發(fā)師題庫
- DZ∕T 0214-2020 礦產(chǎn)地質(zhì)勘查規(guī)范 銅、鉛、鋅、銀、鎳、鉬(正式版)
- 博奧工程量清單計價軟件操作指南
- 2024年度-《醫(yī)療事故處理?xiàng)l例》解讀
- (2024年)面神經(jīng)炎課件完整版
- 宇宙星空基礎(chǔ)知識講座
評論
0/150
提交評論