原文地址：http://www.jianshu.com/p/311141f2047d

问题描述

程序实现

13-15

# coding: utf-8

import numpy as np

import numpy.random as random

import matplotlib.pyplot as plt

def sign(x):

    if(x>=0):

        return 1

    else:

        return -1

def gen_data():

    x1=random.uniform(-1,1,1000)

    x2=random.uniform(-1,1,1000)

    id_array=random.permutation([i for i in range(1000)])

    dataY=np.zeros((1000,1))

    for i in range(1000):

        if(i<1000*0.1):

            i = id_array[i]

            dataY[i][0]=-sign(x1[i]**2+x2[i]**2-0.6)

        else:

            i = id_array[i]

            dataY[i][0]=sign(x1[i]**2+x2[i]**2-0.6)

    dataX=np.concatenate((np.ones((1000,1)),np.array(x1).reshape((1000,1)),np.array(x2).reshape((1000,1))),axis=1)

    return dataX,dataY

def w_lin(dataX,dataY):

    dataX_T=np.transpose(dataX)

    tmp=np.dot(np.linalg.inv(np.dot(dataX_T,dataX)),dataX_T)

    return np.dot(tmp,dataY)

def pred(dataX,wLIN):

    pred=np.dot(dataX,wLIN)

    num_data=dataX.shape[0]

    for i in range(num_data):

        pred[i][0]=sign(pred[i][0])

    return pred

def zero_one_cost(pred,dataY):

    return np.sum(pred!=dataY)/dataY.shape[0]

def feat_transform(dataX):

    num_data=dataX.shape[0]

    tmp1=dataX[:,1]*dataX[:,2]

    tmp2=dataX[:,1]**2

    tmp3=dataX[:,2]**2

    new_dataX=np.concatenate(

        (dataX,tmp1.reshape((num_data,1)),tmp2.reshape((num_data,1)),tmp3.reshape((num_data,1))),axis=1)

    return new_dataX

if __name__=="__main__":

    cost_list=[]

    for i in range(1000):

        dataX,dataY=gen_data()

        wLIN=w_lin(dataX,dataY)

        cost_list.append(zero_one_cost(pred(dataX,wLIN),dataY))

    # show results

    print("the average Ein over 1000 experiments: ",sum(cost_list)/len(cost_list))

    plt.figure()

    plt.hist(cost_list)

    plt.xlabel("zero_one Ein")

    plt.ylabel("frequency")

    plt.title("13")

    plt.savefig("13.png")

    W=[]

    cost_list=[]

    for i in range(1000):

        # train

        dataX,dataY=gen_data()

        dataX=feat_transform(dataX)

        wLIN=w_lin(dataX,dataY)

        W.append(wLIN[:,0].tolist())

        # test

        testX, testY = gen_data()

        testX = feat_transform(testX)

        cost_list.append(zero_one_cost(pred(testX, wLIN), testY))

    min_cost=min(cost_list)

    min_id=cost_list.index(min_cost)

    print(W[min_id])

    W=np.array(W)

    # show w3

    print("the average w3 over 1000 experiments: ",np.average(W,axis=0)[3])

    plt.figure()

    plt.hist(W[:,3].tolist())

    plt.xlabel("w3")

    plt.ylabel("frequency")

    plt.title("14")

    plt.savefig("14.png")

    # show Eout

    print("the average Eout over 1000 experiments: ",sum(cost_list)/len(cost_list))

    plt.figure()

    plt.hist(cost_list)

    plt.xlabel("Eout")

    plt.ylabel("frequency")

    plt.title("15")

    plt.savefig("15.png")

18-20

# coding: utf-8

import numpy as np

def sigmoid(x):

    return 1/(1+np.e**(-x))

def read_data(dataFile):

    with open(dataFile,'r') as f:

        lines=f.readlines()

        data_list=[]

        for line in lines:

            line=line.strip().split()

            data_list.append([1.0] + [float(l) for l in line])

        dataArray=np.array(data_list)

        num_data=dataArray.shape[0]

        num_dim=dataArray.shape[1]-1

        dataX=dataArray[:,:-1].reshape((num_data,num_dim))

        dataY=dataArray[:,-1].reshape((num_data,1))

        return dataX,dataY

def gradient_descent(w,dataX,dataY,eta):

    assert w.shape[0]==dataX.shape[1],"wrong shape!"

    assert w.shape[1]==1,"wrong shape of w!"

    num_data=dataX.shape[0]

    num_dim=dataX.shape[1]

    tmp1=-dataY*dataX

    tmp2=-dataY*np.dot(dataX,w)

    for i in range(num_data):

        tmp2[i][0]=sigmoid(tmp2[i][0])

    tmp3=np.average(tmp1 * tmp2, axis=0)

    new_w=w-eta*tmp3.reshape((num_dim,1))

    return new_w

def s_gradient_descent(w,dataX,dataY,eta):

    assert w.shape[0]==dataX.shape[1],"wrong shape!"

    assert w.shape[1]==1,"wrong shape of w!"

    assert dataX.shape[0]==1,"wrong shape of x!"

    assert dataY.shape[0]==1,"wrong shape of y!"

    num_dim=dataX.shape[1]

    tmp1=-dataY*dataX

    tmp2=-dataY*np.dot(dataX,w)

    tmp2[0][0]=sigmoid(tmp2[0][0])

    tmp3=np.average(tmp1 * tmp2, axis=0)

    new_w=w-eta*tmp3.reshape((num_dim,1))

    return new_w

def pred(wLOG,dataX):

    pred=np.dot(dataX,wLOG)

    num_data=dataX.shape[0]

    for i in range(num_data):

        pred[i][0]=sigmoid(pred[i][0])

        if(pred[i][0]>=0.5):

            pred[i][0]=1

        else:

            pred[i][0]=-1

    return pred

def zero_one_cost(pred,dataY):

    return np.sum(pred!=dataY)/dataY.shape[0]

if __name__=="__main__":

    # train

    dataX,dataY=read_data("hw3_train.dat")

    num_dim=dataX.shape[1]

    w=np.zeros((num_dim,1))

    print("\n18")

    for i in range(2000):

        w=gradient_descent(w,dataX,dataY,eta=0.001)

    print("the weight vector within g: ",w[:,0])

    # test

    testX,testY=read_data("hw3_test.dat")

    Eout=zero_one_cost(pred(w,testX),testY)

    print("the Eout(g) on the test set: ",Eout)

    print("\n18.1")

    w = np.zeros((num_dim, 1))

    for i in range(20000):

        w = gradient_descent(w, dataX, dataY, eta=0.001)

    print("the weight vector within g: ", w[:, 0])

    # test

    Eout = zero_one_cost(pred(w, testX), testY)

    print("the Eout(g) on the test set: ", Eout)

    print("\n19")

    w=np.zeros((num_dim,1))

    for i in range(2000):

        w = gradient_descent(w, dataX, dataY, eta=0.01)

    print("the weight vector within g: ", w[:, 0])

    # test

    Eout = zero_one_cost(pred(w, testX), testY)

    print("the Eout(g) on the test set: ", Eout)

    print("\n20")

    w=np.zeros((num_dim,1))

    num_data=dataX.shape[0]

    for i in range(2000):

        i%=num_data

        x=dataX[i,:].reshape((1,num_dim))

        y=dataY[i,:].reshape((1,1))

        w=s_gradient_descent(w,x,y,eta=0.001)

    print("the weight vector within g: ", w[:, 0])

    # test

    Eout = zero_one_cost(pred(w, testX), testY)

    print("the Eout(g) on the test set: ", Eout)

运行结果及分析

13-15

18-20

对比18和18.1，可知迭代步长较小时，需要较多迭代次数才能达到较优效果。

机器学习基石笔记：Homework #3 LinReg&LogReg相关习题的更多相关文章

机器学习基石笔记：Homework #1 PLA&PA相关习题
原文地址:http://www.jianshu.com/p/5b4a64874650 问题描述程序实现 # coding: utf-8 import numpy as np import matpl ...
机器学习基石笔记：Homework #2 decision stump相关习题
原文地址:http://www.jianshu.com/p/4bc01760ac20 问题描述程序实现 17-18 # coding: utf-8 import numpy as np import ...
机器学习基石笔记：11 Linear Models for Classification、LC vs LinReg vs LogReg、OVA、OVO
原文地址:https://www.jianshu.com/p/6f86290e70f9 一.二元分类的线性模型线性回归后的参数值常用于PLA/PA/Logistic Regression的参数初始化 ...
机器学习基石笔记：Homework #4 Regularization&Validation相关习题
原文地址:https://www.jianshu.com/p/3f7d4aa6a7cf 问题描述程序实现 # coding: utf-8 import numpy as np import math ...
机器学习基石：Homework #0 SVD相关&常用矩阵求导公式
林轩田机器学习基石笔记1—The Learning Problem
机器学习分为四步: When Can Machine Learn? Why Can Machine Learn? How Can Machine Learn? How Can Machine Lear ...
机器学习基石笔记：01 The Learning Problem
原文地址:https://www.jianshu.com/p/bd7cb6c78e5e 什么时候适合用机器学习算法? 存在某种规则/模式,能够使性能提升,比如准确率: 这种规则难以程序化定义,人难以给 ...
机器学习基石笔记：04 Feasibility of Learning
原文地址:https://www.jianshu.com/p/f2f4d509060e 机器学习是设计算法\(A\),在假设集合\(H\)里,根据给定数据集\(D\),选出与实际模式\(f\)最为相近 ...
机器学习基石笔记：03 Types of Learning
原文地址:https://www.jianshu.com/p/86b2a9cef742 一.学习的分类根据输出空间\(Y\):分类(二分类.多分类).回归.结构化(监督学习+输出空间有结构): 根据 ...

随机推荐

QTP学习笔记---datatable应用
DataTable应用1.定位数据行 DataTable.GetSheet() 2.获取当前行 GetCurrentRow3.获取指定行的值 getValueByRow = DataTable.Get ...
NGINX-二级域名
先给二级域名添加到 DNS 解析再配置 nginx server { #侦听80端口 listen 80; #定义使用 www.nginx.cn访问 server_name ~^(?<subdo ...
git 远程库和本地库处理
创建git库的方法第一种方法: 在码云建立一个demo的git库.git clone在本地一个文件夹.之后会出现在.git的目录下方(这是clone而非pull切记分清楚) 而不是在.git的上一层 ...
Notepad++ 连接 FTP 实现编辑 Linux文件
下载并安装插件 github 下载 :https://github.com/ashkulz/NppFTP/releases/ 安装过程将下载后解压的文件夹中的 NppFTP.dll 文件,拷贝到 n ...
Nginx网络架构实战学习笔记（二）：编译PHP并与nginx整合、安装ecshop、商城url重写实战
文章目录编译PHP并与nginx整合安装ecshop(这是一个多年前php的项目貌似,作为java开发的我暂时不去关心) 商城url重写实战编译PHP并与nginx整合安装mysql yum ...
VMware Workstation Pro 15.5.0 官方版本及激活密钥
0x01:下载连接: https://download3.vmware.com/software/wkst/file/VMware-workstation-full-15.5.1-15018445.e ...
原生js星星评分源码
html: <div id="fiveStars"> <div>到场时间:<img v-for="(star,index) in stars ...
Java的枚举类型使用方法详解
1.背景在java语言中还没有引入枚举类型之前,表示枚举类型的常用模式是声明一组具有int常量.之前我们通常利用public final static 方法定义的代码如下,分别用1 表示春天,2表示夏 ...
teb教程10 teb questions
http://wiki.ros.org/teb_local_planner/Tutorials/Frequently%20Asked%20Questions
在linux上安装jdk
一.环境 1. VMware虚拟机 2. Linux系统(centos7) 安装步骤: 1. 先安装VMtools 2. 查看是否有openjdk,若有先将系统内的openJDK删除(采用 r ...

机器学习基石笔记：Homework #3 LinReg&LogReg相关习题

问题描述

程序实现

13-15

18-20

运行结果及分析

13-15

18-20

机器学习基石笔记：Homework #3 LinReg&LogReg相关习题的更多相关文章

随机推荐

热门专题