caffe 利用VGG训练自己的数据

写这个是因为有童鞋在跑VGG的时候遇到各种问题，供参考一下。

网络结构

以VGG16为例，自己跑的细胞数据

solver.prototxt：

net: "/media/dl/source/Experiment/cell/test/vgg/vgg16.prototxt"

test_iter:

test_interval:

base_lr: 0.0001

lr_policy: "step"

gamma: 0.1

stepsize:

display:

max_iter:

momentum: 0.9

weight_decay: 0.0005

snapshot:

snapshot_prefix: "/media/dl/source/Experiment/cell/test/vgg/vgg"

solver_mode: GPU

vgg16.prototxt：

注意，这里的数据层我是用的“ImageData”格式，也就是没有转为LMDB，直接导入图片进去的，因为我用的服务器，为了方便。如果为了更高效，还是使用LMDB数据库的形式。使用LMDB数据库形式的数据层我也写了下，放在这个prototxt后面作为补充。

另外，注意修改最后一个全连接层的num_output为自己的类别数。并修改该层的名字，如我改为了“cellfc8”，是为了finetune vgg时重新训练该层，不使用该层的预训练参数。

 name: "VGG16"

 layer {

   name: "data"

   type: "ImageData"

   top: "data"

   top: "label"

   include {

     phase: TRAIN

   }

   # transform_param {

   #   mirror: true

   #   crop_size:

   #   mean_file: "data/ilsvrc12_shrt_256/imagenet_mean.binaryproto"

   # }

   image_data_param {

     source: "/media/dl/source/Experiment/cell/data/trainnew2_resize/trainnew.txt"

     batch_size:

     shuffle:true

     #is_color: false

     new_height:

     new_width:

   }

 }

 layer {

   name: "data"

   type: "ImageData"

   top: "data"

   top: "label"

   include {

     phase: TEST

   }

   # transform_param {

   #   mirror: false

   #   crop_size:

   #   mean_file: "data/ilsvrc12_shrt_256/imagenet_mean.binaryproto"

   # }

   image_data_param {

     source: "/media/dl/source/Experiment/cell/data/val2_resize/valnew.txt"

     batch_size:

     #is_color: false

     new_height:

     new_width:

   }

 }

 layer {

   bottom: "data"

   top: "conv1_1"

   name: "conv1_1"

   type: "Convolution"

   param {

     lr_mult:

     decay_mult:

   }

   param {

     lr_mult:

     decay_mult:

   }

   convolution_param {

     num_output:

     pad:

     kernel_size:

     weight_filler {

       type: "gaussian"

       std: 0.01

     }

     bias_filler {

       type: "constant"

       value:

     }

   }

 }

 layer {

   bottom: "conv1_1"

   top: "conv1_1"

   name: "relu1_1"

   type: "ReLU"

 }

 layer {

   bottom: "conv1_1"

   top: "conv1_2"

   name: "conv1_2"

   type: "Convolution"

   param {

     lr_mult:

     decay_mult:

   }

   param {

     lr_mult:

     decay_mult:

   }

   convolution_param {

     num_output:

     pad:

     kernel_size:

     weight_filler {

       type: "gaussian"

       std: 0.01

     }

     bias_filler {

       type: "constant"

       value:

     }

   }

 }

 layer {

   bottom: "conv1_2"

   top: "conv1_2"

   name: "relu1_2"

   type: "ReLU"

 }

 layer {

   bottom: "conv1_2"

   top: "pool1"

   name: "pool1"

   type: "Pooling"

   pooling_param {

     pool: MAX

     kernel_size:

     stride:

   }

 }

 layer {

   bottom: "pool1"

   top: "conv2_1"

   name: "conv2_1"

   type: "Convolution"

   param {

     lr_mult:

     decay_mult:

   }

   param {

     lr_mult:

     decay_mult:

   }

   convolution_param {

     num_output:

     pad:

     kernel_size:

     weight_filler {

       type: "gaussian"

       std: 0.01

     }

     bias_filler {

       type: "constant"

       value:

     }

   }

 }

 layer {

   bottom: "conv2_1"

   top: "conv2_1"

   name: "relu2_1"

   type: "ReLU"

 }

 layer {

   bottom: "conv2_1"

   top: "conv2_2"

   name: "conv2_2"

   type: "Convolution"

   param {

     lr_mult:

     decay_mult:

   }

   param {

     lr_mult:

     decay_mult:

   }

   convolution_param {

     num_output:

     pad:

     kernel_size:

     weight_filler {

       type: "gaussian"

       std: 0.01

     }

     bias_filler {

       type: "constant"

       value:

     }

   }

 }

 layer {

   bottom: "conv2_2"

   top: "conv2_2"

   name: "relu2_2"

   type: "ReLU"

 }

 layer {

   bottom: "conv2_2"

   top: "pool2"

   name: "pool2"

   type: "Pooling"

   pooling_param {

     pool: MAX

     kernel_size:

     stride:

   }

 }

 layer {

   bottom: "pool2"

   top: "conv3_1"

   name: "conv3_1"

   type: "Convolution"

   param {

     lr_mult:

     decay_mult:

   }

   param {

     lr_mult:

     decay_mult:

   }

   convolution_param {

     num_output:

     pad:

     kernel_size:

     weight_filler {

       type: "gaussian"

       std: 0.01

     }

     bias_filler {

       type: "constant"

       value:

     }

   }

 }

 layer {

   bottom: "conv3_1"

   top: "conv3_1"

   name: "relu3_1"

   type: "ReLU"

 }

 layer {

   bottom: "conv3_1"

   top: "conv3_2"

   name: "conv3_2"

   type: "Convolution"

   param {

     lr_mult:

     decay_mult:

   }

   param {

     lr_mult:

     decay_mult:

   }

   convolution_param {

     num_output:

     pad:

     kernel_size:

     weight_filler {

       type: "gaussian"

       std: 0.01

     }

     bias_filler {

       type: "constant"

       value:

     }

   }

 }

 layer {

   bottom: "conv3_2"

   top: "conv3_2"

   name: "relu3_2"

   type: "ReLU"

 }

 layer {

   bottom: "conv3_2"

   top: "conv3_3"

   name: "conv3_3"

   type: "Convolution"

   param {

     lr_mult:

     decay_mult:

   }

   param {

     lr_mult:

     decay_mult:

   }

   convolution_param {

     num_output:

     pad:

     kernel_size:

     weight_filler {

       type: "gaussian"

       std: 0.01

     }

     bias_filler {

       type: "constant"

       value:

     }

   }

 }

 layer {

   bottom: "conv3_3"

   top: "conv3_3"

   name: "relu3_3"

   type: "ReLU"

 }

 layer {

   bottom: "conv3_3"

   top: "pool3"

   name: "pool3"

   type: "Pooling"

   pooling_param {

     pool: MAX

     kernel_size:

     stride:

   }

 }

 layer {

   bottom: "pool3"

   top: "conv4_1"

   name: "conv4_1"

   type: "Convolution"

   param {

     lr_mult:

     decay_mult:

   }

   param {

     lr_mult:

     decay_mult:

   }

   convolution_param {

     num_output:

     pad:

     kernel_size:

     weight_filler {

       type: "gaussian"

       std: 0.01

     }

     bias_filler {

       type: "constant"

       value:

     }

   }

 }

 layer {

   bottom: "conv4_1"

   top: "conv4_1"

   name: "relu4_1"

   type: "ReLU"

 }

 layer {

   bottom: "conv4_1"

   top: "conv4_2"

   name: "conv4_2"

   type: "Convolution"

   param {

     lr_mult:

     decay_mult:

   }

   param {

     lr_mult:

     decay_mult:

   }

   convolution_param {

     num_output:

     pad:

     kernel_size:

     weight_filler {

       type: "gaussian"

       std: 0.01

     }

     bias_filler {

       type: "constant"

       value:

     }

   }

 }

 layer {

   bottom: "conv4_2"

   top: "conv4_2"

   name: "relu4_2"

   type: "ReLU"

 }

 layer {

   bottom: "conv4_2"

   top: "conv4_3"

   name: "conv4_3"

   type: "Convolution"

   param {

     lr_mult:

     decay_mult:

   }

   param {

     lr_mult:

     decay_mult:

   }

   convolution_param {

     num_output:

     pad:

     kernel_size:

     weight_filler {

       type: "gaussian"

       std: 0.01

     }

     bias_filler {

       type: "constant"

       value:

     }

   }

 }

 layer {

   bottom: "conv4_3"

   top: "conv4_3"

   name: "relu4_3"

   type: "ReLU"

 }

 layer {

   bottom: "conv4_3"

   top: "pool4"

   name: "pool4"

   type: "Pooling"

   pooling_param {

     pool: MAX

     kernel_size:

     stride:

   }

 }

 layer {

   bottom: "pool4"

   top: "conv5_1"

   name: "conv5_1"

   type: "Convolution"

   param {

     lr_mult:

     decay_mult:

   }

   param {

     lr_mult:

     decay_mult:

   }

   convolution_param {

     num_output:

     pad:

     kernel_size:

     weight_filler {

       type: "gaussian"

       std: 0.01

     }

     bias_filler {

       type: "constant"

       value:

     }

   }

 }

 layer {

   bottom: "conv5_1"

   top: "conv5_1"

   name: "relu5_1"

   type: "ReLU"

 }

 layer {

   bottom: "conv5_1"

   top: "conv5_2"

   name: "conv5_2"

   type: "Convolution"

   param {

     lr_mult:

     decay_mult:

   }

   param {

     lr_mult:

     decay_mult:

   }

   convolution_param {

     num_output:

     pad:

     kernel_size:

     weight_filler {

       type: "gaussian"

       std: 0.01

     }

     bias_filler {

       type: "constant"

       value:

     }

   }

 }

 layer {

   bottom: "conv5_2"

   top: "conv5_2"

   name: "relu5_2"

   type: "ReLU"

 }

 layer {

   bottom: "conv5_2"

   top: "conv5_3"

   name: "conv5_3"

   type: "Convolution"

   param {

     lr_mult:

     decay_mult:

   }

   param {

     lr_mult:

     decay_mult:

   }

   convolution_param {

     num_output:

     pad:

     kernel_size:

     weight_filler {

       type: "gaussian"

       std: 0.01

     }

     bias_filler {

       type: "constant"

       value:

     }

   }

 }

 layer {

   bottom: "conv5_3"

   top: "conv5_3"

   name: "relu5_3"

   type: "ReLU"

 }

 layer {

   bottom: "conv5_3"

   top: "pool5"

   name: "pool5"

   type: "Pooling"

   pooling_param {

     pool: MAX

     kernel_size:

     stride:

   }

 }

 layer {

   bottom: "pool5"

   top: "fc6"

   name: "fc6"

   type: "InnerProduct"

   param {

     lr_mult:

     decay_mult:

   }

   param {

     lr_mult:

     decay_mult:

   }

   inner_product_param {

     num_output:

     weight_filler {

       type: "gaussian"

       std: 0.005

     }

     bias_filler {

       type: "constant"

       value: 0.1

     }

   }

 }

 layer {

   bottom: "fc6"

   top: "fc6"

   name: "relu6"

   type: "ReLU"

 }

 layer {

   bottom: "fc6"

   top: "fc6"

   name: "drop6"

   type: "Dropout"

   dropout_param {

     dropout_ratio: 0.5

   }

 }

 layer {

   bottom: "fc6"

   top: "fc7"

   name: "fc7"

   type: "InnerProduct"

   param {

     lr_mult:

     decay_mult:

   }

   param {

     lr_mult:

     decay_mult:

   }

   inner_product_param {

     num_output:

     weight_filler {

       type: "gaussian"

       std: 0.005

     }

     bias_filler {

       type: "constant"

       value: 0.1

     }

   }

 }

 layer {

   bottom: "fc7"

   top: "fc7"

   name: "relu7"

   type: "ReLU"

 }

 layer {

   bottom: "fc7"

   top: "fc7"

   name: "drop7"

   type: "Dropout"

   dropout_param {

     dropout_ratio: 0.5

   }

 }

 layer {

   bottom: "fc7"

   top: "fc8"

   name: "cellfc8"

   type: "InnerProduct"

   param {

     lr_mult:

     decay_mult:

   }

   param {

     lr_mult:

     decay_mult:

   }

   inner_product_param {

     num_output: 7 #改为自己的类别数

     weight_filler {

       type: "gaussian"

       std: 0.005

     }

     bias_filler {

       type: "constant"

       value: 0.1

     }

   }

 }

 layer {

   name: "accuracy_at_1"

   type: "Accuracy"

   bottom: "fc8"

   bottom: "label"

   top: "accuracy_at_1"

   accuracy_param {

     top_k:

   }

   include {

     phase: TEST

   }

 }

 layer {

   name: "accuracy_at_5"

   type: "Accuracy"

   bottom: "fc8"

   bottom: "label"

   top: "accuracy_at_5"

   accuracy_param {

     top_k:

   }

   include {

     phase: TEST

   }

 }

 layer {

   bottom: "fc8"

   bottom: "label"

   top: "loss"

   name: "loss"

   type: "SoftmaxWithLoss"

 }

如果使用LMDB数据库形式，将前面的数据层改为：

 name: "vgg"

 layer {

   name: "data"

   type: "Data"

   top: "data"

   top: "label"

   include {

     phase: TRAIN

   }

   transform_param {

     mirror: true

     crop_size:

 #如果图片大于224，则使用crop的方式，小于则使用下面的new_height和new_width

    # new_height:

     #new_width:

     mean_file: "vggface/face_mean.binaryproto"

   }

   data_param {

     source: "vggface/face_train_lmdb"

     batch_size:

     backend: LMDB

   }

 }

 layer {

   name: "data"

   type: "Data"

   top: "data"

   top: "label"

   include {

     phase: TEST

   }

   transform_param {

     mirror: false

     crop_size:

 #如果图片大于224，则使用crop的方式，小于则使用下面的new_height和new_width

    # new_height:

     #new_width:

     mean_file: "vggface/face_mean.binaryproto"

   }

   data_param {

     source: "vggface/face_val_lmdb"

     batch_size:

     backend: LMDB

   }

 }

训练

放一个shell命令：

#!/usr/bin/env sh

TOOLS=/home/dl/caffe-jonlong/build/tools

$TOOLS/caffe train \

  -solver=/media/dl/source/Experiment/cell/test/vgg/solver.prototxt \

  -weights=/media/dl/source/Experiment/cell/test/vgg/VGG_ILSVRC_16_layers.caffemodel \

  -gpu=all \

预训练模型VGG_ILSVRC_16_layers.caffemodel的下载地址为

http://www.robots.ox.ac.uk/~vgg/software/very_deep/caffe/VGG_ILSVRC_16_layers.caffemodel

caffe 利用VGG训练自己的数据的更多相关文章

利用YOLOV3训练自己的数据
写在前面:YOLOV3只有修改了源码才需要重新make,而且make之前要先make clean. 一.准备数据在/darknet/VOCdevkit1下建立文件夹VOC2007. voc2007文 ...
caffe 如何训练自己的数据图片
申明:此教程加工于caffe 如何训练自己的数据图片一.准备数据有条件的同学,可以去imagenet的官网http://www.image-net.org/download-images,下载im ...
caffe学习三：使用Faster RCNN训练自己的数据
本文假设你已经完成了安装,并可以运行demo.py 不会安装且用PASCAL VOC数据集的请看另来两篇博客. caffe学习一:ubuntu16.04下跑Faster R-CNN demo (基于c ...
caffe 用faster rcnn 训练自己的数据遇到的问题
1 . 怎么处理那些pyx和.c .h文件在lib下有一些文件为.pyx文件,遇到不能import可以cython 那个文件,然后把lib文件夹重新make一下. 遇到.c 和 .h一样的操作. 2 ...
YOLO2解读，训练自己的数据及相关转载以供学习
https://pjreddie.com/darknet/yolo/ 具体安装及使用可以参考官方文档https://github.com/pjreddie/darknet https://blog.c ...
YOLOv3：训练自己的数据（附优化与问题总结）
环境说明系统:ubuntu16.04 显卡:Tesla k80 12G显存 python环境: 2.7 && 3.6 前提条件:cuda9.0 cudnn7.0 opencv3.4. ...
人脸检测及识别python实现系列（3）——为模型训练准备人脸数据
人脸检测及识别python实现系列(3)——为模型训练准备人脸数据机器学习最本质的地方就是基于海量数据统计的学习,说白了,机器学习其实就是在模拟人类儿童的学习行为.举一个简单的例子,成年人并没有主动 ...
TensorFlow下利用MNIST训练模型识别手写数字
本文将参考TensorFlow中文社区官方文档使用mnist数据集训练一个多层卷积神经网络(LeNet5网络),并利用所训练的模型识别自己手写数字. 训练MNIST数据集,并保存训练模型 # Pyth ...
py-faster-rcnn 训练自己的数据
转载:http://blog.csdn.net/sinat_30071459/article/details/51332084 Faster-RCNN+ZF用自己的数据集训练模型(Python版本) ...

随机推荐

eMMC基础技术6：eMMC data读写
1. 前言 data可以经data线从host发往device,也可以从device发往host 数据线以是1线(DATA0),4线(DATA0~DATA3),8线(DATA0~DATA7) 对每条数 ...
ARMV8 datasheet学习笔记4：AArch64系统级体系结构之编程模型（2）- 寄存器
1. 前言 2. 指令运行与异常处理寄存器 ARM体系结构的寄存器分为两类: (1)系统控制和状态报告寄存器 (2)指令处理寄存器,如累加.异常处理本部分将主要介绍如上第(2)部分的寄存器,分为AA ...
Python3学习笔记10-条件控制
Python条件语句是通过一条或多条语句的执行结果(True或者False)来决定执行的代码块 var1 = 100 if var1: print("1 - if 表达式条件为 true&q ...
svn的常用命令
svn :看log.版本库.增删.提交 (1)svn up //代码更新到最新版本. (2)svn checkout //将代码checkout出来. (3)svn revert -R ./ //将代 ...
java中不同类型的数值占用字节数
在Java中一共有8种基本数据类型,其中有4种整型,2种浮点类型,1种用于表示Unicode编码的字符单元的字符类型和1种用于表示真值的boolean类型.(一个字节等于8个bit) 1.整型类型 ...
nginx配置集群
1.准备两个Tomcat 首先在Linux机器上部署两个Tomcat,端口分别为80和8080 2.分别部署测试应用在两个tomcat下分别部署同一个应用testapp,很简单,就是在页面显示当前系 ...
红黑树与AVL树
概述:本文从排序二叉树作为引子,讲解了红黑树,最后把红黑树和AVL树做了一个比较全面的对比. 1 排序二叉树排序二叉树是一种特殊结构的二叉树,可以非常方便地对树中所有节点进行排序和检索. 排序二叉树 ...
activemq 消息类型
//文本消息 TextMessage textMessage = session.createTextMessage("文本消息"); producer.send(textMess ...
转载：Linux内核参数的优化（1.3.4）《深入理解Nginx》（陶辉）
原文:https://book.2cto.com/201304/19615.html 由于默认的Linux内核参数考虑的是最通用的场景,这明显不符合用于支持高并发访问的Web服务器的定义,所以需要修改 ...
Python-bootstrap
1 引入如果想要用到BootStrap提供的js插件,那么还需要引入jQuery框架,因为BootStrap提供的js插件是依赖于jQuery的 <link type="text/c ...

caffe 利用VGG训练自己的数据

网络结构

训练

caffe 利用VGG训练自己的数据的更多相关文章

随机推荐

热门专题