In [1]:
import torch
import torch.nn as nn
from torchvision import datasets, transforms
from torch.utils.data import DataLoader, random_split
import numpy as np
import os
from collections import namedtuple
from torch import optim

device = torch.device('cuda' if torch.cuda.is_available() else 'cpu')

# ResNet

## 데이터 전처리 설정

In [2]:
def transform_config(resize, phase):
    data_transform = {
        'train': transforms.Compose([
            transforms.RandomResizedCrop(resize, scale=(0.5,1.0)),
            transforms.RandomHorizontalFlip(),
            transforms.RandomRotation(10.),
            transforms.ToTensor()
        ]),
        'val': transforms.Compose([
            transforms.Resize(resize),
            # transforms.CenterCrop(resize),
            transforms.ToTensor()
        ]),
        'test': transforms.Compose([
            transforms.Resize(resize),
            transforms.ToTensor()
        ])
    }
    
    return data_transform[phase]

## param 설정

In [3]:
data_size = 128
epochs = 20
batch_size = 32
learning_rate = 1e-3

## dataset 정의

In [4]:
all_train_data = datasets.FashionMNIST(root = './data/train', train = True, download = True, transform = transform_config(data_size,'train'))
test_data = datasets.FashionMNIST(root = './data/test', train = False, download = True, transform = transform_config(data_size,'test'))

all_train_data_length = len(all_train_data)
train_data_length = int(all_train_data_length*0.9)
val_data_length = int(all_train_data_length - train_data_length)
test_data_length = len(test_data)

train_data, val_data = random_split(all_train_data,[train_data_length, val_data_length])

print('train_data_shape:{} / train_data_length:{}'.format(train_data[0][0].shape, train_data_length))
print('val_data_shape:{} / val_data_length:{}'.format(val_data[0][0].shape, val_data_length))
print('test_data_shape:{} / test_data_length:{}'.format(test_data[0][0].shape, test_data_length))

train_data_shape:torch.Size([1, 128, 128]) / train_data_length:54000
val_data_shape:torch.Size([1, 128, 128]) / val_data_length:6000
test_data_shape:torch.Size([1, 128, 128]) / test_data_length:10000


### dataset의 데이터 DataLoader를 이용하여 메모리로 불러오기

In [5]:
train_iterator = DataLoader(train_data, batch_size=batch_size, shuffle=True)
val_iterator = DataLoader(val_data, batch_size=batch_size, shuffle=False)

## ResNet 모델 구현

In [6]:
class BasicBlock(nn.Module):
    expansion = 1
    
    def __init__(self, in_channels, out_channels, stride):
        super().__init__()
        
        self.conv1 = nn.Conv2d(in_channels, out_channels, kernel_size=3, stride=stride, padding=1, bias=False)
        self.bn1 = nn.BatchNorm2d(out_channels)
        
        self.conv2 = nn.Conv2d(out_channels, out_channels, kernel_size=3, stride=1, padding=1, bias=False)
        self.bn2 = nn.BatchNorm2d(out_channels)
        
        self.relu = nn.ReLU(inplace=True)
        
        self.shortcut = nn.Sequential()
        
        if stride != 1:
            self.shortcut = nn.Sequential(
                nn.Conv2d(in_channels, out_channels, kernel_size=1, stride=stride, bias=False),
                nn.BatchNorm2d(out_channels)
            )
    
    def forward(self, x):
        i = x
        out = self.conv1(x)
        out = self.bn1(out)
        out = self.relu(out)
        out = self.conv2(out)
        out = self.bn2(out)
        out += self.shortcut(i)
        out = relu(out)
        return out

In [7]:
class Bottleneck(nn.Module):
    expansion = 4
    
    def __init__(self, in_channels, out_channels, stride):
        super().__init__()
        
        self.conv1 = nn.Conv2d(in_channels, out_channels, kernel_size=1, stride=1, bias=False)
        self.bn1 = nn.BatchNorm2d(out_channels)
        
        self.conv2 = nn.Conv2d(out_channels, out_channels, kernel_size=3, stride=stride, padding=1, bias=False)
        self.bn2 = nn.BatchNorm2d(out_channels)
        
        self.conv3 = nn.Conv2d(out_channels, self.expansion*out_channels, kernel_size=1, stride=1, bias=False)
        self.bn3 = nn.BatchNorm2d(out_channels)
        
        self.relu = nn.ReLU(inplace=True)
        
        self.shortcut = nn.Sequential()
        
        if stride != 1 or in_channels != out_channels*self.expansion:
            self.shortcut = nn.Sequential(
                nn.Conv2d(in_channels, out_channels, kernel_size=1, stride=stride, bias=False),
                nn.BatchNorm2d(out_channels)
            )
        
    def forwar(self, x):
        i = x
        out = self.conv1(x)
        out = self.bn1(out)
        out = self.relu(out)
        out = self.conv2(out)
        out = self.bn2(out)
        out = self.relu(out)
        out = self.conv3(out)
        out = self.bn3(out)
        out += self.shortcut(i)
        out = relu(out)
        return out

In [8]:
class ResNet(nn.Module):
    def __init__(self, config, num_classes, zero_init_residual=False):
        super().__init__()
        
        block, n_block, channels = config
        self.in_channels = channels[0]
        
        self.conv1 = nn.Conv2d(3, self.in_channels, kernel_size=7, stride=2, padding=3, bias=False)
        self.bn1 = nn.BatchNorm2d(self.in_channels)
        self.maxpool = nn.MaxPool2d(kernel_size=3, stride=2, padding=1)
        
        self.layer1 = self.get_resnet_layer(block, n_block[0], channels[0], stride=1)
        self.layer2 = self.get_resnet_layer(block, n_block[1], channels[1], stride=2)
        self.layer3 = self.get_resnet_layer(block, n_block[2], channels[2], stride=2)
        self.layer4 = self.get_resnet_layer(block, n_block[3], channels[3], stride=2)
        
        self.avgpool = nn.AdaptiveAvgPool2d((1,1))
        self.fc = nn.Linear(self.in_channels, num_classes)
        
        self.relu = nn.ReLU(inplace=True)
        
        if zero_init_residual:
            for m in self.modules():
                if isinstance(m, Bottleneck):
                    nn.init.constant_(m.bn3.weight, 0)
                elif isinstance(m, BasicBlock):
                    nn.init.constant_(m.bn2.weight, 0)
    
    def get_resnet_layer(self, block, n_blocks, channels, stride):
        layer = []
        
        layer.append(block(self.in_channels, channels, stride))
        
        for i in range(1, n_blocks):
            layer.append(block(block.expansion*channels, channels, stride=1))
            
        self.in_channels = block.expansion*channels
        
        return nn.Sequential(*layer)
    
    def forward(self, x):
        x = self.conv1(x)
        x = self.bn1(x)
        x = self.relu(x)
        x = self.maxpool(x)
        x = self.layer1(x)
        x = self.layer2(x)
        x = self.layer3(x)
        x = self.layer4(x)
        x = self.avgpool(x)
        h = x.view(x.shape[0], -1)
        x = self.fc(h)
        
        return x, h

### ResNet model 함수 정의

In [9]:
ResNetConfig = namedtuple('ResNetConfig', ['block', 'n_blocks', 'channels'])

resnet18_config = ResNetConfig(block = BasicBlock,
                               n_blocks = [2,2,2,2],
                               channels = [64, 128, 256, 512])

resnet34_config = ResNetConfig(block = BasicBlock,
                               n_blocks = [3,4,6,3],
                               channels = [64, 128, 256, 512])

resnet50_config = ResNetConfig(block = Bottleneck,
                               n_blocks = [3, 4, 6, 3],
                               channels = [64, 128, 256, 512])

resnet101_config = ResNetConfig(block = Bottleneck,
                                n_blocks = [3, 4, 23, 3],
                                channels = [64, 128, 256, 512])

resnet152_config = ResNetConfig(block = Bottleneck,
                                n_blocks = [3, 8, 36, 3],
                                channels = [64, 128, 256, 512])

def ResNet50():
    return ResNet(resnet50_config, 10, True)

#### 모델 구조

In [10]:
model = ResNet50() #maxpool이전 relu함수 적용 왜 안되는지...?

print(model)

ResNet(
  (conv1): Conv2d(3, 64, kernel_size=(7, 7), stride=(2, 2), padding=(3, 3), bias=False)
  (bn1): BatchNorm2d(64, eps=1e-05, momentum=0.1, affine=True, track_running_stats=True)
  (maxpool): MaxPool2d(kernel_size=3, stride=2, padding=1, dilation=1, ceil_mode=False)
  (layer1): Sequential(
    (0): Bottleneck(
      (conv1): Conv2d(64, 64, kernel_size=(1, 1), stride=(1, 1), bias=False)
      (bn1): BatchNorm2d(64, eps=1e-05, momentum=0.1, affine=True, track_running_stats=True)
      (conv2): Conv2d(64, 64, kernel_size=(3, 3), stride=(1, 1), padding=(1, 1), bias=False)
      (bn2): BatchNorm2d(64, eps=1e-05, momentum=0.1, affine=True, track_running_stats=True)
      (conv3): Conv2d(64, 256, kernel_size=(1, 1), stride=(1, 1), bias=False)
      (bn3): BatchNorm2d(64, eps=1e-05, momentum=0.1, affine=True, track_running_stats=True)
      (relu): ReLU(inplace=True)
      (shortcut): Sequential(
        (0): Conv2d(64, 64, kernel_size=(1, 1), stride=(1, 1), bias=False)
        (1): Batc

#### pytorch에서 제공하는 사전 훈련된 구조

In [11]:
import torchvision.models as models

pretrained_resnet50 = models.resnet50(pretrained = True)

print(pretrained_resnet50)



ResNet(
  (conv1): Conv2d(3, 64, kernel_size=(7, 7), stride=(2, 2), padding=(3, 3), bias=False)
  (bn1): BatchNorm2d(64, eps=1e-05, momentum=0.1, affine=True, track_running_stats=True)
  (relu): ReLU(inplace=True)
  (maxpool): MaxPool2d(kernel_size=3, stride=2, padding=1, dilation=1, ceil_mode=False)
  (layer1): Sequential(
    (0): Bottleneck(
      (conv1): Conv2d(64, 64, kernel_size=(1, 1), stride=(1, 1), bias=False)
      (bn1): BatchNorm2d(64, eps=1e-05, momentum=0.1, affine=True, track_running_stats=True)
      (conv2): Conv2d(64, 64, kernel_size=(3, 3), stride=(1, 1), padding=(1, 1), bias=False)
      (bn2): BatchNorm2d(64, eps=1e-05, momentum=0.1, affine=True, track_running_stats=True)
      (conv3): Conv2d(64, 256, kernel_size=(1, 1), stride=(1, 1), bias=False)
      (bn3): BatchNorm2d(256, eps=1e-05, momentum=0.1, affine=True, track_running_stats=True)
      (relu): ReLU(inplace=True)
      (downsample): Sequential(
        (0): Conv2d(64, 256, kernel_size=(1, 1), stride=(1, 

## optimizer & loos_function 정의

In [12]:
optimizer = optim.Adam(model.parameters(), lr=learning_rate)
criterion = nn.CrossEntropyLoss()

model = model.to(device)
criterion = criterion.to(device)

In [13]:
target_name = {0:'t-shirt', 1:'pants', 2:'sweater', 3:'dress', 4:'coat', 5:'sandal', 6:'shirt', 7:'sneakers', 8:'bag', 9:'ankle_boots'}