In [1]:
import mnist
import torch
import numpy as np
import matplotlib.pyplot as plt

from sklearn.preprocessing import OneHotEncoder

try:
    from sklearnex import patch_sklearn
    patch_sklearn()
except:
    print("No scikit-learn-intelex found.  We go on with the classic implementation.")

import time
from pathlib import Path

from torchvision import datasets, transforms
from torch.utils.data import DataLoader

# Transformaciones para normalizar los datos
transform = transforms.Compose([
    transforms.ToTensor(),
    transforms.Normalize((0.1307,), (0.3081,))
])

# Cargar los datos de entrenamiento y prueba
train_data = datasets.MNIST(root='data', train=True, download=True, transform=transform)
test_data = datasets.MNIST(root='data', train=False, download=True, transform=transform)

# Crear dataloaders para iterar sobre los datos en lotes
train_loader = DataLoader(train_data, batch_size=64, shuffle=True)
test_loader = DataLoader(test_data, batch_size=64, shuffle=True)

import torch.nn as nn
import torch.nn.functional as F

class Net(nn.Module):
    def __init__(self):
        super(Net, self).__init__()
        # Capas de la red
        self.fc1 = nn.Linear(28 * 28, 512)
        self.fc2 = nn.Linear(512, 10)

    def forward(self, x):
        x = x.view(-1, 28 * 28)
        x = F.relu(self.fc1(x))
        x = self.fc2(x)
        return F.log_softmax(x, dim=1)

net = Net()
print(net)

import torch.optim as optim

# Crear una instancia del modelo
model = Net()

# Definir una función de pérdida
criterion = nn.NLLLoss()

# Definir un optimizador
optimizer = optim.SGD(model.parameters(), lr=0.003)

# Entrenar la red neuronal
epochs = 100
for e in range(epochs):
    running_loss = 0
    for images, labels in train_loader:
        # Limpiar los gradientes
        optimizer.zero_grad()
        
        # Calcular la salida del modelo
        output = model(images)
        
        # Calcular la pérdida
        loss = criterion(output, labels)
        
        # Calcular los gradientes
        loss.backward()
        
        # Actualizar los pesos
        optimizer.step()
        
        running_loss += loss.item()
    else:
        print(f"Epoch: {e+1}/{epochs}.. Training loss: {running_loss/len(train_loader)}")
        
# Guardar el estado del modelo
torch.save(model.state_dict(), 'model.pth')

# Crear una nueva instancia del modelo
#model = Net()

# Cargar el estado del modelo
#model.load_state_dict(torch.load('model.pth'))

# Evaluar el modelo
#model.eval()

Intel(R) Extension for Scikit-learn* enabled (https://github.com/intel/scikit-learn-intelex)


Net(
  (fc1): Linear(in_features=784, out_features=512, bias=True)
  (fc2): Linear(in_features=512, out_features=10, bias=True)
)
Epoch: 1/100.. Training loss: 0.9337186587454159
Epoch: 2/100.. Training loss: 0.4194766165160421
Epoch: 3/100.. Training loss: 0.3490758377478829
Epoch: 4/100.. Training loss: 0.31479394089565604
Epoch: 5/100.. Training loss: 0.2917882199687109
Epoch: 6/100.. Training loss: 0.27396028679110473
Epoch: 7/100.. Training loss: 0.2587689061615386
Epoch: 8/100.. Training loss: 0.24577567486493573
Epoch: 9/100.. Training loss: 0.23416705582854844
Epoch: 10/100.. Training loss: 0.22369768120237252
Epoch: 11/100.. Training loss: 0.21389666094041582
Epoch: 12/100.. Training loss: 0.20530382205428346
Epoch: 13/100.. Training loss: 0.19720816490317838
Epoch: 14/100.. Training loss: 0.18976158048234767
Epoch: 15/100.. Training loss: 0.18272143556897255
Epoch: 16/100.. Training loss: 0.17615293463997878
Epoch: 17/100.. Training loss: 0.17003841756551125
Epoch: 18/100.. T