In [1]:
import numpy as np
import torch
import torch.nn as nn
import torch.optim as optim
from torch.autograd import Variable

''' 参考：https://github.com/graykode/nlp-tutorial/blob/master/3-1.TextRNN/TextRNN_Torch.ipynb '''

' 参考：https://github.com/graykode/nlp-tutorial/blob/master/3-1.TextRNN/TextRNN_Torch.ipynb '

In [2]:
dtype = torch.FloatTensor

#原始数据
sentences = [ "i like dog", "i love coffee", "i hate milk"]

In [3]:
#构造词典
word_list = " ".join(sentences).split()
word_list = list(set(word_list))
word_dict = {w: i for i, w in enumerate(word_list)}
number_dict = {i: w for i, w in enumerate(word_list)}
vocab_size = len(word_dict)

In [4]:
# TextRNN Parameter
batch_size = len(sentences)
n_step = 2      # number of cells(= number of Step) 每句话长度
hidden_size = 5    # number of hidden units in one cell

In [5]:
#将数据划分为batch
def make_batch(sentences):
    input_batch, target_batch = [], []
    for sen in sentences:
        word = sen.split()
        input = [word_dict[n] for n in word[:-1]]
        target = word_dict[word[-1]]

        input_batch.append(np.eye(vocab_size)[input])
        target_batch.append(target)
    return input_batch, target_batch

In [6]:
# to Torch.Tensor
input_batch, target_batch = make_batch(sentences)
input_batch = Variable(torch.Tensor(input_batch))           #【batch_size, n_step, vocab_size】前n-1个单词的one-hot向量
target_batch = Variable(torch.LongTensor(target_batch))     #第n个单词index

In [7]:
class TextRNN(nn.Module):
    def __init__(self, vocab_size, hidden_size):
        super().__init__()
        self.rnn = nn.RNN(input_size=vocab_size, hidden_size=hidden_size)
        self.W = nn.Parameter(torch.randn([hidden_size, vocab_size]).type(dtype))
        self.b = nn.Parameter(torch.randn([vocab_size]).type(dtype))

    def forward(self, hidden, X):   #hidden:初始隐状态
        X = X.permute(1, 0, 2)       #0, 1维度互换  X:[n_step, batch_size, vocab_size]
        outputs, hidden = self.rnn(X, hidden)   # outputs : [n_step, batch_size, num_directions(=1) * hidden_size]
                                                # hidden(最后时刻隐状态): [num_layers(=1) * num_directions(=1), batch_size, hidden_size]
        outputs = outputs[-1]   #取最终时间隐藏状态作为输出
        outputs = torch.mm(outputs, self.W) + self.b      #[batch_size, vocab_size]
        return outputs

In [8]:
model = TextRNN(vocab_size, hidden_size)
criterion = nn.CrossEntropyLoss()
optimizer = optim.Adam(model.parameters(), lr=0.001)

# Training
for epoch in range(6000):
    optimizer.zero_grad()
    hidden = Variable(torch.zeros(1, batch_size, hidden_size))     #[num_layers(层数), batch_size, hidden_size]
    output = model(hidden, input_batch)     # input_batch : [batch_size, n_step, vocab_size]
    loss = criterion(output, target_batch)  # output: [batch_size, vocab_size], target_batch:[batch_size] (LongTensor, not one-hot)
    loss.backward()
    optimizer.step()    
    if (epoch + 1) % 1000 == 0:
        print('Epoch:', '%04d' % (epoch + 1), 'cost =', '{:.6f}'.format(loss))

Epoch: 1000 cost = 0.063719
Epoch: 2000 cost = 0.013499
Epoch: 3000 cost = 0.005070
Epoch: 4000 cost = 0.002362
Epoch: 5000 cost = 0.001221
Epoch: 6000 cost = 0.000669


In [9]:
input = [sen.split()[:2] for sen in sentences]

# Predict
hidden = Variable(torch.zeros(1, batch_size, hidden_size))
predict = model(hidden, input_batch).data.max(1, keepdim=True)[1]
print([sen.split()[:2] for sen in sentences], '->', [number_dict[n.item()] for n in predict.squeeze()])

[['i', 'like'], ['i', 'love'], ['i', 'hate']] -> ['dog', 'coffee', 'milk']


In [10]:
'''完成'''

'完成'