In [49]:
import numpy as np
import pandas as pd
import matplotlib.pyplot as plt

In [50]:
import tensorflow as tf

In [51]:
with open("shakespeare.txt") as f:
    doc = f.read()

In [52]:

vocab = sorted(set(doc))
len(vocab)

84

In [53]:

char_to_ind = {u:i for i, u in enumerate(vocab)}

In [54]:
char_to_ind

{'\n': 0,
 ' ': 1,
 '!': 2,
 '"': 3,
 '&': 4,
 "'": 5,
 '(': 6,
 ')': 7,
 ',': 8,
 '-': 9,
 '.': 10,
 '0': 11,
 '1': 12,
 '2': 13,
 '3': 14,
 '4': 15,
 '5': 16,
 '6': 17,
 '7': 18,
 '8': 19,
 '9': 20,
 ':': 21,
 ';': 22,
 '<': 23,
 '>': 24,
 '?': 25,
 'A': 26,
 'B': 27,
 'C': 28,
 'D': 29,
 'E': 30,
 'F': 31,
 'G': 32,
 'H': 33,
 'I': 34,
 'J': 35,
 'K': 36,
 'L': 37,
 'M': 38,
 'N': 39,
 'O': 40,
 'P': 41,
 'Q': 42,
 'R': 43,
 'S': 44,
 'T': 45,
 'U': 46,
 'V': 47,
 'W': 48,
 'X': 49,
 'Y': 50,
 'Z': 51,
 '[': 52,
 ']': 53,
 '_': 54,
 '`': 55,
 'a': 56,
 'b': 57,
 'c': 58,
 'd': 59,
 'e': 60,
 'f': 61,
 'g': 62,
 'h': 63,
 'i': 64,
 'j': 65,
 'k': 66,
 'l': 67,
 'm': 68,
 'n': 69,
 'o': 70,
 'p': 71,
 'q': 72,
 'r': 73,
 's': 74,
 't': 75,
 'u': 76,
 'v': 77,
 'w': 78,
 'x': 79,
 'y': 80,
 'z': 81,
 '|': 82,
 '}': 83}

In [55]:
ind_to_char = np.array(vocab)

In [56]:
ind_to_char

array(['\n', ' ', '!', '"', '&', "'", '(', ')', ',', '-', '.', '0', '1',
       '2', '3', '4', '5', '6', '7', '8', '9', ':', ';', '<', '>', '?',
       'A', 'B', 'C', 'D', 'E', 'F', 'G', 'H', 'I', 'J', 'K', 'L', 'M',
       'N', 'O', 'P', 'Q', 'R', 'S', 'T', 'U', 'V', 'W', 'X', 'Y', 'Z',
       '[', ']', '_', '`', 'a', 'b', 'c', 'd', 'e', 'f', 'g', 'h', 'i',
       'j', 'k', 'l', 'm', 'n', 'o', 'p', 'q', 'r', 's', 't', 'u', 'v',
       'w', 'x', 'y', 'z', '|', '}'], dtype='<U1')

In [57]:
encoded_text = np.array([char_to_ind[c] for c in doc])

In [58]:
encoded_text

array([ 0,  1,  1, ..., 30, 39, 29])

In [59]:
seq_len = 120
char_dataset = tf.data.Dataset.from_tensor_slices(encoded_text)

In [60]:
sequences = char_dataset.batch(seq_len+1,drop_remainder=True)

In [61]:
sequences

<BatchDataset shapes: (121,), types: tf.int64>

In [62]:
def create_seq_target(seq):
    input_text = seq[:-1]
    output_text = seq[1:]
    return input_text,output_text

In [63]:
dataset = sequences.map(create_seq_target)

In [64]:
dataset

<MapDataset shapes: ((120,), (120,)), types: (tf.int64, tf.int64)>

In [65]:
for input_text,output_text in dataset.take(1):
    print(input_text.numpy())
    print("".join(ind_to_char[i] for i in input_text.numpy()))
    print("\n")
    print(output_text.numpy())
    print("".join(ind_to_char[j] for j in output_text.numpy()))

[ 0  1  1  1  1  1  1  1  1  1  1  1  1  1  1  1  1  1  1  1  1  1 12  0
  1  1 31 73 70 68  1 61 56 64 73 60 74 75  1 58 73 60 56 75 76 73 60 74
  1 78 60  1 59 60 74 64 73 60  1 64 69 58 73 60 56 74 60  8  0  1  1 45
 63 56 75  1 75 63 60 73 60 57 80  1 57 60 56 76 75 80  5 74  1 73 70 74
 60  1 68 64 62 63 75  1 69 60 77 60 73  1 59 64 60  8  0  1  1 27 76 75]

                     1
  From fairest creatures we desire increase,
  That thereby beauty's rose might never die,
  But


[ 1  1  1  1  1  1  1  1  1  1  1  1  1  1  1  1  1  1  1  1  1 12  0  1
  1 31 73 70 68  1 61 56 64 73 60 74 75  1 58 73 60 56 75 76 73 60 74  1
 78 60  1 59 60 74 64 73 60  1 64 69 58 73 60 56 74 60  8  0  1  1 45 63
 56 75  1 75 63 60 73 60 57 80  1 57 60 56 76 75 80  5 74  1 73 70 74 60
  1 68 64 62 63 75  1 69 60 77 60 73  1 59 64 60  8  0  1  1 27 76 75  1]
                     1
  From fairest creatures we desire increase,
  That thereby beauty's rose might never die,
  But 


In [66]:
batch_size = 128

In [67]:
buffer_size = 10000
dataset = dataset.shuffle(buffer_size=buffer_size).batch(batch_size,drop_remainder=True)

In [68]:
dataset

<BatchDataset shapes: ((128, 120), (128, 120)), types: (tf.int64, tf.int64)>

In [69]:
embed_dim = 64
rnn_neurons = 1026

In [70]:
from tensorflow.keras.losses import sparse_categorical_crossentropy

In [71]:
def sparse_cat_loss(y_true,y_pred):
    return sparse_categorical_crossentropy(y_true,y_pred,from_logits=True)

In [72]:
from tensorflow.keras.layers import Embedding,GRU, Dense
from tensorflow.keras.models import Sequential

In [73]:
def model(vocab_size,embed_dim,rnn_neurons,batch_size):
    model = Sequential()
    model.add(Embedding(vocab_size,embed_dim,batch_input_shape=[batch_size,None]))
    model.add(GRU(rnn_neurons,return_sequences=True,stateful=True,recurrent_initializer='glorot_uniform'))
    model.add(Dense(vocab_size))
    model.compile('adam',loss = sparse_cat_loss)
    return model

In [74]:
model = model(
  vocab_size = len(vocab),
  embed_dim=embed_dim,
  rnn_neurons=rnn_neurons,
  batch_size=batch_size)

In [75]:
model.summary()

Model: "sequential_1"
_________________________________________________________________
Layer (type)                 Output Shape              Param #   
embedding_1 (Embedding)      (128, None, 64)           5376      
_________________________________________________________________
gru_1 (GRU)                  (128, None, 1026)         3361176   
_________________________________________________________________
dense_1 (Dense)              (128, None, 84)           86268     
Total params: 3,452,820
Trainable params: 3,452,820
Non-trainable params: 0
_________________________________________________________________


In [76]:
model.fit(dataset,epochs=20)

Epoch 1/20
Epoch 2/20
Epoch 3/20
Epoch 4/20
Epoch 5/20
Epoch 6/20
Epoch 7/20
Epoch 8/20
Epoch 9/20
Epoch 10/20
Epoch 11/20
Epoch 12/20
Epoch 13/20
Epoch 14/20
Epoch 15/20
Epoch 16/20
Epoch 17/20
Epoch 18/20
Epoch 19/20
Epoch 20/20


<tensorflow.python.keras.callbacks.History at 0x7f04cfa346d8>

In [77]:
model.save('shakespeare_gen.h5') 

In [78]:
from tensorflow.keras.models import load_model

In [82]:
vocab_size = len(vocab)

model = Sequential()
model.add(Embedding(vocab_size,embed_dim,batch_input_shape=[1,None]))
model.add(GRU(rnn_neurons,return_sequences=True,stateful=True,recurrent_initializer='glorot_uniform'))
model.add(Dense(vocab_size))
model.compile('adam',loss = sparse_cat_loss)


model.load_weights('shakespeare_gen.h5')

model.build(tf.TensorShape([1, None]))

In [83]:
model.summary()

Model: "sequential_3"
_________________________________________________________________
Layer (type)                 Output Shape              Param #   
embedding_3 (Embedding)      (1, None, 64)             5376      
_________________________________________________________________
gru_3 (GRU)                  (1, None, 1026)           3361176   
_________________________________________________________________
dense_3 (Dense)              (1, None, 84)             86268     
Total params: 3,452,820
Trainable params: 3,452,820
Non-trainable params: 0
_________________________________________________________________


In [84]:
def generate_text(model, start_seed,gen_size=100,temp=1.0):
  '''
  model: Trained Model to Generate Text
  start_seed: Intial Seed text in string form
  gen_size: Number of characters to generate

  Basic idea behind this function is to take in some seed text, format it so
  that it is in the correct shape for our network, then loop the sequence as
  we keep adding our own predicted characters. Similar to our work in the RNN
  time series problems.
  '''

  # Number of characters to generate
  num_generate = gen_size

  # Vecotrizing starting seed text
  input_eval = [char_to_ind[s] for s in start_seed]

  # Expand to match batch format shape
  input_eval = tf.expand_dims(input_eval, 0)

  # Empty list to hold resulting generated text
  text_generated = []

  # Temperature effects randomness in our resulting text
  # The term is derived from entropy/thermodynamics.
  # The temperature is used to effect probability of next characters.
  # Higher probability == lesss surprising/ more expected
  # Lower temperature == more surprising / less expected
 
  temperature = temp

  # Here batch size == 1
  model.reset_states()

  for i in range(num_generate):

      # Generate Predictions
      predictions = model(input_eval)

      # Remove the batch shape dimension
      predictions = tf.squeeze(predictions, 0)

      # Use a cateogircal disitribution to select the next character
      predictions = predictions / temperature
      predicted_id = tf.random.categorical(predictions, num_samples=1)[-1,0].numpy()

      # Pass the predicted charracter for the next input
      input_eval = tf.expand_dims([predicted_id], 0)

      # Transform back to character letter
      text_generated.append(ind_to_char[predicted_id])

  return (start_seed + ''.join(text_generated))

In [85]:
print(generate_text(model,"flower",gen_size=1000))

flower ags
    to show your owns. Well, well may stir uprom with two presence
    From the sog of confusion; tesom wave,
    Plead thee belonging my head.
  ANNE. No, dear command,
    I'm death to be or danger fiery throne,
         Dismay y'T  FLUILLENCELLE. An hour, my lord. By my trown,
    Was well well? He loves me of the worm;
    Very good nurse, we but a wretched accuse may you might;
    But reason have theif forth. There dare not tell;
     The tree remember. Softly and edift eternal bed
    Of epltain'd
    Hither rotten droppier, two of fortuous gates;
    You hope better refuse, old message to your  
    Amaze your offence but for ex.
  Mer. Yet I love you fastuade you
    'Tid bite thy old kingdom shall not be.
  YORK. I'll hear a rye with the middle purity-
  VALENTINE. I have a           Exit with ARET and LARCUS

  TIMON. What Herods, what a noble honour, I am ped
    your master's blessing?
  KING RICHARD. Your presence shall I Said they adward to Boorthums.
    The 