# Imports

In [3]:
import os
import torch
import numpy as np
import pandas as pd
from tqdm import tqdm
import matplotlib.pyplot as plt
from typing import Optional
from transformers import AutoTokenizer, AutoModel, AutoModelForCausalLM, pipeline, BitsAndBytesConfig
from collections import Counter
from nltk.tokenize import word_tokenize
from nltk.corpus import stopwords

from datasets import load_dataset
from transformers import AutoModelForCausalLM, AutoTokenizer
from nltk.translate.bleu_score import sentence_bleu
import nltk
import faiss
import torch.nn.functional as F
from torch.utils.data import Dataset

import os
import torch
from torch.utils.data import DataLoader
from torch.optim import AdamW
from tqdm import tqdm

import itertools
import time
import warnings

# Chargement des données

In [None]:
from datasets import load_dataset
questions = load_dataset("rag-datasets/rag-mini-wikipedia", "question-answer")
text = load_dataset("rag-datasets/rag-mini-wikipedia", "text-corpus")

In [5]:
print(questions)
print(text)

DatasetDict({
    test: Dataset({
        features: ['question', 'answer', 'id'],
        num_rows: 918
    })
})
DatasetDict({
    passages: Dataset({
        features: ['passage', 'id'],
        num_rows: 3200
    })
})


In [15]:
from datasets import DatasetDict

questions_all = questions["test"]

# split = questions_all.train_test_split(test_size=0.2, seed=42)
# questions_train = split["train"]
# temp = split["test"].train_test_split(test_size=0.5, seed=42)
# questions_val = temp["train"]
# questions_test = temp["test"]

# questions_split = DatasetDict({
#     "train": questions_train,
#     "validation": questions_val,
#     "test": questions_test
# })

questions_val = questions_all

questions_split = DatasetDict({
    "validation": questions_val
})

In [16]:
print(questions_split)

DatasetDict({
    validation: Dataset({
        features: ['question', 'answer', 'id'],
        num_rows: 918
    })
})


## Constantes

In [17]:
DEVICE = torch.device("cuda" if torch.cuda.is_available() else "cpu")
BATCH_SIZE = 16

## Encodeur

### Chargement de l'encodeur

In [None]:
from transformers import AutoTokenizer, AutoModel
model_name = 'BAAI/bge-base-en-v1.5'
emb_tokenizer = AutoTokenizer.from_pretrained(model_name)
emb_model = AutoModel.from_pretrained(model_name)
emb_model.to(DEVICE)

### Stockage dans FAISS

In [18]:
def encode_sequences(sequences: list, tokenizer, model, device=DEVICE, batch_size=BATCH_SIZE, training=False) -> torch.Tensor:
    model.train() if training else model.eval()
    context = torch.enable_grad() if training else torch.no_grad()

    embeddings = []
    with context:
        for i in range(0, len(sequences), batch_size):
            batch = sequences[i:i + batch_size]

            # Tokenisation des textes avant entrée du modèle
            inputs = tokenizer(
                batch, padding=True, truncation=True, return_tensors="pt"
            ).to(device)

            # Passage dans le modèle encodeur
            output = model(**inputs).last_hidden_state

            # Pooling : moyenne des vecteurs de tous les tokens pour chaque séquence
            pooled = output.mean(dim=1)  # [batch_size, hidden_size]

            embeddings.append(pooled)

    # Concatène tous les embeddings pour obtenir un seul tensor de sortie
    return torch.cat(embeddings, dim=0)

In [19]:
def retrieve_top_k_passages_faiss(questions: list, vector_index: faiss.IndexFlatL2, embedding_model_tokenizer, embedding_model, k: int):
    question_embeddings = encode_sequences(questions, embedding_model_tokenizer, embedding_model)
    question_embeddings = question_embeddings.cpu().numpy()
    faiss.normalize_L2(question_embeddings)
    _, top_k_indices = vector_index.search(question_embeddings, k)
    
    return top_k_indices

In [12]:
passages = text["passages"]["passage"]
passage_embed = encode_sequences(passages, emb_tokenizer, emb_model)
passage_embeddings = passage_embed.cpu().numpy()
faiss.normalize_L2(passage_embeddings)

d = passage_embeddings.shape[1]
index = faiss.IndexFlatIP(d)
index.add(passage_embeddings)

## Generator

### Chargement du générateur

In [None]:
from huggingface_hub import login
login(token = '')

In [None]:
from transformers import AutoTokenizer, AutoModelForCausalLM
import torch

model_name = "mistralai/Mistral-7B-Instruct-v0.1"

gen_tokenizer = AutoTokenizer.from_pretrained(model_name)
gen_model = AutoModelForCausalLM.from_pretrained(
    model_name,
    device_map="auto",           # auto-placement sur GPU si dispo
    torch_dtype=torch.float16,   # allège la mémoire (fp16)
    trust_remote_code=True,       # pour activer les modèles qui utilisent du code perso
)
gen_tokenizer.pad_token = gen_tokenizer.eos_token

### Prompts

In [None]:
def generate_prompts(questions, index, emb_tokenizer, emb_model, k, passages):
    prompts = []
    faiss_passages = retrieve_top_k_passages_faiss(questions, index, emb_tokenizer, emb_model, k)
    help_instruction = "Give an answer to the question using one word if possible."

    for i, passages_ids in enumerate(faiss_passages):
        # Récupérer les textes directement via les indices
        retrieved_passages = [passages[j] for j in passages_ids]

        # Construire le prompt
        context_string = (
            f"{help_instruction}\n" +
            "\n".join(retrieved_passages) +
            "\n" +
            f"Question: {questions[i]}\n"
        )
        prompts.append(context_string)

    return prompts

In [23]:
questions_val = questions_split["validation"]["question"]
passages = text["passages"]["passage"]

k_values = [1, 2, 3, 4]
prompts_k = {}

for k in k_values:
    prompts_k[k] = generate_prompts(questions_val, index, emb_tokenizer, emb_model, k, passages)

# Affichage exemple
print(prompts_k[3][0])

Give an answer to the question using one word if possible.
Abraham Lincoln (February 12, 1809 â April 15, 1865) was the sixteenth President of the United States, serving from March 4, 1861 until his assassination. As an outspoken opponent of the expansion of slavery in the United States, "[I]n his short autobiography written for the 1860 presidential campaign, Lincoln would describe his protest in the Illinois legislature as one that 'briefly defined his position on the slavery question, and so far as it goes, it was then the same that it is now." This was in reference to the anti-expansion sentiments he had then expressed. Doris Kearns Goodwin, Team of Rivals: The Political Genius of Abraham Lincoln (2005) p. 91.  Holzer pg. 232.  Writing of the Cooper Union  speech, Holzer notes, "Cooper Union proved a unique confluence of political culture, rhetorical opportunity, technological innovation, and human genius, and it brought Abraham Lincoln to the center stage of American politics at

### Génération

In [24]:
def process_batch_prompts(batch_prompts, tokenizer, model, device, max_new_tokens=64):
    formatted_prompts = [
        f"<s>[INST] {prompt.strip()} [/INST]" for prompt in batch_prompts
    ]

    print(formatted_prompts)

    inputs = tokenizer(formatted_prompts, return_tensors="pt", padding=True, truncation=True).to(device)

    # Génération
    outputs = model.generate(
        input_ids=inputs["input_ids"],
        attention_mask=inputs["attention_mask"],
        max_new_tokens=max_new_tokens,
        do_sample=True,
        top_p=0.9,
        temperature=0.7,
        eos_token_id=tokenizer.eos_token_id
    )

    # Décodage et extraction de la réponse
    responses = [
        tokenizer.decode(output, skip_special_tokens=True).split('[/INST]')[-1].strip()
        for output in outputs
    ]

    print(responses)

    return responses

def chunked_iterable(iterable, size):
    it = iter(iterable)
    while True:
        chunk = tuple(itertools.islice(it, size))
        if not chunk:
            return
        yield chunk

def generate_responses(gen_model, gen_tokenizer, prompts, batch_size = 32) : 
    batch_responses = []
    
    for batch_index, batch_prompts in enumerate(chunked_iterable(prompts, batch_size)):
        print(f"[k={k} | Batch {batch_index}] Started generation...")
        start_time = time.time()
        
        batch_responses.extend(process_batch_prompts(batch_prompts, gen_tokenizer, gen_model, DEVICE))
        
        end_time = time.time()
        elapsed_time = end_time - start_time
        print(f"[k={k} | Batch {batch_index}] Time for generation = {elapsed_time:.2f} seconds")
    return batch_responses

In [25]:
k_values = [3]
generated_responses = {}
for k in k_values:
    generated_responses[k] = generate_responses(gen_model, gen_tokenizer, prompts_k[k])
    print(f"Done with k = {k}!")

Asking to truncate to max_length but no maximum length is provided and the model has no predefined maximum length. Default to no truncation.


[k=3 | Batch 0] Started generation...
['<s>[INST] Give an answer to the question using one word if possible.\nAbraham Lincoln (February 12, 1809 â\x80\x93 April 15, 1865) was the sixteenth President of the United States, serving from March 4, 1861 until his assassination. As an outspoken opponent of the expansion of slavery in the United States, "[I]n his short autobiography written for the 1860 presidential campaign, Lincoln would describe his protest in the Illinois legislature as one that \'briefly defined his position on the slavery question, and so far as it goes, it was then the same that it is now." This was in reference to the anti-expansion sentiments he had then expressed. Doris Kearns Goodwin, Team of Rivals: The Political Genius of Abraham Lincoln (2005) p. 91.  Holzer pg. 232.  Writing of the Cooper Union  speech, Holzer notes, "Cooper Union proved a unique confluence of political culture, rhetorical opportunity, technological innovation, and human genius, and it brought A

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.
Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Yes.', 'Yes.', 'No.', '18', '1832.', 'Paper currency.', 'Bedell', '1776.', 'No.', 'No. Abraham Lincoln was the 16th President of the United States.', 'Yes.', 'Yes.', 'Kentucky.', 'March 4, 1861.', 'John Wilkes Booth.', 'Yes.', 'McClellan', 'To free slaves in Confederate-held territories.', 'Yes.', 'Yes.', 'Yes.', 'Seven.', '1846.', '23 years.', 'The trail used by Lincoln in the William "Duff" Armstrong murder trial was the one that led to the acquittal of the accused.', 'Yes.', 'No.', 'Avogadro.', 'John Adams', '1811.', '1779', 'No.']
[k=3 | Batch 0] Time for generation = 6.36 seconds
[k=3 | Batch 1] Started generation...
["<s>[INST] Give an answer to the question using one word if possible.\nLorenzo Romano Amedeo Carlo Avogadro, Count of Quaregna and Cerreto (August 9, 1776 July 9, 1856) was an Italian savant.  He is most noted for his contributions to the theory of molarity and molecular weight. As a tribute to him, the number of elementary entities (atoms, molecules, ions or other

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Yes', 'Yes', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'Nothing.', 'Anders Celsius.', 'Tesla', 'Published.', 'No.', 'No.', 'Yes', 'Yes.', 'Yes.', 'Anders Celsius.', 'Celsius', 'Yes.', 'Yes.', 'Yes', 'Yes.', 'Nothing.', 'Nothing.', 'Yes.', 'No.', 'Yes.', 'Head, Thorax, Abdomen.', 'Mimicry.', 'Potato.', 'Head, mouthparts, spiracles.', 'Beetles eat plants, fungi, animals, and decaying organic matter.']
[k=3 | Batch 1] Time for generation = 2.84 seconds
[k=3 | Batch 2] Started generation...
['<s>[INST] Give an answer to the question using one word if possible.\nBeetles have mouthparts similar to those of grasshoppers. Of these parts, the most commonly known are probably the mandibles, which appear as large pincers on the front of some beetles. The mandibles are a pair of hard, often tooth-like structures that move horizontally to grasp, crush, or cut food or enemies (see defence, below). Two pairs of finger-like appendages are found around the mouth in most beetles, serving to move food in

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Mandibles.', 'Yes.', 'No.', 'Yes.', 'Approximately 350,000.', 'During pairing.', 'Coleopterology.', 'Yes.', 'No.', 'No.', 'Beetles.', 'Coleoptera.', 'Antennae vary greatly in form within the Coleoptera.', 'Yes.', 'Yes.', 'Oxygen is acquired.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'Plants and fungi were not mentioned in the provided information.', 'Yes.', 'Yes.', 'Yes.', 'Watt did not drop John from his name.', '1905', 'Roaring Twenties.', 'Became a lawyer.', 'Vermont.', 'Coolidge', 'No.']
[k=3 | Batch 2] Time for generation = 3.09 seconds
[k=3 | Batch 3] Started generation...
["<s>[INST] Give an answer to the question using one word if possible.\nJohn Calvin Coolidge, Jr. (July 4 1872   January 5 1933), more commonly known as Calvin Coolidge, was the thirtieth President of the United States (1923â\x80\x931929).  A lawyer from Vermont, Coolidge worked his way up the ladder of Massachusetts state politics, eventually becoming governor of that state. His actions during the Boston Pol

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Yes.', 'Yes.', 'July 4, 1872.', 'Plymouth.', 'Vermont.', 'No.', 'No.', 'Infection.', 'No.', 'No.', 'Yes.', '1898.', 'Phi Gamma Delta.', 'Grace Anna Goodhue.', 'He did not attend law school, but rather apprenticed with a local firm.', 'No.', '1900', 'No.', 'Yes.', 'No.', 'Chinese', 'Quebec City', 'British North America Act.', 'Britain', 'Significantly.', '27.7%', 'Yes.', 'Yes.', 'No.', 'Ten provinces and three territories.', '5.9%', 'Quebec City-Windsor Corridor']
[k=3 | Batch 3] Time for generation = 4.28 seconds
[k=3 | Batch 4] Started generation...
["<s>[INST] Give an answer to the question using one word if possible.\nAt  ,  World Factbook area rank order  Egypt is the world's 38th-largest country (after Mauritania). It is comparable in size to Tanzania, twice the size of France, four times the size of the United Kingdom, and is more than half the size of the US state of Alaska.\nCanada occupies a major northern portion of North America, sharing land borders with the contiguous Un

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Russia.', 'Catholics.', 'During World War II.', 'Aboriginal Peoples', '1545', 'Fur trade.', 'Yes.', 'Declining.', 'False.', 'The Maple Leaf Flag of Canada was adopted in 1965.', 'Immigration.', 'Yes.', 'Yes.', 'Yes.', 'No.', 'Yes.', 'No.', 'Yes.', 'Confederation.', 'Yes.', 'Yes.', 'No.', 'Smew.', 'Duck.', 'Farming.', 'Daffy Duck is a cartoon character.', 'Weight.', 'Food and sport.', 'Yes.', 'Yes.', 'Meat, eggs, feathers, and hunting.', 'Land ducks.']
[k=3 | Batch 4] Time for generation = 2.94 seconds
[k=3 | Batch 5] Started generation...
['<s>[INST] Give an answer to the question using one word if possible.\nDucks are sometimes confused with several types of unrelated water birds with similar forms, such as loons or divers, grebes, gallinules, and coots.\nDespite widespread misconceptions, most ducks other than female Mallards and domestic ducks do not "quack"; for example, the scaup makes a noise like "scaup", which its name came from.\nDuck is the common name for a number of speci

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Loons, divers, grebes, gallinules, coots.', 'Bird.', 'Many.', 'Push.', 'Yes.', 'Yes.', 'Yes.', 'Loons, divers, grebes, gallinules, and coots.', 'Vulnerable.', 'Duck farming, eggs, meat, feathers, zoos, hunting.', 'Lamellae.', 'Duck test.', 'Hunting.', 'Yes.', 'Yes.', 'Yes.', '1967.', 'Outlawed.', 'Approximately 100 million.', 'Arable land.', 'No.', 'Akhenaten was a Pharaoh during the Eighteenth Dynasty of the New Kingdom period (c.1550-1070 BC).', 'No.', 'Yes.', 'Yes.', '1953', 'Old Kingdom.', 'Libya, Sudan, Gaza Strip, Israel, Mediterranean Sea, Red Sea.', '78 million.', 'Yes.', 'Organized agriculture appeared in the Nile Valley around 6000 BC.', 'Yes.']
[k=3 | Batch 5] Time for generation = 4.82 seconds
[k=3 | Batch 6] Started generation...
['<s>[INST] Give an answer to the question using one word if possible.\nElephants are among the world\'s most potentially dangerous animals, capable of crushing and killing any other land animal, from humans to lions and even rhinoceros. They ca

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Yes.', 'Yes.', 'Two.', '70 years.', 'None.', 'One.', '120 kg.', 'Wise.', 'Yes.', 'No.', 'Yes.', 'Infrasound.', 'Rarely.', '22 months.', 'Two.', 'Regulation.', 'Elephant.', 'Yes.', 'Yes.', 'Yes.', 'Approximately 60,000.', 'Ears.', 'Angola.', 'Yes.', 'There is no land animal larger than an elephant.', 'No.', 'No.', 'Yes.', 'No.', 'Alexander I', '75 years.', '1952']
[k=3 | Batch 6] Time for generation = 3.54 seconds
[k=3 | Batch 7] Started generation...


Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Gulf of Finland.', 'Golden Eagle.', 'No.', 'Yes.', 'No.', 'No.', '75', 'Helsinki.', 'CMX.', 'Diversified.', 'Europe.', 'Hot.', 'No.', 'Yes.', 'Yes.', 'Twelfth century.', 'Capital.', 'Christmas.', 'No.', 'Yes.', 'None.', 'Yes.', 'Ford was a member of the House of Representatives for twenty-four years.', 'No.', 'Yes.', 'Yes.', 'Center and linebacker.', 'Abusive.', 'Agnew', 'No.', 'Intense.', 'Yes.']
[k=3 | Batch 7] Time for generation = 3.61 seconds
[k=3 | Batch 8] Started generation...
['<s>[INST] Give an answer to the question using one word if possible.\nIn November 1963, President Lyndon B. Johnson appointed Ford to the Warren Commission, a special task force set up to investigate the assassination of President John F. Kennedy. Ford was assigned to prepare a biography of Lee Harvey Oswald, the accused assassin. In 1997 the Assassination Records Review Board (ARRB) released a document that revealed that Ford had altered the first draft of the report to read: "A bullet had entered th

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['No.', 'Yes.', 'Yes.', 'Yes.', 'Fifth.', '400.', 'Berwanger.', 'No.', 'Yes.', 'No.', 'They are related through their locations in the Pacific Ocean.', 'Commission.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'No.', 'Yes.', 'Yes.', 'Yes.', 'No.', '1957', 'Ten.', 'President John Agyekum Kufuor.', 'Christianity.', 'Kwa and Gur.', 'Warm.', 'No.', 'No.', 'Yes.', 'Warrior King.', 'Kwame Nkrumah was an anti-colonial leader and the founder of modern Ghana.']
[k=3 | Batch 8] Time for generation = 5.87 seconds
[k=3 | Batch 9] Started generation...


Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Ancestral.', "Ghana's Parliament is made up of the Eduskunta (Finnish) or Riksdag (Swedish), a unicameral parliament consisting of 200 members elected for a four-year term on the basis of proportional representation through open list multi-member districts.", 'Portugal.', 'Yes.', 'Yes.', 'Yes.', 'Blended.', 'All over their bodies.', 'Wolves can weigh between 32 to 62 kg (70 to 135 pounds).', 'Communication.', 'Surplus killing is defined as the killing of several prey animals too numerous to eat at one sitting.', 'No.', 'Yes.', 'Yes.', 'Yes.', 'Kazakhstan.', 'Canis', 'Spring/Summer', 'Dominance.', 'Darting and foot hold traps.', 'Dominance.', 'No.', 'Yes.', 'No.', 'White House.', 'Heart attack.', 'One year.', '"Ma, Ma, where\'s my Pa?"', 'Blaine', 'Fear.', 'No.', 'Yes.']
[k=3 | Batch 9] Time for generation = 6.66 seconds
[k=3 | Batch 10] Started generation...
["<s>[INST] Give an answer to the question using one word if possible.\nGrover Cleveland was the second President married in of

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Yes.', '1870', 'He died on March 8, 1874.', '1884', 'Reverend Richard Cleveland and Anne Neal.', 'Five.', 'Grover Cleveland.', 'Yes.', 'Yes.', 'Yes.', 'Yes', 'Yes.', 'Yes.', 'No.', '1894', 'Discovery of radioactivity.', 'Yes.', 'Yes.', 'He did not marry Louise Désirée Lorieux.', 'Yes.', 'Yes', 'Becquerel', 'No.', 'No.', 'Election.', 'The bromide emulsion in two sheets of very thick black paper was exposed to the sun for a day and became clouded.', 'Yes.', 'Yes.', 'No.', 'Yes', '1999', '16th century.']
[k=3 | Batch 10] Time for generation = 5.32 seconds
[k=3 | Batch 11] Started generation...
["<s>[INST] Give an answer to the question using one word if possible.\nIndonesia's size, tropical climate, and archipelagic geography, support the world's second highest level of biodiversity (after Brazil),    and its flora and fauna is a mixture of Asian and Australasian species.    Once linked to the Asian mainland, the islands of the Sunda Shelf (Sumatra, Borneo, Java, Borneo, and Bali) have 

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['No.', 'No.', "The exact reason for Indonesia breaking off relations with the People's Republic of China is not specified in the provided text. However, it is mentioned that Indonesia temporarily withdrew from the UN on January 20, 1965, in response to the fact that Malaysia was elected as a non-perman", 'Nutmeg', 'City.', 'Indonesia.', 'Javanese.', 'The Srivijaya Kingdom formed trade links with China and other countries, leading to its prosperity and influence.', 'No.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'State.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', '"George"', '1.', '1825', 'Spence', 'Maria', 'Reconstruction.', 'Yes.', 'Proclamation.', 'Yes.', 'Yes.', 'Yes.', 'Westmoreland County, Virginia.']
[k=3 | Batch 11] Time for generation = 7.53 seconds
[k=3 | Batch 12] Started generation...
["<s>[INST] Give an answer to the question using one word if possible.\nThe Presidentâ\x80\x99s parents, father Spence Monroe (ca. 1727 1774), a woodworker and tobacco farmer, and mother Elizabet

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Battle of Trenton.', 'Elizabeth Kortright.', 'Before being elected governor, James Monroe served as Secretary of State.', '1816, 1820', 'John Monroe was not standing behind in the painting of Washington Crossing the Delaware.', '1814.', '1831', '1776.', '1816', 'Monroe sold his Highland Plantation in 1817.', 'December 2, 1823.', 'Daughter', 'Religious beliefs.', 'Rejection of the Jay Treaty led to increased tension between the United States and Great Britain, and ultimately contributed to the outbreak of the War of 1812.', 'Washington', 'The Democratic-Republican Party.', 'No.', 'Tours.', 'No.', 'Yes', 'Agreements.', 'Yes.', 'Woodrow Wilson', 'Publication.', 'Monroe was elected president nearly uncontested in the election of 1816.', 'Margaret Miller', 'Not James Watt or Michael Faraday.', 'Yes.', 'Yes.', 'Yes', 'Four years after opening his shop, Watt began what? Inventing.', 'No.']
[k=3 | Batch 12] Time for generation = 5.54 seconds
[k=3 | Batch 13] Started generation...
['<s>[INST]

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['No.', 'Yes.', 'Yes.', '1755.', 'Thomas Jefferson, Benjamin Franklin, Robert R. Livingston, and Roger Sherman.', '"Colossus of Independence".', 'Federalist.', 'No.', 'No.', 'Yes.', 'No.', 'Yes.', '1789-1797.', 'Federalist Party.', 'Quincy.', 'Adams.', "Adams and Hamilton had a complex relationship, and Hamilton did not openly oppose Adams's presidential campaign in 1796. However, it is true that Hamilton had a grudge against Adams and considered him to be the lesser of the two evils, which may have influenced some of his actions and statements during", "Sherman recorded his military experiences and thoughts during the American Civil War. Stukeley recorded information about Sir Isaac Newton's life.", 'John Adams, Jr.', 'John Adams', 'Quincy.', 'Yes.', 'Yes.', 'John Adams.', 'No.', 'Books.', 'Jefferson and Burr.', 'No.', 'The Stamp Act was declared invalid in 1764.', 'No.', 'Yes.', 'Yes.']
[k=3 | Batch 13] Time for generation = 8.71 seconds
[k=3 | Batch 14] Started generation...
["<s>[I

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Mob.', 'Pouch.', 'Hopping.', 'Specialized teeth.', 'Bacteria', 'Yes.', 'Yes.', 'No.', '46', '1994', 'Roo bars.', 'Yes.', 'Thylacine, Marsupial Lion, Megalania, Wonambi.', 'Yes.', 'Yes.', 'Yes.', 'Hopping.', 'Mob.', 'Marsupial.', 'Roobars.', 'Hopping.', 'Captain James Cook.', 'Kangaroos.', 'Yes.', 'No.', 'Researchers.', 'Different.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'Yes.']
[k=3 | Batch 14] Time for generation = 2.90 seconds
[k=3 | Batch 15] Started generation...
['<s>[INST] Give an answer to the question using one word if possible.\nKangaroos have large, powerful hind legs, large feet adapted for leaping, a long muscular tail for balance, and a small head. Like all marsupials, female kangaroos have a pouch called a marsupium in which joeys complete postnatal development.\nIn general, larger kangaroos have adapted much better to changes wrought to the Australian landscape by humans and though many of their smaller cousins are endangered, they are plentiful. They are not farmed to any e

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Habitat changes.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'Cheetah, jaguar.', 'Pumapard.', 'Trees.', 'Tail length.', 'Black Panthers.', 'Round.', 'Yes.', 'Yes.', 'Yes.', 'Africa, Asia.', 'Long.', '37-90 kg.', 'Yes.', 'Rosettes.', 'Unknown.', 'Leopard.', 'Felis pardus.', 'Leopards.', 'Circular.', 'Yes.', 'No answer possible with one word as "centered" is the correct word, not "centred".', "Felis pardus was a species described in Linnaeus's Systema Naturae.", 'Yes.', 'Yes.', 'No.', 'Yes.', 'Yes.']
[k=3 | Batch 15] Time for generation = 4.25 seconds
[k=3 | Batch 16] Started generation...
['<s>[INST] Give an answer to the question using one word if possible.\nA panther can be any of several species of large felid; in North America, the term refers to cougars, in South America, jaguars, and elsewhere, leopards. Early naturalists distinguished between leopards and panthers not by colour (a common misconception), but by the length of the tail panthers having longer tails than leopards.\nThe leopar

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Yes.', 'Yes.', 'Yes.', 'No.', 'Yes.', '1836.', '11', 'The national currency of Liechtenstein is the Swiss franc.', 'Yes.', "The company that administers Liechtenstein's railways is the Austrian Federal Railways.", 'Two-thirds.', 'No.', 'No.', 'Yes.', 'No.', 'October 1921.', 'German.', 'Switzerland and Austria.', 'Raetia.', 'Five.', 'Liechtenstein', '1861', 'Enjoy', 'Yes.', 'Both have a global presence and employ a large number of workers.', 'No.', 'History.', 'Constitution.', 'Preserves', 'Yes.', 'Yes.', 'Yes.']
[k=3 | Batch 16] Time for generation = 4.28 seconds
[k=3 | Batch 17] Started generation...
["<s>[INST] Give an answer to the question using one word if possible.\nThe Principality of Liechtenstein (  ( , ( ) is a tiny, doubly landlocked alpine country in Western Europe, bordered by Switzerland to its west and by Austria to its east. Mountainous, it is a winter sports resort, although it is perhaps best known as a tax haven. Despite this, it is not heavily urbanized. Many cult

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Yes.', 'Yes.', 'No.', 'These territories were held in fief under other more senior feudal lords.', 'Yes.', 'Yes.', 'Yes.', 'Weed.', 'Japan', 'Fiction.', 'Before.', 'Yes.', '8 years.', 'Yes.', 'Yes.', 'No. Millard Fillmore died on March 8, 1874.', 'Summerhill.', 'Nathaniel and Phoebe Millard Fillmore.', 'Abigail Powers.', 'Slavery.', 'Inheritance.', 'Angry.', 'Fillmore', '1846.', 'John Adams.', 'Yes.', 'Yes.', 'No.', 'Yes.', 'Yes.', 'Yes.', 'Yes.']
[k=3 | Batch 17] Time for generation = 3.22 seconds
[k=3 | Batch 18] Started generation...
['<s>[INST] Give an answer to the question using one word if possible.\nFillmore was born in a log cabin in Summerhill, New York, to Nathaniel and Phoebe Millard Fillmore, as the second of nine children and the eldest son.    Though a Unitarian in later life,   \nAfter the uniquely severe U.S. winter of 1886-1887 wiped out his herd of cattle and his $60,000 investment (together with those of his competitors), he returned to the East, where in 1885, he

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Fillmore was born in a log cabin in Summerhill, New York.', 'Nikola Tesla', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'No.', '13.', 'Fish.', 'Sea otters weigh 30 kg (about 65 pounds).', 'Air.', 'Hunting.', 'Roadkill deaths.', 'Yes.', '13 species and 7 genera.', 'Yes.', 'Sea otters live in the sea.', 'Old English.', 'Give birth.', 'Hunting', 'Culture.', 'Yes.', 'Yes.', 'No.', 'Fish.', 'Holt.', 'Poaching, habitat loss, and use of mercury in gold mining.', 'Air.', 'Playful.', 'Mustelidae', 'Aquatic beetles.', 'Otters.']
[k=3 | Batch 18] Time for generation = 4.07 seconds
[k=3 | Batch 19] Started generation...
['<s>[INST] Give an answer to the question using one word if possible.\nMost otters have fish as the primary item in their diet, supplemented by frogs, crayfish and crabs.  Some are expert at opening shellfish, and others will take any available small mammals or birds.  This prey-dependence leaves otters very vulnerable to prey depletion.\nOtters do not depend on their specialized fur alone 

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Yes.', 'Yes.', 'Yes.', 'Yes.', 'Holt.', 'Sea otters eat shellfish and other invertebrates.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'Emperor Penguin.', 'Yes.', 'Sliding.', 'Camouflage.', 'No.', 'No information provided.', '17-20.', 'Color.', 'Yes.', 'No.', 'Yes.', 'Sphenisciformes.', "Nordenskjoeld's Giant Penguin was nearly 1.80 meters (6 ft) tall.", 'Adaptation.', 'Penguins "tobaggan" to conserve energy while moving quickly.', 'Penguins.', '3 meters (9 ft)', 'Penguin wings.', "Mothers attempt to steal other mothers' chicks.", 'Yes.', 'Yes.']
[k=3 | Batch 19] Time for generation = 3.37 seconds
[k=3 | Batch 20] Started generation...
['<s>[INST] Give an answer to the question using one word if possible.\nPenguins are superbly adapted to an aquatic life. Their wings have become flippers, useless for flight in the air. In the water, however, penguins are astonishingly agile. Within the smooth plumage a layer of air is preserved, ensuring buoyancy. The air layer also helps insulate the 

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Yes.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'Cooling.', '.', 'Yes.', 'Yes.', 'Weighing.', 'Black.', 'Longer.', 'Global warming.', 'Size, fur, adaptation, and unique characteristics.', 'Unclear.', 'Yes.', 'Yes.', 'Yes.', 'Heavy.', '1002 kg.', 'Camouflage and insulation.', 'No.', 'Female polar bears gain approximately 150-250 kg (330-550 lb) during pregnancy.', 'Yes', 'Black.', 'Fur.', 'Database entry.', 'Oppose polar bear habitat protection.']
[k=3 | Batch 20] Time for generation = 4.11 seconds
[k=3 | Batch 21] Started generation...


Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['No.', 'Yes.', 'Camouflage.', 'Consumed.', 'Compete.', 'Yes.', 'Yes.', 'No.', 'Yes.', 'No.', 'Yes.', 'Yes.', '1971.', 'Doha.', 'Doha', 'Yes.', 'The Arabic word for municipalities is "baladiyah".', "Qatar University, Carnegie Mellon University, Georgetown University School of Foreign Service, Texas A&M University, Virginia Commonwealth University, Cornell University's Weill Medical College, Northwestern University.", 'The diplomatic response of the British to the Bahraini aggression on the 1820 Anglo-Bahraini Treaty.', '1971', 'Desert', 'Interconnected.', 'Yes.', 'Yes.', 'Qatari.', 'Qatar', 'Yes.', 'Yes.', 'No.', 'Yes.', 'Yes.', 'Nothing.']
[k=3 | Batch 21] Time for generation = 6.19 seconds
[k=3 | Batch 22] Started generation...
['<s>[INST] Give an answer to the question using one word if possible.\nIn terms of English phonemics, the vowels sound halfway between short u   and broad a  . The q and the t have no direct counterparts, but are closest to the unaspirated allophones of Engli

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['No.', 'No.', 'Yes.', 'Yes.', 'No.', 'Forty-one.', 'Moldoveanu Peak.', 'Romanian', 'Yes.', 'Yes.', 'Four.', 'Yes.', 'Yes', 'No.', 'Nadia ComÄ\x83neci', 'Bucharest.', 'Hungarians.', 'One.', 'Romania is a secular state.', 'Romania was not neutral during World War II because it was pressured by both Moscow and Berlin to join the Axis powers, and it was awarded southern Dobruja and part of Northern Transylvania as a result.', '2007', 'Homo sapiens sapiens.', 'Yes', 'Regulation.', 'Tourism in Romania attracted 400 million euros in investments in 2005.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'Yes.']
[k=3 | Batch 22] Time for generation = 8.51 seconds
[k=3 | Batch 23] Started generation...


Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['No.', "No. The media industry contributed 1.56% to Singapore's GDP in 2001.", 'More than 2.78 million people use the bus network daily.', 'Orang Laut people lived around the coast, rivers, and smaller islands.', 'Mass Rapid Transit.', 'Tropical.', '3%', '6', 'No.', 'Yes', 'No.', 'Lion.', '2nd century AD.', 'Japan.', 'Churchill called the occupation of Singapore "Britain\'s greatest defeat" because it was a significant loss for the British Empire and a symbol of their decline during World War II.', 'Yes.', '1990.', 'Rainfall.', 'English', 'Singapore', 'Geography.', 'Not provided.', 'British East India Company', '1959-1965', 'Capital punishment.', 'The British East India Company established a trading post on the island in 1819.', 'The British East India Company.', 'Singapore', 'Montevideo', 'Reservoirs are used to store water.', 'Yes.', 'Yes.']
[k=3 | Batch 23] Time for generation = 5.10 seconds
[k=3 | Batch 24] Started generation...
['<s>[INST] Give an answer to the question using one

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Yes.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'Gramercy.', 'Exercise.', 'Bullies.', 'Bamie', '"Personal experience."', 'Newspapers responded with praise.', 'No.', 'No.', 'Yes.', 'New York', 'Lincoln.', 'Roosevelt died on January 6, 1919.', 'Carnegie', 'No.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'Spherical.', 'Infrequently.', 'Air.', 'Reproduce.', 'Researchers have discovered that turtle organs do not break down or become less efficient over time.', 'Pleurodira']
[k=3 | Batch 24] Time for generation = 5.86 seconds
[k=3 | Batch 25] Started generation...
['<s>[INST] Give an answer to the question using one word if possible.\nSea turtles lay their eggs on dry sandy beaches. Immature sea turtles are not cared for by the adults. Most are endangered largely as a result of beach development and over hunting.\nTurtles lay eggs, like other reptiles, which are slightly soft and leathery. The eggs of the largest species are spherical, while the eggs of the rest are elongated. Their a

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Yes.', 'Yes.', 'No.', 'Three.', 'Russian Tortoise.', '1. Encyclopedia of Turtles by Pritchard, Pether C H (1979).\n2. Turtles of the World by Pritchard, Pether C H (1979).\n3. Turtles and Crocodiles of Ins', 'Female.', 'Papillae.', 'Neck.', 'No.', 'Yes', 'Yes.', 'Paracryptodira.', 'Yolk.', 'Beaches.', 'Pleurodira.', 'Turtle eggs are covered in sand when they incubate.', 'Reptiles.', 'Humus.', 'All turtles are ectothermic.', 'Yes.', 'Yes.', 'Papillae have various shapes and sizes.', 'Spherical.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'Yes.']
[k=3 | Batch 25] Time for generation = 5.43 seconds
[k=3 | Batch 26] Started generation...
["<s>[INST] Give an answer to the question using one word if possible.\nOn August 22, 1848, Grant married Julia Boggs Dent (1826 1902), the daughter of a slave owner. They had four children: Frederick Dent Grant, Ulysses S. Grant, Jr. (Buck), Ellen Wrenshall Grant (Nellie), and Jesse Root Grant.\nPresident Grant with his wife, Julia, and so

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Julia', 'Point Pleasant, Ohio.', 'Republican.', 'It is not clear from the given information what Grant said in response to the question.', 'Yes.', 'Grant thought the Mexican-American War was unjust because it was designed to gain land open to slavery.', 'Ulysses S. Grant.', '1868.', 'Grant', '1856.', 'Grant', 'Yes.', 'The question is asking about the origin of the phrase "Ma, Ma, where\'s my Pa?" and it was not from Pennsylvania.', 'No.', 'Yes.', 'No.', 'No.', 'Yes.', 'Lincoln delivered his second inaugural address.', 'Crisis.', 'No.', 'Yes.', 'No.', 'Spanish', 'San Carlos.', 'Europeans', 'Cantegriles', 'No.', 'Catholic.', 'Yes', 'Yes.', 'Yes.']
[k=3 | Batch 26] Time for generation = 4.25 seconds
[k=3 | Batch 27] Started generation...
["<s>[INST] Give an answer to the question using one word if possible.\nAt 176,214 square kilometres (68,036 square miles) of continental land and 142,199 square kilometres (54,903 sq mi) of jurisdictional waters and small river islands,  Instituto Naci

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['176,214', '88%', 'Military stronghold.', 'Two.', '65%', 'Yes.', 'No.', 'Yes.', 'Plebiscite.', 'Catholicism, Protestantism, Judaism.', 'Uruguay has been competitive in agriculture by focusing on ecological and natural products, and by developing a strong export-oriented agricultural sector.', 'Partido Colorado, Partido Blanco, Frente Amplio.', 'No.', 'Montevideo', 'Meanings.', 'Import/export.', 'Independence.', 'No.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'No.', 'Yes.', 'No.', 'Paris.', 'December 28, 1856.', 'Wilson was President when Franklin D. Roosevelt was elected.', 'History and political science.']
[k=3 | Batch 27] Time for generation = 4.01 seconds
[k=3 | Batch 28] Started generation...
['<s>[INST] Give an answer to the question using one word if possible.\nWilson spent 1914 through the beginning of 1917 trying to keep America out of the war in Europe. He offered to be a mediator, but neither the Allies nor the Central Powers took his requests seriously. Repub

## Evaluation

In [28]:
questions_val = questions_split["validation"].to_pandas()
df_export = questions_val[['id', 'question', 'answer']].copy()
df_export['generated_answer'] = generated_responses[k]
df_export.to_csv(f"generated_responses_simple.csv", index=False, encoding='utf-8')