# Imports

In [1]:
!pip install -q \
    torch \
    torchvision \
    torchaudio \
    transformers \
    datasets \
    nltk \
    faiss-cpu \
    matplotlib \
    tqdm \
    numpy \
    pandas \
    accelerate


[1m[[0m[34;49mnotice[0m[1;39;49m][0m[39;49m A new release of pip is available: [0m[31;49m24.3.1[0m[39;49m -> [0m[32;49m25.0.1[0m
[1m[[0m[34;49mnotice[0m[1;39;49m][0m[39;49m To update, run: [0m[32;49mpip install --upgrade pip[0m


In [2]:
import torch
torch.cuda.is_available()

True

In [3]:
import os
import torch
import numpy as np
import pandas as pd
from tqdm import tqdm
import matplotlib.pyplot as plt
from typing import Optional
from transformers import AutoTokenizer, AutoModel, AutoModelForCausalLM, pipeline, BitsAndBytesConfig
from collections import Counter
from nltk.tokenize import word_tokenize
from nltk.corpus import stopwords

from datasets import load_dataset
from transformers import AutoModelForCausalLM, AutoTokenizer
from nltk.translate.bleu_score import sentence_bleu
import nltk
import faiss
import torch.nn.functional as F
from torch.utils.data import Dataset

import os
import torch
from torch.utils.data import DataLoader
from torch.optim import AdamW
from tqdm import tqdm

import itertools
import time
import warnings

# Chargement des données

In [4]:
from datasets import load_dataset
questions = load_dataset("rag-datasets/rag-mini-wikipedia", "question-answer")
text = load_dataset("rag-datasets/rag-mini-wikipedia", "text-corpus")

In [5]:
print(questions)
print(text)

DatasetDict({
    test: Dataset({
        features: ['question', 'answer', 'id'],
        num_rows: 918
    })
})
DatasetDict({
    passages: Dataset({
        features: ['passage', 'id'],
        num_rows: 3200
    })
})


In [6]:
from datasets import DatasetDict

questions_all = questions["test"]

# split = questions_all.train_test_split(test_size=0.2, seed=42)
# questions_train = split["train"]
# temp = split["test"].train_test_split(test_size=0.5, seed=42)
# questions_val = temp["train"]
# questions_test = temp["test"]

# questions_split = DatasetDict({
#     "train": questions_train,
#     "validation": questions_val,
#     "test": questions_test
# })

questions_val = questions_all
questions_split = DatasetDict({
    "validation": questions_val
})

In [7]:
print(questions_split)

DatasetDict({
    validation: Dataset({
        features: ['question', 'answer', 'id'],
        num_rows: 918
    })
})


## Constantes

In [8]:
DEVICE = torch.device("cuda" if torch.cuda.is_available() else "cpu")
BATCH_SIZE = 16

## Generator

### Chargement du générateur

In [9]:
from huggingface_hub import login
login(token = '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')

In [11]:
from transformers import AutoTokenizer, AutoModelForCausalLM
import torch

model_name = "mistralai/Mistral-7B-Instruct-v0.1"

gen_tokenizer = AutoTokenizer.from_pretrained(model_name)
gen_model = AutoModelForCausalLM.from_pretrained(
    model_name,
    device_map="auto",           # auto-placement sur GPU si dispo
    torch_dtype=torch.float16,   # allège la mémoire (fp16)
    trust_remote_code=True,       # pour activer les modèles qui utilisent du code perso
)
gen_tokenizer.pad_token = gen_tokenizer.eos_token

Loading checkpoint shards:   0%|          | 0/2 [00:00<?, ?it/s]

### Prompts

In [12]:
def generate_prompts(questions):
    prompts = []
    help_instruction = "Give an answer to the question. Use one word if possible."

    for i, question in enumerate(questions):
        context_string = (
            f"{help_instruction}\n" +
            "\n" +
            f"Question: {question}\n"
        )
        prompts.append(context_string)

    return prompts

In [13]:
questions_val = questions_split["validation"]["question"]
passages = text["passages"]["passage"]

prompts = generate_prompts(questions_val)
print(prompts[0])

Give an answer to the question. Use one word if possible.

Question: Was Abraham Lincoln the sixteenth President of the United States?



### Génération

In [14]:
def process_batch_prompts(batch_prompts, tokenizer, model, device, max_new_tokens=64):
    formatted_prompts = [
        f"<s>[INST] {prompt.strip()} [/INST]" for prompt in batch_prompts
    ]

    print(formatted_prompts)

    inputs = tokenizer(formatted_prompts, return_tensors="pt", padding=True, truncation=True).to(device)

    # Génération
    outputs = model.generate(
        input_ids=inputs["input_ids"],
        attention_mask=inputs["attention_mask"],
        max_new_tokens=max_new_tokens,
        do_sample=True,
        top_p=0.9,
        temperature=0.7,
        eos_token_id=tokenizer.eos_token_id
    )

    # Décodage et extraction de la réponse
    responses = [
        tokenizer.decode(output, skip_special_tokens=True).split('[/INST]')[-1].strip()
        for output in outputs
    ]

    print(responses)

    return responses

def chunked_iterable(iterable, size):
    it = iter(iterable)
    while True:
        chunk = tuple(itertools.islice(it, size))
        if not chunk:
            return
        yield chunk

def generate_responses(gen_model, gen_tokenizer, prompts, batch_size = 32) : 
    batch_responses = []
    
    for batch_index, batch_prompts in enumerate(chunked_iterable(prompts, batch_size)):
        print(f"[Batch {batch_index}] Started generation...")
        start_time = time.time()
        
        batch_responses.extend(process_batch_prompts(batch_prompts, gen_tokenizer, gen_model, DEVICE))
        
        end_time = time.time()
        elapsed_time = end_time - start_time
        print(f"[Batch {batch_index}] Time for generation = {elapsed_time:.2f} seconds")
    return batch_responses

In [15]:
generated_responses = generate_responses(gen_model, gen_tokenizer, prompts)

Asking to truncate to max_length but no maximum length is provided and the model has no predefined maximum length. Default to no truncation.
Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


[Batch 0] Started generation...
['<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Was Abraham Lincoln the sixteenth President of the United States? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Did Lincoln sign the National Banking Act of 1863? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Did his mother die of pneumonia? [/INST]', "<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: How many long was Lincoln's formal education? [/INST]", '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: When did Lincoln begin his political career? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: What did The Legal Tender Act of 1862 establish? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Who suggested Lincoln grow a beard? [/INST]

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Yes.', 'Yes.', 'Yes.', 'Long.', '1832.', 'Legal tender.', 'No one.', '1863', 'Yes.', 'No.', 'Yes.', 'No.', 'Hodgenville.', 'Elected.', 'John Wilkes Booth.', 'Yes.', 'McClellan', 'Freedom.', 'Yes.', 'No.', 'Yes.', '16', '1861', '22', 'Road', 'No.', 'No.', 'Dalton.', 'Thomas.', 'Never.', 'Unknown.', 'No.']
[Batch 0] Time for generation = 1.08 seconds
[Batch 1] Started generation...
['<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Was Lorenzo Romano Amedeo Carlo Avogadro an Italian savant? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Was Amedeo Avogadro born in Turin August 9th 1776 to a noble ancient family of Piedmont, Italy? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Is he most noted for his contributions to the theory of molarity and molecular weight? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Was King Vict

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Yes.', 'No.', 'Yes.', 'No.', 'Yes.', 'No.', 'Yes.', 'Revolution.', 'Boyle.', 'Legacy.', 'Unknown.', 'No.', 'No.', 'No.', 'Yes.', 'No.', 'Kelvin.', 'Impact.', 'No.', 'No.', 'No.', 'Agree.', 'Nothing.', 'Revolution.', 'Yes.', 'No.', 'Yes.', 'Body, wings, antennae.', 'Mimicry.', 'Potato.', 'Color.', 'Leaves.']
[Batch 1] Time for generation = 0.41 seconds
[Batch 2] Started generation...
['<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: What are the similarities between beetles and grasshoppers? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Are certain species of beetles considered pests? [/INST]', "<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Is a beetle's general anatomy uniform? [/INST]", '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Are beetles endopterygotes? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Insects.', 'Yes.', 'No.', 'No.', 'Thousands.', 'Insemination.', 'Entomology.', 'Yes.', 'No.', 'Yes.', 'Food.', 'Aristotle described the hardened shield as "armor."', 'Beetles.', 'Yes.', 'Yes.', 'Essential.', 'Yes.', 'Yes.', 'Yes.', 'No.', 'No.', 'Nothing.', 'Yes.', 'Yes.', 'Yes.', 'Never.', '1909', 'Roaring Twenties.', 'Lawyer', 'Massachusetts', 'Franklin D. Roosevelt.', 'Yes.']
[Batch 2] Time for generation = 0.57 seconds
[Batch 3] Started generation...
['<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Was Calvin Coolidge Republican? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Was Calvin Coolidge a governor of MassachuS08_setts? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: When was Coolidge born? [/INST]', "<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Where did Coolidge's grandfather had government offices? [/INST]", '<s>

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Yes.', 'Yes.', 'Calvin.', 'Massachusetts', 'Vermont', 'No.', 'Yes.', 'Tragedy.', 'No.', 'Yes.', 'Yes.', '1906', 'Calvin', 'Prudence.', 'Answer: Harvard', 'Yes.', '1910, 1912', 'No.', 'Yes.', 'No.', 'Mandarin.', 'Quebec', 'Bill of Rights', 'Britain', 'Contributed.', 'Majority.', 'Yes.', 'Yes.', 'No.', 'Ten.', 'Unemployment.', 'Toronto.']
[Batch 3] Time for generation = 0.51 seconds
[Batch 4] Started generation...
['<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: What is the largest country in the world? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: What is the largest religious group in Canada? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: When did Canada have one of the largest armed forces in the world? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: What have been inhabited for millennia by aboriginal peoples? [/INS

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Russia.', 'Catholic.', 'Never.', 'Australia', '1555', 'Fur.', 'Yes.', 'Declining.', 'True.', 'The flag.', 'Invasion.', 'Yes.', 'Yes.', 'Yes.', 'No.', 'No.', 'No.', 'No.', 'Year.', 'Yes.', 'Yes.', 'Yes.', 'Muscovy.', 'Monkey', 'Duckling.', 'Character.', 'Water.', 'Sport.', 'Yes.', 'Yes.', 'Ducks have none.', 'Wood.']
[Batch 4] Time for generation = 0.33 seconds
[Batch 5] Started generation...
['<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: What unrelated water birds are ducks sometimes confused with? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: What does the word duck mean? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: What are some common predators of ducks? [/INST]', "<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: How does a Mallard's tongue work? [/INST]", '<s>[INST] Give an answer to the question. Use one word if possible

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Geese.', 'Duck.', 'Alligators.', 'Taste.', 'No.', 'No.', 'Yes.', 'Geese.', 'Hatchling.', 'Quack.', 'Filtration.', 'Testing.', 'Origin.', 'Yes.', 'No.', 'Yes.', '1967', 'None.', 'Millions.', 'Water.', 'No.', 'Amarna.', 'No.', 'No.', 'No.', '1952.', 'Ancient.', 'Sudan.', 'Millions.', 'No.', 'Ancient.', 'Yes.']
[Batch 5] Time for generation = 0.32 seconds
[Batch 6] Started generation...
['<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Can an elephant kill a rhinoceros? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Are elephants good swimmers? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: How many living species of African Elephants are there? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: How long may elephants live? [/INST]', "<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Why does th

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Yes.', 'No.', 'Two.', 'Long.', 'Forget.', 'Two.', 'Massive.', 'Unknown.', 'Fragmented.', 'Yes.', 'Yes.', 'Radio.', 'Never.', 'Long.', 'Two.', 'Hearing.', 'Elephant.', 'Yes.', 'Yes.', 'Yes.', 'Unknown.', 'Ears.', 'Botswana.', 'No.', 'None.', 'No.', 'No.', 'Yes.', 'No.', 'Alexander.', '80', '1920']
[Batch 6] Time for generation = 0.29 seconds
[Batch 7] Started generation...
['<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: What body of water lies to the south of Finland? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Name an animal that is growing in number due to recent conservation efforts [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Today, are there more cell phone subscriptions than people in Finland? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Is Matti Vanhanen the Prime Minister of Finland? [/INST]', '<s>[INST

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Sea.', 'Giant Panda.', 'Yes.', 'No.', 'No.', 'Yes.', 'Variable.', 'Helsinki.', 'Helsinki.', 'Strong.', 'North.', 'Sunny.', 'Yes.', 'Yes.', 'Yes.', 'Cuneiform.', 'London', 'Christmas.', 'No.', 'No.', 'Russia.', 'No.', 'Ford served in the House for 14 years.', 'Yes.', 'Yes.', 'No.', 'Ford was a player on the school football team.', 'Secret.', 'Spiro.', 'No.', 'Difficult.', 'No.']
[Batch 7] Time for generation = 0.48 seconds
[Batch 8] Started generation...
['<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Is Ford related with the assassination of John F. Kennedy? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Was Gerald Ford the 38th President of the United States? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Did Gerald Ford serve as the Republican Minority Leader of the House of Representatives? [/INST]', '<s>[INST] Give an answer to the question. Use one word if

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['No.', 'Yes.', 'No.', 'No.', 'Dearborn.', 'Unknown.', 'Smith', 'Yes.', 'No.', 'No.', 'Islands.', 'Birthday.', 'Yes.', 'No.', 'No.', 'No.', 'No.', 'Yes.', 'No.', 'No.', 'Yes', '1957.', '10.', 'President.', 'Christianity.', 'Akan and Ewe.', 'Sunny.', 'No.', 'No.', 'No.', 'Meaning.', 'Pan-Africanist.']
[Batch 8] Time for generation = 0.38 seconds
[Batch 9] Started generation...
['<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Where is old Ghana in relation to present Ghana? [/INST]', "<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Who makes up Ghana's Parliament? [/INST]", '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: What European countries established states in Ghana? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Does the Gray Wolf share an ancestry with the domestic dog? [/INST]', '<s>[INST] Give an answer to the question. Use one word i

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Past.', 'Members.', 'Colonizers.', 'Yes.', 'No.', 'No.', 'Fur.', 'Face.', 'Varies.', 'Communication.', 'Excessive.', 'Yes.', 'Yes.', 'No.', 'Endurance.', 'Canada.', 'Canis.', 'Wolves molt annually.', 'Territorial.', 'Capture.', 'Alertness', 'No.', 'Yes.', 'No.', 'Married.', 'Heart attack.', '18', 'Corrupt.', 'Benjamin Harrison.', 'Privacy.', 'Yes.', 'No.']
[Batch 9] Time for generation = 0.32 seconds
[Batch 10] Started generation...
['<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Was Grover Cleveland married in the whitehouse? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: When was he elected sheriff of Erire County, New York? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: When did he die? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Which election did Grover Cleveland win? [/INST]', "<s>[INST] Give an answer to th

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['No.', 'Elected.', 'Deceased.', 'Grover Cleveland won the presidential election of 1884.', 'Neither.', 'None.', 'William McKinley.', 'No.', 'No.', 'Yes.', 'No.', 'No.', 'Yes.', 'Yes.', '1995.', 'Curiosity', 'False', 'No.', 'Never.', 'No.', 'Yes.', 'Curie', 'No.', 'Yes.', 'Edison.', 'Nothing.', 'No.', 'No.', 'Yes.', 'No.', '1999.', 'Islamicization.']
[Batch 10] Time for generation = 0.63 seconds
[Batch 11] Started generation...
["<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Does Indonesia have the world's hightest level of biodiversity? [/INST]", '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Was Indonesia named after it became an independent country? [/INST]', "<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Why did Indonesia break off relations with the People's Republic of China? [/INST]", '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Which spic

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['No.', 'No.', 'Balance', 'Cloves.', 'City.', 'Borders.', 'Majority.', 'Prosperity.', 'Yes.', 'Yes.', 'No.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'State.', 'Yes.', 'Yes.', 'Yes.', 'No.', 'Alphabet.', 'One.', '1825.', 'Monkeypaw', 'Marilyn', 'Epoch.', 'Yes.', 'Interventionism.', 'Yes.', 'No.', 'No.', 'Monroe']
[Batch 11] Time for generation = 0.33 seconds
[Batch 12] Started generation...
['<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Where was James Monroe shot? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Who did James Monroe marry? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: What did James Monroe do before being elected governor? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: In which years were John Monroe elected as President? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestio

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Monroe.', 'Monroe.', 'Lawyer', 'Elected.', 'Washington', 'Monroe.', 'Monroe died in 1838.', 'Never.', 'Elected.', 'Home', '1823', 'Monroe.', 'Secrets.', 'War.', 'Washington.', 'Compromise', 'No.', 'Constitution', 'No.', 'Yes.', 'Nothing.', 'Yes.', 'Unity.', 'War.', 'Result.', 'Marie', 'Andrew Jackson.', 'No.', 'Yes.', 'Yes.', 'Expanded.', 'No.']
[Batch 12] Time for generation = 0.46 seconds
[Batch 13] Started generation...
['<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Did John Adams represent the Continental Congress in Europe? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Was Adams raised Congregationalist? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Was Adams an opponent of the Stamp Act? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: When did Adams graduate from college? [/INST]', '<s>[INST] Give an answer t

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['No.', 'No.', 'Yes.', 'Never.', 'Jefferson', 'Friend.', 'Republican.', 'No.', 'No.', 'Yes.', 'No.', 'No.', '1789-1797', 'Federalist.', 'Dead.', 'None.', 'Ideologically.', 'Details.', 'John Adams', 'Thomas Jefferson', 'Independence Hall', 'False', 'No.', 'Tea.', 'No.', 'Books.', 'Democrats.', 'Yes.', 'Revolution.', 'No.', 'No.', 'No.']
[Batch 13] Time for generation = 0.48 seconds
[Batch 14] Started generation...
['<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: What is a collective noun for kangaroos? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Where do joeys complete postnatal development? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: What do kangaroos use for "crawl-walking?" [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Why do kangaroos have a wide bite? [/INST]', '<s>[INST] Give an answer to the question. Use o

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Mob.', 'Joey.', 'Paws.', 'Power.', 'Acetogen.', 'No.', 'Thrived.', 'Few.', '12', 'Unknown', 'Four-wheel drive.', 'Yes.', 'Tasmanian.', 'Yes.', 'Yes.', 'No.', 'Hopping.', 'Mob', 'Roof.', 'Steel.', 'Hopping.', 'Curious', 'Kangaroos.', 'No.', 'No.', 'Few.', 'Various.', 'Yes.', 'Yes.', 'No.', 'No.', 'Yes.']
[Batch 14] Time for generation = 0.34 seconds
[Batch 15] Started generation...
['<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: What kinds of changes have larger kangaroos adapted much better to? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Have larger kangaroos adapted much better to changes? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Is the leopard smaller than the other members of Panthera? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Is a leopard larger and less lanky than a cheetah? [/INST]', '<s>[INST] Gi

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Adaptation.', 'Adapted.', 'No.', 'No.', 'Yes.', 'Jaguars.', 'Pum Leopard', 'Trees.', 'Spots.', 'Panthers.', 'Spotted.', 'Yes.', 'Yes.', 'Yes.', 'Africa.', 'Long.', 'Heavy.', 'Yes.', 'Spots.', 'Smaller.', 'Puma.', 'Classification.', 'Independent.', 'No.', 'Yes.', 'Mountain', 'Leopard.', 'Yes.', 'No.', 'Yes.', 'No.', 'Yes.']
[Batch 15] Time for generation = 0.36 seconds
[Batch 16] Started generation...
['<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Can a panther be any of several species of large felid? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Is the leopard an agile and graceful predator? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Is Liechtenstein bordered by Switzerland? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Is Liechtenstein the smallest German-speaking country in the world? [/INST]', '<s>[INST] G

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Yes.', 'Yes.', 'Yes.', 'Yes.', 'No.', '1764.', 'Ten.', 'Euro', 'Yes.', 'PrivatBahn.', '17.5', 'No.', 'No.', 'Yes.', 'No.', 'Adopted.', 'German.', 'Switzerland.', 'Rhine.', 'Numerous.', 'Vatican.', '1793', 'Switzerland', 'Yes.', 'Brands.', 'No.', 'Month.', 'Jurisdiction.', 'Music.', 'No.', 'Yes.', 'No.']
[Batch 16] Time for generation = 0.33 seconds
[Batch 17] Started generation...
['<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Is it a winter sports resort , although it is perhaps best known as a tax haven ? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Is it a winter sports resort? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Is it the smallest german-speaking country in the world? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: What happened in moravia , lower austria , silesia , and styria , though in all cases 

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['No.', 'No.', 'No.', 'Feudalism.', 'No.', 'Yes.', 'Yes.', 'Lincoln', 'Japan.', 'President.', 'After.', 'Yes.', 'Immediately.', 'Yes.', 'No.', 'No.', 'New York', 'Parents.', 'Love.', 'Different', 'Elected.', 'Opposed', 'First.', '1871.', 'Abigail.', 'No.', 'Yes.', 'No.', 'Yes.', 'No.', 'No.', 'No.']
[Batch 17] Time for generation = 0.38 seconds
[Batch 18] Started generation...
['<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: What happened in a log cabin in summerhill , new york? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Who was born precisely at midnight during an electrical storm , to a Serbian family in the village of Smiljan near Gospić , in the Lika region of the Croatian Krajina in Military Frontier ( part of the Austrian Empire ) , in the present-day Croatia . ? [/INST]', "<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Are there at least two films describing Tesla

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Fire.', 'Nikola Tesla', 'Yes.', 'No.', 'No.', 'Yes.', 'No.', 'Four.', 'Fish.', 'Otters.', 'Fur.', 'Habitat.', 'Habitat.', 'Yes.', 'Species: 13. Genera: 2.', 'Yes.', 'Coast.', 'River.', 'Give.', 'Predation.', 'Symbolism.', 'Yes.', 'No.', 'No.', 'Fish.', 'A otter\'s den is called a "lodge."', 'Habitat.', 'Fur.', 'Playful.', 'Otters are related to weasels.', 'Bubbles.', 'People.']
[Batch 18] Time for generation = 0.67 seconds
[Batch 19] Started generation...
['<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Is it true that otters eat a variety of fish? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Are otters playful animals? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Are otters very active? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Are male otters dog-otters? [/INST]', "<s>[INST] Give an answer to the question. 

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Yes.', 'Yes.', 'Yes.', 'No.', 'Burrow.', 'Fish.', 'No.', 'No.', 'Yes.', 'No.', 'Yes.', 'Emperor.', 'Yes.', 'Sledding.', 'Camouflage.', 'Fearful.', 'Little.', 'Six.', 'Slender.', 'Yes.', 'No.', 'No.', 'Penguins belong to the order Sphenisciformes.', 'Gigantic.', 'Adaptation.', 'Waddle.', 'Seahorses.', 'Distance.', 'Wings.', 'Chick dies.', 'No.', 'No.']
[Batch 19] Time for generation = 0.57 seconds
[Batch 20] Started generation...
['<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Are penguins astonishingly agile? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Are penguins at risk? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Is it also possible that penguin comes from the Latin pinguis , fat ? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Have penguins an average sense of hearing for birds ? [/INST]', '<s>[INST] Give 

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['No.', 'Endangered.', 'No.', 'No.', 'Yes.', 'Yes', 'No.', 'No.', 'Yes.', 'Volcanic activity.', '"End."', 'Yes.', 'Yes.', 'Record.', 'White.', 'Long.', 'Pollution.', 'Unique.', 'Convergence.', 'Yes.', 'Yes.', 'Yes.', 'Massive.', 'Massive.', 'Warmth.', 'No.', 'Pounds.', 'Yes.', 'Coal.', 'Blanket.', 'Justification.', 'Environmentalism.']
[Batch 20] Time for generation = 0.42 seconds
[Batch 21] Started generation...
['<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Are studies insufficient evidence for global protection? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Are polar bears excellent swimmers? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: A semi-aquatic marine mammal , the polar bear has what? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Garbage is what? [/INST]', '<s>[INST] Give an answer to the question. Use o

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['No.', 'Yes.', 'Fur.', 'Trash.', 'Eat.', 'No.', 'Yes.', 'Yes.', 'No.', 'Yes.', 'Yes.', 'No.', '1971.', 'Doha.', 'Baghdad', 'No.', 'Baladiyah', 'Universities.', 'Recognition.', '1971', 'Desert.', 'Connected.', 'Yes.', 'Yes.', 'Jordan', 'Kuwait', 'Yes.', 'Yes.', 'No.', 'Yes.', 'Yes.', 'Nothing.']
[Batch 21] Time for generation = 0.40 seconds
[Batch 22] Started generation...
['<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Can these allophones not occur in these positions in english? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Is -RRB- , officially the state of qatar -LRB- arabic (: : transliterated as dawlat qatar -RRB- , an arab emirate in southwest asia? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Does Romania border Hungary? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Is Romania a secular state? [/INST]', '<s

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['No.', 'Yes.', 'Yes.', 'No.', 'No.', 'Counties.', 'Carpathians.', 'Romanian.', 'Thriving.', 'Yes.', 'Seven.', 'Yes.', 'No.', 'No.', 'Olga Korbut.', 'Bucharest.', 'Hungarians.', 'Few.', 'Secular', 'Fascism', '2017.', 'Mammoths.', 'No.', 'Regulation.', 'Tourists.', 'Yes.', 'Yes.', 'No.', 'No.', 'Yes.', 'Yes.', 'Yes.']
[Batch 22] Time for generation = 0.40 seconds
[Batch 23] Started generation...
["<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Did Singapore's standard of living decrease since independence? [/INST]", "<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Did the media contribute 15.6% to Singapore's GDP in 2001? [/INST]", '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: How many people use the bus network daily? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Where did Orang Laut people live? [/INST]', '<s>[INST] Give an answer to the 

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['No.', 'No.', 'Unknown.', 'Orang Laut people lived in Borneo.', 'Mass Rapid Transit.', 'Tropical.', 'Increase.', '3', 'No.', 'Yes.', 'No.', 'Singa means "sign language."', 'Never.', 'Japan', 'Defeat.', 'Yes.', 'Succeeded.', 'Tap.', 'English.', 'Singapore.', 'Vital.', 'Monaco', 'British.', '1965-1990', 'Death.', 'Fur.', 'Chinese.', 'London', 'Denver', 'Contain.', 'Disputes.', 'Yes.']
[Batch 23] Time for generation = 0.48 seconds
[Batch 24] Started generation...
['<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Has Singapore a network of reservoirs and water catchment areas ? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Has Urbanization eliminated many areas of former primary rainforest , with the only remaining area of primary rainforest being Bukit Timah Nature Reserve ? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Does the economy not depend heavily on export

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Yes.', 'No.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'No.', 'No.', 'New York', 'Exercised.', 'Roosevelt started boxing lessons for fitness.', 'Daisy', 'Equality.', 'Euphoric.', 'Yes.', 'Yes.', 'Yes.', 'New York', 'Theodore Roosevelt.', '1945.', 'Roosevelt.', 'Yes.', 'Affluent.', 'Yes.', 'Yes.', 'No.', 'Oval.', 'Infrequently.', 'Air.', 'Eggs.', 'Discovered.', 'Draws.']
[Batch 24] Time for generation = 0.47 seconds
[Batch 25] Started generation...
['<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Do sea turtles lay eggs on dry sandy beaches? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Are Pleurodira known as side-necked turtles? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Are turtles a part of the Flying Spaghetti Monster? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: How many suborders are turtles divided into? [/INST]'

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['No.', 'Yes.', 'No.', 'Suborders.', 'Sheldon.', 'Turtles, Habitat, Diet, Conservation.', 'Turtle.', 'Gills.', 'Neck', 'No.', 'No.', 'Yes.', 'Extinct.', 'Protein.', 'Beaches.', 'Testudines.', 'Shells.', 'Mammals.', 'Footprints.', 'All.', 'No.', 'Yes.', 'Papillae have teeth.', 'Ostrich.', 'Inspired.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'Yes.', 'No.', 'No.']
[Batch 25] Time for generation = 0.57 seconds
[Batch 26] Started generation...
["<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Who was Grant's brother in law? [/INST]", '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Where was Grant born? [/INST]', "<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: What was Grant's political affiliation? [/INST]", '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Why did Grant say "Damn, I had nothing to do with this batte."? [/INST]', '<s>[INST] Give an answer to the quest

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['Julia', 'Ohio', 'Republican.', 'Innocence.', 'Yes.', 'Unjust.', 'Grant', '1969.', 'Law enforcement.', 'Never.', 'Ulysses', 'No.', 'Liberty.', 'No.', 'No.', 'No.', 'No.', 'Yes.', 'Civil War ended.', 'Developments.', 'No.', 'No.', 'No.', 'Founded.', 'Colonia.', 'Uruguayan', 'Ghettos.', 'No.', 'Catholic.', 'Yes.', 'Yes.', 'No.']
[Batch 26] Time for generation = 0.35 seconds
[Batch 27] Started generation...
['<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: How many square kilometres of continental land is Uruguay? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: How much of the population is of white European descent? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Why was the capital of Uruguay founded? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: How many times has Uruguay won the World Cup? [/INST]', '<s>[INST] Give an a

Setting `pad_token_id` to `eos_token_id`:None for open-end generation.


['188.000', 'Unknown.', 'Colonia', 'Twice.', 'Less.', 'Yes.', 'No.', 'No.', 'Petition.', 'Christianity.', 'Diversification.', 'Parties.', 'No.', 'Fort', 'Ambiguous', 'Culturally.', 'Independence.', 'No.', 'Yes.', 'Yes.', 'No.', 'Yes.', 'No.', 'No.', 'No.', 'No.', 'Yes.', 'Yes.', 'Geneva.', '1856', 'Woodrow.', 'Politics']
[Batch 27] Time for generation = 0.42 seconds
[Batch 28] Started generation...
['<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: What caused Wilson to ask Congress to declare war on the Central Powers? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: What was more damaging than moving students into colleges? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Was Wilson a member of the Phi Kappa Psi fraternity? [/INST]', '<s>[INST] Give an answer to the question. Use one word if possible.\n\nQuestion: Was Wilson an automobile enthusiast? [/INST]', "<s>[IN

## Evaluation

In [16]:
def evaluate_bleu(df_true: pd.DataFrame, df_pred: pd.DataFrame, bleu_type: int):
    df_merged = pd.merge(df_true[['id', 'answer']], df_pred[['id', 'answer']], on='id', suffixes=('_true', '_pred'))

    weights = (1.0, 0.0, 0.0, 0.0) if bleu_type == 1 else (0.5, 0.5, 0.0, 0.0)
    
    bleu_scores = []

    with warnings.catch_warnings():
        warnings.simplefilter("ignore")
        for _, row in df_merged.iterrows():
            reference = row['answer_true'].split()
            candidate = row['answer_pred'].split() 

            bleu_score = sentence_bleu([reference], candidate, weights=weights)
            bleu_scores.append(bleu_score)
        
    mean_bleu = sum(bleu_scores) / len(bleu_scores)
    return mean_bleu

In [17]:
questions_val = questions_split["validation"].to_pandas()

df_true_faiss = questions_val[['id', 'answer']].copy()
df_pred_faiss = questions_val[['id']].copy()
df_pred_faiss['answer'] = generated_responses

bleu_1 = evaluate_bleu(df_true_faiss, df_pred_faiss, 1)
bleu_2 = evaluate_bleu(df_true_faiss, df_pred_faiss, 2)

print(f"Le score BLEU-1 est {bleu_1:.3f}.")
print(f"Le score BLEU-2 est {bleu_2:.3f}.\n")

Le score BLEU-1 est 0.048.
Le score BLEU-2 est 0.003.



In [18]:
questions_val = questions_split["validation"].to_pandas()
df_export = questions_val[['id', 'question', 'answer']].copy()
df_export['generated_answer'] = generated_responses
df_export.to_csv(f"generated_responses_simple_Base.csv", index=False, encoding='utf-8')