In [1]:
import torch
import transformers
from transformers import TrainingArguments
from datasets import load_dataset, DatasetDict, Dataset

from unsloth import FastLanguageModel
from trl import SFTTrainer
from unsloth.chat_templates import get_chat_template
from unsloth import is_bfloat16_supported

from tqdm import tqdm
import random

import os
os.environ["HF_HOME"] = "~/scratch/hf-cache"

sat_math_dict = load_dataset("knazarali3/group2_processed_sat_math_cot")
test_dataset = sat_math_dict["test"]


ðŸ¦¥ Unsloth: Will patch your computer to enable 2x faster free finetuning.


In [2]:

max_seq_length = 2048 # supports RoPE Scaling internally
dtype = None # None for auto detection
load_in_4bit = True # Use 4bit quantization to reduce memory usage.

random.seed(42)

In [3]:
model, tokenizer = FastLanguageModel.from_pretrained(
    model_name = "knazarali3/llama3_sat_math", # YOUR MODEL YOU USED FOR TRAINING
    max_seq_length = max_seq_length,
    dtype = dtype,
    load_in_4bit = load_in_4bit,
    token = "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"
)
FastLanguageModel.for_inference(model) # Enable native 2x faster inference

==((====))==  Unsloth 2024.11.5: Fast Llama patching. Transformers = 4.47.0.dev0.
   \\   /|    GPU: Tesla V100-PCIE-32GB. Max memory: 31.733 GB. Platform = Linux.
O^O/ \_/ \    Pytorch: 2.5.0+cu121. CUDA = 7.0. CUDA Toolkit = 12.1.
\        /    Bfloat16 = FALSE. FA [Xformers = 0.0.28.post2. FA2 = False]
 "-____-"     Free Apache license: http://github.com/unslothai/unsloth
Unsloth: Fast downloading is enabled - ignore downloading bars which are red colored!


Unsloth 2024.11.5 patched 32 layers with 32 QKV layers, 32 O layers and 32 MLP layers.


PeftModelForCausalLM(
  (base_model): LoraModel(
    (model): LlamaForCausalLM(
      (model): LlamaModel(
        (embed_tokens): Embedding(128256, 4096, padding_idx=128004)
        (layers): ModuleList(
          (0-31): 32 x LlamaDecoderLayer(
            (self_attn): LlamaAttention(
              (q_proj): lora.Linear4bit(
                (base_layer): Linear4bit(in_features=4096, out_features=4096, bias=False)
                (lora_dropout): ModuleDict(
                  (default): Identity()
                )
                (lora_A): ModuleDict(
                  (default): Linear(in_features=4096, out_features=16, bias=False)
                )
                (lora_B): ModuleDict(
                  (default): Linear(in_features=16, out_features=4096, bias=False)
                )
                (lora_embedding_A): ParameterDict()
                (lora_embedding_B): ParameterDict()
                (lora_magnitude_vector): ModuleDict()
              )
              (k_proj): lor

In [4]:
def format_sat_math(test_dataset, tokenizer):
    alpaca_prompt = """Below is an instruction that describes a task, paired with an input that provides further context. Write a response that appropriately completes the request.

    ### Instruction:
    Please provide the answer in the following format: A final_answer key with the direct answer. A steps key containing a list of step-by-step calculations and explanations that should lead to the final answer. Each step should have an explanation key describing the calculation. Please pay attention to the input because it is the question you have to answer.
    
    ### Input:
    {}

    ### Response:
    """

    def formatting_prompts_func(examples):
        return {"text": alpaca_prompt.format(examples["question"]) + tokenizer.eos_token }

    return test_dataset.map(formatting_prompts_func)  # Process examples individually

In [5]:
dataset = format_sat_math(test_dataset, tokenizer)

Map:   0%|          | 0/6290 [00:00<?, ? examples/s]

In [9]:
import re
for i in tqdm(range(5), desc="Generating Responses"):  # No need for batch_size here
    print(dataset["text"][i])
    inputs = tokenizer(dataset["text"][i], return_tensors = "pt").to("cuda")
    with torch.no_grad():
        outputs = model.generate(**inputs, use_cache = True, pad_token_id = tokenizer.eos_token_id)
        output = tokenizer.batch_decode(outputs)
        output = output[0]

    answer = re.search(r"'final_answer':\s*'([^']*)'", output)

    if answer:
        final_answer = answer.group(1)
    
    print("output: ", output)
    
    print("final answer: ", final_answer)
    print("actual answer: ", dataset["answer"][i])

Generating Responses:   0%|          | 0/5 [00:00<?, ?it/s]

Below is an instruction that describes a task, paired with an input that provides further context. Write a response that appropriately completes the request.

    ### Instruction:
    Please provide the answer in the following format: A final_answer key with the direct answer. A steps key containing a list of step-by-step calculations and explanations that should lead to the final answer. Each step should have an explanation key describing the calculation. Please pay attention to the input because it is the question you have to answer.
    
    ### Input:
    A few friends are making a small movie. They hire a few actors and that cost $1200. They also get food for some people there. Each person gets $3 worth of food. Equipment rental costs twice as much as food and actors combined. They sold the movie for $10,000 and made a profit of $5950. How many people were there?

    ### Response:
    <|end_of_text|>


Generating Responses:  20%|â–ˆâ–ˆ        | 1/5 [00:46<03:07, 46.93s/it]

output:  <|begin_of_text|>Below is an instruction that describes a task, paired with an input that provides further context. Write a response that appropriately completes the request.

    ### Instruction:
    Please provide the answer in the following format: A final_answer key with the direct answer. A steps key containing a list of step-by-step calculations and explanations that should lead to the final answer. Each step should have an explanation key describing the calculation. Please pay attention to the input because it is the question you have to answer.
    
    ### Input:
    A few friends are making a small movie. They hire a few actors and that cost $1200. They also get food for some people there. Each person gets $3 worth of food. Equipment rental costs twice as much as food and actors combined. They sold the movie for $10,000 and made a profit of $5950. How many people were there?

    ### Response:
    <|end_of_text|><|begin_of_text|>Below is an instruction that describes

Generating Responses:  40%|â–ˆâ–ˆâ–ˆâ–ˆ      | 2/5 [01:15<01:47, 35.97s/it]

output:  <|begin_of_text|>Below is an instruction that describes a task, paired with an input that provides further context. Write a response that appropriately completes the request.

    ### Instruction:
    Please provide the answer in the following format: A final_answer key with the direct answer. A steps key containing a list of step-by-step calculations and explanations that should lead to the final answer. Each step should have an explanation key describing the calculation. Please pay attention to the input because it is the question you have to answer.
    
    ### Input:
    Subtracting 36 from a number and multipling by 22 gets you 748. Find the number.

    ### Response:
    <|end_of_text|><|begin_of_text|>Below is an instruction that describes a task, paired with an input that provides further context. Write a response that appropriately completes the request.

    ### Instruction:
    Please provide the answer in the following format: A final_answer key with the direct an

Generating Responses:  60%|â–ˆâ–ˆâ–ˆâ–ˆâ–ˆâ–ˆ    | 3/5 [01:32<00:54, 27.49s/it]

output:  <|begin_of_text|>Below is an instruction that describes a task, paired with an input that provides further context. Write a response that appropriately completes the request.

    ### Instruction:
    Please provide the answer in the following format: A final_answer key with the direct answer. A steps key containing a list of step-by-step calculations and explanations that should lead to the final answer. Each step should have an explanation key describing the calculation. Please pay attention to the input because it is the question you have to answer.
    
    ### Input:
    Steven's teacher sends the class an assignment to collect 60 different fruit seeds. Apples average 6 seeds, pears average 2 seeds, and grapes average a certain number of seeds. Steven has set aside 4 apples, 3 pears, and 9 grapes to extract their seeds. He needs 3 more seeds to fulfill his assignment. How many seeds do grapes average?

    ### Response:
    <|end_of_text|><|begin_of_text|>:// The number o

Generating Responses:  80%|â–ˆâ–ˆâ–ˆâ–ˆâ–ˆâ–ˆâ–ˆâ–ˆ  | 4/5 [01:47<00:22, 22.50s/it]

output:  <|begin_of_text|>Below is an instruction that describes a task, paired with an input that provides further context. Write a response that appropriately completes the request.

    ### Instruction:
    Please provide the answer in the following format: A final_answer key with the direct answer. A steps key containing a list of step-by-step calculations and explanations that should lead to the final answer. Each step should have an explanation key describing the calculation. Please pay attention to the input because it is the question you have to answer.
    
    ### Input:
    In a week, Rayman works out for half the number of hours Junior takes to work out. Wolverine, their friend, works out for some ratio of the combined total number of hours both Rayman and Junior take during their workouts. Wolverine takes 60 hours doing her workouts in a week. What is the ratio of the number of hours Wolverine works out to the combined total number of hours both Rayman and Junior take duri

Generating Responses: 100%|â–ˆâ–ˆâ–ˆâ–ˆâ–ˆâ–ˆâ–ˆâ–ˆâ–ˆâ–ˆ| 5/5 [02:18<00:00, 27.67s/it]

output:  <|begin_of_text|>Below is an instruction that describes a task, paired with an input that provides further context. Write a response that appropriately completes the request.

    ### Instruction:
    Please provide the answer in the following format: A final_answer key with the direct answer. A steps key containing a list of step-by-step calculations and explanations that should lead to the final answer. Each step should have an explanation key describing the calculation. Please pay attention to the input because it is the question you have to answer.
    
    ### Input:
    Find A from the 4-digit equation A2B2+1C1D=3333.

    ### Response:
    <|end_of_text|><|begin_of_text|>Below is an instruction that describes a task, paired with an input that provides further context. Write a response that appropriately completes the request.

    ### Instruction:
    Please provide the answer in the following format: A final_answer key with the direct answer. A steps key containing a l


