## Running prompts on OpenAI

In [1]:
# %pip install markdown
# %pip install annoy
# %pip install openai

In [2]:
from openai import OpenAI # OpenAI API
import json
import requests # to download some resources
import os # file operations
import numpy as np # linear algebra
import pandas as pd # data processing
from markdown import markdown # to render markdown
from IPython.display import Markdown
import annoy # Approximate Nearest Neighbors Oh Yeah for fast searching
import pickle
from annoy import AnnoyIndex

In [3]:
# read in our api key
with open('../api-keys/our_api_key.txt', 'r') as file:
    api_key = file.read().replace('\n', '')

# read in finns api key (the one we'll use for testing)
with open('../api-keys/finns_api_key.txt', 'r') as file:
    finns_api_key = file.read().replace('\n', '')

In [4]:
# Setting the OpenAI
client = OpenAI(api_key=finns_api_key) # using finns for testing

#### Sample of countries

In [10]:
# Sample of countries
sample_countries = [
    # 'Finland','Netherlands','Denmark','Norway', 'Chile', 'Bangladesh', 'Canada', 'Sweden', 
    'Libya', 'Malta', 'Myanmar', 'Oman', 'Australia_2021-01-01']

# Folder with raw PDFs
folder_mds = "../data/3-naps-md"

# Importing MDs files
for file in os.listdir(folder_mds):
    for country in sample_countries:
        if country in file:  
            with open(os.path.join(folder_mds, file), "r", encoding="utf-8") as md_file:
                content = md_file.read()
            globals()[country.lower()] = content # Saving the MD file in lowercase


### 1. Using the WHOLE document

In [44]:
country_test = malta
country_name = "malta"

In [45]:
system_prompt = """
You are an assistant that extracts and summarizes national action plans for antimicrobial resistance (NAPs). Your task is to answer the following questions concisely while relying strictly on the document content. 

For each question, provide:
1. A 'yes' or 'no' answer.
2. The specific text from the document that supports the answer.

Here are the questions:

1. What is the period considered for the NAP? (Provide the exact years or date range)
2. Does the NAP adopt a One Health approach by addressing multiple priority sectors (such as human health, animal health, environment, and agriculture/food security) through strategies, policies, goals, or actions? (yes/no + supporting text)  
3. Does the NAP include any distinct strategy/policy/goal/action explicitly targeting the priority sector of 'Human Health'? (yes/no + supporting text)  
4. Does the NAP include any distinct strategy/policy/goal/action explicitly targeting the priority sector of 'Animal Health'? (yes/no + supporting text)  
5. Does the NAP include any distinct strategy/policy/goal/action explicitly targeting the priority sector of 'Environment'? (yes/no + supporting text)  
6. Does the NAP include any distinct strategy/policy/goal/action explicitly targeting the priority sector of 'Agriculture/Food Security'? (yes/no + supporting text)  
7. Does the NAP include any mechanisms for progress reporting to track how its objectives are being met, such as an annual or semi-annual progress report, a dashboard displaying the status, or similar documents/tools? (yes/no + supporting text)  
8. Does the NAP include any specific, time-bound targets to track progress? (yes/no + supporting text)  
9. Does the NAP outline quantitative targets for AMR prevalence and antimicrobial use? (yes/no + supporting text)  
10. Does the NAP specify any budget allocation assigned for the strategies/policies/goals/actions directed at tackling AMR? (yes/no + supporting text)  
11. Does the NAP contain an assessment of future budget requirements for different activities listed in the document? (yes/no + supporting text)  
12. Does the NAP establish a multisectoral committee or task force for AMR coordination? (yes/no + supporting text)  
13. Does the NAP include any strategy/policy/goal/action related to strengthening the legislative and regulatory framework for AMR, including laws and regulations for antimicrobial use? (yes/no + supporting text)  
14. Does the NAP include any strategy/policy/goal/action related to training and educational initiatives to increase awareness about AMR in school curricula or professional training programs for doctors and pharmacists? (yes/no + supporting text)  
15. Does the NAP include any strategy/policy/goal/action related to public awareness campaigns in combating AMR, such as media involvement and community engagement? (yes/no + supporting text)  
16. Does the NAP include any strategy/policy/goal/action related to international collaboration and support for AMR initiatives, such as partnerships with organizations like World Health Organization (WHO), Food and Agriculture Organization (FAO), and World Organisation for Animal Health (OIE)? (yes/no + supporting text)  
17. Does the NAP include any strategy/policy/goal/action related to data collection and analysis for AMR surveillance, such as developing surveillance systems, data collection tools, setting up laboratories for AMR monitoring, or other similar actions? (yes/no + supporting text)  
18. Does the NAP include any strategy/policy/goal/action related to developing the capacity to detect and report newly emerged resistance that may constitute a public health emergency of international concern, as required by the International Health Regulations (2005)? (yes/no + supporting text)  
19. Does the NAP include any strategy/policy/goal/action related to specific prevention strategies such as sanitation, hygiene measures, and infection prevention and control (IPC)? (yes/no + supporting text)  
20. Does the NAP include any strategy/policy/goal/action related to implementing international standards and guidelines for antimicrobial resistance (AMR) control, such as (but not limited to) the OIE Terrestrial and Aquatic Animal Health Codes or the FAO/WHO Codex Alimentarius Code of Practice to Minimize and Contain AMR? (yes/no + supporting text)  
21. Does the NAP include any strategy/policy/goal/action related to regulating the distribution, prescription, and dispensation of antibiotics, such as developing or maintaining essential medicine lists for antibiotics? (yes/no + supporting text)  
22. Does the NAP include any strategy/policy/goal/action related to requiring the rational use of antibiotics through regulations/laws? (yes/no + supporting text)  
23. Does the NAP include any strategy/policy/goal/action related to participating in international collaborative research to support the development of new medicines, diagnostic tools, and vaccines? (yes/no + supporting text)

Format the response as a JSON object:
{
  "period_start": "XXXX",
  "period_end":"YYYY",
  "questions": [
    {
      "question": "Does the NAP include any strategy/policy/goal/action related to the priority sector of 'Human Health'?",
      "answer": "yes",
      "supporting_text": "The NAP mentions specific actions in the human health sector, including..."
    },
    ...
  ]
}
"""

In [46]:
response = client.chat.completions.create(
    model="gpt-4o-mini",
    temperature=0,  # Ensures deterministic output
    top_p=0.5,  # Controls diversity
    response_format={"type": "json_object"},
    messages=[
        {"role": "system", "content": system_prompt},
        {"role": "user", "content": country_test}
    ]
)

# Print the structured JSON response
print(response.choices[0].message.content)

{
  "period_start": "2020",
  "period_end": "2028",
  "questions": [
    {
      "question": "Does the NAP adopt a One Health approach by addressing multiple priority sectors (such as human health, animal health, environment, and agriculture/food security) through strategies, policies, goals, or actions?",
      "answer": "yes",
      "supporting_text": "The Strategy adopts a whole of Government, a whole of society and a comprehensive One Health Approach to AMR, recognizing the complexity of the issue."
    },
    {
      "question": "Does the NAP include any distinct strategy/policy/goal/action explicitly targeting the priority sector of 'Human Health'?",
      "answer": "yes",
      "supporting_text": "The Strategy includes specific actions in the human health sector, including improved surveillance, stewardship measures for the proper use of antibiotics, and infection prevention and control."
    },
    {
      "question": "Does the NAP include any distinct strategy/policy/goal/acti

In [47]:
# Formatting the JSON

c = json.loads(response.choices[0].message.content)
period_data = pd.DataFrame([
    {"question": "period_start", "answer": c["period_start"], "supporting_text": None},
    {"question": "period_end", "answer": c["period_end"], "supporting_text": None}
])

df = pd.DataFrame(c["questions"])
df = pd.concat([period_data, df], ignore_index=True)
df['country']=country_name
df = df.rename(columns={"question": "question_text", "answer": "answer_full"})
df['question']=df['question_text']
df.loc[2:, 'question'] = range(1, 1 + len(df) - 2)
df = df.drop(columns=['question_text'])
output_json = df.to_dict(orient="records")

# Saving the JSON
output_path = f"../data/4-summary-responses-json/version_yesno_updated/summary_response_{country_name}.json"
with open(output_path, "w") as f:
    json.dump(output_json, f, indent=4)

print(f"Saved response to {output_path}")


Saved response to ../data/4-summary-responses-json/version_yesno_updated/summary_response_malta.json


In [48]:
df.head()

Unnamed: 0,answer_full,supporting_text,country,question
0,2020,,malta,period_start
1,2028,,malta,period_end
2,yes,"The Strategy adopts a whole of Government, a w...",malta,1
3,yes,The Strategy includes specific actions in the ...,malta,2
4,yes,The Strategy outlines interventions to address...,malta,3


In [49]:
# Load the annotations
annotations_df = pd.read_excel('../data/5-annotations/annotations_new.xlsx')
annotations_df['question'] = annotations_df['question'].astype(str)
annotations_df['answer'] = annotations_df['answer'].astype(str)

# Filter for the current country
annotations_df = annotations_df[annotations_df['country'] == country_name]


In [50]:
annotations_df

Unnamed: 0,country,question,answer,comment_input_for_protocol,supporting_text_when_true_answer_is_yes_but_prediction_is_no
168,malta,period_start,2020,,
169,malta,period_end,2028,,
170,malta,1,yes,,
171,malta,2,yes,,
172,malta,3,yes,,
173,malta,4,yes,,
174,malta,5,yes,,
175,malta,6,yes,,
176,malta,7,yes,,
177,malta,8,no,,


In [51]:
# Extract period and questions from the model response
model_answers = [{"country": country_name, "question": "period_start", "answer": c["period_start"]}]  # Corrected
model_answers.append({"country": country_name, "question": "period_end", "answer": c["period_end"]})  # Corrected

# Add answers for the other questions
for q in c["questions"]:
    model_answers.append({
        "country": country_name,
        "question": q["question"],
        "answer": q["answer"]
    })

# Convert to DataFrame
model_answers_df = pd.DataFrame(model_answers)
model_answers_df['question'] = model_answers_df['question'].astype(str)

# Mapping the questions to numbers
question_mapping = {
   "Does the NAP adopt a One Health approach by addressing multiple priority sectors (such as human health, animal health, environment, and agriculture/food security) through strategies, policies, goals, or actions?": "1",
    "Does the NAP include any distinct strategy/policy/goal/action explicitly targeting the priority sector of 'Human Health'?": "2",
    "Does the NAP include any distinct strategy/policy/goal/action explicitly targeting the priority sector of 'Animal Health'?": "3",
    "Does the NAP include any distinct strategy/policy/goal/action explicitly targeting the priority sector of 'Environment'?": "4",
    "Does the NAP include any distinct strategy/policy/goal/action explicitly targeting the priority sector of 'Agriculture/Food Security'?": "5",
    "Does the NAP include any mechanisms for progress reporting to track how its objectives are being met, such as an annual or semi-annual progress report, a dashboard displaying the status, or similar documents/tools?": "6",
    "Does the NAP include any specific, time-bound targets to track progress?": "7",
    "Does the NAP outline quantitative targets for AMR prevalence and antimicrobial use?": "8",
    "Does the NAP specify any budget allocation assigned for the strategies/policies/goals/actions directed at tackling AMR?": "9",
    "Does the NAP contain an assessment of future budget requirements for different activities listed in the document?": "10",
    "Does the NAP establish a multisectoral committee or task force for AMR coordination?": "11",
    "Does the NAP include any strategy/policy/goal/action related to strengthening the legislative and regulatory framework for AMR, including laws and regulations for antimicrobial use?": "12",
    "Does the NAP include any strategy/policy/goal/action related to training and educational initiatives to increase awareness about AMR in school curricula or professional training programs for doctors and pharmacists?": "13",
    "Does the NAP include any strategy/policy/goal/action related to public awareness campaigns in combating AMR, such as media involvement and community engagement?": "14",
    "Does the NAP include any strategy/policy/goal/action related to international collaboration and support for AMR initiatives, such as partnerships with organizations like World Health Organization (WHO), Food and Agriculture Organization (FAO), and World Organisation for Animal Health (OIE)?": "15",
    "Does the NAP include any strategy/policy/goal/action related to data collection and analysis for AMR surveillance, such as developing surveillance systems, data collection tools, setting up laboratories for AMR monitoring, or other similar actions?": "16",
    "Does the NAP include any strategy/policy/goal/action related to developing the capacity to detect and report newly emerged resistance that may constitute a public health emergency of international concern, as required by the International Health Regulations (2005)?": "17",
    "Does the NAP include any strategy/policy/goal/action related to specific prevention strategies such as sanitation, hygiene measures, and infection prevention and control (IPC)?": "18",
    "Does the NAP include any strategy/policy/goal/action related to implementing international standards and guidelines for antimicrobial resistance (AMR) control, such as (but not limited to) the OIE Terrestrial and Aquatic Animal Health Codes or the FAO/WHO Codex Alimentarius Code of Practice to Minimize and Contain AMR?": "19",
    "Does the NAP include any strategy/policy/goal/action related to regulating the distribution, prescription, and dispensation of antibiotics, such as developing or maintaining essential medicine lists for antibiotics?": "20",
    "Does the NAP include any strategy/policy/goal/action related to requiring the rational use of antibiotics through regulations/laws?": "21",
    "Does the NAP include any strategy/policy/goal/action related to participating in international collaborative research to support the development of new medicines, diagnostic tools, and vaccines?": "22"
}

# Replacing the questions with numbers
model_answers_df["question"] = model_answers_df["question"].replace(question_mapping)

# Merge based on country and question
merged_df = pd.merge(annotations_df, model_answers_df, on=["country", "question"], suffixes=('_gs', '_llm'))

# Now you should have the full merged dataframe


In [52]:
merged_df

Unnamed: 0,country,question,answer_gs,comment_input_for_protocol,supporting_text_when_true_answer_is_yes_but_prediction_is_no,answer_llm
0,malta,period_start,2020,,,2020
1,malta,period_end,2028,,,2028
2,malta,1,yes,,,yes
3,malta,2,yes,,,yes
4,malta,3,yes,,,yes
5,malta,4,yes,,,yes
6,malta,5,yes,,,yes
7,malta,6,yes,,,yes
8,malta,7,yes,,,yes
9,malta,8,no,,,yes


In [53]:
from sklearn.metrics import accuracy_score, precision_score, recall_score, f1_score

# Convert categorical yes/no answers to binary (1 for yes, 0 for no)
merged_df['gs_binary'] = merged_df['answer_gs'].apply(lambda x: 1 if x == 'yes' else (0 if x == 'no' else None))
merged_df['llm_binary'] = merged_df['answer_llm'].apply(lambda x: 1 if x == 'yes' else (0 if x == 'no' else None))

# Drop rows with NaN values in either gs_binary or llm_binary for accurate calculations
yesno_df = merged_df[~merged_df['question'].isin(["period_start", "period_end"])]
yesno_df = yesno_df.dropna(subset=['gs_binary', 'llm_binary'])

# Performance for yes/no questions
accuracy = accuracy_score(yesno_df['gs_binary'], yesno_df['llm_binary'])
precision = precision_score(yesno_df['gs_binary'], yesno_df['llm_binary'])
recall = recall_score(yesno_df['gs_binary'], yesno_df['llm_binary'])
f1 = f1_score(yesno_df['gs_binary'], yesno_df['llm_binary'])

# Performance for period extraction (period_start, period_end)
period_matches = merged_df[merged_df['question'].isin(["period_start", "period_end"])]

# For period extraction, we don't need binary conversion, so we directly compare the answers
period_matches['period_match'] = period_matches['answer_gs'] == period_matches['answer_llm']

# Calculate period extraction accuracy as the percentage of exact matches
period_accuracy = period_matches['period_match'].mean()  # Exact match percentage

# Print results
print(f"Accuracy (Yes/No Questions): {accuracy:.2f}")
print(f"Precision (Yes/No Questions): {precision:.2f}")
print(f"Recall (Yes/No Questions): {recall:.2f}")
print(f"F1 Score (Yes/No Questions): {f1:.2f}")
print(f"Period Extraction Accuracy: {period_accuracy:.2f}")

Accuracy (Yes/No Questions): 0.91
Precision (Yes/No Questions): 0.95
Recall (Yes/No Questions): 0.95
F1 Score (Yes/No Questions): 0.95
Period Extraction Accuracy: 1.00


A value is trying to be set on a copy of a slice from a DataFrame.
Try using .loc[row_indexer,col_indexer] = value instead

See the caveats in the documentation: https://pandas.pydata.org/pandas-docs/stable/user_guide/indexing.html#returning-a-view-versus-a-copy
  period_matches['period_match'] = period_matches['answer_gs'] == period_matches['answer_llm']


---