In [1]:
from pyserini.search.lucene import LuceneSearcher
from datasets import load_dataset
from dotenv import load_dotenv
import os
from langchain.chat_models import ChatOpenAI, ChatCohere, ChatAnyscale

from main import get_query_expansion_dataset, run_search

  from .autonotebook import tqdm as notebook_tqdm


In [3]:
# set up LLMs, assumes that API keys are stored in .env file
load_dotenv(override=True)
chats = {}

# OpenAI
try: 
    openai_api_key = os.getenv("OPENAI_API_KEY")
    chat_openai = ChatOpenAI(openai_api_key=openai_api_key, model="gpt-3.5-turbo")
    chats["OpenAI"] = chat_openai
except:
    print("Could not initialize OpenAI chatbot. Please store a valid API key in a .env file.")

# Cohere
try:
    cohere_api_key = os.getenv("COHERE_API_KEY")
    chat_cohere = ChatCohere(cohere_api_key=cohere_api_key)
    chats["Cohere"] = chat_cohere
except:
    print("Could not initialize Cohere chatbot. Please store a valid API key in a .env file.")

# Llama 2 (Anyscale)
try: 
    anyscale_api_key = os.getenv("ANYSCALE_API_KEY")
    chat_llama = ChatAnyscale(model_name="meta-llama/Llama-2-7b-chat-hf", anyscale_api_key=anyscale_api_key)
    chats["Llama 2"] = chat_llama
except:
    print("Could not initialize Anyscale chatbot. Please store a valid API key in a .env file.")

## Experiments Set 1: All prompts, english language, all models

In [4]:
# Load the Miracl dataset (english version)
dataset = load_dataset("Cohere/miracl-en-queries-22-12", split="dev")
dataset = dataset.to_pandas().to_dict(orient='records')

# set up searcher
searcher = LuceneSearcher.from_prebuilt_index('miracl-v1.0-en')

## 1) Baseline

In [4]:
recall_baseline, ndcg_baseline = run_search(searcher, dataset)

print(f'BM25 Recall@100: {recall_baseline:.4f}')
print(f'BM25 nDCG@10: {ndcg_baseline:.4f}')
print("\n")

## 2) Query2Doc Zero Shot Prompting

In [6]:
for chat_name, chat in chats.items():
    expanded_dataset = get_query_expansion_dataset(dataset, chat, chat_name, lang='en', prompt='q2d-zs')

    recall, ndcg = run_search(searcher, expanded_dataset)
    print(f'{chat_name} Recall@100: {recall:.4f}') 
    print(f'{chat_name} nDCG@10: {ndcg:.4f}')
    print("\n")

Expanding queries:   0%|          | 0/5 [00:00<?, ?it/s]

Expanding queries: 100%|██████████| 5/5 [01:10<00:00, 14.11s/it]
Searching: 100%|██████████| 5/5 [00:14<00:00,  2.83s/it]


OpenAI Recall@100: 0.8333
OpenAI nDCG@10: 0.2921




Expanding queries: 100%|██████████| 5/5 [01:06<00:00, 13.30s/it]
Searching: 100%|██████████| 5/5 [00:17<00:00,  3.51s/it]


Cohere Recall@100: 0.7667
Cohere nDCG@10: 0.3049




Expanding queries: 100%|██████████| 5/5 [00:27<00:00,  5.52s/it]
Searching: 100%|██████████| 5/5 [00:08<00:00,  1.77s/it]

Llama 2 Recall@100: 0.7667
Llama 2 nDCG@10: 0.3563







## 3) Q2D zero shot prompting with pseudo-relevant feedback

In [None]:
for chat_name, chat in chats.items():
    expanded_dataset = get_query_expansion_dataset(dataset, chat, chat_name, lang='en', prompt='q2d-zs-prf')

    recall, ndcg = run_search(searcher, expanded_dataset)
    print(f'{chat_name} Recall@100: {recall:.4f}') 
    print(f'{chat_name} nDCG@10: {ndcg:.4f}')
    print("\n")

Getting prf documents...


Searching: 100%|██████████| 5/5 [00:00<00:00, 72.71it/s]


Done.


Expanding queries: 100%|██████████| 5/5 [00:52<00:00, 10.54s/it]
Searching: 100%|██████████| 5/5 [00:08<00:00,  1.60s/it]

BM25 Recall@100: 0.9000
BM25 nDCG@10: 0.3924







## 4) Keywords with zero-shot prompting

In [8]:
for chat_name, chat in chats.items():
    expanded_dataset = get_query_expansion_dataset(dataset, chat, chat_name, lang='en', prompt='q2e-zs')

    recall, ndcg = run_search(searcher, expanded_dataset)
    print(f'{chat_name} Recall@100: {recall:.4f}') 
    print(f'{chat_name} nDCG@10: {ndcg:.4f}')
    print("\n")

Expanding queries:   0%|          | 0/799 [00:00<?, ?it/s]

Expanding queries: 100%|██████████| 799/799 [54:50<00:00,  4.12s/it]  
Searching: 100%|██████████| 799/799 [01:25<00:00,  9.38it/s]

BM25 Recall@100: 0.8433
BM25 nDCG@10: 0.3745







## 5) Keywords with pseudo-relevant feedback

In [9]:
for chat_name, chat in chats.items():
    expanded_dataset = get_query_expansion_dataset(dataset, chat, chat_name, lang='en', prompt='q2e-zs-prf')

    recall, ndcg = run_search(searcher, expanded_dataset)
    print(f'{chat_name} Recall@100: {recall:.4f}') 
    print(f'{chat_name} nDCG@10: {ndcg:.4f}')
    print("\n")

Getting prf documents...


Searching:  11%|█         | 89/799 [00:02<00:28, 25.03it/s]

Searching: 100%|██████████| 799/799 [00:26<00:00, 30.04it/s]


Done.


Expanding queries: 100%|██████████| 799/799 [1:02:36<00:00,  4.70s/it]
Searching: 100%|██████████| 799/799 [02:56<00:00,  4.53it/s]

BM25 Recall@100: 0.8337
BM25 nDCG@10: 0.3722







## 6) Chain of thought prompting

In [15]:
for chat_name, chat in chats.items():
    expanded_dataset = get_query_expansion_dataset(dataset, chat, chat_name, lang='en', prompt='chain-of-thought')

    recall, ndcg = run_search(searcher, expanded_dataset)
    print(f'{chat_name} Recall@100: {recall:.4f}') 
    print(f'{chat_name} nDCG@10: {ndcg:.4f}')
    print("\n")

Expanding queries:  33%|███▎      | 1/3 [00:16<00:33, 16.59s/it]

Is Creole a pidgin of French? Is Creole a pidgin of French? Is Creole a pidgin of French? Is Creole a pidgin of French? Is Creole a pidgin of French? Yes, Creole can be considered a pidgin of French. 

A pidgin is a simplified language that develops as a means of communication between two groups of people who do not share a common language. It typically arises in situations of trade or colonization, where people with different native languages need to communicate.

Creole languages, including Haitian Creole and Louisiana Creole, emerged as pidgins when French-speaking colonizers interacted with African slaves and other non-French speakers. These pidgins gradually evolved and became the native languages of the communities, gaining complexity and distinct features over time.

Therefore, Creole languages can be seen as pidgins that have undergone further development and become fully-formed languages in their own right, while still maintaining influences from the French language.


Expanding queries:  67%|██████▋   | 2/3 [00:28<00:13, 13.94s/it]

What percentage of the Earth's atmosphere is oxygen? What percentage of the Earth's atmosphere is oxygen? What percentage of the Earth's atmosphere is oxygen? What percentage of the Earth's atmosphere is oxygen? What percentage of the Earth's atmosphere is oxygen? Approximately 21% of the Earth's atmosphere is oxygen. This percentage is based on the composition of the atmosphere which consists of various gases such as nitrogen, oxygen, carbon dioxide, and others. Oxygen is the second most abundant gas in the atmosphere after nitrogen. The rational behind this answer is that oxygen is essential for supporting life, and its presence in the atmosphere allows for respiration and combustion processes. However, it is important to note that the precise percentage may vary slightly depending on location and altitude.


Expanding queries: 100%|██████████| 3/3 [00:59<00:00, 19.76s/it]


When did Marxism develop? When did Marxism develop? When did Marxism develop? When did Marxism develop? When did Marxism develop? Marxism developed in the mid-19th century during a period of significant social and economic changes. The rational behind its development can be traced back to several factors:

1. Industrial Revolution: The rapid industrialization in Europe during the 18th and 19th centuries led to the rise of capitalism, which resulted in stark social inequalities and exploitation of the working class. This created a fertile ground for the development of Marxist ideas.

2. Class struggle: Karl Marx and Friedrich Engels, the founders of Marxism, observed the struggles and conflicts between the bourgeoisie (capitalist class) and the proletariat (working class). They believed that these class struggles were inherent to the capitalist system and would eventually lead to its downfall.

3. Hegelian philosophy: Marx was influenced by the German philosopher Georg Wilhelm Friedrich

Searching: 100%|██████████| 3/3 [00:04<00:00,  1.51s/it]


OpenAI Recall@100: 0.6111
OpenAI nDCG@10: 0.0952




Expanding queries:  33%|███▎      | 1/3 [00:04<00:08,  4.14s/it]

Is Creole a pidgin of French? Is Creole a pidgin of French? Is Creole a pidgin of French? Is Creole a pidgin of French? Is Creole a pidgin of French? A pidgin language is a language that develops when two or more languages come into contact and is used for basic communication among people who do not have a common language. It is characterized by a simplified grammar and limited vocabulary, often with words borrowed from different languages. 

 Creole is a term used to describe a variety of languages that have developed from French. These languages are characterized by a mix of French vocabulary and grammar rules adapted to local conditions and languages that were spoken in the region. 

Therefore, the answer is "Yes". 

Would you like to know how Creole came to be?


Expanding queries:  67%|██████▋   | 2/3 [00:14<00:08,  8.06s/it]

What percentage of the Earth's atmosphere is oxygen? What percentage of the Earth's atmosphere is oxygen? What percentage of the Earth's atmosphere is oxygen? What percentage of the Earth's atmosphere is oxygen? What percentage of the Earth's atmosphere is oxygen? To begin, it is crucial to acknowledge that the Earth's atmosphere is a intricate mixture of various gases, including nitrogen, oxygen, argon, carbon dioxide, and hydrogen. The precise composition of the atmosphere varies slightly depending on factors such as altitude and location, but the following are the primary components:
1. Nitrogen: Nitrogen gas makes up approximately 78% of the Earth's atmosphere. It is an inert gas that is vital for life on Earth, as it is a major component of air and plays a crucial role in the nitrogen cycle.
2. Oxygen: Oxygen gas makes up approximately 21% of the Earth's atmosphere. Oxygen is essential for human respiration and is required by organisms for the process of respiration.
3. Argon: Arg

Expanding queries: 100%|██████████| 3/3 [00:17<00:00,  5.67s/it]


When did Marxism develop? When did Marxism develop? When did Marxism develop? When did Marxism develop? When did Marxism develop? To answer the question, consider the following: Marxism is a philosophical, socioeconomic, and political theory originated by Karl Marx and Friedrich Engels in the mid-19th century. 

Therefore, the final answer is mid-19th century. 

Would you like to know more about Marxism?


Searching: 100%|██████████| 3/3 [00:02<00:00,  1.14it/s]


Cohere Recall@100: 0.8333
Cohere nDCG@10: 0.1016




Expanding queries:  33%|███▎      | 1/3 [00:15<00:31, 15.58s/it]

Is Creole a pidgin of French? Is Creole a pidgin of French? Is Creole a pidgin of French? Is Creole a pidgin of French? Is Creole a pidgin of French?   Creole is a language that has been the subject of much debate and discussion among linguists and scholars. While some argue that Creole is a pidgin of French, others maintain that it is a distinct language with its own grammar and syntax. In this answer, we will explore the arguments for and against the idea that Creole is a pidgin of French, and examine the linguistic evidence that supports or challenges this claim.

Arguments for Creole being a pidgin of French:

1. Historical origins: Creole is believed to have originated in the French colonies of the Americas and the Caribbean in the 18th and 19th centuries. As a result, many linguists assume that it must have developed from French as a means of communication between French speakers and enslaved Africans or other non-French speakers.
2. Lexical similarity: Creole shares many cognate

Expanding queries:  67%|██████▋   | 2/3 [00:20<00:09,  9.04s/it]

What percentage of the Earth's atmosphere is oxygen? What percentage of the Earth's atmosphere is oxygen? What percentage of the Earth's atmosphere is oxygen? What percentage of the Earth's atmosphere is oxygen? What percentage of the Earth's atmosphere is oxygen?   The Earth's atmosphere is composed of several gases, with oxygen (O2) being the most abundant. According to the National Oceanic and Atmospheric Administration (NOAA), oxygen makes up approximately 21% of the Earth's atmosphere.

To calculate the percentage of oxygen in the atmosphere, we need to know the total volume of the atmosphere. The Earth's atmosphere is approximately 5.14 x 10^18 meters cubed, according to NASA.

Now, let's calculate the percentage of oxygen in the atmosphere:

Percentage of oxygen in the atmosphere = Mass of oxygen / Total mass of atmosphere

The mass of oxygen in the atmosphere is approximately 0.21 x 10^18 kilograms (based on the average atmospheric pressure of 1013 mbar and the molar mass of ox

Expanding queries: 100%|██████████| 3/3 [00:28<00:00,  9.47s/it]


When did Marxism develop? When did Marxism develop? When did Marxism develop? When did Marxism develop? When did Marxism develop?   Marxism, as a political and economic ideology, developed over a period of time, with key contributions from Karl Marx and Friedrich Engels in the 19th century. Here's a brief timeline of the major milestones in the development of Marxism:

1. Early influences (1815-1830s): Marx was influenced by the ideas of the French Revolution, the Enlightenment, and the young Hegelian school. He also read the works of Adam Smith, David Ricardo, and Thomas Malthus, which laid the groundwork for his ideas on capitalism and economic inequality.
2. Young Hegelian period (1830s-1840s): Marx was a member of the Young Hegelian circle in Berlin, where he developed his early ideas on dialectical materialism and the role of class struggle in history.
3. Parisian period (1843-1848): Marx moved to Paris, where he became familiar with the works of Charles Fourier, Louis Blanc, and 

Searching: 100%|██████████| 3/3 [00:09<00:00,  3.16s/it]

Llama 2 Recall@100: 0.6111
Llama 2 nDCG@10: 0.2592







## 7) Chain of thought prompting with psuedo-relevant feedback

In [7]:
for chat_name, chat in chats.items():
    expanded_dataset = get_query_expansion_dataset(dataset, chat, chat_name, lang='en', prompt='chain-of-thought-prf')

    recall, ndcg = run_search(searcher, expanded_dataset)
    print(f'{chat_name} Recall@100: {recall:.4f}') 
    print(f'{chat_name} nDCG@10: {ndcg:.4f}')
    print("\n")

Getting prf documents...


Searching: 100%|██████████| 799/799 [00:30<00:00, 25.81it/s]


Done.


Expanding queries: 100%|██████████| 799/799 [1:41:54<00:00,  7.65s/it]  
Searching: 100%|██████████| 799/799 [09:45<00:00,  1.36it/s]

BM25 Recall@100: 0.8486
BM25 nDCG@10: 0.4343







## 8) Chain of thought prompting with short answer

In [5]:
for chat_name, chat in chats.items():
    expanded_dataset = get_query_expansion_dataset(dataset, chat, chat_name, lang='en', prompt='chain-of-thought-short')

    recall, ndcg = run_search(searcher, expanded_dataset)
    print(f'{chat_name} Recall@100: {recall:.4f}') 
    print(f'{chat_name} nDCG@10: {ndcg:.4f}')
    print("\n")

# openai time: 2:04:00
# cohere time: 
# llama2 time: 

Searching:   0%|          | 0/799 [00:00<?, ?it/s]

Searching:  33%|███▎      | 262/799 [03:28<05:23,  1.66it/s]

# Experiments Set 2: Chain-of-thought prompt, various languages, all models

## 1) French

In [None]:
# Load the Miracl dataset (english version)
dataset = load_dataset("Cohere/miracl-fr-queries-22-12", split="dev")
dataset = dataset.to_pandas().to_dict(orient='records')

# set up searcher
searcher = LuceneSearcher.from_prebuilt_index('miracl-v1.0-fr')
searcher.set_language('fr')

# exclude cohere model as it does not support french
chats_fr = {k: v for k, v in chats.items() if k != "Cohere"}

In [None]:
recall_baseline, ndcg_baseline = run_search(searcher, dataset)

print(f'BM25 Recall@100: {recall_baseline:.4f}')
print(f'BM25 nDCG@10: {ndcg_baseline:.4f}')
print("\n")

Searching: 100%|██████████| 343/343 [00:11<00:00, 29.70it/s]

BM25 Recall@100: 0.6528
BM25 nDCG@10: 0.1832







In [None]:
for chat_name, chat in chats_fr.items():
    expanded_dataset = get_query_expansion_dataset(dataset, chat, chat_name, lang='fr', prompt='chain-of-thought-fr')

    recall, ndcg = run_search(searcher, expanded_dataset)
    print(f'{chat_name} Recall@100: {recall:.4f}') 
    print(f'{chat_name} nDCG@10: {ndcg:.4f}')
    print("\n")

# openai time: 
# llama 2 time: 25:12 min

Searching:   0%|          | 0/343 [00:00<?, ?it/s]

Searching: 100%|██████████| 343/343 [01:09<00:00,  4.92it/s]


Llama 2 Recall@100: 0.3979
Llama 2 nDCG@10: 0.1341




Expanding queries:   4%|▍         | 14/343 [17:39<6:55:00, 75.68s/it]


KeyboardInterrupt: 

## 2) German

In [None]:
# Load the Miracl dataset (english version)
dataset = load_dataset("Cohere/miracl-de-queries-22-12", split="dev")
dataset = dataset.to_pandas().to_dict(orient='records')

# set up searcher
searcher = LuceneSearcher.from_prebuilt_index('miracl-v1.0-de')
searcher.set_language('de')

# exclude cohere model as it does not support french
chats_de = {k: v for k, v in chats.items() if k != "Cohere"}

In [None]:
recall_baseline, ndcg_baseline = run_search(searcher, dataset)

print(f'BM25 Recall@100: {recall_baseline:.4f}')
print(f'BM25 nDCG@10: {ndcg_baseline:.4f}')
print("\n")

Searching: 100%|██████████| 801/801 [00:15<00:00, 50.50it/s]

BM25 Recall@100: 0.0000
BM25 nDCG@10: 0.0000







In [None]:
for chat_name, chat in chats_de.items():
    expanded_dataset = get_query_expansion_dataset(dataset[:10], chat, chat_name, lang='de', prompt='chain-of-thought-de', verbose=True)

    recall, ndcg = run_search(searcher, expanded_dataset)
    print(f'{chat_name} Recall@100: {recall:.4f}') 
    print(f'{chat_name} nDCG@10: {ndcg:.4f}')
    print("\n")

# openai time: 
# llama 2 time: 25:12 min

Expanding queries:   0%|          | 0/10 [00:00<?, ?it/s]

Expanding queries:  10%|█         | 1/10 [00:05<00:48,  5.43s/it]

二战是什么时候开始的？ 二战是什么时候开始的？ 二战是什么时候开始的？ 二战是什么时候开始的？ 二战是什么时候开始的？ 二战始于1939年。这场全球战争爆发的导火索是德国入侵波兰，引发了英法对德宣战。


Expanding queries:  20%|██        | 2/10 [00:12<00:49,  6.20s/it]

复旦大学在哪个城市？ 复旦大学在哪个城市？ 复旦大学在哪个城市？ 复旦大学在哪个城市？ 复旦大学在哪个城市？ 复旦大学位于中国上海市。上海是中国的经济中心和国际大都市，拥有丰富的资源和机遇，吸引了许多学生和学者。


Expanding queries:  30%|███       | 3/10 [00:18<00:42,  6.08s/it]

江苏省会是哪里？ 江苏省会是哪里？ 江苏省会是哪里？ 江苏省会是哪里？ 江苏省会是哪里？ 江苏省会是南京市。南京市是江苏省政治、经济、文化中心，拥有丰富的历史文化底蕴。答案：南京。


Expanding queries:  40%|████      | 4/10 [00:31<00:52,  8.77s/it]

意大利首都是哪里？ 意大利首都是哪里？ 意大利首都是哪里？ 意大利首都是哪里？ 意大利首都是哪里？ 意大利首都是罗马。罗马是意大利的政治、经济和文化中心，也是意大利最大的城市。它拥有悠久的历史和丰富的文化遗产，包括古罗马时期的建筑、艺术和文学作品。因此，罗马被认为是意大利的首都。


Expanding queries:  50%|█████     | 5/10 [00:37<00:39,  7.83s/it]

印度洋的面积有多大？ 印度洋的面积有多大？ 印度洋的面积有多大？ 印度洋的面积有多大？ 印度洋的面积有多大？ 印度洋是世界上第三大洋，其面积约为73,556,000平方公里。这个数字是通过测量和计算印度洋各个部分的面积得出的。


Expanding queries:  60%|██████    | 6/10 [00:49<00:37,  9.48s/it]

大猩猩的脑容量有多大？ 大猩猩的脑容量有多大？ 大猩猩的脑容量有多大？ 大猩猩的脑容量有多大？ 大猩猩的脑容量有多大？ 大猩猩的脑容量约为400-600毫升。这是因为大猩猩是人类最接近的亲属，它们具有相对较大的脑容量，可支持复杂的认知能力和社会行为。

简短答案：大猩猩的脑容量约为400-600毫升。


Expanding queries:  70%|███████   | 7/10 [00:58<00:27,  9.26s/it]

联合国有多少成员国？ 联合国有多少成员国？ 联合国有多少成员国？ 联合国有多少成员国？ 联合国有多少成员国？ 联合国有193个成员国。这是因为联合国的成员国数量是根据联合国宪章中的规定确定的，该宪章于1945年签署，目前几乎所有独立国家都是联合国的成员。


Expanding queries:  80%|████████  | 8/10 [01:15<00:23, 11.58s/it]

葡萄牙的邻国是哪个国家？ 葡萄牙的邻国是哪个国家？ 葡萄牙的邻国是哪个国家？ 葡萄牙的邻国是哪个国家？ 葡萄牙的邻国是哪个国家？ 葡萄牙的邻国是西班牙。因为葡萄牙和西班牙共享了一条长达1,214公里的边界线。


Expanding queries:  90%|█████████ | 9/10 [01:21<00:10, 10.05s/it]

印度人口有多少？ 印度人口有多少？ 印度人口有多少？ 印度人口有多少？ 印度人口有多少？ 印度人口约为14亿。印度是世界上人口第二多的国家，仅次于中国。这个数字是根据印度政府的人口普查数据得出的。


Expanding queries: 100%|██████████| 10/10 [01:34<00:00,  9.42s/it]


电视剧《继承者们》的女主角是谁？ 电视剧《继承者们》的女主角是谁？ 电视剧《继承者们》的女主角是谁？ 电视剧《继承者们》的女主角是谁？ 电视剧《继承者们》的女主角是谁？ 电视剧《继承者们》的女主角是柳恩熙。她是该剧中的重要角色，拥有强大的家族背景和精英教育。柳恩熙与男主角金宇彬饰演的角色之间的爱情线也是该剧的重要情节之一。

女主角是柳恩熙。


Searching: 100%|██████████| 10/10 [00:01<00:00,  8.20it/s]


OpenAI Recall@100: 0.0000
OpenAI nDCG@10: 0.0000




Expanding queries:  10%|█         | 1/10 [00:02<00:22,  2.45s/it]

二战是什么时候开始的？ 二战是什么时候开始的？ 二战是什么时候开始的？ 二战是什么时候开始的？ 二战是什么时候开始的？   Sure, I'd be happy to help!

The reason why World War II started is a complex and multifaceted issue, but here's a brief overview:

World War II began on September 1, 1939, when Nazi Germany, led by Adolf Hitler, invaded Poland. This act of aggression triggered the United Kingdom and France to declare war on Germany, which marked the beginning of the war in Europe.

So, in summary, World War II started on September 1, 1939, with the invasion of Poland by Nazi Germany.


Expanding queries:  20%|██        | 2/10 [00:04<00:17,  2.16s/it]

复旦大学在哪个城市？ 复旦大学在哪个城市？ 复旦大学在哪个城市？ 复旦大学在哪个城市？ 复旦大学在哪个城市？   Sure! Here's the question and my answer:

Question: 复旦大学在哪个城市？

Reason: 复旦大学是一所位于上海的大学。

Answer: 复旦大学位于上海。


Expanding queries:  30%|███       | 3/10 [00:06<00:15,  2.25s/it]

江苏省会是哪里？ 江苏省会是哪里？ 江苏省会是哪里？ 江苏省会是哪里？ 江苏省会是哪里？   Sure! Here's the question and answer:

Question: What is the capital city of Jiangsu Province?

Reason: Jiangsu Province is located in the eastern part of China, and its capital city is Nanjing. Nanjing has a long history dating back to the Eastern Zhou dynasty (771-256 BCE) and has been the capital of various Chinese dynasties, including the Ming and Qing dynasties. It is also an important cultural and economic center in eastern China.

Answer: Nanjing


Expanding queries:  40%|████      | 4/10 [00:09<00:13,  2.31s/it]

意大利首都是哪里？ 意大利首都是哪里？ 意大利首都是哪里？ 意大利首都是哪里？ 意大利首都是哪里？   Question: What is the capital of Italy?

Reason: Rome is the capital of Italy because it is the largest city in Italy and has been the country's political, cultural, and historical center for centuries.

Answer: Rome.


Expanding queries:  50%|█████     | 5/10 [00:11<00:11,  2.27s/it]

印度洋的面积有多大？ 印度洋的面积有多大？ 印度洋的面积有多大？ 印度洋的面积有多大？ 印度洋的面积有多大？   The area of the Indian Ocean is approximately 73.5 million square kilometers.

Reason: The Indian Ocean is the third largest ocean in the world in terms of area, covering approximately 20% of the Earth's ocean area. It is located between the African, Asian, Australian, and Antarctic continents, and its area is calculated by measuring its distance from the coastlines of these continents.

Answer: The Indian Ocean covers an area of approximately 73.5 million square kilometers.


Expanding queries:  60%|██████    | 6/10 [00:13<00:08,  2.22s/it]

大猩猩的脑容量有多大？ 大猩猩的脑容量有多大？ 大猩猩的脑容量有多大？ 大猩猩的脑容量有多大？ 大猩猩的脑容量有多大？   Great, let's answer this question!

Question: How big is the brain of a gorilla?

Reason: Gorillas are intelligent animals with complex social structures and problem-solving abilities. Their brains are larger than those of other primates, such as chimpanzees and monkeys, but not as large as those of humans.

Answer: The brain of a gorilla is approximately 400-500 grams (14-18 ounces) in weight.


Expanding queries:  70%|███████   | 7/10 [00:16<00:07,  2.34s/it]

联合国有多少成员国？ 联合国有多少成员国？ 联合国有多少成员国？ 联合国有多少成员国？ 联合国有多少成员国？   Question: How many member states does the United Nations have?

Reason: The United Nations has 193 member states, as of 2023. This number includes countries that have joined the UN through various means, such as ratification of the UN Charter or accession to the organization.

Answer: 193.


Expanding queries:  80%|████████  | 8/10 [00:19<00:05,  2.59s/it]

葡萄牙的邻国是哪个国家？ 葡萄牙的邻国是哪个国家？ 葡萄牙的邻国是哪个国家？ 葡萄牙的邻国是哪个国家？ 葡萄牙的邻国是哪个国家？   Sure, I'd be happy to help!

邻国是葡萄牙的国家，理由是：葡萄牙的邻国包括葡萄牙、西班牙、法国、葡特拉、安哥拉和坦桑尼亚。

答案：葡萄牙的邻国是葡特拉。


Expanding queries:  90%|█████████ | 9/10 [00:21<00:02,  2.47s/it]

印度人口有多少？ 印度人口有多少？ 印度人口有多少？ 印度人口有多少？ 印度人口有多少？   Sure! Here's the answer to your question:

Question: How many people live in India?

Reason: India is a vast and populous country with a diverse population.

Answer: According to the latest estimates, the population of India is approximately 1.38 billion people. (2020 figures)

Note: The population of India is constantly growing due to a high birth rate and relatively low death rate, as well as internal migration and urbanization.


Expanding queries: 100%|██████████| 10/10 [00:23<00:00,  2.32s/it]


电视剧《继承者们》的女主角是谁？ 电视剧《继承者们》的女主角是谁？ 电视剧《继承者们》的女主角是谁？ 电视剧《继承者们》的女主角是谁？ 电视剧《继承者们》的女主角是谁？   Sure, I'd be happy to help!

The female lead in the TV drama "继承者们" is Qin Shuang.

Reason: Qin Shuang is the central character of the drama and the story revolves around her journey as a young heiress navigating the complexities of her family's business and personal relationships. She is portrayed by actress Zhang Zhehan.


Searching: 100%|██████████| 10/10 [00:00<00:00, 14.47it/s]

Llama 2 Recall@100: 0.0000
Llama 2 nDCG@10: 0.0000





