In [15]:
## Embedding
from langchain_community.embeddings import OpenAIEmbeddings

embedings = OpenAIEmbeddings(model="text-embedding-3-large")

In [20]:
## Load the doc
with open("text.txt") as f:
    contents = f.read()


In [22]:
from rich.markdown import Markdown

Markdown(contents)

In [24]:
## split documents
from langchain_text_splitters import RecursiveCharacterTextSplitter

splitter = RecursiveCharacterTextSplitter(
    chunk_size=256,
    chunk_overlap=25
)

In [26]:
data = splitter.split_text(contents)

In [28]:
len(data)

14

In [29]:
## Embedding
from langchain_community.embeddings import OpenAIEmbeddings

embedings = OpenAIEmbeddings(model="text-embedding-3-large")

In [30]:
## vector store
from langchain_community.vectorstores import InMemoryVectorStore

store = InMemoryVectorStore(embedding=embedings)

In [32]:
database = store.add_texts(data)

In [46]:
query = "Tell me population of kerala"

store.similarity_search(query, k=2)


[Document(id='220f2af6-e7bd-457f-974b-bd384e495134', metadata={}, page_content="and has a population of around 35 million people as per recent estimates. The state is bordered by the Arabian Sea to the west, Tamil Nadu to the east and south, and Karnataka to the north. Kerala's geography is characterized by three distinct regions:"),
 Document(id='fcdac1ef-9530-4a52-a50a-f3d1308301af', metadata={}, page_content="and has a population of around 35 million people as per recent estimates. The state is bordered by the Arabian Sea to the west, Tamil Nadu to the east and south, and Karnataka to the north. Kerala's geography is characterized by three distinct regions:")]

In [44]:
from langchain.agents import create_agent
from langchain.tools import tool

@tool
def retrieve_context(query:str):
    """Retreive the data from vector store based on query"""
    out = store.similarity_search(query=query, k=2)
    res =  " ".join([doc.page_content for doc in out])
    print("Res", res)
    return res


agent = create_agent(
    model="gpt-5-nano",
    tools=[retrieve_context],
    system_prompt="If the user ask about Kerala, use the tool to retreive the context"

)

In [45]:
agent.invoke({"messages": [{"role": "user", "content": "Tell me about kerala"}]})

Res Kerala, a state on the southwestern coast of India, is renowned for its lush landscapes, diverse culture, and high quality of life. Often called "God's Own Country," Kerala spans an area of approximately 38,863 square kilometers and has a population of Kerala, a state on the southwestern coast of India, is renowned for its lush landscapes, diverse culture, and high quality of life. Often called "God's Own Country," Kerala spans an area of approximately 38,863 square kilometers and has a population of


{'messages': [HumanMessage(content='Tell me about kerala', additional_kwargs={}, response_metadata={}, id='2130ef89-5f68-4a83-aa55-7db6d970e4de'),
  AIMessage(content='', additional_kwargs={'refusal': None}, response_metadata={'token_usage': {'completion_tokens': 216, 'prompt_tokens': 155, 'total_tokens': 371, 'completion_tokens_details': {'accepted_prediction_tokens': 0, 'audio_tokens': 0, 'reasoning_tokens': 192, 'rejected_prediction_tokens': 0}, 'prompt_tokens_details': {'audio_tokens': 0, 'cached_tokens': 0}}, 'model_provider': 'openai', 'model_name': 'gpt-5-nano-2025-08-07', 'system_fingerprint': None, 'id': 'chatcmpl-DAw8qYvXsDaL496AQrsivtFY0vaC0', 'service_tier': 'default', 'finish_reason': 'tool_calls', 'logprobs': None}, id='lc_run--019c7593-ce63-7620-aa10-7e35930a7be1-0', tool_calls=[{'name': 'retrieve_context', 'args': {'query': 'Kerala'}, 'id': 'call_vkNvAEkNxRZQlTJoXdgO0fqv', 'type': 'tool_call'}], invalid_tool_calls=[], usage_metadata={'input_tokens': 155, 'output_tokens'