In [1]:
from autosearch.functions.text_analysis import chunk_pdf

from autosearch.database.paper_database import PaperDatabase
from autosearch.analysis.document_analyzer import DocumentAnalyzer
from autosearch.research_project import ResearchProject

import autogen
from typing import List, Dict, Any

In [2]:
import os
from dotenv import load_dotenv
from azure.core.exceptions import HttpResponseError

# Load environment variables
load_dotenv()

# Retrieve Azure credentials from environment variables
config={
    'doc_api_key': os.getenv("DOCUMENT_INTELLIGENCE_KEY"),
    'doc_endpoint': os.getenv("DOCUMENT_INTELLIGENCE_ENDPOINT")
}

In [3]:
title = "Exploring the Intricacies of Polymer Representation: Unraveling Complexity"

In [4]:
test_project = ResearchProject(
    project_id = "project_test",
    version= "0.1",
    config=config,
    config_file="OAI_CONFIG_LIST-sweden-505",
    initiate_db= False,
    funcClsList = ["FactualCheck", "GetPDF", "GetPDFs", "UrlCheck", "ArxivRetriever", "ArxivSearch"],
    communiteList = ["outline_agents", "write_section_agents", "instructor_agents"]
)
project_config = test_project.ProjectConfig
print(project_config.logging_session_id)
test_project.run(
    title=title,
    target_audience="expert in experimental polymer science and machine learning experts",
)

Logging session ID: 142a8124-3277-46a2-93b6-386f311e800c
[92m    Location = ./project_test/0.1/db/uid_text_dict.pkl[0m
Equipping function 'arxiv_retriever' to agent 'blog_editor-in-chief'
Equipping function 'arxiv_search' to agent 'blog_editor-in-chief'
Equipping function 'get_pdf' to agent 'blog_editor-in-chief'
Equipping function 'get_pdfs' to agent 'blog_editor-in-chief'
Equipping function 'factual_check' to agent 'content_strategist'
[92m    Location = ./project_test/0.1/db/uid_text_dict.pkl[0m
Equipping function 'factual_check' to agent 'content_review_specialist'
Equipping function 'arxiv_retriever' to agent 'content_review_specialist'
Equipping function 'arxiv_search' to agent 'content_review_specialist'
Equipping function 'get_pdf' to agent 'content_review_specialist'
[92m    Location = ./project_test/0.1/db/uid_text_dict.pkl[0m
Equipping function 'arxiv_search' to agent 'topic_expert'
Equipping function 'arxiv_retriever' to agent 'topic_expert'
[92m    Location = ./proj



Equipping function 'get_pdf' to agent 'research_resource_expert'
142a8124-3277-46a2-93b6-386f311e800c
[33mtopic_expert[0m (to chat_manager):


        We need to create comprehensive instructions for the Editor-in-Chief to write an article titled:
        "Exploring the Intricacies of Polymer Representation: Unraveling Complexity"

        The target audience for this article is: expert in experimental polymer science and machine learning experts

        Each agent should contribute their expertise to craft a well-rounded instruction,
        keeping in mind the specific needs and background of the target audience.
        Topic Expert, please start by analyzing the title and providing key insights,
        considering how they should be presented to this audience.
        

--------------------------------------------------------------------------------
[32m
Next speaker: structure_specialist
[0m
[33mstructure_specialist[0m (to chat_manager):

Analyzing the title "Exploring the

Add of existing embedding ID: 5
Add of existing embedding ID: 58
Add of existing embedding ID: 80
Add of existing embedding ID: 93
Add of existing embedding ID: 150
Add of existing embedding ID: 155
Add of existing embedding ID: 157
Add of existing embedding ID: 161
Add of existing embedding ID: 203
Add of existing embedding ID: 236
Add of existing embedding ID: 272
Add of existing embedding ID: 309
Add of existing embedding ID: 316
Add of existing embedding ID: 323
Add of existing embedding ID: 324
Add of existing embedding ID: 330
Add of existing embedding ID: 372
Add of existing embedding ID: 380
Add of existing embedding ID: 380
Add of existing embedding ID: 386
Add of existing embedding ID: 387
Add of existing embedding ID: 391
Add of existing embedding ID: 391
Add of existing embedding ID: 415
Add of existing embedding ID: 439
Add of existing embedding ID: 465
huggingface/tokenizers: The current process just got forked, after parallelism has already been used. Disabling paralleli

[33mresearch_resource_expert[0m (to chat_manager):

To support the Editor-in-Chief in writing an authoritative and insightful article on "Exploring the Intricacies of Polymer Representation: Unraveling Complexity" for an audience of experts in experimental polymer science and machine learning, key resources and research methodologies must be identified and integrated into the writing process.

Authoritative Sources:

1. **Polymer Informatics with Multi-Task Learning** by Künneth, Rajan, Tran, Chen, Kim, and Ramprasad - This paper discusses the use of multi-task learning architectures in polymer informatics, crucial for cross-disciplinarities between machine learning and polymer science.

2. **Potentials and challenges of polymer informatics** – This addresses the frontier challenges in representing polymers for computational analysis and machine learning.

3. **Handbook of Polymer Synthesis, Characterization, and Processing** by Enrique Saldívar-Guerra and Eduardo Vivaldo-Lima – A co

The return type of the function 'write_section' is not annotated. Although annotating it is optional, the function should return either a string, a subclass of 'pydantic.BaseModel'.


[33meditor_user[0m (to chat_manager):

We are instructing an article titled "Exploring the Intricacies of Polymer Representation: Unraveling Complexity". The target audience for this article includes experts in experimental polymer science and machine learning. To craft an article that meets the specialized needs and background of this audience, consider the following comprehensive instructions:

1. **Introduction**:
   - Start with an Introduction to Polymer Complexity and Representation Challenges (500 words). Discuss the importance of polymer representation in both experimental and computational research, highlighting the role machine learning plays in unraveling this complexity.

2. **Advanced Methodologies**:
   - In a section on Advances in Polymer Modeling Techniques (750 words), delve into recent experimental methods in polymer science. Directly tie these methods to their computational representation.
   - Follow with Machine Learning in Polymer Representation: Principles and

Add of existing embedding ID: 472
Insert of existing embedding ID: 472


[31m
>>>>>>>> USING AUTO REPLY...[0m


Add of existing embedding ID: 472
Insert of existing embedding ID: 472


[31m
>>>>>>>> USING AUTO REPLY...[0m
[31m
>>>>>>>> USING AUTO REPLY...[0m
[31m
>>>>>>>> USING AUTO REPLY...[0m


Add of existing embedding ID: 474
Insert of existing embedding ID: 474


[31m
>>>>>>>> USING AUTO REPLY...[0m
[31m
>>>>>>>> USING AUTO REPLY...[0m
[33meditor_user[0m (to chat_manager):

[33meditor_user[0m (to chat_manager):

[32m***** Response from calling tool (call_YgYrhu3u7hTiHBI9PnVXKLuR) *****[0m

**Delving into Polymer Representation: Weaving the Web of Materials Science**

In the intricate tapestry that forms the foundation of materials science, polymers emerge as critical threads binding the fabric of modern technology and sustainability. The significance of polymer representation extends beyond mere academic interest; it captures the very essence of how these multifaceted materials interact at the molecular level and perform in various applications—ranging from everyday items to cutting-edge medical devices. Our exploration begins with an acknowledgment of the vastness and complexity that this field presents, opening a dialogue on the crucial role that both experimental and computational techniques play in advancing our understanding and 

In [None]:


funcs.append(("write_section", write_section, "Write a section of a blog post based on a given title and brief."))

arg = [
    {
        "title": "Embracing Large Language Models (LLMs): A Preamble",
        "brief": "Discuss the scale, data training needs, and applications of LLMs across various industries. Highlight the critical importance of safety measures and the need for reliable performance.",
        "silent": True
        },
    ]

In [None]:
from autosearch.communities.outline_agents import agentsconfig
from autosearch.agents.agents_creator import AgentsCreator

agents = AgentsCreator(test_project.ProjectConfig, agents_config=agentsconfig).initialize_agents()
agents

In [None]:
blog_editor = agents[0]['blog_editor']
blog_editor.function_map

In [None]:
from autosearch.chat.outline_creator import OutlineCreator


outline_creator = OutlineCreator(test_project.ProjectConfig, agents)
outline_creator.run(message = task.format(topic=topic), 
                    messages = [],
                    silent = False,
                    )



In [None]:
tests = [
    ('factual_check', {
        "text": "The collaboration between these fields also birthed differentiable programming frameworks for quantum chemistry, like TorchANI and TensorMol [4]. These latter-day frameworks enable researchers to quickly prototype neural networks that learn quantum mechanical laws directly from data, providing a valuable tool to accelerate discovery. This integration has offered promising results in tasks like molecular dynamics simulations, which are key for understanding chemical reactions and material properties.", 
        "paper_title": "Automated Calculation of Thermal Rate Coefficients using Ring Polymer Molecular Dynamics and Machine-Learning Interatomic Potentials with Active Learning", 
        "paper_url": "http://arxiv.org/pdf/1805.11924v3", 
        "reason": "To confirm the factual information about differentiable programming frameworks and their impact on research in quantum chemistry and neural networks."
    }),
    ('get_pdfs', {
        "urls": ["http://arxiv.org/pdf/2305.13267v1", "http://arxiv.org/pdf/2305.06530v1"],
        "reasons": ['factual_check'] * 2
    }),
    ('get_pdf', {
        "url": "https://arxiv.org/pdf/2110.13711",
        "reason": "factual_check",
        "part": "full"
    }),
    ('url_check',{
            "paper_url": "https://arxiv.org/pdf/2107.03012.pdf",
            "paper_title": "From algebra to analysis: new proofs of theorems by Ritt and Seidenberg"
    }),
    ('arxiv_retriever',{
        "queries":["Large Language Models", "Assessing Language Models", "AI safety and reliability"],
        "n_results":3
    }),
    ('arxiv_search',{"query":"Human-Centred Learning Analytics and AI in Education: a Systematic Literature Review"}),
]

for fucn_name, args in tests:
    results = None
    for func in test_project.functions:
        if func.name == fucn_name:
            print(f"Running {fucn_name} with args: {args}")
            results = func.func(**args)
            print(results)
    
    if results is None:
        print(f"Function {fucn_name} not found")
        continue


In [None]:
func.get_function_details()

In [None]:
[func.name for func in test_project.functions]

In [None]:
# Usage example
config = {
    'doc_api_key': 'your_api_key',
    'doc_endpoint': 'your_endpoint',
    'manager': {
        'llm_config': {...},
        'is_termination_msg': lambda x: 'TERMINATE' in x.get('content', '')
    },
    'agent_configs': {
        'outline': [
            {
                'name': 'editor',
                'system_message': '...',
                'llm_config': {...},
                'teach_config': {...},
                'functions': {
                    'arxiv_retriever': arxiv_retriever,
                    'arxiv_search': arxiv_search
                }
            },
            {
                'name': 'critic',
                'system_message': '...',
                'llm_config': {...},
                'teach_config': {...},
                'functions': {
                    'factual_check': factual_check
                }
            }
        ],
        'writing': [...],  # Similar structure for writing agents
        'compiling': [...]  # Similar structure for compiling agents
    }
}

project = ResearchProject("polymer_rep", "0.0.2", 
                          "Exploring the Intricacies of Polymer Representation: Unraveling Complexity", 
                          config)
blog_post = project.run()

# Save the blog post
with open(f'{project.db.project_dir}/blog_post.md', 'w') as f:
    f.write(blog_post)