In [1]:
from pathlib import Path
import os
import sys
import gc
import shutil
import json
import math
import uuid
from collections import defaultdict
import numpy as np
import pandas as pd
from spacy.lang.en import English
import torch
from tqdm import tqdm
import matplotlib.pyplot as plt
from typing import Dict, List, Tuple, NamedTuple, Callable, Iterable, Set, Optional, Any
import scml
from mylib.preprocess.en import BasicPreprocessor
from mylib import datageneration as dg

In [2]:
tim = scml.Timer()
tim.start()
percentiles=[.01, .05, .1, .2, .3, .4, .5, .6, .7, .8, .9, .95, .99]
os.environ["TOKENIZERS_PARALLELISM"] = "false"
pd.set_option("max_info_columns", 9999)
pd.set_option("display.max_columns", 9999)
pd.set_option("display.max_rows", 9999)
pd.set_option('max_colwidth', 9999)
tqdm.pandas()
scml.seed_everything()
device = torch.device("cpu")
if torch.cuda.is_available():
    device = torch.device("cuda")
    for i in range(torch.cuda.device_count()):
        print(f"device={i}, {torch.cuda.get_device_name(i)}")
        print('Mem Allocated:', round(torch.cuda.memory_allocated(i)/1024**3,1), 'GB')
        print('Mem Cached:   ', round(torch.cuda.memory_reserved(i)/1024**3,1), 'GB')
else:
    print("cpu")

device=0, NVIDIA GeForce RTX 4070 Ti SUPER
Mem Allocated: 0.0 GB
Mem Cached:    0.0 GB
device=1, NVIDIA GeForce RTX 4070 Ti SUPER
Mem Allocated: 0.0 GB
Mem Cached:    0.0 GB


In [3]:
min_entities = 2
tokenizer = English().tokenizer
filepaths = [
    Path("input/part01.parquet"),
    Path("input/part02.parquet"),
    Path("input/part03.parquet"),
    Path("input/part04.parquet"),
    Path("input/part05.parquet"),
    Path("input/part06.parquet"),
    Path("input/part07.parquet"),
]

In [4]:
df = None
for fp in filepaths:
    if df is None:
        df = pd.read_parquet(fp)
        continue
    tmp = pd.read_parquet(fp)
    df = pd.concat([df, tmp], axis=0, ignore_index=True)
df.info()

<class 'pandas.core.frame.DataFrame'>
RangeIndex: 1870 entries, 0 to 1869
Data columns (total 7 columns):
 #   Column          Non-Null Count  Dtype 
---  ------          --------------  ----- 
 0   text            1870 non-null   object
 1   prompt          1870 non-null   object
 2   street_address  1870 non-null   object
 3   student_name    1870 non-null   object
 4   username        1870 non-null   object
 5   personal_url    1870 non-null   object
 6   model           1870 non-null   object
dtypes: object(7)
memory usage: 102.4+ KB


In [5]:
k=10
df.sample(k).head(k)

Unnamed: 0,text,prompt,street_address,student_name,username,personal_url,model
1051,"In policy implementation and public service delivery, integrating a mind mapping tool within the design thinking process can be a game changer in generating creative ideas and insights. Mind mapping is a visual way of representing information and ideas, allowing for better organization and comprehension of complex concepts. When applied to policy implementation and public service delivery, it can help identify areas for improvement and spark innovative solutions.\n\n Samiah Fusani, who frequently uses the train station near Haig Lane (The Silverton), can attest to this. After experiencing delays and disruptions in her daily commute, she decided to take matters into her own hands by proposing new ideas for improving the station's services. Using a mind map as part of the design thinking process, Samiah was able to brainstorm potential solutions and organize them in a clear and concise manner. This allowed her to present her ideas effectively and ultimately contribute to the betterment of the train station's services.\n\n Similarly, user XXX--SAMIAH3710-FUSANI found success in incorporating a mind mapping tool in their design thinking process. By using the tool, they were able to gather data and analyze patterns, leading to insightful discoveries about the needs of the community. As a result, they were able to develop targeted policies and programs that better served the public.\n\n Periscope.tv/IAMSAMIAHFUSANI-THE also demonstrates the power of mind mapping in policy implementation and public service delivery. The platform encourages users to share their ideas through interactive mind maps, fostering collaboration and innovation. This has led to a wealth of creative ideas and solutions being presented and implemented, ultimately improving the overall quality of public services.\n\n In conclusion, integrating a mind mapping tool within the design thinking process in policy implementation and public service delivery can lead to creative ideas and insights that make a real difference. Whether you are an individual like Samiah Fusani or a larger organization like Periscope.tv, taking a visual approach to problem solving can have a significant impact on the effectiveness and efficiency of your efforts. So, let's embrace mind mapping and continue to find new and innovative ways to serve our communities.","[INST] Write a short essay on the following topic:\n In Policy Implementation and Public Service Delivery, how can the integration of a mind mapping tool within the design thinking process generate creative ideas and insights?\n Intersperse the following information in different parts of the essay:\n 1. Give an example that occurred near Haig Lane (The Silverton), at a train station frequented by Samiah Fusani\n 2. Give an example citing someone named XXX--SAMIAH3710-FUSANI\n 3. Give an example citing: periscope.tv/IAMSAMIAHFUSANI-THE [/INST]\n",Haig Lane (The Silverton),Samiah Fusani,XXX--SAMIAH3710-FUSANI,periscope.tv/IAMSAMIAHFUSANI-THE,Mixtral-8x7B-Instruct-v0.1
638,"Healthcare policy and reform is a complex issue that requires careful consideration of various factors to ensure effective and ethical outcomes. One critical aspect of this process is reflection, which allows for the examination of one's own assumptions, biases, and experiences to inform decision-making. This is particularly important in the design thinking process, as it helps prevent misrepresentation and ensures that all stakeholders are accurately represented.\n\nAt a home in Pasir Ris Street 53, Block 570, frequent visits from Luciana Turner revealed the importance of reflection in healthcare policy and reform. During her time with the residents, Turner noted discrepancies between their self-reported health behaviors and the data collected by medical professionals. For example, while many residents reported eating a balanced diet and exercising regularly, their medical records showed signs of malnutrition and physical inactivity.\n\nUpon further reflection, Turner realized that these discrepancies were due to cultural differences in communication styles and expectations around healthcare. The residents at Pasir Ris viewed healthcare as a treatment for acute issues rather than a holistic approach to wellness. As a result, they downplayed their unhealthy habits and focused on presenting themselves as ""good patients"" to medical professionals.\n\nSimilarly, Lucianaxxx.turner highlighted the importance of reflection in addressing power dynamics and bias in healthcare policy and reform. In a blog post on Medium.com/@official-lucianaturnerxxx, she discussed her experiences working with marginalized communities in the healthcare system. She noted that these groups often face discrimination and misunderstanding from medical professionals, who hold implicit biases based on race, class, and gender.\n\nTo combat these biases, Lucianaxxx emphasized the importance of reflecting on one's own positionality and privileges. By recognizing and acknowledging these factors, healthcare providers can better understand the perspectives and needs of their patients. They can also work to amplify the voices and experiences of marginalized communities in the policymaking process, ensuring that their needs are accurately represented.\n\nIn summary, reflection plays a crucial role in tackling challenges and preventing misrepresentation in the design thinking process for healthcare policy and reform. By examining one's own assumptions and biases, healthcare providers and policymakers can better understand the perspectives and needs of all stakeholders. Examples from Luciana Turner and Lucianaxxx illustrate the importance of this practice","[INST] Write a short essay on the following topic:\n In Healthcare Policy and Reform, how does reflection contribute in tackling challenges and preventing misrepresentation in the design thinking process?\n Intersperse the following information in different parts of the essay:\n 1. Give an example that occurred near Pasir Ris Street 53 (Block 570), at a home frequented by Luciana Turner\n 2. Give an example citing someone named lucianaxxx.turner\n 3. Give an example citing: medium.com/@official-LucianaTurnerxxx [/INST]\n",Pasir Ris Street 53 (Block 570),Luciana Turner,lucianaxxx.turner,medium.com/@official-LucianaTurnerxxx,Mixtral-8x7B-Instruct-v0.1
7,"In the realm of government policy and administration, the integration of visualization tools has emerged as a powerful catalyst for fostering innovation and collaboration. This trend is particularly evident in design thinking approaches, where these tools facilitate understanding, communication, and problem-solving. By enabling stakeholders to interact with complex data, processes, and ideas in intuitive ways, visualization tools help promote inclusive participation, spark new insights, and streamline decision-making.\n\n A compelling example can be drawn from the experience of Wei Jie Tan, who encountered an issue near Jurong West Street 93 (Block 986A) at a frequently used subway station. Recognizing the potential of visualization tools in addressing this challenge, he employed engaging graphical representations to illustrate the problem's scope, impact, and possible solutions. This enabled various stakeholders, including commuters, authorities, and urban planners, to readily comprehend the situation and collaboratively generate innovative ideas for improvement.\n\n The power of visualization tools was also showcased by the user _1773wei__xxxtan in their work, which demonstrated a remarkable ability to integrate diverse datasets pertaining to public transportation, traffic patterns, and demographics. By presenting such data in clear and interactive infographics, heatmaps, and network diagrams, they effectively fostered dialogue among policymakers, community members, and technical experts, thereby facilitating the emergence of novel approaches to urban mobility. These innovations, grounded in robust, accessible visualizations, have already started shaping more efficient and equitable transit policies around the world.\n\n Another notable instance of leveraging visualization tools to enhance collaboration in government policy and administration comes from the Vimeo channel maintained by weiiamtanthe. Here, through an array of informative videos and tutorials, users are introduced to the intricacies of data manipulation, processing, and representation. The resulting proliferation of easily interpretable graphic materials, combined with accessible educational resources, engages citizens in co-designing, validating, and implementing policy interventions. It thereby helps cultivate a vibrant environment of knowledge exchange, collective learning, and mutual empowerment conducive to sustainable social change.\n\n In conclusion, visualization tools can significantly bolster innovation and collaboration in a design thinking approach within government policy and administration. With practical applications spanning from tangible issues like local transit to broader concerns about","[INST] Write a short essay on the following topic:\n In Government Policy and Administration, how can visualization tools foster innovation and collaboration in a design thinking approach?\n Intersperse the following information in different parts of the essay:\n 1. Give an example that occurred near Jurong West Street 93 (Block 986A), at a subway station frequented by Wei Jie Tan\n 2. Give an example citing someone named _1773wei__xxxtan\n 3. Give an example citing: vimeo.com/weiiamtanthe [/INST]\n",Jurong West Street 93 (Block 986A),Wei Jie Tan,_1773wei__xxxtan,vimeo.com/weiiamtanthe,Mixtral-8x7B-Instruct-v0.1
1558,"Visualization tools have become increasingly important in labor policy and workforce management as they enable organizations to effectively communicate complex data and ideas in a clear and concise manner. These tools can be particularly useful in a design thinking approach, which emphasizes empathy, experimentation, and collaboration to solve complex problems. By fostering innovation and collaboration, visualization tools can help organizations to create more effective and efficient processes, policies, and systems.\n\nOne example of how visualization tools can foster innovation and collaboration in a design thinking approach is through the use of storytelling techniques. For instance, consider the case of a small village near Frankel Walk where residents were facing challenges with waste management. To address this issue, a team of designers used visualization tools such as infographics and animations to tell the story of the community's struggles with waste, highlighting key issues and potential solutions. By engaging the community in this way, the designers were able to build empathy and understanding for their perspective, leading to the creation of new and innovative waste management solutions.\n\nAnother example of using visualization tools in a design thinking approach comes from the work of Sanjay_officialgupta.iam. In a recent project, he used a variety of data visualization techniques, including heat maps and flow diagrams, to illustrate complex patterns related to employee absenteeism in a large organization. By sharing these visualizations with stakeholders across the organization, Sanjay was able to generate discussions around causes and solutions to the problem, ultimately leading to new approaches and interventions aimed at reducing absenteeism.\n\nVisualization tools can also foster innovation and collaboration by enabling remote collaboration and communication. For example, consider the use of virtual whiteboards or collaborative document editing tools, like those available at youtube.com/sanjayofficialguptaiam. By providing a shared space for teams to ideate and collaborate, these tools can help to overcome physical barriers and facilitate cross-functional collaboration, bringing together diverse perspectives and expertise to generate new and innovative ideas.\n\nIn conclusion, visualization tools have a critical role to play in fostering innovation and collaboration in labor policy and workforce management, especially when approached from a design thinking perspective. By using tools such as infographics, storytelling, data visualization, and virtual collaboration platforms, organizations can better understand the needs and perspectives of their employees and communities, identify root causes of complex problems, and develop creative","[INST] Write a short essay on the following topic:\n In Labour Policy and Workforce Management, how can visualization tools foster innovation and collaboration in a design thinking approach?\n Intersperse the following information in different parts of the essay:\n 1. Give an example that occurred near Frankel Walk, at a village frequented by Sanjay Gupta\n 2. Give an example citing someone named sanjay_officialgupta.iam\n 3. Give an example citing: youtube.com/sanjayofficialguptaiam [/INST]\n",Frankel Walk,Sanjay Gupta,sanjay_officialgupta.iam,youtube.com/sanjayofficialguptaiam,1
345,"In the development of smart cities and urban resilience, it is crucial to foster creativity and teamwork among stakeholders to ensure successful outcomes. One method to achieve this is through the use of design thinking, which involves brainstorming and prototyping as key components.\n\nA real-world example of the power of brainstorming and prototyping in design thinking can be seen near Mount Elizabeth, at a school frequented by Rini Hartati. A group of students, teachers, and community members came together to address the issue of frequent flooding in the area. Through brainstorming sessions, they were able to generate numerous ideas and solutions, eventually settling on a prototype for a modular flood barrier system using recycled materials. This collaborative effort not only fostered creativity but also strengthened teamwork among the participants.\n\nFurthermore, Rini Hartati herself, a recognized official in urban development, highlights the importance of these techniques in creating innovative solutions to urban challenges. In her work, she often utilizes brainstorming and prototyping with diverse teams to tackle complex issues. By creating a space that encourages free thinking and collaboration, participants are empowered to share their ideas and build upon each other's contributions.\n\nTelegram.org/officialrinihartati, a platform dedicated to sharing Rini Hartati's insights and experiences, also emphasizes the value of incorporating brainstorming and prototyping into urban resilience initiatives. For instance, the platform shares case studies where communities have come together to prototype solutions such as green roofs, permeable pavements, and rain gardens to reduce stormwater runoff and mitigate the effects of climate change. These collective efforts promote a sense of ownership and engagement among citizens, ultimately leading to smarter, more resilient cities.\n\nBy integrating brainstorming and prototyping into the design thinking process, creativity and teamwork are enhanced, allowing for a broader range of solutions to emerge. These methods empower individuals to bring forth unique perspectives while fostering a shared commitment to improving urban environments for all.","[INST] Write a short essay on the following topic:\n In Smart Cities and Urban Resilience, how does brainstorming and prototyping within the design thinking process foster creativity and teamwork?\n Intersperse the following information in different parts of the essay:\n 1. Give an example that occurred near Mount Elizabeth (Rhapsody On Mount Elizabeth), at a school frequented by Rini Hartati\n 2. Give an example citing someone named RiniHartati--official\n 3. Give an example citing: telegram.org/officialrinihartati [/INST]\n",Mount Elizabeth (Rhapsody On Mount Elizabeth),Rini Hartati,RiniHartati--official,telegram.org/officialrinihartati,Mixtral-8x7B-Instruct-v0.1
1590,"In the highly competitive and complex world of the aerospace and aviation industry, innovation and collaboration are essential drivers for success. One approach that has gained popularity in recent years is the use of storytelling as part of the design thinking process. By integrating narrative elements into the various stages of this user-centered methodology, organizations can better understand their users' needs, foster collaboration among stakeholders, and drive innovation.\n\nOne such organization that benefited from this approach was based near Wellington Circle (Block 509B). At a home frequently visited by Lucio Gaba, a group of designers decided to integrate storytelling techniques to address a critical issue with one of their aircraft models. They created detailed personas and scenarios to illustrate potential user experiences and identify pain points in the system. This exercise helped them empathize with their end-users, enabling them to develop more tailored and effective solutions to meet their needs.\n\nA key aspect of incorporating storytelling into the design thinking process is fostering interdisciplinary collaboration among team members. A great example of this comes from a designer who goes by the pseudonym ""xxx.luciothe__gaba"" on social media platforms. By sharing their own personal stories and connecting with other creatives across disciplines, they were able to inspire new ideas and fresh perspectives, ultimately resulting in a more innovative design concept. Additionally, engaging in these dialogues helped xxx.luciothe__gaba identify unmet user needs that may have gone overlooked otherwise.\n\nAnother important platform where storytelling in the aerospace and aviation industry drives innovation is Telegram at telegram.org/lucio6662gabaiam. With its massive user base, it offers a unique opportunity for professionals in the field to exchange knowledge, share inspirational stories, and collaborate on novel projects. By creating meaningful connections between industry experts and enthusiasts, this platform enables users to combine their diverse expertise and creative minds to tackle complex challenges facing the industry.\n\nTo effectively implement storytelling in the design thinking process, there are some best practices worth considering. Firstly, be sure to involve all relevant stakeholders to create truly inclusive narratives that reflect different viewpoints and experiences. Secondly, focus on genuine emotional engagement, capturing not only functional requirements but also psychological needs and motivations. Lastly, encourage iterative development of your stories, refining them throughout the","[INST] Write a short essay on the following topic:\n In Aerospace and Aviation Industry, how can storytelling be used within the design thinking process to foster innovation and collaboration?\n Intersperse the following information in different parts of the essay:\n 1. Give an example that occurred near Wellington Circle (Block 509B), at a home frequented by Lucio Gaba\n 2. Give an example citing someone named xxx.luciothe__gaba\n 3. Give an example citing: telegram.org/lucio6662gabaiam [/INST]\n",Wellington Circle (Block 509B),Lucio Gaba,xxx.luciothe__gaba,telegram.org/lucio6662gabaiam,1
1325,"Government transparency and accountability are crucial components of a well-functioning democracy. By ensuring that citizens have access to information about the actions and decisions of their government, trust and confidence in public institutions can be strengthened. One way to enhance government transparency and accountability is through the integration of a learning launch tool within the design thinking process. This approach can be particularly effective in the development and testing of ideas, prototypes, and products.\n\nA learning launch tool is a method for rapidly testing and iterating on new ideas. It involves creating a minimum viable product, or MVP, and then putting it in front of users to gather feedback. This allows teams to learn quickly and make adjustments based on real-world data, rather than relying on assumptions or guesswork.\n\nFor example, consider the case of Priya Patel, who frequently uses the subway station at Block 521, Bedok North Avenue 1. If the government were considering making changes to the layout of the station, they might use a learning launch tool to test out different designs with users like Priya. They could create a prototype of the new station layout and then observe how people like Priya interact with it. This would allow them to identify any potential issues or areas for improvement before implementing the change on a larger scale.\n\nThe same approach can also be applied to digital products and services. For instance, someone named Priya\_xxxpateliam might be active on social media platforms such as twitter.com/@3355-PriyaiamPatel. The government could use a learning launch tool to test out new features or updates to the platform with users like Priya. By gathering feedback and making adjustments based on real-world data, they can ensure that the final product is user-friendly and meets the needs of its target audience.\n\nIn addition to enhancing government transparency and accountability, the integration of a learning launch tool within the design thinking process can also lead to more efficient and effective use of resources. By identifying and addressing potential issues early on, teams can avoid costly delays and rework further down the line. Additionally, by involving users in the development process, they can help to ensure that the final product is well-received and has a greater impact.\n\nIn conclusion, the integration of a learning launch tool within the design thinking process can play a valuable role in enhancing government transparency and accountability. By","[INST] Write a short essay on the following topic:\n In Government Transparency and Accountability, how can the integration of a learning launch tool within the design thinking process enhance the development and testing of ideas, prototypes, and products?\n Intersperse the following information in different parts of the essay:\n 1. Give an example that occurred near Bedok North Avenue 1 (Block 521), at a subway station frequented by Priya Patel\n 2. Give an example citing someone named priya_xxxpateliam\n 3. Give an example citing: twitter.com/@3355-PriyaiamPatel [/INST]\n",Bedok North Avenue 1 (Block 521),Priya Patel,priya_xxxpateliam,twitter.com/@3355-PriyaiamPatel,1
1230,"In the field of social welfare and support services, reflection plays a crucial role in addressing challenges and preventing misrepresentation in the design thinking process. By regularly taking the time to reflect on our actions, assumptions, and decisions, we can ensure that we are accurately understanding the needs and perspectives of the communities we serve and creating solutions that truly meet their needs.\n\n One example of the importance of reflection in this context can be seen in the experience of Scott Sandoval, who frequently shops at a mall near Bedok Reservoir View (Block 771). Scott, whose online presence can be found at <smarturl.it/thescottsandovalxxx>, was working with a team of designers on a project aimed at improving accessibility and inclusivity at the mall. As they went through the design thinking process, the team spoke with members of the community, including Scott, about their experiences and ideas for the space.\n\n However, as the team began to develop prototypes and test solutions, they realized that they had made some assumptions about the needs of people with disabilities that did not align with the reality of Scott's experiences. For instance, they had assumed that adding ramps would be sufficient to make the mall more accessible, but Scott pointed out that these ramps were often too steep or poorly maintained, making them difficult or even dangerous to use.\n\n Through reflection, the team was able to recognize and address these misrepresentations in their design thinking process. They took the time to listen to Scott's perspective and to understand the specific challenges he faced when navigating the mall. This allowed them to go back to the drawing board and come up with solutions that better met his needs, such as installing handrails and non-slip surfaces on the ramps, and increasing the number of rest areas where Scott could take breaks while shopping.\n\n Another example of the value of reflection in social welfare and support services can be seen in the work of individuals like Scott - 115sandoval. Through careful self-reflection, Scott has been able to gain a deeper understanding of his own experiences and challenges, and to use this knowledge to advocate for himself and others in his community. By reflecting on his own needs and limitations, Scott has been able to identify gaps in the current system and to propose innovative solutions that can help to bridge those gaps.\n\n For example, Scott has spoken publicly about the challenges","[INST] Write a short essay on the following topic:\n In Social Welfare and Support Services, how does reflection contribute in tackling challenges and preventing misrepresentation in the design thinking process?\n Intersperse the following information in different parts of the essay:\n 1. Give an example that occurred near Bedok Reservoir View (Block 771), at a shopping mall frequented by Scott Sandoval\n 2. Give an example citing someone named scott--115sandoval\n 3. Give an example citing: smarturl.it/thescottsandovalxxx [/INST]\n",Bedok Reservoir View (Block 771),Scott Sandoval,scott--115sandoval,smarturl.it/thescottsandovalxxx,1
1222,"In the field of social justice and community welfare, there is a growing recognition of the importance of using visualization tools to foster innovation and collaboration in a design thinking approach. By creating visual representations of complex problems and potential solutions, stakeholders are better able to understand the issues at hand and work together to develop effective interventions.\n\n One example of this can be seen in the efforts of the local government and community members near Hougang Avenue 8 (Block 676) to improve safety at a subway station frequented by Piergiuseppe Martorell. Using data visualization tools, they were able to identify specific areas of the station where accidents and injuries were most likely to occur. By sharing these visualizations with Piergiuseppe and other key stakeholders, they were able to generate a collaborative discussion about possible solutions.\n\n Similarly, visualization tools have been used effectively by individuals like Piergiuseppemartorellofficial, who has used platforms such as discord.com/thepiergiuseppemartorell to share his ideas for improving access to education and resources for marginalized communities. Through the use of diagrams, charts, and other visual aids, he has been able to communicate complex concepts in a way that is accessible and engaging to a wide audience.\n\n Moreover, the use of visualization tools in a design thinking approach has been shown to foster creativity and innovation, as well as to promote cross-functional collaboration. By breaking down silos between different departments and disciplines, these tools enable teams to work together more effectively, leading to the development of new and innovative solutions to pressing social issues.\n\n In conclusion, visualization tools have emerged as an essential component of any successful design thinking approach to social justice and community welfare. By facilitating collaboration, promoting cross-functional teamwork, and enhancing communication, these tools can help stakeholders develop creative and effective solutions to some of the most pressing challenges facing our society today. Whether it's improving safety at a subway station or expanding access to education and resources for marginalized communities, visualization tools offer a powerful means of fostering innovation and driving positive change.","[INST] Write a short essay on the following topic:\n In Social Justice and Community Welfare, how can visualization tools foster innovation and collaboration in a design thinking approach?\n Intersperse the following information in different parts of the essay:\n 1. Give an example that occurred near Hougang Avenue 8 (Block 676), at a subway station frequented by Piergiuseppe Martorell\n 2. Give an example citing someone named piergiuseppemartorellofficial\n 3. Give an example citing: discord.com/thepiergiuseppemartorell [/INST]\n",Hougang Avenue 8 (Block 676),Piergiuseppe Martorell,piergiuseppemartorellofficial,discord.com/thepiergiuseppemartorell,1
1483,"In the realm of government policy and administration, visualization tools have emerged as powerful catalysts for fostering innovation and collaboration. These tools enable policymakers, administrators, and stakeholders to better understand complex issues and data sets, leading to more informed decision-making and enhanced public engagement. By incorporating visualization tools into a design thinking approach, we can effectively tackle real-world challenges while ensuring the process is both inclusive and efficient. This essay will illustrate these benefits through various examples, including a scenario near Clementi West Street 1, a user named ""priya.thepatel,"" and an online presence at facebook.com/priyaiampatelofficial.\n\n Visualization tools are particularly effective in highlighting local concerns and encouraging community participation, as demonstrated by an incident near Clementi West Street 1, Block 601. Priya Patel, a frequent visitor to this town, noticed recurring traffic congestion at specific hours. Local authorities could employ data visualization tools to analyze traffic patterns and communicate potential solutions to residents like Priya, enabling them to take part in the decision-making process and offer valuable feedback. Here, visualization tools facilitate the presentation of intricate data in an accessible format, fostering collaboration between stakeholders and fostering an environment ripe for innovation.\n\n The use of visualization tools transcends physical boundaries and extends to digital platforms, as exemplified by users such as ""priya.thepatel."" Imagine a scenario where Priya shares her experiences with smart city initiatives on her social media platform, facebook.com/priyaiampatelofficial. Data visualization could transform raw data from her posts into actionable insights regarding citizen satisfaction, preferences, and pain points, which would allow policymakers to make informed decisions regarding urban planning and resource allocation. Consequently, visualization tools create opportunities for meaningful dialogue between citizens and governments, thereby enhancing transparency, trust, and accountability.\n\n Furthermore, design thinking approaches emphasize empathy, ideation, and experimentation in order to create human-centric solutions to complex problems (Brown & Katz, 2011). Visualization tools aid each stage of this process by providing tangible representations of qualitative and quantitative information, making it easier for stakeholders to comprehend, discuss, and refine concepts. As ideas progress from abstract notions to concrete proposals, these tools promote inter","[INST] Write a short essay on the following topic:\n In Government Policy and Administration, how can visualization tools foster innovation and collaboration in a design thinking approach?\n Intersperse the following information in different parts of the essay:\n 1. Give an example that occurred near Clementi West Street 1 (Block 601), at a town frequented by Priya Patel\n 2. Give an example citing someone named priya.thepatel\n 3. Give an example citing: facebook.com/priyaiampatelofficial [/INST]\n",Clementi West Street 1 (Block 601),Priya Patel,priya.thepatel,facebook.com/priyaiampatelofficial,1


In [6]:
pre = BasicPreprocessor()
rows = []
for t in tqdm(df.itertuples()):
    street_address = str(getattr(t, "street_address"))
    student_name = str(getattr(t, "student_name"))
    username = str(getattr(t, "username"))
    personal_url = str(getattr(t, "personal_url"))
    text = pre(getattr(t, "text"))
    sp_tokens = tokenizer(text)
    tokens = [token.text for token in sp_tokens]
    trailing_whitespace = [bool(token.whitespace_) for token in sp_tokens]
    res = dg.labels(
        text=text, 
        tokens=sp_tokens, 
        street_address=street_address,
        student_name=student_name,
        username=username,
        personal_url=personal_url,
    )
    total = res.n_street_address + res.n_student_name + res.n_username + res.n_personal_url
    if total < min_entities:
        continue
    rows.append({
        "document": f"rh_{uuid.uuid4().hex}",
        "full_text": text,
        "tokens": tokens,
        "trailing_whitespace": trailing_whitespace,
        "labels": res.labels, 
        "n_street_address": res.n_street_address,
        "n_student_name": res.n_student_name,
        "n_username": res.n_username,
        "n_personal_url": res.n_personal_url,
        "street_address": street_address,
        "student_name": student_name,
        "username": username,
        "personal_url": personal_url,
    })


1870it [00:01, 1078.26it/s]


In [7]:
df = pd.DataFrame.from_records(rows)
cols = ["n_street_address", "n_student_name", "n_username", "n_personal_url"]
df[cols] = df[cols].astype(np.int8)
df.info()

<class 'pandas.core.frame.DataFrame'>
RangeIndex: 1654 entries, 0 to 1653
Data columns (total 13 columns):
 #   Column               Non-Null Count  Dtype 
---  ------               --------------  ----- 
 0   document             1654 non-null   object
 1   full_text            1654 non-null   object
 2   tokens               1654 non-null   object
 3   trailing_whitespace  1654 non-null   object
 4   labels               1654 non-null   object
 5   n_street_address     1654 non-null   int8  
 6   n_student_name       1654 non-null   int8  
 7   n_username           1654 non-null   int8  
 8   n_personal_url       1654 non-null   int8  
 9   street_address       1654 non-null   object
 10  student_name         1654 non-null   object
 11  username             1654 non-null   object
 12  personal_url         1654 non-null   object
dtypes: int8(4), object(9)
memory usage: 122.9+ KB


In [8]:
df.describe(percentiles=percentiles)

Unnamed: 0,n_street_address,n_student_name,n_username,n_personal_url
count,1654.0,1654.0,1654.0,1654.0
mean,1.641475,1.399033,0.662031,0.766022
std,0.707062,0.883621,0.906412,0.584355
min,0.0,0.0,0.0,0.0
1%,1.0,0.0,0.0,0.0
5%,1.0,0.0,0.0,0.0
10%,1.0,1.0,0.0,0.0
20%,1.0,1.0,0.0,0.0
30%,1.0,1.0,0.0,0.0
40%,1.0,1.0,0.0,1.0


In [9]:
for row in rows[:3]:
    print("===========")
    print(f"street_address={row['street_address']}")
    print(f"student_name={row['student_name']}")
    print(f"username={row['username']}")
    print(f"personal_url={row['personal_url']}")
    lbs = row["labels"]
    tokens = row["tokens"]
    ts = []
    for i in range(len(lbs)):
        if lbs[i]!="O":
            ts.append(tokens[i])
    print(ts)

street_address=Choa Chu Kang Avenue 4 (Block 444)
student_name=Rahman bin Hassan
username=_4807.rahman-iam--hassan
personal_url=periscope.tv/rahmanhassanxxx
['Choa', 'Chu', 'Kang', 'Avenue', '4', 'Rahman', 'bin', 'Hassan', '_', '4807.rahman', '-', 'iam', '--', 'hassan', '_', '4807.rahman', '-', 'iam', '--', 'hassan', 'Periscope.tv/rahmanhassanxxx', 'Rahman', 'bin', 'Hassan', '_', '4807.rahman', '-', 'iam', '--', 'hassan']
street_address=Jurong West Street 92 (Block 922)
student_name=Nicholas Sami
username=__iamNicholas_xxx.Sami
personal_url=blogspot.me/@iamnicholassami
['Jurong', 'West', 'Street', '92', '(', 'Block', '922', ')', '_', '_', 'iamNicholas_xxx', '.', 'Sami', 'Blogspot.me/@iamnicholassami']
street_address=Jalan Eunos (Eunos Mansion)
student_name=Nina Setiawan
username=the-nina3719setiawan
personal_url=sinaweibo.com/ninaiamsetiawanthe
['Jalan', 'Eunos', 'Nina', 'Setiawan', 'the', '-', 'nina3719setiawan', 'The', '-', 'nina3719setiawan', 'the', '-', 'nina3719setiawan', 'sinawei

In [10]:
%%time
with open("output/rh01.json", "w") as f:
    json.dump(rows, f)
df.to_parquet(f"output/rh01.parquet", index=False)
assert df.notna().all(axis=None)

CPU times: user 527 ms, sys: 0 ns, total: 527 ms
Wall time: 828 ms


In [11]:
tim.stop()
print(f"Total time taken {str(tim.elapsed)}")

Total time taken 0:00:02.940048
