# Fix pathing

In [1]:
import sys


sys.path.append("../..")


In [2]:
import constants

import os


constants.PROJECT_DIRECTORY_PATH = os.path.dirname(os.path.dirname(os.path.dirname(os.path.dirname(constants.PROJECT_DIRECTORY_PATH))))


# Imports

In [3]:
import plotter
import datahandler

import matplotlib.pyplot
import numpy as np
import pandas as pd
import seaborn as sns
import IPython.display


# Constants

In [4]:
FOLDER_NAME = "exp_1_avg_active"

FOLDER_PATH = os.path.join(os.path.dirname(constants.PROJECT_DIRECTORY_PATH), "Simulator", "data", FOLDER_NAME)


# Methods

In [5]:
def load_csv(fileName):
    df = pd.read_csv(os.path.join(FOLDER_PATH, fileName + ".csv"))

    response_time_cols = [
        'duration_incident_creation',
        'duration_resource_appointment',
        'duration_resource_preparing_departure',
        'duration_dispatching_to_scene'
    ]
    df['total_response_time'] = df[response_time_cols].sum(axis=1)

    return df


In [6]:
def print_results(df: pd.DataFrame):
    # Define the criteria for response times
    criteria = {
        ('A', True): 12 * 60,
        ('A', False): 25 * 60,
        ('H', True): 30 * 60,
        ('H', False): 40 * 60
    }

    # Function to calculate compliance for each group
    def calculate_compliance(group, triage, urban):
        limit = criteria.get((triage, urban))
        if limit is not None:
            return (group['total_response_time'] < limit).mean()
        return None

    # Calculate statistics and compliance for each group
    results = []
    for (triage, urban), group in df.groupby(['triage_impression_during_call', 'urban']):
        mean = group['total_response_time'].mean()
        median = group['total_response_time'].median()
        compliance = calculate_compliance(group, triage, urban)
        results.append({
            'Triage': triage,
            'Urban': urban,
            'Mean (sec)': mean,
            'Median (sec)': median,
            'Compliance': compliance
        })

    stats = pd.DataFrame(results)

    # Convert mean and median to minutes
    stats['Mean (min)'] = (stats['Mean (sec)'] / 60)
    stats['Median (min)'] = (stats['Median (sec)'] / 60)
    stats.drop(columns=['Mean (sec)', 'Median (sec)'], inplace=True)

    # Map urban values to Yes/No
    stats['Urban'] = stats['Urban'].map({True: 'Yes', False: 'No'})
    
    # Sort values
    stats.sort_values(by=["Urban", "Triage"], ascending=[False, True], inplace=True)
    
    # Display the DataFrame
    formatted_stats = stats.style.format({
        'Mean (min)': "{:.2f}",
        'Median (min)': "{:.2f}",
        'Compliance': "{:.2%}"
    }).hide(axis='index')
    IPython.display.display(formatted_stats)


In [7]:
def boxplot_time_at_steps_modified(
    dataframe: pd.DataFrame,
    triage_impression: str = None
):
    title = "Time Taken At Each Step of the Incident"

    if triage_impression is not None:
        # Filter the dataframe without overwriting the original one
        temp_df = dataframe[dataframe["triage_impression_during_call"] == triage_impression].copy()
        title += f" ({triage_impression})"
    else:
        # Use the original dataframe if no triage_impression filter is applied
        temp_df = dataframe.copy()

    steps = {
        "Creating Incident": "duration_incident_creation",
        "Appointing Resource": "duration_resource_appointment",
        "Resource to Start Task": "duration_resource_preparing_departure",
        "Dispatching to Scene": "duration_dispatching_to_scene",
        "At Scene": "duration_at_scene",
        "Dispatching to Hospital": "duration_dispatching_to_hospital",
        "At Hospital": "duration_at_hospital"
    }

    # Calculating durations for each step
    plot_data = [(temp_df[duration_column][temp_df[duration_column] > 0] / 60) for step, duration_column in steps.items()]

    # Plotting
    matplotlib.pyplot.figure(figsize=(8, 4))
    matplotlib.pyplot.boxplot(plot_data[::-1], labels=list(steps.keys())[::-1], vert=False, patch_artist=True, showfliers=False)
    matplotlib.pyplot.title(title)
    matplotlib.pyplot.xlabel("Time in Minutes")
    matplotlib.pyplot.xticks()
    matplotlib.pyplot.show()


In [8]:
def boxplot_time_at_steps(
    historic_dataframe: pd.DataFrame,
    simulated_dataframe: pd.DataFrame,
    bounds: tuple[str, str] = None,
    triage_impressions: list[str] = ["A", "H", "V1"],
):
    historic_df = historic_dataframe.copy()
    simulated_df = simulated_dataframe.copy()
    
    # filter historic data by time frame
    if bounds is not None:
        start_bound, end_bound = pd.to_datetime(bounds[0]), pd.to_datetime(bounds[1])
        historic_df = historic_df[(historic_df['time_call_received'] >= start_bound) & (historic_df['time_call_received'] <= end_bound)]

    # calculate duration at each stage in minutes (simulated dataframe already has this calculated)
    historic_steps = {
        "duration_incident_creation": ("time_call_received", "time_incident_created"),
        "duration_resource_appointment": ("time_incident_created", "time_resource_appointed"),
        "duration_resource_preparing_departure": ("time_resource_appointed", "time_ambulance_dispatch_to_scene"),
        "duration_dispatching_to_scene": ("time_ambulance_dispatch_to_scene", "time_ambulance_arrived_at_scene"),
        "duration_at_scene": ("time_ambulance_arrived_at_scene", "time_ambulance_dispatch_to_hospital", "time_ambulance_available"),
        "duration_dispatching_to_hospital": ("time_ambulance_dispatch_to_hospital", "time_ambulance_arrived_at_hospital"),
        "duration_at_hospital": ("time_ambulance_arrived_at_hospital", "time_ambulance_available")
    }

    for step, times in historic_steps.items():
        if len(times) == 3:
            historic_df.loc[historic_df[times[1]].isna(), step] = (historic_df[times[2]] - historic_df[times[0]]).dt.total_seconds() / 60
            historic_df.loc[~historic_df[times[1]].isna(), step] = (historic_df[times[1]] - historic_df[times[0]]).dt.total_seconds() / 60
        else:
            historic_df[step] = (historic_df[times[1]] - historic_df[times[0]]).dt.total_seconds() / 60
    
    # convert seconds to minutes and replace zeros in simulated data to nan
    for step in historic_steps.keys():
        simulated_df[step] /= 60

    simulated_df[["duration_dispatching_to_hospital", "duration_at_hospital"]] = simulated_df[["duration_dispatching_to_hospital", "duration_at_hospital"]].replace(0, np.nan)

    # plot data
    matplotlib.pyplot.figure(figsize=(8, 8))

    data = []
    positions = []
    labels = []
    colors = []
    stripes = []

    i = 0
    for step in historic_steps.keys():
        labels.append(f"{step}")
        for triage in triage_impressions:
            for df, stripe, alpha in zip([historic_df, simulated_df], [False, True], [1.00, 0.65]):
                data.append(df[df['triage_impression_during_call'] == triage][step].dropna())

                labels.append("")

                if (triage == "A"):
                    colors.append([1.00, 0.37, 0.28, alpha])
                elif (triage == "H"):
                    colors.append([0.12, 0.56, 1.00, alpha])
                else:
                    colors.append([0.20, 0.80, 0.20, alpha])

                stripes.append(stripe)

                positions.append(i)
                i += 0.75
        labels.pop()
        i += 1

    bplot = matplotlib.pyplot.boxplot(
        data[::-1],
        labels=labels[::-1],
        positions=positions,
        vert=False,
        patch_artist=True,
        showfliers=True
    )

    for patch, color, stripe in zip(bplot["boxes"], colors[::-1], stripes[::-1]):
        patch.set_facecolor(color)

    matplotlib.pyplot.title("Time Taken At Each Step of the Incident")
    matplotlib.pyplot.xlabel("Time in Minutes")
    matplotlib.pyplot.xticks()
    matplotlib.pyplot.show()


In [9]:
def table_results(filter_urban: bool = None, strategies = ["closest", "random"]):
    def calculate_compliance(df):
        if df is None:
            return None

        # Criteria for both triage types in both urban and rural settings
        criteria = {
            ('A', True): 12 * 60,  # 12 minutes for triage 'A' in urban areas
            ('A', False): 25 * 60, # 25 minutes for triage 'A' in rural areas
            ('H', True): 30 * 60,  # 30 minutes for triage 'H' in urban areas
            ('H', False): 40 * 60  # 40 minutes for triage 'H' in rural areas
        }

        # Initialize counters
        total_compliant_cases = 0
        total_cases = 0
        
        # Check necessary columns exist
        for urban in [False, True]:
            if filter_urban is not None and urban != filter_urban:
                continue
            for triage in ['A', 'H']:
                filtered_df = df[(df['triage_impression_during_call'] == triage) & (df['urban'] == urban)]
                limit = criteria.get((triage, urban))
                if not filtered_df.empty:
                    # Count compliant cases for this triage type
                    compliant_cases = filtered_df['total_response_time'] < limit
                    total_compliant_cases += compliant_cases.sum()
                    total_cases += len(filtered_df)
        if total_cases > 0:
            # Calculate overall compliance rate across both triage types
            overall_compliance = total_compliant_cases / total_cases
            return overall_compliance
        else:
            return None  # Return None if there are no cases to evaluate

    # Settings combinations
    prioritize_triages = [False, True]
    response_restricteds = [False, True]
    schedule_breaks = [False, True]

    results = []

    for strategy in strategies:
        for prioritize_triage in prioritize_triages:
            for response_restricted in response_restricteds:
                for schedule_break in schedule_breaks:
                    filename = f"events_strategy={strategy}_prioritizeTriage={'true' if prioritize_triage else 'false'}_responseRestricted={'true' if response_restricted else 'false'}_scheduleBreaks={'true' if schedule_break else 'false'}"
                    df = load_csv(filename)
                    compliance = calculate_compliance(df)
                    if compliance is not None:
                        results.append({
                            "Strategy": strategy,
                            "Prioritize Triage": prioritize_triage,
                            "Response Restricted": response_restricted,
                            "Schedule Breaks": schedule_break,
                            "Compliance": compliance
                        })

    # Create DataFrame
    results_df = pd.DataFrame(results)
    # Pivot Table for visualization
    pivot_table = results_df.pivot_table(index=["Schedule Breaks", "Prioritize Triage", "Response Restricted"],
                                        columns=["Strategy"], 
                                        values="Compliance")
    # Color coding from green to red
    cm = sns.light_palette("green", as_cmap=True, n_colors=8)
    styled_pivot = pivot_table.style.background_gradient(cmap=cm).format("{:.2%}")

    return styled_pivot


# Main

In [10]:
table_results()


Unnamed: 0_level_0,Unnamed: 1_level_0,Strategy,closest,random
Schedule Breaks,Prioritize Triage,Response Restricted,Unnamed: 3_level_1,Unnamed: 4_level_1
False,False,False,74.80%,21.95%
False,False,True,75.61%,24.39%
False,True,False,75.61%,22.76%
False,True,True,73.98%,16.26%
True,False,False,73.17%,20.33%
True,False,True,70.73%,26.02%
True,True,False,74.80%,18.70%
True,True,True,73.17%,18.70%


In [11]:
table_results(filter_urban=True)


Unnamed: 0_level_0,Unnamed: 1_level_0,Strategy,closest,random
Schedule Breaks,Prioritize Triage,Response Restricted,Unnamed: 3_level_1,Unnamed: 4_level_1
False,False,False,73.45%,21.24%
False,False,True,74.34%,23.89%
False,True,False,75.22%,22.12%
False,True,True,73.45%,15.93%
True,False,False,72.57%,19.47%
True,False,True,69.91%,26.55%
True,True,False,74.34%,17.70%
True,True,True,71.68%,18.58%


In [12]:
table_results(filter_urban=False)


Unnamed: 0_level_0,Unnamed: 1_level_0,Strategy,closest,random
Schedule Breaks,Prioritize Triage,Response Restricted,Unnamed: 3_level_1,Unnamed: 4_level_1
False,False,False,90.00%,30.00%
False,False,True,90.00%,30.00%
False,True,False,80.00%,30.00%
False,True,True,80.00%,20.00%
True,False,False,80.00%,30.00%
True,False,True,80.00%,20.00%
True,True,False,80.00%,30.00%
True,True,True,90.00%,20.00%


In [13]:
print_results(load_csv("events_strategy=closest_prioritizeTriage=false_responseRestricted=false_scheduleBreaks=false"))


Triage,Urban,Compliance,Mean (min),Median (min)
A,Yes,63.46%,11.35,10.47
H,Yes,81.97%,22.2,20.75
V1,Yes,nan%,83.28,66.2
A,No,100.00%,14.37,15.55
H,No,83.33%,26.49,26.3
V1,No,nan%,66.96,66.0


In [14]:
print_results(load_csv("events_strategy=closest_prioritizeTriage=false_responseRestricted=false_scheduleBreaks=true"))


Triage,Urban,Compliance,Mean (min),Median (min)
A,Yes,65.38%,11.56,10.29
H,Yes,78.69%,22.31,20.58
V1,Yes,nan%,84.17,66.46
A,No,100.00%,14.12,15.11
H,No,66.67%,30.41,28.84
V1,No,nan%,70.51,70.47


In [15]:
print_results(load_csv("events_strategy=closest_prioritizeTriage=false_responseRestricted=true_scheduleBreaks=false"))


Triage,Urban,Compliance,Mean (min),Median (min)
A,Yes,63.46%,12.12,10.32
H,Yes,83.61%,21.72,20.13
V1,Yes,nan%,83.09,66.7
A,No,100.00%,14.49,15.43
H,No,83.33%,26.17,25.95
V1,No,nan%,69.13,69.07


In [16]:
print_results(load_csv("events_strategy=closest_prioritizeTriage=false_responseRestricted=true_scheduleBreaks=true"))


Triage,Urban,Compliance,Mean (min),Median (min)
A,Yes,59.62%,13.03,10.59
H,Yes,78.69%,22.44,20.33
V1,Yes,nan%,84.52,66.5
A,No,100.00%,14.66,15.72
H,No,66.67%,28.19,26.9
V1,No,nan%,69.49,69.08


In [17]:
print_results(load_csv("events_strategy=closest_prioritizeTriage=true_responseRestricted=false_scheduleBreaks=false"))


Triage,Urban,Compliance,Mean (min),Median (min)
A,Yes,67.31%,11.37,10.28
H,Yes,81.97%,22.02,20.75
V1,Yes,nan%,84.6,67.22
A,No,100.00%,14.35,15.37
H,No,66.67%,29.56,26.68
V1,No,nan%,66.92,66.13


In [18]:
print_results(load_csv("events_strategy=closest_prioritizeTriage=true_responseRestricted=false_scheduleBreaks=true"))


Triage,Urban,Compliance,Mean (min),Median (min)
A,Yes,67.31%,11.4,10.42
H,Yes,80.33%,22.49,20.67
V1,Yes,nan%,85.97,67.71
A,No,100.00%,14.46,15.85
H,No,66.67%,29.43,25.38
V1,No,nan%,70.05,69.08


In [19]:
print_results(load_csv("events_strategy=closest_prioritizeTriage=true_responseRestricted=true_scheduleBreaks=false"))


Triage,Urban,Compliance,Mean (min),Median (min)
A,Yes,63.46%,12.1,10.42
H,Yes,81.97%,21.49,20.7
V1,Yes,nan%,84.73,67.45
A,No,100.00%,14.61,15.67
H,No,66.67%,28.72,25.52
V1,No,nan%,66.97,66.06


In [20]:
print_results(load_csv("events_strategy=closest_prioritizeTriage=true_responseRestricted=true_scheduleBreaks=true"))


Triage,Urban,Compliance,Mean (min),Median (min)
A,Yes,59.62%,12.61,10.65
H,Yes,81.97%,22.02,20.57
V1,Yes,nan%,86.87,67.18
A,No,100.00%,14.7,15.5
H,No,83.33%,26.84,26.76
V1,No,nan%,69.4,69.12


In [21]:
print_results(load_csv("events_strategy=random_prioritizeTriage=false_responseRestricted=false_scheduleBreaks=false"))


Triage,Urban,Compliance,Mean (min),Median (min)
A,Yes,7.69%,26.9,27.19
H,Yes,32.79%,40.67,38.77
V1,Yes,nan%,105.07,89.02
A,No,0.00%,34.86,33.17
H,No,50.00%,43.21,44.43
V1,No,nan%,104.96,105.02


In [22]:
print_results(load_csv("events_strategy=random_prioritizeTriage=false_responseRestricted=false_scheduleBreaks=true"))


Triage,Urban,Compliance,Mean (min),Median (min)
A,Yes,3.85%,28.6,26.16
H,Yes,32.79%,41.34,36.23
V1,Yes,nan%,102.73,88.73
A,No,0.00%,43.14,45.3
H,No,50.00%,47.48,44.92
V1,No,nan%,107.91,102.71


In [23]:
print_results(load_csv("events_strategy=random_prioritizeTriage=false_responseRestricted=true_scheduleBreaks=false"))


Triage,Urban,Compliance,Mean (min),Median (min)
A,Yes,7.69%,26.76,22.14
H,Yes,37.70%,40.5,37.43
V1,Yes,nan%,103.36,89.62
A,No,25.00%,38.47,43.35
H,No,33.33%,51.87,53.78
V1,No,nan%,95.44,87.26


In [24]:
print_results(load_csv("events_strategy=random_prioritizeTriage=false_responseRestricted=true_scheduleBreaks=true"))


Triage,Urban,Compliance,Mean (min),Median (min)
A,Yes,15.38%,25.7,22.87
H,Yes,36.07%,39.91,37.37
V1,Yes,nan%,107.06,85.55
A,No,0.00%,48.45,49.92
H,No,33.33%,50.67,48.26
V1,No,nan%,101.53,88.16


In [25]:
print_results(load_csv("events_strategy=random_prioritizeTriage=true_responseRestricted=false_scheduleBreaks=false"))


Triage,Urban,Compliance,Mean (min),Median (min)
A,Yes,9.62%,27.34,25.25
H,Yes,32.79%,42.11,38.82
V1,Yes,nan%,106.23,85.09
A,No,50.00%,27.75,24.95
H,No,16.67%,54.11,55.23
V1,No,nan%,109.26,108.9


In [26]:
print_results(load_csv("events_strategy=random_prioritizeTriage=true_responseRestricted=false_scheduleBreaks=true"))


Triage,Urban,Compliance,Mean (min),Median (min)
A,Yes,9.62%,27.4,24.33
H,Yes,24.59%,40.95,35.87
V1,Yes,nan%,113.55,94.39
A,No,25.00%,37.83,38.83
H,No,33.33%,44.53,45.36
V1,No,nan%,101.34,98.41


In [27]:
print_results(load_csv("events_strategy=random_prioritizeTriage=true_responseRestricted=true_scheduleBreaks=false"))


Triage,Urban,Compliance,Mean (min),Median (min)
A,Yes,7.69%,23.72,21.17
H,Yes,22.95%,45.31,41.72
V1,Yes,nan%,111.32,91.05
A,No,0.00%,40.48,42.03
H,No,33.33%,44.11,48.08
V1,No,nan%,103.76,96.55


In [28]:
print_results(load_csv("events_strategy=random_prioritizeTriage=true_responseRestricted=true_scheduleBreaks=true"))


Triage,Urban,Compliance,Mean (min),Median (min)
A,Yes,13.46%,25.46,23.02
H,Yes,22.95%,44.27,42.53
V1,Yes,nan%,113.2,100.17
A,No,25.00%,35.71,38.02
H,No,16.67%,45.99,46.39
V1,No,nan%,97.5,89.12
