# Fix pathing

In [1]:
import sys


sys.path.append("../..")


In [2]:
import constants

import os


constants.PROJECT_DIRECTORY_PATH = os.path.dirname(os.path.dirname(os.path.dirname(os.path.dirname(constants.PROJECT_DIRECTORY_PATH))))


# Imports

In [3]:
import plotter
import datahandler

import matplotlib.pyplot
import numpy as np
import pandas as pd
import seaborn as sns
import IPython.display


# Constants

In [104]:
FOLDER_NAME = "ex_1_most_NONE"

FOLDER_PATH = os.path.join(os.path.dirname(constants.PROJECT_DIRECTORY_PATH), "Simulator", "data", FOLDER_NAME)


# Methods

In [105]:
def load_csv(fileName):
    df = pd.read_csv(os.path.join(FOLDER_PATH, fileName + ".csv"))

    response_time_cols = [
        'duration_incident_creation',
        'duration_resource_appointment',
        'duration_resource_preparing_departure',
        'duration_dispatching_to_scene'
    ]
    df['total_response_time'] = df[response_time_cols].sum(axis=1)

    return df


In [106]:
def print_results(df: pd.DataFrame):
    # Define the criteria for response times
    criteria = {
        ('A', True): 12 * 60,
        ('A', False): 25 * 60,
        ('H', True): 30 * 60,
        ('H', False): 40 * 60
    }

    # Function to calculate compliance for each group
    def calculate_compliance(group, triage, urban):
        limit = criteria.get((triage, urban))
        if limit is not None:
            return (group['total_response_time'] < limit).mean()
        return None

    # Calculate statistics and compliance for each group
    results = []
    for (triage, urban), group in df.groupby(['triage_impression_during_call', 'urban']):
        mean = group['total_response_time'].mean()
        median = group['total_response_time'].median()
        compliance = calculate_compliance(group, triage, urban)
        results.append({
            'Triage': triage,
            'Urban': urban,
            'Mean (sec)': mean,
            'Median (sec)': median,
            'Compliance': compliance
        })

    stats = pd.DataFrame(results)

    # Convert mean and median to minutes
    stats['Mean (min)'] = (stats['Mean (sec)'] / 60)
    stats['Median (min)'] = (stats['Median (sec)'] / 60)
    stats.drop(columns=['Mean (sec)', 'Median (sec)'], inplace=True)

    # Map urban values to Yes/No
    stats['Urban'] = stats['Urban'].map({True: 'Yes', False: 'No'})
    
    # Sort values
    stats.sort_values(by=["Urban", "Triage"], ascending=[False, True], inplace=True)
    
    # Display the DataFrame
    formatted_stats = stats.style.format({
        'Mean (min)': "{:.2f}",
        'Median (min)': "{:.2f}",
        'Compliance': "{:.2%}"
    }).hide(axis='index')
    IPython.display.display(formatted_stats)


In [107]:
def boxplot_time_at_steps_modified(
    dataframe: pd.DataFrame,
    triage_impression: str = None
):
    title = "Time Taken At Each Step of the Incident"

    if triage_impression is not None:
        # Filter the dataframe without overwriting the original one
        temp_df = dataframe[dataframe["triage_impression_during_call"] == triage_impression].copy()
        title += f" ({triage_impression})"
    else:
        # Use the original dataframe if no triage_impression filter is applied
        temp_df = dataframe.copy()

    steps = {
        "Creating Incident": "duration_incident_creation",
        "Appointing Resource": "duration_resource_appointment",
        "Resource to Start Task": "duration_resource_preparing_departure",
        "Dispatching to Scene": "duration_dispatching_to_scene",
        "At Scene": "duration_at_scene",
        "Dispatching to Hospital": "duration_dispatching_to_hospital",
        "At Hospital": "duration_at_hospital"
    }

    # Calculating durations for each step
    plot_data = [(temp_df[duration_column][temp_df[duration_column] > 0] / 60) for step, duration_column in steps.items()]

    # Plotting
    matplotlib.pyplot.figure(figsize=(8, 4))
    matplotlib.pyplot.boxplot(plot_data[::-1], labels=list(steps.keys())[::-1], vert=False, patch_artist=True, showfliers=False)
    matplotlib.pyplot.title(title)
    matplotlib.pyplot.xlabel("Time in Minutes")
    matplotlib.pyplot.xticks()
    matplotlib.pyplot.show()


In [108]:
def boxplot_time_at_steps(
    historic_dataframe: pd.DataFrame,
    simulated_dataframe: pd.DataFrame,
    bounds: tuple[str, str] = None,
    triage_impressions: list[str] = ["A", "H", "V1"],
):
    historic_df = historic_dataframe.copy()
    simulated_df = simulated_dataframe.copy()
    
    # filter historic data by time frame
    if bounds is not None:
        start_bound, end_bound = pd.to_datetime(bounds[0]), pd.to_datetime(bounds[1])
        historic_df = historic_df[(historic_df['time_call_received'] >= start_bound) & (historic_df['time_call_received'] <= end_bound)]

    # calculate duration at each stage in minutes (simulated dataframe already has this calculated)
    historic_steps = {
        "duration_incident_creation": ("time_call_received", "time_incident_created"),
        "duration_resource_appointment": ("time_incident_created", "time_resource_appointed"),
        "duration_resource_preparing_departure": ("time_resource_appointed", "time_ambulance_dispatch_to_scene"),
        "duration_dispatching_to_scene": ("time_ambulance_dispatch_to_scene", "time_ambulance_arrived_at_scene"),
        "duration_at_scene": ("time_ambulance_arrived_at_scene", "time_ambulance_dispatch_to_hospital", "time_ambulance_available"),
        "duration_dispatching_to_hospital": ("time_ambulance_dispatch_to_hospital", "time_ambulance_arrived_at_hospital"),
        "duration_at_hospital": ("time_ambulance_arrived_at_hospital", "time_ambulance_available")
    }

    for step, times in historic_steps.items():
        if len(times) == 3:
            historic_df.loc[historic_df[times[1]].isna(), step] = (historic_df[times[2]] - historic_df[times[0]]).dt.total_seconds() / 60
            historic_df.loc[~historic_df[times[1]].isna(), step] = (historic_df[times[1]] - historic_df[times[0]]).dt.total_seconds() / 60
        else:
            historic_df[step] = (historic_df[times[1]] - historic_df[times[0]]).dt.total_seconds() / 60
    
    # convert seconds to minutes and replace zeros in simulated data to nan
    for step in historic_steps.keys():
        simulated_df[step] /= 60

    simulated_df[["duration_dispatching_to_hospital", "duration_at_hospital"]] = simulated_df[["duration_dispatching_to_hospital", "duration_at_hospital"]].replace(0, np.nan)

    # plot data
    matplotlib.pyplot.figure(figsize=(8, 8))

    data = []
    positions = []
    labels = []
    colors = []
    stripes = []

    i = 0
    for step in historic_steps.keys():
        labels.append(f"{step}")
        for triage in triage_impressions:
            for df, stripe, alpha in zip([historic_df, simulated_df], [False, True], [1.00, 0.65]):
                data.append(df[df['triage_impression_during_call'] == triage][step].dropna())

                labels.append("")

                if (triage == "A"):
                    colors.append([1.00, 0.37, 0.28, alpha])
                elif (triage == "H"):
                    colors.append([0.12, 0.56, 1.00, alpha])
                else:
                    colors.append([0.20, 0.80, 0.20, alpha])

                stripes.append(stripe)

                positions.append(i)
                i += 0.75
        labels.pop()
        i += 1

    bplot = matplotlib.pyplot.boxplot(
        data[::-1],
        labels=labels[::-1],
        positions=positions,
        vert=False,
        patch_artist=True,
        showfliers=True
    )

    for patch, color, stripe in zip(bplot["boxes"], colors[::-1], stripes[::-1]):
        patch.set_facecolor(color)

    matplotlib.pyplot.title("Time Taken At Each Step of the Incident")
    matplotlib.pyplot.xlabel("Time in Minutes")
    matplotlib.pyplot.xticks()
    matplotlib.pyplot.show()


In [109]:
def table_results(filter_urban: bool = None, strategies = ["closest", "random"]):
    def calculate_compliance(df):
        if df is None:
            return None

        # Criteria for both triage types in both urban and rural settings
        criteria = {
            ('A', True): 12 * 60,  # 12 minutes for triage 'A' in urban areas
            ('A', False): 25 * 60, # 25 minutes for triage 'A' in rural areas
            ('H', True): 30 * 60,  # 30 minutes for triage 'H' in urban areas
            ('H', False): 40 * 60  # 40 minutes for triage 'H' in rural areas
        }

        # Initialize counters
        total_compliant_cases = 0
        total_cases = 0
        
        # Check necessary columns exist
        for urban in [False, True]:
            if filter_urban is not None and urban != filter_urban:
                continue
            for triage in ['A', 'H']:
                filtered_df = df[(df['triage_impression_during_call'] == triage) & (df['urban'] == urban)]
                limit = criteria.get((triage, urban))
                if not filtered_df.empty:
                    # Count compliant cases for this triage type
                    compliant_cases = filtered_df['total_response_time'] < limit
                    total_compliant_cases += compliant_cases.sum()
                    total_cases += len(filtered_df)
        if total_cases > 0:
            # Calculate overall compliance rate across both triage types
            overall_compliance = total_compliant_cases / total_cases
            return overall_compliance
        else:
            return None  # Return None if there are no cases to evaluate

    # Settings combinations
    prioritize_triages = [False, True]
    response_restricteds = [False, True]
    schedule_breaks = [False, True]

    results = []

    for strategy in strategies:
        for prioritize_triage in prioritize_triages:
            for response_restricted in response_restricteds:
                for schedule_break in schedule_breaks:
                    filename = f"events_strategy={strategy}_prioritizeTriage={'true' if prioritize_triage else 'false'}_responseRestricted={'true' if response_restricted else 'false'}_scheduleBreaks={'true' if schedule_break else 'false'}"
                    df = load_csv(filename)
                    compliance = calculate_compliance(df)
                    if compliance is not None:
                        results.append({
                            "Strategy": strategy,
                            "Prioritize Triage": prioritize_triage,
                            "Response Restricted": response_restricted,
                            "Schedule Breaks": schedule_break,
                            "Compliance": compliance
                        })

    # Create DataFrame
    results_df = pd.DataFrame(results)
    # Pivot Table for visualization
    pivot_table = results_df.pivot_table(index=["Schedule Breaks", "Prioritize Triage", "Response Restricted"],
                                        columns=["Strategy"], 
                                        values="Compliance")
    # Color coding from green to red
    cm = sns.light_palette("green", as_cmap=True, n_colors=8)
    styled_pivot = pivot_table.style.background_gradient(cmap=cm).format("{:.2%}")

    return styled_pivot


# Main

In [110]:
table_results()


Unnamed: 0_level_0,Unnamed: 1_level_0,Strategy,closest,random
Schedule Breaks,Prioritize Triage,Response Restricted,Unnamed: 3_level_1,Unnamed: 4_level_1
False,False,False,79.23%,31.69%
False,False,True,73.77%,24.04%
False,True,False,82.51%,19.67%
False,True,True,74.86%,26.23%
True,False,False,67.21%,16.94%
True,False,True,66.67%,15.30%
True,True,False,73.77%,21.86%
True,True,True,69.95%,19.67%


In [111]:
table_results(filter_urban=True)


Unnamed: 0_level_0,Unnamed: 1_level_0,Strategy,closest,random
Schedule Breaks,Prioritize Triage,Response Restricted,Unnamed: 3_level_1,Unnamed: 4_level_1
False,False,False,77.50%,33.75%
False,False,True,71.88%,25.62%
False,True,False,80.00%,22.50%
False,True,True,71.88%,26.25%
True,False,False,66.25%,15.62%
True,False,True,65.00%,15.00%
True,True,False,73.75%,21.88%
True,True,True,67.50%,20.62%


In [112]:
table_results(filter_urban=False)


Unnamed: 0_level_0,Unnamed: 1_level_0,Strategy,closest,random
Schedule Breaks,Prioritize Triage,Response Restricted,Unnamed: 3_level_1,Unnamed: 4_level_1
False,False,False,91.30%,17.39%
False,False,True,86.96%,13.04%
False,True,False,100.00%,0.00%
False,True,True,95.65%,26.09%
True,False,False,73.91%,26.09%
True,False,True,78.26%,17.39%
True,True,False,73.91%,21.74%
True,True,True,86.96%,13.04%


In [113]:
print_results(load_csv("events_strategy=closest_prioritizeTriage=false_responseRestricted=false_scheduleBreaks=false"))


Triage,Urban,Compliance,Mean (min),Median (min)
A,Yes,62.16%,11.94,9.71
H,Yes,90.70%,18.65,17.45
V1,Yes,nan%,49.78,41.2
A,No,93.33%,16.31,14.95
H,No,87.50%,23.59,18.56
V1,No,nan%,51.44,55.54


In [114]:
print_results(load_csv("events_strategy=closest_prioritizeTriage=false_responseRestricted=false_scheduleBreaks=true"))


Triage,Urban,Compliance,Mean (min),Median (min)
A,Yes,48.65%,14.21,12.59
H,Yes,81.40%,22.09,19.47
V1,Yes,nan%,58.58,53.26
A,No,73.33%,18.98,14.78
H,No,75.00%,28.64,26.84
V1,No,nan%,73.71,76.75


In [115]:
print_results(load_csv("events_strategy=closest_prioritizeTriage=false_responseRestricted=true_scheduleBreaks=false"))


Triage,Urban,Compliance,Mean (min),Median (min)
A,Yes,51.35%,13.63,11.19
H,Yes,89.53%,18.68,17.48
V1,Yes,nan%,49.3,40.71
A,No,80.00%,17.97,18.03
H,No,100.00%,20.87,18.29
V1,No,nan%,46.88,51.27


In [116]:
print_results(load_csv("events_strategy=closest_prioritizeTriage=false_responseRestricted=true_scheduleBreaks=true"))


Triage,Urban,Compliance,Mean (min),Median (min)
A,Yes,41.89%,15.52,12.88
H,Yes,84.88%,20.06,17.39
V1,Yes,nan%,54.24,46.34
A,No,73.33%,21.51,19.2
H,No,87.50%,23.79,18.73
V1,No,nan%,63.46,64.39


In [117]:
print_results(load_csv("events_strategy=closest_prioritizeTriage=true_responseRestricted=false_scheduleBreaks=false"))


Triage,Urban,Compliance,Mean (min),Median (min)
A,Yes,68.92%,10.6,9.39
H,Yes,89.53%,19.32,17.52
V1,Yes,nan%,56.01,51.52
A,No,100.00%,15.83,14.78
H,No,100.00%,21.55,18.2
V1,No,nan%,47.02,51.35


In [118]:
print_results(load_csv("events_strategy=closest_prioritizeTriage=true_responseRestricted=false_scheduleBreaks=true"))


Triage,Urban,Compliance,Mean (min),Median (min)
A,Yes,58.11%,11.49,10.09
H,Yes,87.21%,20.65,18.27
V1,Yes,nan%,63.21,54.95
A,No,73.33%,18.91,18.05
H,No,75.00%,24.13,18.15
V1,No,nan%,75.63,59.66


In [119]:
print_results(load_csv("events_strategy=closest_prioritizeTriage=true_responseRestricted=true_scheduleBreaks=false"))


Triage,Urban,Compliance,Mean (min),Median (min)
A,Yes,52.70%,12.28,10.84
H,Yes,88.37%,19.8,17.88
V1,Yes,nan%,53.8,41.42
A,No,93.33%,15.45,14.87
H,No,100.00%,21.62,18.82
V1,No,nan%,46.88,51.23


In [120]:
print_results(load_csv("events_strategy=closest_prioritizeTriage=true_responseRestricted=true_scheduleBreaks=true"))


Triage,Urban,Compliance,Mean (min),Median (min)
A,Yes,52.70%,13.02,11.38
H,Yes,80.23%,22.33,20.27
V1,Yes,nan%,66.89,54.7
A,No,80.00%,17.04,16.73
H,No,100.00%,20.85,18.98
V1,No,nan%,74.25,63.09


In [121]:
print_results(load_csv("events_strategy=random_prioritizeTriage=false_responseRestricted=false_scheduleBreaks=false"))


Triage,Urban,Compliance,Mean (min),Median (min)
A,Yes,20.27%,24.09,22.82
H,Yes,45.35%,34.46,30.87
V1,Yes,nan%,77.64,67.59
A,No,20.00%,36.62,36.33
H,No,12.50%,63.5,66.67
V1,No,nan%,98.67,91.08


In [122]:
print_results(load_csv("events_strategy=random_prioritizeTriage=false_responseRestricted=false_scheduleBreaks=true"))


Triage,Urban,Compliance,Mean (min),Median (min)
A,Yes,5.41%,66.13,31.49
H,Yes,24.42%,63.51,39.56
V1,Yes,nan%,106.3,79.12
A,No,26.67%,38.33,35.1
H,No,25.00%,73.3,58.77
V1,No,nan%,124.21,138.51


In [123]:
print_results(load_csv("events_strategy=random_prioritizeTriage=false_responseRestricted=true_scheduleBreaks=false"))


Triage,Urban,Compliance,Mean (min),Median (min)
A,Yes,16.22%,25.99,25.3
H,Yes,33.72%,36.8,33.71
V1,Yes,nan%,77.51,69.08
A,No,6.67%,38.15,35.43
H,No,25.00%,51.12,54.17
V1,No,nan%,103.69,96.86


In [124]:
print_results(load_csv("events_strategy=random_prioritizeTriage=false_responseRestricted=true_scheduleBreaks=true"))


Triage,Urban,Compliance,Mean (min),Median (min)
A,Yes,6.76%,109.57,31.93
H,Yes,22.09%,56.07,40.06
V1,Yes,nan%,95.72,82.85
A,No,20.00%,81.36,36.18
H,No,12.50%,74.69,50.48
V1,No,nan%,134.56,119.98


In [125]:
print_results(load_csv("events_strategy=random_prioritizeTriage=true_responseRestricted=false_scheduleBreaks=false"))


Triage,Urban,Compliance,Mean (min),Median (min)
A,Yes,6.76%,24.84,22.96
H,Yes,36.05%,42.67,38.23
V1,Yes,nan%,119.79,91.14
A,No,0.00%,43.54,40.93
H,No,0.00%,56.63,53.41
V1,No,nan%,235.04,244.53


In [126]:
print_results(load_csv("events_strategy=random_prioritizeTriage=true_responseRestricted=false_scheduleBreaks=true"))


Triage,Urban,Compliance,Mean (min),Median (min)
A,Yes,9.46%,23.09,21.03
H,Yes,32.56%,44.9,41.48
V1,Yes,nan%,234.76,177.39
A,No,26.67%,39.45,39.38
H,No,12.50%,73.97,59.32
V1,No,nan%,304.01,264.61


In [127]:
print_results(load_csv("events_strategy=random_prioritizeTriage=true_responseRestricted=true_scheduleBreaks=false"))


Triage,Urban,Compliance,Mean (min),Median (min)
A,Yes,14.86%,25.77,23.52
H,Yes,36.05%,40.47,35.23
V1,Yes,nan%,110.82,83.79
A,No,20.00%,35.02,37.82
H,No,37.50%,47.5,52.84
V1,No,nan%,142.58,113.97


In [128]:
print_results(load_csv("events_strategy=random_prioritizeTriage=true_responseRestricted=true_scheduleBreaks=true"))


Triage,Urban,Compliance,Mean (min),Median (min)
A,Yes,12.16%,24.88,24.61
H,Yes,27.91%,48.21,43.32
V1,Yes,nan%,201.66,139.62
A,No,6.67%,45.11,47.78
H,No,25.00%,65.67,60.52
V1,No,nan%,311.93,319.45
