<?xml version="1.0" encoding="UTF-8"?><!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.0 20040830//EN" "journalpublishing.dtd"><article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" dtd-version="2.0" xml:lang="en" article-type="research-article"><front><journal-meta><journal-id journal-id-type="nlm-ta">J Med Internet Res</journal-id><journal-id journal-id-type="publisher-id">jmir</journal-id><journal-id journal-id-type="index">1</journal-id><journal-title>Journal of Medical Internet Research</journal-title><abbrev-journal-title>J Med Internet Res</abbrev-journal-title><issn pub-type="epub">1438-8871</issn><publisher><publisher-name>JMIR Publications</publisher-name><publisher-loc>Toronto, Canada</publisher-loc></publisher></journal-meta><article-meta><article-id pub-id-type="publisher-id">v28i1e97672</article-id><article-id pub-id-type="doi">10.2196/97672</article-id><article-categories><subj-group subj-group-type="heading"><subject>Original Paper</subject></subj-group></article-categories><title-group><article-title>Predicting Critical Outcomes in Suspected Cardiopulmonary Emergencies Using Dispatch Narratives: Temporal Validation Study</article-title></title-group><contrib-group><contrib contrib-type="author" equal-contrib="yes"><name name-style="western"><surname>Li</surname><given-names>Zhe</given-names></name><degrees>PhD</degrees><xref ref-type="aff" rid="aff1">1</xref><xref ref-type="fn" rid="equal-contrib1">*</xref></contrib><contrib contrib-type="author" equal-contrib="yes"><name name-style="western"><surname>Shi</surname><given-names>Lei</given-names></name><degrees>MD</degrees><xref ref-type="aff" rid="aff1">1</xref><xref ref-type="fn" rid="equal-contrib1">*</xref></contrib><contrib contrib-type="author" equal-contrib="yes"><name name-style="western"><surname>Luo</surname><given-names>Chunting</given-names></name><degrees>MD</degrees><xref ref-type="aff" rid="aff2">2</xref><xref ref-type="fn" rid="equal-contrib1">*</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Huang</surname><given-names>Siqi</given-names></name><degrees>MD</degrees><xref ref-type="aff" rid="aff1">1</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Qin</surname><given-names>Jianmin</given-names></name><degrees>MD</degrees><xref ref-type="aff" rid="aff2">2</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Yao</surname><given-names>Min</given-names></name><degrees>MD</degrees><xref ref-type="aff" rid="aff2">2</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Zhu</surname><given-names>Sanshan</given-names></name><degrees>MD</degrees><xref ref-type="aff" rid="aff1">1</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Huang</surname><given-names>Zhengzhuang</given-names></name><degrees>MD</degrees><xref ref-type="aff" rid="aff1">1</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Nong</surname><given-names>Yinghua</given-names></name><degrees>MD</degrees><xref ref-type="aff" rid="aff2">2</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Qiu</surname><given-names>Guozheng</given-names></name><degrees>MD</degrees><xref ref-type="aff" rid="aff1">1</xref></contrib><contrib contrib-type="author" corresp="yes"><name name-style="western"><surname>Lyu</surname><given-names>Liwen</given-names></name><degrees>MD, PhD</degrees><xref ref-type="aff" rid="aff1">1</xref></contrib></contrib-group><aff id="aff1"><institution>The People's Hospital of Guangxi Zhuang Autonomous Region</institution><addr-line>No. 6 Taoyuan Road</addr-line><addr-line>Nanning</addr-line><addr-line>Guangxi</addr-line><country>China</country></aff><aff id="aff2"><institution>Nanning Emergency Medical Center</institution><addr-line>Nanning</addr-line><addr-line>Guangxi</addr-line><country>China</country></aff><contrib-group><contrib contrib-type="editor"><name name-style="western"><surname>Steenstra</surname><given-names>Ivan</given-names></name></contrib></contrib-group><contrib-group><contrib contrib-type="reviewer"><name name-style="western"><surname>Bethanabatla</surname><given-names>Amruthavalli</given-names></name></contrib><contrib contrib-type="reviewer"><name name-style="western"><surname>Mahmoud</surname><given-names>Mahmoud Badee Rokaya</given-names></name></contrib><contrib contrib-type="reviewer"><name name-style="western"><surname>Ogunbowale</surname><given-names>Oluwatobilola</given-names></name></contrib></contrib-group><author-notes><corresp>Correspondence to Liwen Lyu, MD, PhD, The People's Hospital of Guangxi Zhuang Autonomous Region, No. 6 Taoyuan Road, Nanning, Guangxi, 530021, China, +86 15978194424; <email>iculvliwen@163.com</email></corresp><fn fn-type="equal" id="equal-contrib1"><label>*</label><p>these authors contributed equally</p></fn></author-notes><pub-date pub-type="collection"><year>2026</year></pub-date><pub-date pub-type="epub"><day>14</day><month>8</month><year>2026</year></pub-date><volume>28</volume><elocation-id>e97672</elocation-id><history><date date-type="received"><day>09</day><month>04</month><year>2026</year></date><date date-type="rev-recd"><day>22</day><month>07</month><year>2026</year></date><date date-type="accepted"><day>22</day><month>07</month><year>2026</year></date></history><copyright-statement>&#x00A9; Zhe Li, Lei Shi, Chunting Luo, Siqi Huang, Jianmin Qin, Min Yao, Sanshan Zhu, Zhengzhuang Huang, Yinghua Nong, Guozheng Qiu, Liwen Lyu. Originally published in the Journal of Medical Internet Research (<ext-link ext-link-type="uri" xlink:href="https://www.jmir.org">https://www.jmir.org</ext-link>), 14.8.2026. </copyright-statement><copyright-year>2026</copyright-year><license license-type="open-access" xlink:href="https://creativecommons.org/licenses/by/4.0/"><p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (<ext-link ext-link-type="uri" xlink:href="https://creativecommons.org/licenses/by/4.0/">https://creativecommons.org/licenses/by/4.0/</ext-link>), which permits unrestricted use, distribution, and reproduction in any medium, provided the original work, first published in the Journal of Medical Internet Research (ISSN 1438-8871), is properly cited. The complete bibliographic information, a link to the original publication on <ext-link ext-link-type="uri" xlink:href="https://www.jmir.org/">https://www.jmir.org/</ext-link>, as well as this copyright and license information must be included.</p></license><self-uri xlink:type="simple" xlink:href="https://www.jmir.org/2026/1/e97672"/><abstract><sec><title>Background</title><p>Early risk stratification in emergency medical services (EMS) is essential for patients presenting with acute cardiopulmonary symptoms, yet prehospital decision-making at the dispatch stage is often based on limited structured information. Free-text dispatch narratives may contain additional clinical signals, but their role in early risk assessment remains insufficiently characterized.</p></sec><sec><title>Objective</title><p>This study aims to develop and temporally validate a natural language processing&#x2013;assisted machine learning framework for early risk stratification using free-text EMS dispatch narratives and to evaluate its incremental value beyond conventional structured dispatch information.</p></sec><sec sec-type="methods"><title>Methods</title><p>We conducted a population-based retrospective cohort study using EMS dispatch records from Nanning, China, between 2021 and 2025. Adult patients with suspected cardiopulmonary symptoms were identified based on predefined complaint keywords. After excluding nonmedical and incomplete records, 38,523 cases with available free-text narratives were included. To simulate real-world deployment, data from 2021 to 2024 (n=28,332) were used for model development, and 2025 data (n=10,191) served as an independent temporal test cohort. Dispatch narratives were processed using a natural language processing pipeline based on character-level n-grams and combined with structured variables (age, sex, call time) in a multimodal machine learning framework. The primary outcome was a composite prehospital critical outcome comprising death, clinical deterioration, or lack of response to initial treatment. Model performance was evaluated using area under the receiver operating characteristic curve (AUROC), area under the precision-recall curve, calibration, and decision curve analysis.</p></sec><sec sec-type="results"><title>Results</title><p>Among the 38,523 included patients, 12,476 (32.4%) experienced the primary composite outcome. The median age was 68.0 (IQR 54.0&#x2010;79.0) years, and 23,019 (59.8%) patients were male. In the temporally independent 2025 cohort, an expanded structured baseline model achieved an AUROC of 0.681 (95% CI 0.669&#x2010;0.693). Incorporation of narrative features improved performance (AUROC 0.803, 95% CI 0.792&#x2010;0.814), with marginal additional gain from multimodal integration (AUROC 0.808, 95% CI 0.798&#x2010;0.818). The multimodal model achieved an area under the precision-recall curve of 0.630 (95% CI 0.611&#x2010;0.651), substantially exceeding the no-skill baseline defined by the outcome prevalence in the temporal test cohort (2464/10,191, 24.2%). Model performance remained consistent across age and sex subgroups, and calibration was acceptable (Brier score 0.1878). In a risk enrichment analysis, the top 10% (n=1019) of predicted high-risk cases accounted for 32.1% (n=791) of all critical outcomes, representing a 3.2-fold enrichment. Decision curve analysis indicated a higher net benefit compared with treat-all and treat-none strategies across a range of threshold probabilities.</p></sec><sec sec-type="conclusions"><title>Conclusions</title><p>Free-text dispatch narratives contain clinically relevant information associated with early risk stratification in patients with suspected cardiopulmonary emergencies. Incorporating narrative-derived features into a structured modeling framework may complement existing EMS dispatch systems and support more informed decision-making prior to patient contact. Further external validation and prospective evaluation are warranted.</p></sec></abstract><kwd-group><kwd>digital health</kwd><kwd>emergency medical dispatch</kwd><kwd>natural language processing</kwd><kwd>machine learning</kwd><kwd>clinical decision support</kwd><kwd>prehospital care</kwd><kwd>risk stratification</kwd><kwd>artificial intelligence</kwd></kwd-group></article-meta></front><body><sec id="s1" sec-type="intro"><title>Introduction</title><p>Acute cardiopulmonary emergencies represent one of the most time-critical categories encountered by emergency medical services (EMS), frequently presenting with symptoms such as chest pain, dyspnea, syncope, or sudden loss of consciousness [<xref ref-type="bibr" rid="ref1">1</xref>]. These conditions&#x2014;including acute coronary syndromes, malignant arrhythmias, and cardiogenic shock&#x2014;are associated with rapid clinical deterioration and high mortality if not promptly recognized and managed [<xref ref-type="bibr" rid="ref2">2</xref>]. For EMS teams, the period before patient contact is particularly crucial. Decisions made at the time of dispatch&#x2014;such as mobilizing advanced life support (ALS) resources, preparing defibrillation or airway equipment, and predetermining the most appropriate destination hospital&#x2014;can significantly influence the timeliness and effectiveness of subsequent care [<xref ref-type="bibr" rid="ref3">3</xref>]. In this context, early risk stratification based solely on information obtained during the emergency call has the potential to optimize prehospital readiness and reduce delays in definitive treatment [<xref ref-type="bibr" rid="ref4">4</xref>].</p><p>In current practice, early assessment of patients with suspected acute cardiovascular conditions is largely indirect and experience-driven [<xref ref-type="bibr" rid="ref5">5</xref>]. Before arriving on scene, EMS personnel and dispatchers must rely on limited structured information, such as patient age, reported symptoms, and predefined chief complaint categories, to infer severity. However, such structured variables often lack granularity and fail to reflect the dynamic and heterogeneous nature of cardiovascular emergencies [<xref ref-type="bibr" rid="ref6">6</xref>]. As a result, prehospital decision-making frequently depends on subjective interpretation rather than standardized, quantitative risk assessment. This reliance on experience may lead to variability in triage decisions, with potential consequences including delayed escalation of care for high-risk patients or unnecessary allocation of advanced resources to lower-risk cases [<xref ref-type="bibr" rid="ref7">7</xref>]. These challenges underscore a critical gap in current dispatch systems: the absence of objective tools capable of translating early call information into actionable risk stratification [<xref ref-type="bibr" rid="ref8">8</xref>].</p><p>During emergency calls, particularly in cases involving cardiovascular complaints, callers often provide detailed narrative descriptions that extend beyond structured categories [<xref ref-type="bibr" rid="ref9">9</xref>]. Expressions such as &#x201C;no response,&#x201D; &#x201C;sudden collapse,&#x201D; &#x201C;gasping,&#x201D; or &#x201C;progressively worsening chest discomfort&#x201D; may convey important clinical cues related to hemodynamic instability or impending cardiac arrest [<xref ref-type="bibr" rid="ref10">10</xref>]. These free-text narratives, whether documented by dispatchers or generated through real-time transcription systems, contain rich contextual information that is not captured in conventional triage frameworks [<xref ref-type="bibr" rid="ref11">11</xref>]. With the rapid development of natural language processing (NLP) techniques, it has become increasingly feasible to systematically analyze such unstructured data and extract clinically meaningful patterns [<xref ref-type="bibr" rid="ref12">12</xref>,<xref ref-type="bibr" rid="ref13">13</xref>]. However, many previous clinical NLP studies have focused on longer clinical narratives or electronic health records, while text-based analysis of ultrashort, time-sensitive dispatch communications remains limited. Moreover, previous EMS NLP models have primarily targeted single high-acuity conditions, such as out-of-hospital cardiac arrest. In contrast, comprehensive early risk stratification across heterogeneous cardiopulmonary presentations encountered at dispatch remains largely underexplored. Crucially, unlike models that rely on prehospital vital signs or clinician assessments after EMS arrival, our framework performs risk stratification using unstructured information available at the time of the emergency call, enabling decision support at the earliest stage of the EMS pathway&#x2014;prior to patient contact. Finally, because previous machine learning models in EMS triage frequently rely on random data splits that may overestimate performance, the generalizability of these approaches under real-world temporal variation has not been fully established [<xref ref-type="bibr" rid="ref14">14</xref>]. Therefore, this study advances prior work by developing a lightweight, interpretable NLP framework for emergency dispatch and evaluating its real-world robustness through a rigorous temporal validation strategy.</p><p>In light of these considerations, rather than focusing on algorithmic innovation, the present study aimed to develop and temporally validate an automated risk stratification framework for patients presenting with cardiovascular-related complaints during emergency calls, with a primary focus on evaluating its real-world clinical applicability. Leveraging a large real-world EMS dataset, we applied NLP methods to extract features from dispatch narratives and integrated them with structured variables within a multimodal machine learning framework. Model performance was evaluated using a temporally independent test cohort to assess robustness under evolving case mix and usage patterns. In addition, we examined model interpretability and clinical usefulness through feature attribution and risk enrichment analyses. By focusing on early-stage information available prior to patient contact, this study seeks to explore the potential of NLP-assisted models to support prehospital decision-making and optimize resource allocation in time-sensitive cardiopulmonary emergencies.</p></sec><sec id="s2" sec-type="methods"><title>Methods</title><sec id="s2-1"><title>Study Design and Data Source</title><p>This study was designed as a population-based, real-world retrospective cohort analysis using data derived from the EMS system in Nanning, Guangxi Zhuang Autonomous Region, China. All emergency calls are centrally coordinated through the Nanning Emergency Medical Center, which operates the unified municipal dispatch system and provides prehospital emergency care across both urban and surrounding suburban areas.</p><p>The EMS database integrates dispatch information with standardized electronic prehospital medical records completed by on-scene physicians, allowing linkage between call-level information and clinical outcomes. This structure enabled comprehensive capture of both structured variables and free-text narratives generated during the dispatch process and subsequent patient assessment.</p></sec><sec id="s2-2"><title>Study Population and Temporal Validation Strategy</title><p>All EMS dispatch records between January 1, 2021, and December 31, 2025, were screened. Patients were eligible for inclusion if the emergency call was associated with symptoms suggestive of acute cardiopulmonary conditions, identified through predefined dispatch complaint categories and keyword-based filtering of call narratives (based on standard EMS dispatch guidelines). These included presentations such as chest pain, dyspnea, syncope, and altered consciousness. We excluded records related to trauma, nonmedical events, interfacility transfers, and cases with missing key variables, including dispatch narratives or outcome labels. Given the extremely low rate of missing data (&#x003C;1%), a complete case analysis approach was used. After applying these criteria, a total of 38,523 cases were included in the final analysis.</p><p>To simulate real-world prospective deployment and minimize information leakage, a temporal validation strategy was adopted. Data from January 2021 to December 2024 were used as the training cohort (n=28,332) for model development and internal validation, while data from January to December 2025 were reserved as an independent temporal test cohort (n=10,191). This approach allowed the evaluation of model performance under natural variations in case mix and EMS utilization over time.</p></sec><sec id="s2-3"><title>Definition of the Clinical Outcome</title><p>The primary outcome of interest was a composite end point uniformly defined as a &#x201C;critical outcome.&#x201D; This label was derived from standardized on-scene medical records completed by EMS physicians and reflected patients experiencing a critical outcome requiring urgent intervention. Operationally, these outcomes were recorded by attending EMS physicians using electronic prehospital care reports immediately following the mission. Data standardization was enforced by municipal EMS protocols, and validity was supported by routine electronic logic checks and administrative quality assurance within the EMS database. To improve objectivity and reproducibility, the composite end point was anchored to operational EMS records, including documented on-scene death, requirement for ALS interventions (eg, airway management, cardiopulmonary resuscitation), or persistent instability necessitating immediate emergency department resuscitation upon hospital arrival. These criteria were based on structured EMS documentation fields to reduce interphysician variability.</p><p>Specifically, outcome labels were generated using predefined operational criteria and determined using standardized structured fields in the prehospital electronic medical record rather than subjective free-text descriptions, thereby improving interrater consistency. Critical outcomes included (1) death at the scene, (2) persistent hemodynamic or respiratory instability (eg, systolic blood pressure persistently &#x003C;90 mm Hg requiring vasopressors, or peripheral capillary oxygen saturation [SpO2] &#x003C;90% despite high-flow oxygen) despite initial resuscitative efforts, and (3) refractory or rapidly worsening clinical status necessitating immediate escalation of care. This composite end point was chosen to capture clinically meaningful high-risk cases that would benefit from early identification and optimized prehospital preparation. Although not formally validated as a standardized end point, this composite outcome reflects the occurrence of operationally relevant critical outcomes and has been widely adopted in EMS-based risk stratification studies focusing on early decision-making.</p></sec><sec id="s2-4"><title>NLP Pipeline</title><p>To prevent information leakage, the dataset was temporally partitioned into training (2021&#x2010;2024) and test (2025) cohorts before vocabulary construction, feature selection, and model development. Free-text data, specifically the dispatch narratives documented by call-takers during the emergency call, were processed using a structured NLP pipeline. Prior to feature extraction, all text entries underwent standardized preprocessing, including removal of punctuation, special characters, and noninformative stop-words (using a customized Chinese medical stop-word dictionary), to reduce noise and improve signal consistency. Crucially, the NLP pipeline, including vocabulary construction and selection of the top 300 most frequent n-grams, was developed exclusively using the training cohort. The independent temporal test cohort remained isolated throughout model development and was subsequently transformed using only the predefined vocabulary and preprocessing rules derived from the training data. No information from the temporal test cohort was used during feature engineering, model training, or hyperparameter tuning.</p><p>The cleaned text was then transformed into character-level n-grams, specifically bigrams, to capture local semantic patterns within short phrases commonly used in emergency communication. Given that the dispatch narratives were recorded in Chinese, this character-level tokenization strategy was deliberately used. It effectively bypasses the inaccuracies commonly associated with Chinese word segmentation tools when parsing highly colloquial, fragmented, or abbreviated emergency medical texts, allowing for direct capture of native semantic structures. To ensure computational efficiency and reduce dimensionality, feature selection was performed based on term frequency across the training corpus. The top 300 most frequent n-grams were retained as the core feature set. These selected features were subsequently encoded into a structured numerical matrix. Given the ultra-short nature of emergency dispatch narratives, where term recurrence within a single record is exceedingly rare, a binary encoding scheme (presence vs absence) was used. In this specific short-text context, binary encoding effectively captures the essential semantic signal (functioning as a Boolean term frequency) without the added complexity of global weighting schemes such as term frequency-inverse document frequency (TF-IDF).</p><p>While advanced deep learning architectures such as transformer-based large language models (eg, bidirectional encoder representations from transformers, BERT) are increasingly used in medical informatics, a lightweight, explainable, n-gram-based approach was deliberately selected for this study. Dispatch narratives are inherently highly fragmented, abbreviated, and lack complex grammatical structures, rendering the contextual advantages of deep learning models marginal. Furthermore, the explicit n-gram matrix provides greater semantic transparency, which enhances the interpretability of downstream machine learning predictions&#x2014;a mandatory prerequisite for building algorithmic trust in high-stakes emergency triage. Importantly, feature selection (top 300 n-grams) was performed exclusively within the training dataset to prevent information leakage into the temporal test cohort.</p></sec><sec id="s2-5"><title>Machine Learning Development and Multimodal Fusion</title><p>To leverage both structured and unstructured information, a multimodal modeling framework was constructed by combining the NLP-derived features with baseline structured variables, including age, sex, and call time. Given the class imbalance in the training cohort, with a relatively lower proportion of critical outcomes, class weighting was incorporated into the model training process to enhance sensitivity toward high-risk cases. Specifically, the scale_pos_weight hyperparameter in the Extreme Gradient Boosting (XGBoost) algorithm was defined as the ratio of negative (noncritical) to positive (critical outcome) in the training set.</p><p>The primary model was developed using XGBoost, selected for its ability to handle nonlinear relationships and interactions. Hyperparameter tuning (including parameters such as maximum depth, learning rate, and subsample ratio) was conducted via grid search using 5-fold cross-validation exclusively within the training cohort to optimize the final model parameters. For benchmarking purposes, additional models, including random forest and penalized logistic regression, were also trained under the same framework. To ensure a fair comparison and reproducibility, all model hyperparameters were optimized within the training set only, without access to the temporal test cohort.</p></sec><sec id="s2-6"><title>Model Evaluation, Explainability, and Clinical Usefulness</title><p>Model performance was evaluated on the independent 2025 temporal test cohort. Discrimination was assessed using the area under the receiver operating characteristic curve (AUROC). Given the class imbalance, additional metrics, including the area under the precision-recall curve (AUPRC) and <italic>F</italic><sub>1</sub>-score, were calculated to provide a more comprehensive assessment of model performance. Calibration was evaluated using the Brier score.</p><p>To enhance interpretability, Shapley Additive Explanations (SHAP) were applied to quantify the contribution of individual features to model predictions, allowing the identification of key textual patterns associated with increased risk.</p><p>To assess potential clinical applicability, a risk enrichment analysis was performed by ranking patients according to predicted risk and evaluating the proportion of true critical outcomes captured within predefined high-risk strata (top 5%, 10%, and 20%). In addition, decision curve analysis (DCA) was conducted to estimate the net benefit across a range of threshold probabilities, comparing the model against default strategies of treating all or no patients as high risk.</p></sec><sec id="s2-7"><title>Statistical Analysis</title><p>Continuous variables were presented as medians with IQRs and compared using the Mann-Whitney <italic>U</italic> test, given their nonnormal distributions. Categorical variables were summarized as counts (percentages) and compared using the Pearson chi-square test or Fisher exact test, as appropriate.</p><p>For model performance evaluation, 95% CIs for the AUROC were computed. Statistical comparisons of discriminative performance between the baseline structured model, the NLP-only model, and the multimodal fusion model were conducted using the DeLong test. Given the inherent class imbalance in critical outcomes, the AUPRC and <italic>F</italic><sub>1</sub>-scores were additionally calculated as robust evaluation metrics. The clinical usefulness of the predictive models was assessed using DCA to quantify the net benefit across a range of threshold probabilities.</p><p>All data preprocessing, machine learning development, and statistical analyses were performed using R software (version 4.3.2; R Foundation for Statistical Computing). Key R packages used included xgboost for model training, pROC for ROC analysis and DeLong tests, dcurves for DCA, and relevant libraries for SHAP value computation and visualization. A 2-sided <italic>P</italic> value &#x003C;.05 was considered statistically significant.</p></sec><sec id="s2-8"><title>Ethical Considerations</title><p>This study was approved by the Institutional Review Board of Guangxi Zhuang Autonomous Region People&#x2019;s Hospital (approval number KY-KJT-2024&#x2010;239). The requirement for informed consent was waived because of the retrospective study design and the use of anonymized EMS data.</p><p>Additional methodological details and model specifications are provided in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>. The completed TRIPOD (Transparent Reporting of a Multivariable Prediction Model for Individual Prognosis or Diagnosis) reporting checklist is available in <xref ref-type="supplementary-material" rid="app2">Checklist 1</xref>.</p></sec></sec><sec id="s3" sec-type="results"><title>Results</title><sec id="s3-1"><title>Study Population and Temporal Dataset Shift</title><p>A total of 38,523 EMS dispatch records with available free-text narratives were included in the final analytic cohort. Historical data from 2021 to 2024 (n=28,332) were used for model development, while an independent temporal cohort from 2025 (n=10,191) served as the temporal validation dataset (<xref ref-type="fig" rid="figure1">Figure 1</xref>).</p><p>The flow diagram illustrates the patient selection process and the dual-track methodological framework. From an initial pool of 38,562 dispatch records screened via cardiopulmonary symptom keywords, 38,523 eligible cases were included after excluding records with missing dispatch narratives or unknown clinical outcomes. To simulate real-world prospective deployment and prevent data leakage, the cohort was temporally split into a historical training dataset (2021&#x2010;2024, n=28,332) and an independent temporal test dataset (2025, n=10,191). The diagram outlines the NLP pipeline, machine learning development, and comprehensive model evaluation strategies.</p><p>Baseline demographic characteristics were highly comparable between the 2 cohorts (<xref ref-type="table" rid="table1">Table 1</xref>). The median patient age was 68.0 (IQR 54.0&#x2010;79.0; <italic>P</italic>=.52) years in both groups, and the proportion of male patients was similar (n=16,872, 59.6% vs n=6147, 60.3%; <italic>P</italic>=.18). The temporal distribution of EMS calls was also consistent, with a median call hour of 13:00 (IQR 8.0-18.0) in both cohorts (<italic>P</italic>=.14). Structured major complaint categories, including dyspnea (n=10,265, 36.2% vs n=3599, 35.3%; <italic>P</italic>=.11), syncope (n=6692, 23.6% vs n=2367, 23.2%; <italic>P</italic>=.44), chest tightness (n=5082, 17.9% vs n=1877, 18.4%; <italic>P</italic>=.27), and chest pain (n=2387, 8.4% vs n=894, 8.8%; <italic>P</italic>=.28), showed stable distributions across the 2 periods.</p><p>However, temporal dataset shifts were observed in the underlying clinical severity. The prevalence of critical outcomes decreased from 35.3% (n=10,012) in the historical training cohort to 24.2% (n=2464) in the 2025 temporal validation cohort (<italic>P</italic>&#x003C;.001; <xref ref-type="fig" rid="figure2">Figure 2</xref>). Furthermore, the prevalence of certain NLP-extracted latent semantic signatures, such as second-hand reporting (n=11,070, 39.1% vs n=4199, 41.2%; <italic>P</italic>&#x003C;.001) and convulsions (n=9129, 32.2% vs n=3651, 35.8%; <italic>P</italic>&#x003C;.001), exhibited significant temporal fluctuations. The presence of this natural dataset shift provided a stringent real-world environment for evaluating model robustness under changing clinical distributions.</p><fig position="float" id="figure1"><label>Figure 1.</label><caption><p>Study flowchart and temporal validation design. AUPRC: area under the precision-recall curve; AUROC: area under the receiver operating characteristic curve; CV: cross-validation; DCA: decision curve analysis; EMS: emergency medical services; NLP: natural language processing; SHAP: Shapley Additive Explanations; XGBoost: Extreme Gradient Boosting.</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="jmir_v28i1e97672_fig01.png"/></fig><table-wrap id="t1" position="float"><label>Table 1.</label><caption><p>Baseline characteristics, structured complaints, and natural language processing (NLP)&#x2013;extracted signatures across temporal cohorts<sup><xref ref-type="table-fn" rid="table1fn1">a</xref></sup>.</p></caption><table id="table1" frame="hsides" rules="groups"><thead><tr><td align="left" valign="bottom">Variables<sup><xref ref-type="table-fn" rid="table1fn2">b</xref></sup></td><td align="left" valign="bottom">Overall cohort (N=38,523)</td><td align="left" valign="bottom">Training cohort (2021&#x2010;2024) (n=28,332)</td><td align="left" valign="bottom">Temporal test cohort (2025) (n=10,191)</td><td align="left" valign="bottom"><italic>P</italic> value</td></tr></thead><tbody><tr><td align="left" valign="top" colspan="5">Demographics and call metrics</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Age (y), median (IQR)</td><td align="left" valign="top">68.0 (54.0&#x2010;79.0)</td><td align="left" valign="top">68.0 (54.0&#x2010;79.0)</td><td align="left" valign="top">68.0 (53.0&#x2010;79.0)</td><td align="left" valign="top">.52</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Male sex, n (%)</td><td align="left" valign="top">23,019 (59.8)</td><td align="left" valign="top">16,872 (59.6)</td><td align="left" valign="top">6147 (60.3)</td><td align="left" valign="top">.18</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Call hour, median (IQR)</td><td align="left" valign="top">13.0 (8.0&#x2010;18.0)</td><td align="left" valign="top">13.0 (8.0&#x2010;18.0)</td><td align="left" valign="top">13.0 (8.0&#x2010;18.0)</td><td align="left" valign="top">.14</td></tr><tr><td align="left" valign="top" colspan="3">Structured complaint category, n (%)</td><td align="left" valign="top"/><td align="left" valign="top"/></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Dyspnea or breathing difficulty</td><td align="left" valign="top">13,864 (36.0)</td><td align="left" valign="top">10,265 (36.2)</td><td align="left" valign="top">3599 (35.3)</td><td align="left" valign="top">.19</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Syncope or unconsciousness</td><td align="left" valign="top">9059 (23.5)</td><td align="left" valign="top">6692 (23.6)</td><td align="left" valign="top">2367 (23.2)</td><td align="left" valign="top">.44</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Chest tightness</td><td align="left" valign="top">6959 (18.1)</td><td align="left" valign="top">5082 (17.9)</td><td align="left" valign="top">1877 (18.4)</td><td align="left" valign="top">.28</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Chest pain or angina</td><td align="left" valign="top">3281 (8.5)</td><td align="left" valign="top">2387 (8.4)</td><td align="left" valign="top">894 (8.8)</td><td align="left" valign="top">.28</td></tr><tr><td align="left" valign="top" colspan="3">NLP-extracted high-risk signatures, n (%)</td><td align="left" valign="top"/><td align="left" valign="top"/></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Second-hand reporting (bystander)</td><td align="left" valign="top">15,269 (39.6)</td><td align="left" valign="top">11,070 (39.1)</td><td align="left" valign="top">4199 (41.2)</td><td align="left" valign="top">&#x003C;.001</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Convulsions or seizure-like movements</td><td align="left" valign="top">12,780 (33.2)</td><td align="left" valign="top">9129 (32.2)</td><td align="left" valign="top">3651 (35.8)</td><td align="left" valign="top">&#x003C;.001</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Coma or altered mental status</td><td align="left" valign="top">8306 (21.6)</td><td align="left" valign="top">6033 (21.3)</td><td align="left" valign="top">2273 (22.3)</td><td align="left" valign="top">.04</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Unresponsive to verbal stimuli</td><td align="left" valign="top">7897 (20.5)</td><td align="left" valign="top">5785 (20.4)</td><td align="left" valign="top">2112 (20.7)</td><td align="left" valign="top">.52</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Diaphoresis or cold sweat</td><td align="left" valign="top">3965 (10.3)</td><td align="left" valign="top">2889 (10.2)</td><td align="left" valign="top">1076 (10.6)</td><td align="left" valign="top">.31</td></tr><tr><td align="left" valign="top" colspan="5">Clinical outcome</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Critical outcome, n (%)</td><td align="left" valign="top">12,476 (32.4)</td><td align="left" valign="top">10,012 (35.3)</td><td align="left" valign="top">2464 (24.2)</td><td align="left" valign="top">&#x003C;.001</td></tr></tbody></table><table-wrap-foot><fn id="table1fn1"><p><sup>a</sup>Categorical variables are presented as counts (percentages) and compared using the Pearson <italic>&#x03C7;</italic><sup>2</sup> test. While standard demographics and complaint categories remained stable, temporal dataset shifts were observed in the prevalence of critical outcomes and latent semantic signatures (eg, second-hand reporting, convulsions) in the 2025 test cohort. Percentages for categories do not sum to 100% due to multiple concurrent symptoms.</p></fn><fn id="table1fn2"><p><sup>b</sup>Continuous variables are presented as median (IQR) and compared using the Mann-Whitney <italic>U</italic> test. </p></fn></table-wrap-foot></table-wrap><fig position="float" id="figure2"><label>Figure 2.</label><caption><p>Temporal dataset shift and baseline stability between the historical training and prospective validation cohorts. (A) Bar chart demonstrating a significant decrease in the prevalence of critical outcomes from 35.3% (n=10,012) in the 2021&#x2010;2024 training cohort to 24.2% (n=2464) in the 2025 test cohort (<italic>P</italic>&#x003C;.001), reflecting a natural dataset shift in real-world emergency medical services (EMS) use. (B) Box plots illustrating the stable distribution of patient age across the 2 cohorts (median 68.0, IQR 54.0&#x2010;79.0 y, <italic>P</italic>=.52). (C) Density plots showing highly consistent bimodal circadian rhythms in emergency call hours between the training and test cohorts (<italic>P</italic>=.14).</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="jmir_v28i1e97672_fig02.png"/></fig></sec><sec id="s3-2"><title>Incremental Value of Dispatch Narratives</title><p>Structured triage variables alone demonstrated improved but still limited predictive ability for identifying critical outcomes in the temporal validation cohort. An expanded structured baseline model incorporating age, sex, call hour, and routinely available structured complaint categories (dyspnea, syncope, chest pain, and chest tightness) achieved an AUROC of 0.681 (95% CI 0.669&#x2010;0.693; <xref ref-type="table" rid="table2">Table 2</xref>). In contrast, models incorporating unstructured dispatch narratives substantially improved predictive discrimination beyond the expanded structured baseline model. The NLP-only model achieved an AUROC of 0.803 (95% CI 0.792&#x2010;0.814). The best performance was observed in the multimodal model integrating both structured variables and NLP-derived textual features, which achieved an AUROC of 0.808 (95% CI 0.798&#x2010;0.818; <xref ref-type="table" rid="table2">Table 2</xref>, <xref ref-type="fig" rid="figure3">Figure 3A</xref>).</p><p>Given the outcome prevalence of 24.2% in the temporal test cohort, precision-recall analysis was additionally performed to assess model performance under class imbalance. The multimodal model achieved an AUPRC of 0.630 (95% CI 0.611&#x2010;0.651), markedly exceeding the baseline prevalence of 0.242 (<xref ref-type="fig" rid="figure3">Figure 3B</xref>). Across classification metrics, the multimodal model consistently outperformed other approaches, achieving the highest <italic>F</italic><sub>1</sub>-score (0.637) and an improved positive predictive value (PPV 0.571) compared with the expanded structured baseline model (PPV 0.341) (<xref ref-type="table" rid="table2">Table 2</xref>, <xref ref-type="fig" rid="figure3">Figure 3C</xref>).</p><table-wrap id="t2" position="float"><label>Table 2.</label><caption><p>Incremental predictive value of automated triage models on the 2025 temporal test cohort<sup><xref ref-type="table-fn" rid="table2fn1">a</xref></sup>.</p></caption><table id="table2" frame="hsides" rules="groups"><thead><tr><td align="left" valign="bottom">Model or features included</td><td align="left" valign="bottom">AUROC<sup><xref ref-type="table-fn" rid="table2fn2">b</xref></sup> (95% CI)</td><td align="left" valign="bottom">AUPRC<sup><xref ref-type="table-fn" rid="table2fn3">c</xref></sup></td><td align="left" valign="bottom">Sensitivity</td><td align="left" valign="bottom">Specificity</td><td align="left" valign="bottom">PPV<sup><xref ref-type="table-fn" rid="table2fn4">d</xref></sup></td><td align="left" valign="bottom">NPV<sup><xref ref-type="table-fn" rid="table2fn5">e</xref></sup></td><td align="left" valign="bottom"><italic>F</italic><sub>1</sub>-score</td></tr></thead><tbody><tr><td align="left" valign="top">Model A: expanded structured baseline (age, sex, call hour, and complaint categories)</td><td align="left" valign="top">0.681 (0.669&#x2010;0.693)</td><td align="left" valign="top">0.373</td><td align="left" valign="top">0.711</td><td align="left" valign="top">0.562</td><td align="left" valign="top">0.341</td><td align="left" valign="top">0.859</td><td align="left" valign="top">0.461</td></tr><tr><td align="left" valign="top">Model B: NLP<sup><xref ref-type="table-fn" rid="table2fn6">f</xref></sup> free-text narratives</td><td align="left" valign="top">0.803 (0.792&#x2010;0.814)</td><td align="left" valign="top">0.621</td><td align="left" valign="top">0.704</td><td align="left" valign="top">0.751</td><td align="left" valign="top">0.474</td><td align="left" valign="top">0.890</td><td align="left" valign="top">0.567</td></tr><tr><td align="left" valign="top">Model C: multimodal fusion (baseline + NLP Narratives)</td><td align="left" valign="top">0.808 (0.798&#x2010;0.818)</td><td align="left" valign="top">0.630</td><td align="left" valign="top">0.720</td><td align="left" valign="top">0.738</td><td align="left" valign="top">0.571</td><td align="left" valign="top">0.845</td><td align="left" valign="top">0.637</td></tr></tbody></table><table-wrap-foot><fn id="table2fn1"><p><sup>a</sup>Performance metrics were evaluated on the independent 2025 temporal test cohort (n=10,191). The optimal classification threshold for metrics (sensitivity, specificity, PPV, NPV, <italic>F</italic><sub>1</sub>-score) was determined using the Youden index exclusively on the training set and subsequently locked for application to the test cohort. Positive predictive value is inherently influenced by the lower outcome prevalence (24.2%) in the test set. The substantial improvement across all metrics, particularly the area under the receiver operating characteristic curve and <italic>F</italic><sub>1</sub>-score, demonstrates the meaningful incremental prognostic value embedded within unstructured dispatch narratives.</p></fn><fn id="table2fn2"><p><sup>b</sup>AUROC: area under the receiver operating characteristic curve.</p></fn><fn id="table2fn3"><p><sup>c</sup>AUPRC: area under the precision-recall curve.</p></fn><fn id="table2fn4"><p><sup>d</sup>PPV: positive predictive value. </p></fn><fn id="table2fn5"><p><sup>e</sup>NPV: negative predictive value.</p></fn><fn id="table2fn6"><p><sup>f</sup>NLP: natural language processing.</p></fn></table-wrap-foot></table-wrap><fig position="float" id="figure3"><label>Figure 3.</label><caption><p>Incremental predictive value of unstructured dispatch narratives. (A) Receiver operating characteristic (ROC) curves comparing the discriminative performance of the expanded structured baseline model (incorporating age, sex, call hour, and complaint categories; red dashed line), and the multimodal fusion model (teal thick line) on the 2025 temporal test cohort. (B) Precision-recall (PR) curves for the 3 models, evaluating performance under class imbalance. The horizontal dashed line represents the baseline prevalence of critical outcomes (0.242). The multimodal model achieved the highest area under the precision-recall curve (AUPRC; 0.630). (C) Bar plot summarizing the area under the receiver operating characteristic curve (AUROC), AUPRC, and <italic>F</italic><sub>1</sub>-score across the 3 models. Error bars indicate 95% CIs, demonstrating the meaningful incremental prognostic value provided by free-text dispatch narratives. AUC: area under the curve; NLP: natural language processing.</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="jmir_v28i1e97672_fig03.png"/></fig></sec><sec id="s3-3"><title>Algorithm Robustness and Comparative Performance</title><p>Comparative evaluation across 3 machine learning algorithms&#x2014;XGBoost, random forest, and penalized logistic regression&#x2014;demonstrated consistent predictive performance (<xref ref-type="fig" rid="figure4">Figure 4A</xref>). The XGBoost model achieved the highest discrimination (AUROC 0.808), closely followed by the random forest model (AUROC 0.807), while penalized logistic regression yielded slightly lower performance (AUROC 0.794).</p><p>These similar performances across different modeling approaches suggest that the predictive signal primarily originated from the textual features themselves rather than from a specific algorithmic architecture. Subgroup analyses further demonstrated stable model performance across demographic and spatial groups (<xref ref-type="fig" rid="figure4">Figure 4B</xref>). The multimodal model achieved AUROC values of 0.825 in patients younger than 65 years and 0.793 in those aged 65 years or older. Performance was also comparable between male (AUROC 0.803) and female (AUROC 0.811) patients, as well as between urban (AUROC 0.805) and suburban or rural (AUROC 0.812) dispatch locations.</p><fig position="float" id="figure4"><label>Figure 4.</label><caption><p>Algorithm robustness, subgroup fairness, and semantic feature saturation. (A) Jitter plot combined with point-ranges showing the stable area under the receiver operating characteristic curve (AUROC) variance across 5-fold cross-validation for penalized logistic regression, random forest, and the Extreme Gradient Boosting (XGBoost) model. (B) Forest plot demonstrating algorithmic fairness across key demographic and spatiotemporal subgroups in the test cohort. The vertical dashed red line represents the overall AUROC (0.808), and the gray shaded area indicates a &#x00B1;2% tolerance margin. All subgroups fall within this margin, indicating unbiased predictive performance. (C) Area-line chart illustrating feature size stability. The predictive performance reached a saturation plateau at the top 300 NLP n-gram features, with no further incremental benefit observed upon expanding to 500 features. NLP: natural language processing.</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="jmir_v28i1e97672_fig04.png"/></fig></sec><sec id="s3-4"><title>Semantic Signatures of Criticality</title><p>Model interpretability was examined using SHAP analysis to identify the most influential textual features contributing to predictions. The SHAP summary plot revealed several clinically meaningful linguistic patterns associated with critical outcomes (<xref ref-type="fig" rid="figure5">Figure 5A</xref>).</p><p>Among the most influential features were terms describing impaired consciousness or respiratory distress, including &#x201C;unresponsive,&#x201D; &#x201C;breathing,&#x201D; &#x201C;coma,&#x201D; &#x201C;no response,&#x201D; &#x201C;heartbeat,&#x201D; and &#x201C;dizziness.&#x201D; Because isolated n-gram tokens can sometimes lack clinical context, reviewing the original source phrases provides better interpretability. For example, the high-risk token &#x201C;heartbeat&#x201D; predominantly appeared in dire phrases such as &#x201C;has no heartbeat&#x201D; or &#x201C;heartbeat stopped.&#x201D; Similarly, &#x201C;breathing&#x201D; was often extracted from urgent reports such as &#x201C;can&#x2019;t breathe&#x201D; or &#x201C;stopped breathing.&#x201D; Conversely, features associated with lower risk, such as &#x201C;conscious,&#x201D; typically appeared in reassuring statements such as &#x201C;the patient is still conscious.&#x201D; These descriptors correspond to symptoms frequently reported by callers or bystanders when patients present with severe cardiopulmonary compromise. Notably, several influential features reflected second-hand descriptions from witnesses, suggesting that the inability of the patient to communicate directly may itself function as an important contextual signal of clinical severity. The semantic patterns identified by the model were consistent with established clinical indicators of critical illness and support the interpretability of the NLP-based approach (<xref ref-type="fig" rid="figure5">Figure 5B</xref>).</p><fig position="float" id="figure5"><label>Figure 5.</label><caption><p>Semantic signatures of criticality derived from Shapley Additive Explanations (SHAP) analysis. (A) SHAP summary beeswarm plot illustrating the impact of the top 15 extracted features on the model&#x2019;s output. Each dot represents an individual patient. The color gradient indicates the original feature value (red=high or present, blue=low or absent), while the position on the <italic>x</italic>-axis represents the SHAP value (positive values drive the prediction toward a critical outcome). Notably, implicit contextual cues, such as &#x201C;Second-hand reporting&#x201D; (symptoms reported by a bystander), were identified as high-risk indicators, alongside explicit symptoms such as &#x201C;Unresponsive&#x201D; and &#x201C;Coma.&#x201D; To enhance clinical interpretability, it is important to note that these isolated natural language processing (NLP) tokens represent longer dispatch phrases (eg, the high-risk feature &#x201C;Heartbeat&#x201D; typically corresponds to the caller stating &#x201C;no heartbeat,&#x201D; whereas the low-risk feature &#x201C;Conscious&#x201D; reflects &#x201C;the patient is still conscious&#x201D;). (B) Bar plot showing the mean absolute SHAP values, ranking the overall magnitude of each feature&#x2019;s contribution to the automated triage decision.</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="jmir_v28i1e97672_fig05.png"/></fig></sec><sec id="s3-5"><title>Early-Warning Simulation and Clinical Usefulness</title><p>To evaluate the potential operational value of automated narrative-based triage, a risk enrichment analysis was conducted using the 2025 temporal validation cohort. When the model prioritized the top 10% (n=1019) of dispatch calls with the highest predicted risk scores, 32.1% (n=791) of all true critical outcomes were captured, representing a 3.2-fold enrichment compared with random allocation (<xref ref-type="fig" rid="figure6">Figure 6A</xref>).</p><p>Expanding the prioritization threshold to the top 20% (n=2038) of calls captured 51.9% (n=1279) of the critical outcomes, corresponding to a 2.6-fold enrichment. These findings suggest that even modest prioritization thresholds could substantially concentrate high-risk patients within a limited subset of EMS dispatch calls.</p><p>Based on the outcome prevalence of 24.2% in the temporal validation cohort, the null model yielded a Brier score of 0.183. The multimodal model achieved a Brier score of 0.188 (95% CI 0.185&#x2010;0.191, estimated via 1000 bootstrap resamples), with a calibration intercept of &#x2212;1.27 and a calibration slope of 1.30. These findings indicate imperfect probability calibration (<xref ref-type="fig" rid="figure6">Figure 6B</xref>). The observed calibration deviation is likely attributable to the substantial prevalence shift between the development and temporal validation cohorts, together with the intentional class weighting strategy used during model development. Because the primary objective of this study was to evaluate temporal transportability and risk ranking of the original model, post hoc recalibration was not applied to the temporal validation cohort. Furthermore, DCA showed that the NLP-enabled multimodal model consistently provided greater net benefit than both treat-all and treat-none strategies across clinically relevant threshold probabilities (<xref ref-type="fig" rid="figure6">Figure 6C</xref>). Together, these results indicate that automated analysis of dispatch narratives may offer practical value for early risk stratification and optimized allocation of prehospital emergency resources.</p><fig position="float" id="figure6"><label>Figure 6.</label><caption><p>Early-warning clinical usefulness and decision curve analysis (DCA). (A) Early-warning risk enrichment curve demonstrating the proportion of true critical outcomes captured when prioritizing dispatch resources based on AI-predicted risk thresholds. Prioritizing the top 10% (n=1019) highest-risk calls successfully captured 32.1% (n=791) of all critical outcomes (a 3.2-fold risk enrichment compared to random allocation, dashed line). (B) Calibration plot assessing the agreement between predicted probabilities and observed critical outcome proportions across risk deciles. The blue error bars indicate 95% CIs for the observed proportions. The model yielded a Brier score of 0.188 (95% CI 0.185&#x2010;0.191). (C) DCA demonstrating that the natural language processing (NLP)&#x2013;enabled multimodal model (red line) provides a higher net benefit than both the &#x201C;treat-all&#x201D; (gray dashed line) and &#x201C;treat-none&#x201D; (black solid line) strategies across all clinically relevant threshold probabilities for advanced life support (ALS) dispatch.</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="jmir_v28i1e97672_fig06.png"/></fig></sec></sec><sec id="s4" sec-type="discussion"><title>Discussion</title><sec id="s4-1"><title>Principal Findings</title><p>In this population-based study of 38,523 EMS dispatches for suspected acute cardiopulmonary symptoms, we evaluated a narrative-based approach to early risk stratification using routinely collected dispatch texts. The findings suggest that information embedded in free-text narratives is associated with clinically relevant risk signals that are not captured by conventional structured variables. When incorporated into a multimodal framework, these semantic features were able to improve risk discrimination and remained relatively stable under temporal dataset shift. Rather than representing a stand-alone solution, these results indicate that narrative-derived features may serve as a complementary component within existing prehospital triage systems. Although the underlying NLP and machine-learning methods are well established, the primary contribution of this study lies in the rigorous temporal validation of this framework and in demonstrating its potential clinical applicability within a real-world EMS dispatch setting. These findings support the value of adapting established, interpretable machine learning methods to clinically relevant implementation challenges rather than introducing a novel prediction algorithm.</p></sec><sec id="s4-2"><title>Comparison With Prior Work</title><p>Early risk stratification is a central objective of prehospital care, particularly for acute cardiopulmonary presentations where delays in recognition may directly affect downstream management [<xref ref-type="bibr" rid="ref15">15</xref>]. Existing dispatch systems largely rely on predefined symptom categories, structured checklists, or dispatcher experience. While such approaches provide operational consistency, they are inherently limited in capturing the complexity and evolution of acute illness prior to patient contact [<xref ref-type="bibr" rid="ref16">16</xref>]. In parallel, there has been increasing interest in applying machine learning and NLP in clinical settings, including electronic health records, radiology reports, and clinical documentation, where unstructured text has been shown to contain meaningful prognostic information [<xref ref-type="bibr" rid="ref17">17</xref>]. However, the use of narrative data at the dispatch level remains relatively underexplored, despite being the earliest available clinical description in the care pathway [<xref ref-type="bibr" rid="ref18">18</xref>].</p><p>Previous work in EMS has primarily focused on structured data, such as response intervals, demographic characteristics, or predefined complaint codes, to support triage and system optimization [<xref ref-type="bibr" rid="ref19">19</xref>]. Some studies have explored prediction models for cardiac arrest recognition or dispatch prioritization, but these approaches often rely on limited input features or are developed within static datasets [<xref ref-type="bibr" rid="ref20">20</xref>,<xref ref-type="bibr" rid="ref21">21</xref>]. In contrast, narrative-based approaches leverage the variability and richness of real-world communication to provide complementary prognostic perspectives on patient status.</p><p>The multimodal model remained superior even after strengthening the structured comparator by incorporating routinely available complaint categories, indicating that narrative-derived features provide clinically meaningful information beyond conventional structured dispatch variables. The observed association between narrative features and clinical outcomes likely reflects the way acute illness is described and perceived during emergency calls. Unlike structured variables, which represent relatively static patient characteristics (eg, age and sex), narrative descriptions capture dynamic information regarding symptom onset, progression, and contextual observations [<xref ref-type="bibr" rid="ref22">22</xref>]. The modest incremental improvement achieved by multimodal integration suggests that narrative-derived features may already capture much of the clinically relevant information required for early risk stratification, thereby limiting the additional predictive contribution of basic structured variables when dispatch narratives are sufficiently informative. Nevertheless, structured variables may still provide complementary value in situations where narrative information is limited or incomplete, for example, because of communication difficulties or very brief emergency calls [<xref ref-type="bibr" rid="ref23">23</xref>,<xref ref-type="bibr" rid="ref24">24</xref>]. Taken together, these findings suggest that incorporating free-text dispatch narratives alongside structured information may improve early risk assessment and support more informed allocation of EMS resources. Crucially, transformer-based architectures have demonstrated excellent performance across a wide range of clinical NLP tasks by capturing complex contextual relationships. However, emergency dispatch narratives are typically short, fragmented, and highly abbreviated, and the incremental benefit of large language models in this setting remains uncertain. In addition, lightweight n-gram&#x2013;based models offer greater interpretability through explicit feature attribution and generally require substantially fewer computational resources, which may facilitate future implementation in resource-constrained EMS dispatch environments. Future studies directly comparing lightweight and transformer-based approaches across diverse EMS datasets will be valuable to better define their respective strengths and optimal application scenarios [<xref ref-type="bibr" rid="ref25">25</xref>,<xref ref-type="bibr" rid="ref26">26</xref>].</p><p>From a clinical and operational perspective, the potential value of narrative-informed risk stratification lies in its timing [<xref ref-type="bibr" rid="ref27">27</xref>]. Decisions regarding dispatch priority, resource allocation, and destination planning are often made before direct patient assessment. In this context, even modest improvements in early risk identification may influence preparedness at multiple levels, including equipment readiness, crew configuration, and prenotification of receiving facilities [<xref ref-type="bibr" rid="ref28">28</xref>]. Rather than replacing existing triage protocols, a narrative-based risk signal could be integrated as an additional layer of information to support decision-making under uncertainty [<xref ref-type="bibr" rid="ref29">29</xref>]. This may be particularly relevant in settings with constrained ALS resources, where aligning resource intensity with patient risk remains a persistent challenge [<xref ref-type="bibr" rid="ref30">30</xref>].</p></sec><sec id="s4-3"><title>Limitations</title><p>Several limitations should be considered when interpreting these findings. First, the analysis depends on the accuracy and completeness of dispatch narratives, which may vary according to dispatcher training, caller communication ability, and situational stress. Misclassification or information loss at the transcription stage cannot be excluded. Second, although temporal validation was performed, the data were derived from a single EMS system in China. Differences in language, dispatch workflows, documentation practices, and health care system organization may limit the applicability of the proposed model to other regions and countries. Third, the outcome definition is based on prehospital assessments and does not capture in-hospital end points such as definitive diagnosis, survival, or neurological recovery. Furthermore, our composite &#x201C;critical outcome&#x201D; relies on local EMS operational criteria and intervention thresholds. Although the use of standardized structured fields likely improves internal consistency within our system, these operational practices may vary across different regional EMS networks, potentially affecting the external generalizability of our outcome definition. However, from an EMS operational perspective, the occurrence of a prehospital critical outcome represents the most direct and actionable proxy for dispatch triage because it directly informs decisions regarding on-scene ALS and transport prioritization, independent of downstream in-hospital management. Fourth, our NLP pipeline relied on character-level n-grams and tree-based ensemble models. Although these methods offer advantages in interpretability and computational efficiency, they may not fully capture complex semantic relationships compared with modern pretrained language models. Finally, although the lightweight architecture suggests favorable computational efficiency, deployment-oriented characteristics were not formally evaluated in this retrospective study. Therefore, our findings support the potential for future implementation rather than confirmed operational deployment. Future multicenter studies incorporating external validation, prospective workflow evaluation, and implementation-oriented engineering assessments will be essential to determine the feasibility of integrating this framework into routine EMS dispatch practice.</p></sec><sec id="s4-4"><title>Conclusions</title><p>Free-text dispatch narratives contain clinically relevant information that is associated with early risk stratification in patients with suspected cardiopulmonary emergencies. Incorporating these narrative features into a structured modeling framework may provide additional insight beyond conventional variables and support decision-making at the dispatch stage. Further validation and prospective evaluation are needed to determine how such approaches can be integrated into routine prehospital care. These findings support the potential value of incorporating explainable, NLP-derived information into EMS dispatch systems, providing a foundation for future multicenter validation and prospective implementation studies.</p></sec></sec></body><back><ack><p>The authors would like to thank the staff of the Nanning Emergency Medical Center for their support in data collection and for their continued dedication to prehospital emergency care.</p><p>During the preparation of this work, the authors used ChatGPT (OpenAI) to assist with language editing and formatting of the manuscript. After using this tool, the authors reviewed and edited the content as needed and take full responsibility for the content of the published article.</p></ack><notes><sec><title>Funding</title><p>This work was supported by the Guangxi Key Research and Development Program (grant numbers AB24010174 and AB25069064), the Guangxi Health Commission Self-funded Project (grant number Z-A20241212), and the Guangxi Center for Disease Control and Prevention Science and Technology Project (grant number GXJKKJ2026YB010). The funding bodies had no role in the design of the study; collection, analysis, and interpretation of data; or in writing of the manuscript.</p></sec><sec><title>Data Availability</title><p>The datasets generated and analyzed during the current study are not publicly available due to institutional and data protection policies but are available from the corresponding authors on reasonable request and with permission from the Nanning Emergency Medical Center.</p></sec></notes><fn-group><fn fn-type="con"><p>Conceptualization: ZL, LL</p><p>Data curation: CL, SH, JQ, MY, YN</p><p>Formal analysis: ZL, GQ</p><p>Funding acquisition: LL</p><p>Investigation: CL, SH, SZ, ZH</p><p>Methodology: ZL, LS</p><p>Project administration: LL</p><p>Resources: JQ, MY, YN</p><p>Supervision: GQ, LL</p><p>Validation: LS, CL, SH, SZ, ZH</p><p>Writing &#x2013; original draft: ZL</p><p>Writing &#x2013; review &#x0026; editing: LS, CL, GQ, LL</p><p>ZL, LS, and CL contributed equally to this work.</p><p>All authors read and approved the final manuscript.</p></fn><fn fn-type="conflict"><p>None declared.</p></fn></fn-group><glossary><title>Abbreviations</title><def-list><def-item><term id="abb1">ALS</term><def><p>advanced life support</p></def></def-item><def-item><term id="abb2">AUPRC</term><def><p>area under the precision-recall curve</p></def></def-item><def-item><term id="abb3">AUROC</term><def><p>area under the receiver operating characteristic curve</p></def></def-item><def-item><term id="abb4">BERT</term><def><p>bidirectional encoder representations from transformers</p></def></def-item><def-item><term id="abb5">DCA</term><def><p>decision curve analysis</p></def></def-item><def-item><term id="abb6">EMS</term><def><p>emergency medical services</p></def></def-item><def-item><term id="abb7">NLP</term><def><p>natural language processing</p></def></def-item><def-item><term id="abb8">PPV</term><def><p>positive predictive value</p></def></def-item><def-item><term id="abb9">SHAP</term><def><p>Shapley Additive Explanations</p></def></def-item><def-item><term id="abb10">SpO<sub>2</sub></term><def><p>oxygen saturation</p></def></def-item><def-item><term id="abb11">TF-IDF</term><def><p>frequency-inverse document frequency</p></def></def-item><def-item><term id="abb12">TRIPOD</term><def><p>Transparent Reporting of a Multivariable Prediction Model for Individual Prognosis or Diagnosis</p></def></def-item><def-item><term id="abb13">XGBoost</term><def><p>Extreme Gradient Boosting</p></def></def-item></def-list></glossary><ref-list><title>References</title><ref id="ref1"><label>1</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Rao</surname><given-names>SV</given-names> </name><name name-style="western"><surname>O&#x2019;Donoghue</surname><given-names>ML</given-names> </name><name name-style="western"><surname>Ruel</surname><given-names>M</given-names> </name><etal/></person-group><article-title>ACC/AHA/ACEP/NAEMSP/SCAI guideline for the management of patients with acute coronary syndromes: a report of the American College of Cardiology/American Heart Association Joint Committee on Clinical Practice Guidelines</article-title><source>J Am Coll Cardiol</source><year>2025</year><month>06</month><day>10</day><volume>85</volume><issue>22</issue><fpage>2135</fpage><lpage>2237</lpage><pub-id pub-id-type="doi">10.1016/j.jacc.2024.11.009</pub-id><pub-id pub-id-type="medline">40013746</pub-id></nlm-citation></ref><ref id="ref2"><label>2</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Ganter</surname><given-names>J</given-names> </name><name name-style="western"><surname>Busch</surname><given-names>HJ</given-names> </name><name name-style="western"><surname>Henis</surname><given-names>A</given-names> </name><name name-style="western"><surname>Reifferscheid</surname><given-names>F</given-names> </name><name name-style="western"><surname>Braun</surname><given-names>J</given-names> </name><name name-style="western"><surname>Heinrich</surname><given-names>S</given-names> </name></person-group><article-title>Major medical events in patients with acute coronary syndrome during helicopter emergency medical service operations</article-title><source>BMC Emerg Med</source><year>2025</year><month>08</month><day>2</day><volume>25</volume><issue>1</issue><fpage>145</fpage><pub-id pub-id-type="doi">10.1186/s12873-025-01308-7</pub-id><pub-id pub-id-type="medline">40753408</pub-id></nlm-citation></ref><ref id="ref3"><label>3</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Gamberini</surname><given-names>L</given-names> </name><name name-style="western"><surname>Picoco</surname><given-names>C</given-names> </name><name name-style="western"><surname>Del Giudice</surname><given-names>D</given-names> </name><etal/></person-group><article-title>Improving the appropriateness of advanced life support teams&#x2019; dispatch: a before-after study</article-title><source>Prehosp Disaster Med</source><year>2021</year><month>04</month><volume>36</volume><issue>2</issue><fpage>195</fpage><lpage>201</lpage><pub-id pub-id-type="doi">10.1017/S1049023X21000030</pub-id><pub-id pub-id-type="medline">33517934</pub-id></nlm-citation></ref><ref id="ref4"><label>4</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Stamate</surname><given-names>E</given-names> </name><name name-style="western"><surname>Culea-Florescu</surname><given-names>AL</given-names> </name><name name-style="western"><surname>Miron</surname><given-names>M</given-names> </name><etal/></person-group><article-title>AI-based predictive models for cardiogenic shock in STEMI: real-world data for early risk assessment and prognostic insights</article-title><source>J Clin Med</source><year>2025</year><month>05</month><day>25</day><volume>14</volume><issue>11</issue><fpage>3698</fpage><pub-id pub-id-type="doi">10.3390/jcm14113698</pub-id><pub-id pub-id-type="medline">40507459</pub-id></nlm-citation></ref><ref id="ref5"><label>5</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Arrigo</surname><given-names>M</given-names> </name><name name-style="western"><surname>Price</surname><given-names>S</given-names> </name><name name-style="western"><surname>Harjola</surname><given-names>VP</given-names> </name><etal/></person-group><article-title>Diagnosis and treatment of right ventricular failure secondary to acutely increased right ventricular afterload (acute cor pulmonale): a clinical consensus statement of the association for acute cardiovascular care of the European Society of Cardiology</article-title><source>Eur Heart J Acute Cardiovasc Care</source><year>2024</year><month>03</month><day>11</day><volume>13</volume><issue>3</issue><fpage>304</fpage><lpage>312</lpage><pub-id pub-id-type="doi">10.1093/ehjacc/zuad157</pub-id><pub-id pub-id-type="medline">38135288</pub-id></nlm-citation></ref><ref id="ref6"><label>6</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Chong</surname><given-names>B</given-names> </name><name name-style="western"><surname>Jayabaskaran</surname><given-names>J</given-names> </name><name name-style="western"><surname>Jauhari</surname><given-names>SM</given-names> </name><etal/></person-group><article-title>Global burden of cardiovascular diseases: projections from 2025 to 2050</article-title><source>Eur J Prev Cardiol</source><year>2025</year><month>08</month><day>25</day><volume>32</volume><issue>11</issue><fpage>1001</fpage><lpage>1015</lpage><pub-id pub-id-type="doi">10.1093/eurjpc/zwae281</pub-id><pub-id pub-id-type="medline">39270739</pub-id></nlm-citation></ref><ref id="ref7"><label>7</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Khera</surname><given-names>R</given-names> </name><name name-style="western"><surname>Haimovich</surname><given-names>J</given-names> </name><name name-style="western"><surname>Hurley</surname><given-names>NC</given-names> </name><etal/></person-group><article-title>Use of machine learning models to predict death after acute myocardial infarction</article-title><source>JAMA Cardiol</source><year>2021</year><month>06</month><day>1</day><volume>6</volume><issue>6</issue><fpage>633</fpage><lpage>641</lpage><pub-id pub-id-type="doi">10.1001/jamacardio.2021.0122</pub-id><pub-id pub-id-type="medline">33688915</pub-id></nlm-citation></ref><ref id="ref8"><label>8</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Okyere</surname><given-names>D</given-names> </name><name name-style="western"><surname>Nehme</surname><given-names>E</given-names> </name><name name-style="western"><surname>Mahony</surname><given-names>E</given-names> </name><etal/></person-group><article-title>Incidence, diagnoses, and outcomes of pediatric nontraumatic chest pain attended by ambulance</article-title><source>JAMA Netw Open</source><year>2025</year><month>09</month><day>2</day><volume>8</volume><issue>9</issue><fpage>e2533962</fpage><pub-id pub-id-type="doi">10.1001/jamanetworkopen.2025.33962</pub-id><pub-id pub-id-type="medline">41004146</pub-id></nlm-citation></ref><ref id="ref9"><label>9</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Ahmed</surname><given-names>S</given-names> </name><name name-style="western"><surname>Gnesin</surname><given-names>F</given-names> </name><name name-style="western"><surname>Christensen</surname><given-names>HC</given-names> </name><etal/></person-group><article-title>Prehospital management and outcomes of patients calling with chest pain as the main complaint</article-title><source>Int J Emerg Med</source><year>2024</year><month>10</month><day>18</day><volume>17</volume><issue>1</issue><fpage>158</fpage><pub-id pub-id-type="doi">10.1186/s12245-024-00745-8</pub-id><pub-id pub-id-type="medline">39425037</pub-id></nlm-citation></ref><ref id="ref10"><label>10</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Murasaka</surname><given-names>K</given-names> </name><name name-style="western"><surname>Takada</surname><given-names>K</given-names> </name><name name-style="western"><surname>Yamashita</surname><given-names>A</given-names> </name><name name-style="western"><surname>Ushimoto</surname><given-names>T</given-names> </name><name name-style="western"><surname>Wato</surname><given-names>Y</given-names> </name><name name-style="western"><surname>Inaba</surname><given-names>H</given-names> </name></person-group><article-title>Seizure-like activity at the onset of emergency medical service-witnessed out-of-hospital cardiac arrest: an observational study</article-title><source>Resusc Plus</source><year>2021</year><volume>8</volume><fpage>100168</fpage><pub-id pub-id-type="doi">10.1016/j.resplu.2021.100168</pub-id><pub-id pub-id-type="medline">34661179</pub-id></nlm-citation></ref><ref id="ref11"><label>11</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Jaffe</surname><given-names>E</given-names> </name><name name-style="western"><surname>Bitan</surname><given-names>Y</given-names> </name></person-group><article-title>Israeli dispatchers&#x2019; response time to out-of-hospital cardiac arrest emergency calls</article-title><source>Resuscitation</source><year>2022</year><month>09</month><volume>178</volume><fpage>36</fpage><lpage>37</lpage><pub-id pub-id-type="doi">10.1016/j.resuscitation.2022.07.014</pub-id><pub-id pub-id-type="medline">35842189</pub-id></nlm-citation></ref><ref id="ref12"><label>12</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Harris</surname><given-names>M</given-names> </name><name name-style="western"><surname>Crowe</surname><given-names>RP</given-names> </name><name name-style="western"><surname>Anders</surname><given-names>J</given-names> </name><name name-style="western"><surname>D&#x2019;Acunto</surname><given-names>S</given-names> </name><name name-style="western"><surname>Adelgais</surname><given-names>KM</given-names> </name><name name-style="western"><surname>Fishe</surname><given-names>JN</given-names> </name></person-group><article-title>Identification of factors associated with return of spontaneous circulation after pediatric out-of-hospital cardiac arrest using natural language processing</article-title><source>Prehosp Emerg Care</source><year>2023</year><volume>27</volume><issue>5</issue><fpage>687</fpage><lpage>694</lpage><pub-id pub-id-type="doi">10.1080/10903127.2022.2074180</pub-id><pub-id pub-id-type="medline">35510881</pub-id></nlm-citation></ref><ref id="ref13"><label>13</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Jones</surname><given-names>KA</given-names> </name><name name-style="western"><surname>Jani</surname><given-names>KH</given-names> </name><name name-style="western"><surname>Jones</surname><given-names>GW</given-names> </name><etal/></person-group><article-title>Using natural language processing to compare task&#x2010;specific verbal cues in coached versus noncoached cardiac arrest teams during simulated pediatrics resuscitation</article-title><source>AEM Educ Train</source><year>2021</year><month>08</month><volume>5</volume><issue>4</issue><fpage>e10707</fpage><pub-id pub-id-type="doi">10.1002/aet2.10707</pub-id><pub-id pub-id-type="medline">34926971</pub-id></nlm-citation></ref><ref id="ref14"><label>14</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Lorenzoni</surname><given-names>G</given-names> </name><name name-style="western"><surname>Bressan</surname><given-names>S</given-names> </name><name name-style="western"><surname>Lanera</surname><given-names>C</given-names> </name><name name-style="western"><surname>Azzolina</surname><given-names>D</given-names> </name><name name-style="western"><surname>Da Dalt</surname><given-names>L</given-names> </name><name name-style="western"><surname>Gregori</surname><given-names>D</given-names> </name></person-group><article-title>Analysis of unstructured text-based data using machine learning techniques: the case of pediatric emergency department records in Nicaragua</article-title><source>Med Care Res Rev</source><year>2021</year><month>04</month><volume>78</volume><issue>2</issue><fpage>138</fpage><lpage>145</lpage><pub-id pub-id-type="doi">10.1177/1077558719844123</pub-id><pub-id pub-id-type="medline">31030615</pub-id></nlm-citation></ref><ref id="ref15"><label>15</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Siddiqui</surname><given-names>FJ</given-names> </name><name name-style="western"><surname>Fook-Chong</surname><given-names>S</given-names> </name><name name-style="western"><surname>Shahidah</surname><given-names>N</given-names> </name><etal/></person-group><article-title>Dispatcher-assisted cardiopulmonary resuscitation for out-of-hospital cardiac arrest patients: a site-level analysis of the PAROS trial</article-title><source>Resusc Plus</source><year>2026</year><volume>27</volume><fpage>101193</fpage><pub-id pub-id-type="doi">10.1016/j.resplu.2025.101193</pub-id><pub-id pub-id-type="medline">41583935</pub-id></nlm-citation></ref><ref id="ref16"><label>16</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Imbriaco</surname><given-names>G</given-names> </name><name name-style="western"><surname>Galazzi</surname><given-names>A</given-names> </name><name name-style="western"><surname>Semeraro</surname><given-names>F</given-names> </name><name name-style="western"><surname>Ramacciati</surname><given-names>N</given-names> </name></person-group><article-title>Experiences, challenges, and best practices of dispatcher-assisted cardiopulmonary resuscitation: a scoping review</article-title><source>Intern Emerg Med</source><year>2025</year><month>09</month><volume>20</volume><issue>6</issue><fpage>1869</fpage><lpage>1900</lpage><pub-id pub-id-type="doi">10.1007/s11739-025-03991-7</pub-id><pub-id pub-id-type="medline">40789972</pub-id></nlm-citation></ref><ref id="ref17"><label>17</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Kachman</surname><given-names>MM</given-names> </name><name name-style="western"><surname>Brennan</surname><given-names>I</given-names> </name><name name-style="western"><surname>Oskvarek</surname><given-names>JJ</given-names> </name><name name-style="western"><surname>Waseem</surname><given-names>T</given-names> </name><name name-style="western"><surname>Pines</surname><given-names>JM</given-names> </name></person-group><article-title>How artificial intelligence could transform emergency care</article-title><source>Am J Emerg Med</source><year>2024</year><month>07</month><volume>81</volume><fpage>40</fpage><lpage>46</lpage><pub-id pub-id-type="doi">10.1016/j.ajem.2024.04.024</pub-id><pub-id pub-id-type="medline">38663302</pub-id></nlm-citation></ref><ref id="ref18"><label>18</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Raff</surname><given-names>D</given-names> </name><name name-style="western"><surname>Stewart</surname><given-names>K</given-names> </name><name name-style="western"><surname>Yang</surname><given-names>MC</given-names> </name><etal/></person-group><article-title>Improving triage accuracy in prehospital emergency telemedicine: scoping review of machine learning-enhanced approaches</article-title><source>Interact J Med Res</source><year>2024</year><month>09</month><day>11</day><volume>13</volume><fpage>e56729</fpage><pub-id pub-id-type="doi">10.2196/56729</pub-id><pub-id pub-id-type="medline">39259967</pub-id></nlm-citation></ref><ref id="ref19"><label>19</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Rupp</surname><given-names>D</given-names> </name><name name-style="western"><surname>Heuser</surname><given-names>N</given-names> </name><name name-style="western"><surname>Sassen</surname><given-names>MC</given-names> </name><name name-style="western"><surname>Betz</surname><given-names>S</given-names> </name><name name-style="western"><surname>Volberg</surname><given-names>C</given-names> </name><name name-style="western"><surname>Glass</surname><given-names>S</given-names> </name></person-group><article-title>Resuscitation (un-)wanted: does anyone care? A retrospective real data analysis</article-title><source>Resuscitation</source><year>2024</year><month>05</month><volume>198</volume><fpage>110189</fpage><pub-id pub-id-type="doi">10.1016/j.resuscitation.2024.110189</pub-id><pub-id pub-id-type="medline">38522733</pub-id></nlm-citation></ref><ref id="ref20"><label>20</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Chin</surname><given-names>KC</given-names> </name><name name-style="western"><surname>Hsieh</surname><given-names>TC</given-names> </name><name name-style="western"><surname>Chiang</surname><given-names>WC</given-names> </name><etal/></person-group><article-title>Early recognition of a caller&#x2019;s emotion in out-of-hospital cardiac arrest dispatching: an artificial intelligence approach</article-title><source>Resuscitation</source><year>2021</year><month>10</month><volume>167</volume><fpage>144</fpage><lpage>150</lpage><pub-id pub-id-type="doi">10.1016/j.resuscitation.2021.08.032</pub-id><pub-id pub-id-type="medline">34461203</pub-id></nlm-citation></ref><ref id="ref21"><label>21</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Wibring</surname><given-names>K</given-names> </name><name name-style="western"><surname>Lingman</surname><given-names>M</given-names> </name><name name-style="western"><surname>Herlitz</surname><given-names>J</given-names> </name><name name-style="western"><surname>B&#x00E5;ng</surname><given-names>A</given-names> </name></person-group><article-title>The potential of new prediction models for emergency medical dispatch prioritisation of patients with chest pain: a cohort study</article-title><source>Scand J Trauma Resusc Emerg Med</source><year>2022</year><month>05</month><day>8</day><volume>30</volume><issue>1</issue><fpage>34</fpage><pub-id pub-id-type="doi">10.1186/s13049-022-01021-5</pub-id><pub-id pub-id-type="medline">35527302</pub-id></nlm-citation></ref><ref id="ref22"><label>22</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Zhang</surname><given-names>X</given-names> </name><name name-style="western"><surname>Zhao</surname><given-names>H</given-names> </name><name name-style="western"><surname>Zhang</surname><given-names>L</given-names> </name><name name-style="western"><surname>Xiong</surname><given-names>Y</given-names> </name></person-group><article-title>Leveraging multi-text joint prompts in SAM for robust medical image segmentation</article-title><source>IEEE J Biomed Health Inform</source><year>2026</year><month>03</month><volume>30</volume><issue>3</issue><fpage>2512</fpage><lpage>2523</lpage><pub-id pub-id-type="doi">10.1109/JBHI.2025.3607023</pub-id><pub-id pub-id-type="medline">40920522</pub-id></nlm-citation></ref><ref id="ref23"><label>23</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Palatinus</surname><given-names>HN</given-names> </name><name name-style="western"><surname>Johnson</surname><given-names>MA</given-names> </name><name name-style="western"><surname>Wang</surname><given-names>HE</given-names> </name><name name-style="western"><surname>Hoareau</surname><given-names>GL</given-names> </name><name name-style="western"><surname>Youngquist</surname><given-names>ST</given-names> </name></person-group><article-title>Early intramuscular adrenaline administration is associated with improved survival from out-of-hospital cardiac arrest</article-title><source>Resuscitation</source><year>2024</year><month>08</month><volume>201</volume><fpage>110266</fpage><pub-id pub-id-type="doi">10.1016/j.resuscitation.2024.110266</pub-id><pub-id pub-id-type="medline">38857847</pub-id></nlm-citation></ref><ref id="ref24"><label>24</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Krohn</surname><given-names>JN</given-names> </name><name name-style="western"><surname>Barrett</surname><given-names>J</given-names> </name><name name-style="western"><surname>Heeren</surname><given-names>P</given-names> </name><etal/></person-group><article-title>A European paramedic curriculum for geriatric emergency medicine developed via a modified Delphi technique</article-title><source>Scand J Trauma Resusc Emerg Med</source><year>2026</year><month>01</month><day>12</day><volume>34</volume><issue>1</issue><fpage>14</fpage><pub-id pub-id-type="doi">10.1186/s13049-026-01550-3</pub-id><pub-id pub-id-type="medline">41521317</pub-id></nlm-citation></ref><ref id="ref25"><label>25</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Saeed</surname><given-names>N</given-names> </name><name name-style="western"><surname>Naveed</surname><given-names>H</given-names> </name></person-group><article-title>Medical terminology-based computing system: a lightweight post-processing solution for out-of-vocabulary multi-word terms</article-title><source>Front Mol Biosci</source><year>2022</year><volume>9</volume><fpage>928530</fpage><pub-id pub-id-type="doi">10.3389/fmolb.2022.928530</pub-id><pub-id pub-id-type="medline">36032678</pub-id></nlm-citation></ref><ref id="ref26"><label>26</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Charrin</surname><given-names>L</given-names> </name><name name-style="western"><surname>Romain-Scelle</surname><given-names>N</given-names> </name><name name-style="western"><surname>Di-Filippo</surname><given-names>C</given-names> </name><etal/></person-group><article-title>Impact of delayed mobile medical team dispatch for respiratory distress calls: a propensity score matched study from a French emergency communication center</article-title><source>Scand J Trauma Resusc Emerg Med</source><year>2024</year><month>04</month><day>12</day><volume>32</volume><issue>1</issue><fpage>27</fpage><pub-id pub-id-type="doi">10.1186/s13049-024-01201-5</pub-id><pub-id pub-id-type="medline">38609957</pub-id></nlm-citation></ref><ref id="ref27"><label>27</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Lupton</surname><given-names>JR</given-names> </name><name name-style="western"><surname>Jui</surname><given-names>J</given-names> </name><name name-style="western"><surname>Neth</surname><given-names>MR</given-names> </name><name name-style="western"><surname>Sahni</surname><given-names>R</given-names> </name><name name-style="western"><surname>Daya</surname><given-names>MR</given-names> </name><name name-style="western"><surname>Newgard</surname><given-names>CD</given-names> </name></person-group><article-title>Development of a clinical decision rule for the early prediction of shock-refractory out-of-hospital cardiac arrest</article-title><source>Resuscitation</source><year>2022</year><month>12</month><volume>181</volume><fpage>60</fpage><lpage>67</lpage><pub-id pub-id-type="doi">10.1016/j.resuscitation.2022.10.010</pub-id><pub-id pub-id-type="medline">36280216</pub-id></nlm-citation></ref><ref id="ref28"><label>28</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Andresen</surname><given-names>&#x00C5;EL</given-names> </name><name name-style="western"><surname>Kramer-Johansen</surname><given-names>J</given-names> </name><name name-style="western"><surname>Kristiansen</surname><given-names>T</given-names> </name></person-group><article-title>Emergency cricothyroidotomy in difficult airway simulation&#x2014;a national observational study of Air Ambulance crew performance</article-title><source>BMC Emerg Med</source><year>2022</year><month>04</month><day>9</day><volume>22</volume><issue>1</issue><fpage>64</fpage><pub-id pub-id-type="doi">10.1186/s12873-022-00624-6</pub-id><pub-id pub-id-type="medline">35397493</pub-id></nlm-citation></ref><ref id="ref29"><label>29</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Imbriaco</surname><given-names>G</given-names> </name><name name-style="western"><surname>D&#x2019;arrigo</surname><given-names>S</given-names> </name><name name-style="western"><surname>Limonti</surname><given-names>F</given-names> </name><name name-style="western"><surname>Rehn</surname><given-names>M</given-names> </name><name name-style="western"><surname>Cucino</surname><given-names>A</given-names> </name></person-group><article-title>Managing apnoea in trauma victims: the potential for dispatcher-assisted airway handling&#x2014;a narrative review</article-title><source>Scand J Trauma Resusc Emerg Med</source><year>2025</year><month>12</month><day>18</day><volume>34</volume><issue>1</issue><fpage>11</fpage><pub-id pub-id-type="doi">10.1186/s13049-025-01533-w</pub-id><pub-id pub-id-type="medline">41413606</pub-id></nlm-citation></ref><ref id="ref30"><label>30</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Dainty</surname><given-names>KN</given-names> </name><name name-style="western"><surname>Debaty</surname><given-names>G</given-names> </name><name name-style="western"><surname>Waddick</surname><given-names>J</given-names> </name><etal/></person-group><article-title>Interventions to optimize dispatcher-assisted CPR instructions: a scoping review</article-title><source>Resusc Plus</source><year>2024</year><volume>19</volume><fpage>100715</fpage><pub-id pub-id-type="doi">10.1016/j.resplu.2024.100715</pub-id><pub-id pub-id-type="medline">39135732</pub-id></nlm-citation></ref></ref-list><app-group><supplementary-material id="app1"><label>Multimedia Appendix 1</label><p>Methodological details and model specifications.</p><media xlink:href="jmir_v28i1e97672_app1.docx" xlink:title="DOCX File, 18 KB"/></supplementary-material><supplementary-material id="app2"><label>Checklist 1</label><p>TRIPOD checklist.</p><media xlink:href="jmir_v28i1e97672_app2.pdf" xlink:title="PDF File, 113 KB"/></supplementary-material></app-group></back></article>