<?xml version="1.0" encoding="UTF-8"?><!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.0 20040830//EN" "journalpublishing.dtd"><article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" dtd-version="2.0" xml:lang="en" article-type="research-article"><front><journal-meta><journal-id journal-id-type="nlm-ta">JMIR Med Educ</journal-id><journal-id journal-id-type="publisher-id">mededu</journal-id><journal-id journal-id-type="index">20</journal-id><journal-title>JMIR Medical Education</journal-title><abbrev-journal-title>JMIR Med Educ</abbrev-journal-title><issn pub-type="epub">2369-3762</issn><publisher><publisher-name>JMIR Publications</publisher-name><publisher-loc>Toronto, Canada</publisher-loc></publisher></journal-meta><article-meta><article-id pub-id-type="publisher-id">v12i1e96628</article-id><article-id pub-id-type="doi">10.2196/96628</article-id><article-categories><subj-group subj-group-type="heading"><subject>Original Paper</subject></subj-group></article-categories><title-group><article-title>Speech- and Text-Based Emotion Recognition in Anesthesiology Residents During Critical Incident Simulation Training: Exploratory Observational Study</article-title></title-group><contrib-group><contrib contrib-type="author" corresp="yes"><name name-style="western"><surname>Gershov</surname><given-names>Sapir</given-names></name><degrees>PhD</degrees><xref ref-type="aff" rid="aff1">1</xref><xref ref-type="aff" rid="aff2">2</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Bentov</surname><given-names>Itay</given-names></name><degrees>MD, PhD</degrees><xref ref-type="aff" rid="aff3">3</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Mahameed</surname><given-names>Fadi</given-names></name><degrees>MD</degrees><xref ref-type="aff" rid="aff4">4</xref><xref ref-type="aff" rid="aff5">5</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Raz</surname><given-names>Aeyal</given-names></name><degrees>MD, PhD</degrees><xref ref-type="aff" rid="aff4">4</xref><xref ref-type="aff" rid="aff6">6</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Laufer</surname><given-names>Shlomi</given-names></name><degrees>PhD</degrees><xref ref-type="aff" rid="aff2">2</xref><xref ref-type="aff" rid="aff5">5</xref></contrib></contrib-group><aff id="aff1"><institution>Department of Psychiatry, NYU Grossman School of Medicine</institution><addr-line>1 Park Ave., 8th Floor</addr-line><addr-line>New York</addr-line><addr-line>NY</addr-line><country>United States</country></aff><aff id="aff2"><institution>Technion Autonomous Systems Program, Technion &#x2013; Israel Institute of Technology</institution><addr-line>Haifa</addr-line><country>Israel</country></aff><aff id="aff3"><institution>Department of Anesthesiology and Pain Medicine, Harborview Medical Center</institution><addr-line>Seattle</addr-line><addr-line>WA</addr-line><country>United States</country></aff><aff id="aff4"><institution>Department of Anesthesiology, Rambam Health Care Campus</institution><addr-line>Haifa</addr-line><country>Israel</country></aff><aff id="aff5"><institution>Faculty of Data and Decision Sciences, Technion &#x2013; Israel Institute of Technology</institution><addr-line>Haifa</addr-line><country>Israel</country></aff><aff id="aff6"><institution>Rappaport Faculty of Medicine, Technion &#x2013; Israel Institute of Technology</institution><addr-line>Haifa</addr-line><country>Israel</country></aff><contrib-group><contrib contrib-type="editor"><name name-style="western"><surname>Car</surname><given-names>Lorainne Tudor</given-names></name></contrib></contrib-group><contrib-group><contrib contrib-type="reviewer"><name name-style="western"><surname>Turky</surname><given-names>Ayad</given-names></name></contrib><contrib contrib-type="reviewer"><name name-style="western"><surname>Winterton</surname><given-names>Dario</given-names></name></contrib></contrib-group><author-notes><corresp>Correspondence to Sapir Gershov, PhD, Department of Psychiatry, NYU Grossman School of Medicine, 1 Park Ave., 8th Floor, New York, NY, 10016, United States, 1 646-934-4738; <email>Sapir.Gershov@nyulangone.org</email></corresp></author-notes><pub-date pub-type="collection"><year>2026</year></pub-date><pub-date pub-type="epub"><day>21</day><month>8</month><year>2026</year></pub-date><volume>12</volume><elocation-id>e96628</elocation-id><history><date date-type="received"><day>30</day><month>03</month><year>2026</year></date><date date-type="rev-recd"><day>05</day><month>06</month><year>2026</year></date><date date-type="accepted"><day>22</day><month>07</month><year>2026</year></date></history><copyright-statement>&#x00A9; Sapir Gershov, Itay Bentov, Fadi Mahameed, Aeyal Raz, Shlomi Laufer. Originally published in JMIR Medical Education (<ext-link ext-link-type="uri" xlink:href="https://mededu.jmir.org">https://mededu.jmir.org</ext-link>), 21.8.2026. </copyright-statement><copyright-year>2026</copyright-year><license license-type="open-access" xlink:href="https://creativecommons.org/licenses/by/4.0/"><p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (<ext-link ext-link-type="uri" xlink:href="https://creativecommons.org/licenses/by/4.0/">https://creativecommons.org/licenses/by/4.0/</ext-link>), which permits unrestricted use, distribution, and reproduction in any medium, provided the original work, first published in JMIR Medical Education, is properly cited. The complete bibliographic information, a link to the original publication on <ext-link ext-link-type="uri" xlink:href="https://mededu.jmir.org/">https://mededu.jmir.org/</ext-link>, as well as this copyright and license information must be included.</p></license><self-uri xlink:type="simple" xlink:href="https://mededu.jmir.org/2026/1/e96628"/><abstract><sec><title>Background</title><p>Communication during crisis management involves not only the exchange of clinical information but also the expression and regulation of emotional tone and sentiment, which may reflect clinician performance during a critical incident. In simulation-based medical education, these emotional dynamics are rarely measured objectively, limiting the ability to capture aspects of performance relevant to competency assessment.</p></sec><sec><title>Objective</title><p>This study aimed to examine whether automatically extracted speech- and text-based emotion signals varied across predefined phases of critical incident simulation training and whether these signals were associated with expert-rated performance among anesthesiology residents.</p></sec><sec sec-type="methods"><title>Methods</title><p>In this exploratory observational study, we analyzed 164 simulated crisis scenarios performed by 90 anesthesiology residents from 17 hospitals during a national board preparation workshop. Each high-fidelity simulation comprised 4 predefined phases: initial condition, extreme deterioration, advanced cardiovascular life support (ACLS) management, and recovery. Resident speech was isolated and analyzed using a speech-based emotion recognition (SER) model that estimated 8 emotion categories per utterance. Speech was also automatically transcribed and analyzed using text-based emotion recognition (TER), which estimated 9 emotion categories. Expert anesthesiologists rated performance on a 1-5 scale. ANOVA tested associations between emotion categories and performance, as well as variation across simulation phases. Exploratory mixed-effects sensitivity analyses accounted for repeated simulations and scenario-level clustering.</p></sec><sec sec-type="results"><title>Results</title><p>After 7 recordings were excluded because of technical issues, 65 of 164 (39.6%) simulations were classified as poor performance (score &#x2264;2), 72 (43.9%) as intermediate, and 27 (16.5%) as excellent (score=5). SER categories differed across performance groups (<italic>F</italic><sub>2,161</sub>=4.61; <italic>P</italic>=.02): poor performance showed more fear and surprise, whereas excellent performance showed more calm, happiness, and neutrality. In mixed-effects analyses, performance group remained associated with happy, angry, fearful, surprise, calm, and neutral speech-emotion proportions after false discovery rate correction. The largest excellent vs poor differences were observed for neutral (+4.24 percentage points), happy (+4.11), and surprise (&#x2212;3.17) speech affect. TER showed a directionally similar but nonsignificant omnibus association (<italic>F</italic><sub>2,161</sub>=2.57; <italic>P</italic>=.08); mixed-effects analyses identified associations with happy and sad text-emotion proportions after correction. During extreme deterioration and ACLS management, fear and surprise together accounted for 64.1% and 64.6% of classified speech emotions, respectively, among poor performances. During ACLS management, neutral and calm emotions accounted for 49.1% and 37.7%, respectively, among excellent performances.</p></sec><sec sec-type="conclusions"><title>Conclusions</title><p>Automated speech- and text-based emotion recognition captured phase-dependent affective communication patterns during critical incident simulations. Speech-derived features were more consistently associated with expert-rated performance, whereas text-derived findings were more exploratory and emotion-specific. These markers may complement simulation debriefing and formative feedback, but validation using multirater outcomes, preregistered adjusted models, and larger multicenter datasets is needed before they inform competency assessment.</p></sec></abstract><kwd-group><kwd>simulation-based medical education</kwd><kwd>critical incident simulation</kwd><kwd>affective communication</kwd><kwd>speech-based emotion recognition</kwd><kwd>text-based emotion recognition</kwd><kwd>anesthesiology residents</kwd><kwd>formative feedback</kwd></kwd-group></article-meta></front><body><sec id="s1" sec-type="intro"><title>Introduction</title><sec id="s1-1"><title>Affective Communication in Critical Care</title><p>Effective communication is central to safe care in high acuity clinical environments, such as in the operating room (OR) and intensive care unit (ICU). During critical events, clinicians must rapidly interpret evolving information, issue clear instructions, coordinate team responses, and maintain situational awareness under pressure. These exchanges are not purely informational. They also convey affective signals, including urgency, confidence, calmness, hesitation, and distress, which may influence team coordination, decision-making, and clinical execution [<xref ref-type="bibr" rid="ref1">1</xref>,<xref ref-type="bibr" rid="ref2">2</xref>]. Prior work has shown that both communication quality [<xref ref-type="bibr" rid="ref3">3</xref>] and emotional processes [<xref ref-type="bibr" rid="ref4">4</xref>] are relevant to performance in acute care settings.</p><p>In anesthesiology, these issues are especially complex because anesthesiologists operate under sustained cognitive load. Stress and arousal are expected components of medical emergencies, but when poorly regulated, they may impair attention, communication, and decision-making [<xref ref-type="bibr" rid="ref5">5</xref>-<xref ref-type="bibr" rid="ref7">7</xref>]. This is particularly challenging in perioperative crises, which often elicit strong emotional responses and a need for structured debriefing or recovery time [<xref ref-type="bibr" rid="ref6">6</xref>].</p><p>Emotion and communication are intertwined rather than interchangeable. Well-regulated affect can facilitate clarity, closed-loop confirmation, and team coordination; dysregulated affect may coincide with interruptions, hesitations, or missed confirmations [<xref ref-type="bibr" rid="ref7">7</xref>-<xref ref-type="bibr" rid="ref9">9</xref>]. The key question is therefore not whether emotion is present, but how affective patterns unfold during critical situations and whether they align with effective communication behaviors and clinical performance.</p></sec><sec id="s1-2"><title>Simulation-Based Assessment and the Need for Objective Communication Markers</title><p>Simulation-based crisis incident management training has been extensively used by anesthesiologists for over half a century [<xref ref-type="bibr" rid="ref10">10</xref>,<xref ref-type="bibr" rid="ref11">11</xref>], as it is effective for skill acquisition [<xref ref-type="bibr" rid="ref12">12</xref>] and knowledge retention [<xref ref-type="bibr" rid="ref13">13</xref>]. Critical incident simulation training is commonly used in high acuity training [<xref ref-type="bibr" rid="ref14">14</xref>-<xref ref-type="bibr" rid="ref16">16</xref>] and can also provide an opportunity to improve patient safety by exploring providers&#x2019; responses to stressful situations [<xref ref-type="bibr" rid="ref17">17</xref>,<xref ref-type="bibr" rid="ref18">18</xref>]. Simulation-based assessment (SBA) is now a widely accepted method for evaluating clinical competency and is incorporated into board certification examinations by professional societies, including the Israel Society of Anesthesiologists [<xref ref-type="bibr" rid="ref18">18</xref>]. However, current SBA approaches emphasize observable behavior scored by expert evaluators, which can be vulnerable to bias [<xref ref-type="bibr" rid="ref19">19</xref>,<xref ref-type="bibr" rid="ref20">20</xref>], while the emotional dimensions of communication remain largely unexamined. This creates a need for complementary, objective measures that can enrich traditional assessments without replacing expert evaluation.</p></sec><sec id="s1-3"><title>Speech- and Text-Based Emotion Recognition</title><p>Recent advances in natural language processing (NLP) and speech analysis have created new opportunities to quantify emotional content in spoken clinical communication. Text-based emotion recognition (TER) [<xref ref-type="bibr" rid="ref21">21</xref>,<xref ref-type="bibr" rid="ref22">22</xref>] can infer emotional meaning from transcribed language, whereas speech-based emotion recognition (SER) [<xref ref-type="bibr" rid="ref23">23</xref>,<xref ref-type="bibr" rid="ref24">24</xref>] captures paralinguistic features such as tone, intensity, and vocal affect. The integration of these techniques provides a multidimensional view of communication, extending analysis beyond spoken words to encompass tone, urgency, and affect [<xref ref-type="bibr" rid="ref25">25</xref>,<xref ref-type="bibr" rid="ref26">26</xref>]. At the same time, applying TER and SER in healthcare poses challenges due to the specialized vocabulary and context-dependent nature of clinical communication [<xref ref-type="bibr" rid="ref27">27</xref>]. Recent work has also emphasized the broader movement toward technology-enhanced assessment of nontechnical skills in high acuity procedural environments. For example, digital approaches to assessing nontechnical skills in the operating room have been proposed to provide more objective, scalable, and timely feedback, although validation and interpretability remain major challenges [<xref ref-type="bibr" rid="ref7">7</xref>]. Similarly, recent reviews of SER [<xref ref-type="bibr" rid="ref25">25</xref>] and textual emotion detection [<xref ref-type="bibr" rid="ref28">28</xref>] in healthcare highlight the promise of automated affective analysis while emphasizing the need for careful domain adaptation, transparent evaluation, and validation in clinically meaningful settings. These considerations are particularly relevant in simulation-based anesthesiology training, where communication is brief, protocol-driven, emotionally charged, and shaped by clinical context. Accordingly, SER and TER should be viewed as potential complementary markers of affective communication rather than direct measures of clinical competency.</p><p>The contribution of the present study is the integration of existing speech- and text-based emotion recognition methods within a phase-structured, high-fidelity anesthesiology simulation setting. Specifically, we aligned resident-specific affective communication signals with predefined crisis phases and expert-rated simulation performance. This design allowed us to examine whether automatically derived emotion markers provide complementary information about affective communication during simulation-based training.</p></sec><sec id="s1-4"><title>Study Aim and Hypotheses</title><p>In this study, we conceptualized affective communication as a behavioral signal that may reflect several overlapping processes relevant to crisis management, including cognitive load, stress regulation, communication clarity, situational awareness, and individual communication style. These processes are related to, but not equivalent to, clinical competency. Consistent with this conceptual framing, emotion recognition outputs were treated as exploratory markers of affective communication rather than direct measures of clinical competency. Specifically, our aim was to automatically extract emotion signals from speech and transcribe text associated with expert-rated performance during simulation-based critical incident training among anesthesiology residents. Using audio recordings from simulated ICU scenarios, we analyzed SER and text-based emotion recognition outputs across predefined phases of each simulation. We hypothesized that higher-performing residents would display a greater level of regulated affective patterns, reflected in calmer or more positive emotional signals, whereas lower-performing residents would display a greater level of negative and high arousal emotional patterns. We also examined whether these emotional patterns differed across simulation phases, particularly during the most demanding portions of the scenario.</p></sec></sec><sec id="s2" sec-type="methods"><title>Methods</title><sec id="s2-1"><title>Study Design</title><p>This exploratory observational study examined whether automatically extracted emotion signals from residents&#x2019; speech and transcribed language were associated with expert-rated performance during simulation-based critical incident training in anesthesiology.</p></sec><sec id="s2-2"><title>Setting and Participants</title><p>Anesthesiology residency in Israel, as in other countries [<xref ref-type="bibr" rid="ref29">29</xref>], includes staged qualifying examinations throughout training. In the first stage, residents must pass a national written board exam before advancing to the senior phase of residency. To become an attending anesthesiologist, residents must demonstrate knowledge and competency through oral board examinations and simulated scenarios. The Israel Society of Anesthesiologists organizes a biannual, 2-day workshop to prepare senior anesthesiology residents for their final board examinations. This workshop includes hands-on practice with high-fidelity simulations followed by expert feedback from board-certified anesthesiologists.</p><p>For this study, we analyzed all the residents who participated in the workshops conducted in 2022 and 2023. During enrollment, residents completed a demographic questionnaire that included their level of clinical experience, measured in years of residency, and the number of prior SBAs they had completed.</p></sec><sec id="s2-3"><title>Ethical Considerations</title><p>The Rambam Health Care Campus Institutional Review Board approved this study (0482&#x2010;20-RMB). All participants provided written informed consent before participation.</p></sec><sec id="s2-4"><title>Simulation Procedures and Performance Ratings</title><p>An experienced anesthesiologist and a medical simulation specialist developed 7 clinical simulation scenarios in accordance with advanced cardiovascular life support (ACLS) and advanced trauma life support (ATLS) principles: (1) severe anaphylaxis reaction, (2) postoperative severe bradycardia, (3) postoperative opioid overdose, (4) postoperative hypoglycemia, (5) postoperative residual paralysis, (6) postoperative supraventricular tachycardia, and (7) patient with head trauma.</p><p>Each resident was randomly assigned to 2 simulations. During each simulation, 2 members of the research team played the roles of a nurse and a medical intern, respectively. A board-certified anesthesiologist evaluated the resident&#x2019;s clinical performance using a task-specific checklist. The simulation setup included a Laerdal full-body manikin to create a high-fidelity, realistic environment (see <xref ref-type="fig" rid="figure1">Figure 1</xref>). The simulation area was recorded from multiple angles, and to ensure clear audio recordings, all participants (residents and team members) were equipped with wireless lavalier microphones. The data was recorded and synchronized using StreamPix digital video recording software from NorPix Inc.</p><p>To better capture participants&#x2019; expressed emotions, we divided the simulation timeline into 4 chronologically distinct phases:</p><list list-type="order"><list-item><p>Initial condition: Participants understand the patient&#x2019;s status by examining the patient and communicating with the nurse.</p></list-item><list-item><p>Extreme deterioration: The patient&#x2019;s condition has deteriorated significantly (eg, dropping blood pressure, rising heart rate). Participants must manage this crisis and respond appropriately.</p></list-item><list-item><p>ACLS management: Participants must execute procedures quickly and accurately. Additionally, participants must demonstrate confidence and familiarity with the protocols. In the case of simulation 7 (patient with a head trauma), residents performed and managed ATLS.</p></list-item><list-item><p>Recovery: Participants update the team on the patient&#x2019;s status and coordinate postcrisis care.</p></list-item></list><p>This phase-based structure enabled us to examine whether emotional expression varied across different levels of clinical demand within the scenario. Phases were controlled by the simulator operator, each bounded by start and end timestamps and observable clinical cues. These included changes displayed on the patient monitor, such as alterations in vital signs and alarm activations, and, in some scenarios, physical indicators on the manikin itself (eg, breathing sounds, airway obstruction, chest movement, pupil size, etc).</p><p>At the end of each simulation, the anesthesiologist evaluator provided an overall performance rating on a Likert scale of 1 to 5, with 1 indicating poor performance and 5 indicating excellent performance. Based on these scores, we divided our sample into the following groups: poor performance (evaluation scores &#x2264;2), intermediate performance (evaluation scores 3&#x2010;4), and excellent performance (evaluation scores=5).</p><fig position="float" id="figure1"><label>Figure 1.</label><caption><p>Simulation environment setup. (A) General overview of the simulation area. The faces have been manually blurred; (B) patient monitor display; (C) illustration of the lavalier microphone audio levels.</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="mededu_v12i1e96628_fig01.png"/></fig></sec><sec id="s2-5"><title>Audio Collection and Preprocessing</title><p>To minimize the effect of mixed audio recordings, participants (resident, nurse, and intern) were recorded on separate audio channels. However, participants&#x2019; speech overlap and background noise (eg, patient monitor alarms) disrupted speech coherence, prompting us to integrate a speech-cleaning framework called end-to-end neural diarization with speaker subspace (EEND-SS) [<xref ref-type="bibr" rid="ref30">30</xref>]. This framework performs several speech processing procedures (speaker diarization, speech separation, and speaker counting) and produces 2 types of audio files from each raw audio channel, (1) single-speaker files for each participant and (2) background noise files devoid of any traces of speech. These procedures eliminate disturbances that could have affected speech analysis. Additionally, the EEND-SS framework generates segments of &#x201C;silence&#x201D; areas in all processed audio files; silence thresholds are fine-tuned to optimize the results (see <xref ref-type="fig" rid="figure2">Figure 2</xref>A).</p><fig position="float" id="figure2"><label>Figure 2.</label><caption><p>Illustration of the text-based emotion recognition (TER) and speech-based emotion recognition (SER) pipeline. (A) Preprocessing of the raw audio channels via ESPnet framework to extract intervals of the resident verbal expressions; (B) based on the extracted resident speech intervals, we executed SER for every sentence; (C) based on the extracted resident speech intervals, we transcribed each sentence and executed TER.</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="mededu_v12i1e96628_fig02.png"/></fig></sec><sec id="s2-6"><title>Speech-Based Emotion Recognition</title><p>Each EEND-SS single-speaker file is labeled according to the audio source (ie, anesthesia resident, nurse, and medical intern). The following stages apply only to the anesthesia resident&#x2019;s single-speaker audio file. First, feature extraction was performed using the SpeechBrain toolkit [<xref ref-type="bibr" rid="ref31">31</xref>]. Specifically, we applied the wav2vec 2.0 model [<xref ref-type="bibr" rid="ref32">32</xref>], which was pretrained via self supervision. Then, we utilized Pepino et al&#x2019;s [<xref ref-type="bibr" rid="ref33">33</xref>] model, which generates emotion category probabilities for 8 classes: happy, sad, angry, fearful, surprised, disgusted, calm, and neutral (see <xref ref-type="fig" rid="figure2">Figure 2</xref>B). During inference, the predicted emotion category for each segment was defined as the class with the highest probability. Segment-level predictions were then aggregated within each simulation, performance group, and simulation phase to estimate the relative frequency of each emotion category. These aggregated proportions were used in the statistical analyses.</p></sec><sec id="s2-7"><title>Text-Based Emotion Recognition</title><p>The processed anesthesia resident audio file was transcribed to Hebrew using a state-of-the-art, large-scale, weakly supervised speech recognition model, Whisper [<xref ref-type="bibr" rid="ref34">34</xref>]. In addition, since Whisper only provides timestamps for each speech utterance, we used the work of Bain et al [<xref ref-type="bibr" rid="ref35">35</xref>] to generate a word-level timestamped transcription. Afterward, we performed word- and sentence-level TER using 2 language models trained specifically on Hebrew text: AlephBERT [<xref ref-type="bibr" rid="ref36">36</xref>] and HebEMO [<xref ref-type="bibr" rid="ref37">37</xref>] (see <xref ref-type="fig" rid="figure2">Figure 2</xref>C). As with the speech-based model, the dominant text-based emotion category was defined as the class with the highest model score for a given utterance. To improve compatibility with clinical terminology, we applied the medical domain adaptation procedure described in our previous work [<xref ref-type="bibr" rid="ref38">38</xref>], which involved adapting the text-processing pipeline to medical language. We did not fine-tune the emotion recognition models on the current performance labels, thereby avoiding leakage between performance ratings and emotion classification outputs.</p></sec><sec id="s2-8"><title>Implementation Details</title><p>All analyses were conducted on a computing cluster with 2 NVIDIA RTX A6000 48 GB graphics processing units (GPUs). The software environment consisted of Ubuntu 20.04 LTS, Python (version 3.9), and PyTorch (version 2.00). The EEND-SS framework used 6 transformer encoders and Conv-TasNet [<xref ref-type="bibr" rid="ref39">39</xref>], which incorporates 8 temporal convolutional network blocks. The wav2vec 2.0 model had a feature encoder with 7 convolutional neural network layers and a transformer with 12 encoder layers. Both AlephBERT and HebEMO were used with their default parameters.</p></sec><sec id="s2-9"><title>Statistical Analysis</title><p>To evaluate whether the distribution of emotions differed across performance groups, we first used one-way ANOVA as an exploratory omnibus test. All statistical tests were 2-sided. Statistical significance was defined as <italic>P</italic>&#x003C;.05, and exact <italic>P</italic> values are reported unless <italic>P</italic>&#x003C;.001. Because residents could contribute more than 1 simulation and simulations differed by scenario, we then conducted exploratory mixed-effects sensitivity analyses. Emotion proportions were modeled separately for each emotion category. The performance group was included as the primary fixed effect, with random intercepts for resident and simulation scenario. Residency year and sex were included as covariates when available. Global performance-group effects were evaluated using likelihood ratio tests comparing models with and without performance group. False discovery rate (FDR) correction was applied across emotion-level tests within each modality. These sensitivity analyses were intended to assess whether the observed associations persisted after accounting for repeated observations and scenario-level clustering.</p><p>We also performed phase-stratified analyses to examine whether the relationship between emotional expression and performance varied across phases of the simulation scenario. All statistical analyses were executed using SciPy (version 1.16.1) and statsmodels (version 0.14.6).</p></sec></sec><sec id="s3" sec-type="results"><title>Results</title><sec id="s3-1"><title>Sample Characteristics</title><p>We recruited 90 senior anesthesiology residents from 17 different hospitals, representing a diverse range of clinical backgrounds (see <xref ref-type="table" rid="table1">Table 1</xref>). On average, each resident participated in 2 different SBAs, yielding 171 recordings overall. Seven simulation recordings were excluded due to technical issues, including incomplete audio or video capture and severe audio artifacts, resulting in the final analytic sample of 164 simulations (see <xref ref-type="table" rid="table2">Table 2</xref>).</p><p>Based on expert assessors&#x2019; ratings, 65 out of 164 simulations (39.6%) were classified as poor performance (scores &#x2264;2), 72 out of 164 (43.9%) as intermediate performance (scores 3&#x2010;4), and 27 out of 164 (16.5%) as excellent performance (score=5). The final dataset included recordings from all 7 simulation scenarios, with the largest proportions contributed by postoperative residual paralysis (42 out of 164 simulations, 25.6%) and postoperative severe hypoglycemia (36 out of 164 simulations, 22.0%) scenarios. It is worth noting that although participants communicated in Hebrew (with some English terms), the EEND-SS framework successfully processed our participants&#x2019; speech.</p><table-wrap id="t1" position="float"><label>Table 1.</label><caption><p>Demographic questionnaire summary.</p></caption><table id="table1" frame="hsides" rules="groups"><thead><tr><td align="left" valign="bottom">Characteristic</td><td align="left" valign="bottom">Residents (N=90), n (%)</td><td align="left" valign="bottom">Years of residency, mean (SD)</td></tr></thead><tbody><tr><td align="left" valign="top" colspan="3">Sex</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Male</td><td align="left" valign="top">58 (64)</td><td align="left" valign="top">5.48 (1.32)</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Female</td><td align="left" valign="top">32 (36)</td><td align="left" valign="top">5.06 (3.51)</td></tr><tr><td align="left" valign="top" colspan="3">Residency year</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>4</td><td align="left" valign="top">26 (29)</td><td align="left" valign="top">&#x2014;<sup><xref ref-type="table-fn" rid="table1fn1">a</xref></sup></td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>5</td><td align="left" valign="top">32 (36)</td><td align="left" valign="top">&#x2014;</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>6</td><td align="left" valign="top">19 (21)</td><td align="left" valign="top">&#x2014;</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>&#x2265;7</td><td align="left" valign="top">13 (14)</td><td align="left" valign="top">&#x2014;</td></tr></tbody></table><table-wrap-foot><fn id="table1fn1"><p><sup>a</sup>Not applicable.</p></fn></table-wrap-foot></table-wrap><table-wrap id="t2" position="float"><label>Table 2.</label><caption><p>Simulation dataset summary.</p></caption><table id="table2" frame="hsides" rules="groups"><thead><tr><td align="left" valign="bottom">Simulation type</td><td align="left" valign="bottom">Simulations (N=164), n (%)</td><td align="left" valign="bottom" colspan="5">Performance score, n</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom"/><td align="left" valign="bottom">1</td><td align="left" valign="bottom">2</td><td align="left" valign="bottom">3</td><td align="left" valign="bottom">4</td><td align="left" valign="bottom">5</td></tr></thead><tbody><tr><td align="left" valign="top">Patient with a severe anaphylaxis reaction</td><td align="left" valign="top">24 (14.6)</td><td align="left" valign="top">2</td><td align="left" valign="top">5</td><td align="left" valign="top">7</td><td align="left" valign="top">6</td><td align="left" valign="top">4</td></tr><tr><td align="left" valign="top">Postoperative patient with severe bradycardia</td><td align="left" valign="top">26 (15.9)</td><td align="left" valign="top">5</td><td align="left" valign="top">4</td><td align="left" valign="top">4</td><td align="left" valign="top">10</td><td align="left" valign="top">3</td></tr><tr><td align="left" valign="top">Postoperative patient with opioid overdose</td><td align="left" valign="top">12 (7.3)</td><td align="left" valign="top">1</td><td align="left" valign="top">2</td><td align="left" valign="top">2</td><td align="left" valign="top">4</td><td align="left" valign="top">3</td></tr><tr><td align="left" valign="top">Postoperative patient with severe hypoglycemia</td><td align="left" valign="top">36 (22.0)</td><td align="left" valign="top">7</td><td align="left" valign="top">7</td><td align="left" valign="top">6</td><td align="left" valign="top">8</td><td align="left" valign="top">8</td></tr><tr><td align="left" valign="top">Postoperative patient with residual paralysis</td><td align="left" valign="top">42 (25.6)</td><td align="left" valign="top">7</td><td align="left" valign="top">12</td><td align="left" valign="top">9</td><td align="left" valign="top">7</td><td align="left" valign="top">7</td></tr><tr><td align="left" valign="top">Postoperative patient with severe supraventricular tachycardia</td><td align="left" valign="top">14 (8.5)</td><td align="left" valign="top">5</td><td align="left" valign="top">2</td><td align="left" valign="top">4</td><td align="left" valign="top">2</td><td align="left" valign="top">1</td></tr><tr><td align="left" valign="top">Patient with head trauma</td><td align="left" valign="top">10 (6.1)</td><td align="left" valign="top">4</td><td align="left" valign="top">2</td><td align="left" valign="top">2</td><td align="left" valign="top">1</td><td align="left" valign="top">1</td></tr></tbody></table></table-wrap></sec><sec id="s3-2"><title>Affective Communication Across Performance Groups</title><p>Across performance groups, SER emotion proportions differed significantly (<italic>F</italic><sub>2,161</sub>=4.61; <italic>P</italic>=.02; <xref ref-type="fig" rid="figure3">Figure 3</xref>). Poor performances (scores &#x2264;2) showed higher levels of fear and surprise, whereas excellent performances (score=5) showed higher levels of calm and happiness. TER emotion proportion showed a directionally similar but nonsignificant omnibus effect (<italic>F</italic><sub>2,161</sub>=2.57; <italic>P</italic>=.08; <xref ref-type="fig" rid="figure4">Figure 4</xref>), with poorer performance skewing toward more negative emotions and excellent performance toward more positive and neutral emotions. Cross-modal correspondence is evident: groups with higher negative emotions in SER also tended to have lower positive emotions in TER.</p><p>Across both modalities, the poor and excellent performance groups differed significantly for nearly all emotions.</p><fig position="float" id="figure3"><label>Figure 3.</label><caption><p>Speech-based emotion recognition (SER): emotion classification of participants&#x2019; speech and transcription according to their performance score group. The comparison evaluation is conducted between the poor and excellent groups.</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="mededu_v12i1e96628_fig03.png"/></fig><fig position="float" id="figure4"><label>Figure 4.</label><caption><p>Text-based emotion recognition (TER): emotion classification of participants&#x2019; speech and transcription according to their performance group. The comparison evaluation is conducted between the poor and excellent groups.</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="mededu_v12i1e96628_fig04.png"/></fig></sec><sec id="s3-3"><title>Sensitivity Analysis Using Mixed-Effects Models</title><p>Mixed-effects model results for speech- and text-based emotion proportions across performance groups are presented in <xref ref-type="table" rid="table3">Table 3</xref>. For SER, performance group remained significantly associated with happy, angry, fearful, surprise, calm, and neutral speech-emotion proportions after FDR correction. Sad and disgust were not significantly associated with performance group after correction. The largest excellent vs poor differences were observed for neutral speech affect, which was 4.24 percentage points higher in excellent performances; happy speech affect, which was 4.11 percentage points higher in excellent performances; and surprise speech affect, which was 3.17 percentage points lower in excellent performances.</p><table-wrap id="t3" position="float"><label>Table 3.</label><caption><p>Mixed-effects sensitivity analysis of speech- and text-based emotion proportions across performance groups. Models were estimated separately for each emotion category, and emotion proportions were logit-transformed and modeled as the outcome, with performance group as the primary fixed effect. All models included random intercepts for resident and scenario and were adjusted for residency year and sex. Global performance group effects were evaluated using likelihood ratio tests comparing models with and without performance group. False discovery rate correction was applied across emotion-level tests within each modality.</p></caption><table id="table3" frame="hsides" rules="groups"><thead><tr><td align="left" valign="bottom">Modality and emotion</td><td align="left" valign="bottom">LR<sup><xref ref-type="table-fn" rid="table3fn1">a</xref></sup> &#x03C7;&#x00B2;(2)</td><td align="left" valign="bottom"><italic>P</italic> value</td><td align="left" valign="bottom">FDR<sup><xref ref-type="table-fn" rid="table3fn2">b</xref></sup>&#x2013;adjusted <italic>P</italic></td><td align="left" valign="bottom">Excellent-poor difference</td></tr></thead><tbody><tr><td align="left" valign="top" colspan="5">SER<sup><xref ref-type="table-fn" rid="table3fn3">c</xref></sup></td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Happy</td><td align="left" valign="top">22.03</td><td align="left" valign="top">&#x003C;.001</td><td align="left" valign="top">&#x003C;.001</td><td align="left" valign="top">4.11</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Sad</td><td align="left" valign="top">2.19</td><td align="left" valign="top">.34</td><td align="left" valign="top">.38</td><td align="left" valign="top">&#x2212;0.83</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Angry</td><td align="left" valign="top">11.32</td><td align="left" valign="top">.003</td><td align="left" valign="top">.005</td><td align="left" valign="top">&#x2212;1.66</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Fearful</td><td align="left" valign="top">19.76</td><td align="left" valign="top">&#x003C;.001</td><td align="left" valign="top">&#x003C;.001</td><td align="left" valign="top">&#x2212;2.31</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Surprise</td><td align="left" valign="top">21.88</td><td align="left" valign="top">&#x003C;.001</td><td align="left" valign="top">&#x003C;.001</td><td align="left" valign="top">&#x2212;3.17</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Disgust</td><td align="left" valign="top">0.67</td><td align="left" valign="top">.72</td><td align="left" valign="top">.72</td><td align="left" valign="top">0.00</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Calm</td><td align="left" valign="top">15.53</td><td align="left" valign="top">&#x003C;.001</td><td align="left" valign="top">&#x003C;.001</td><td align="left" valign="top">2.21</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Neutral</td><td align="left" valign="top">25.71</td><td align="left" valign="top">&#x003C;.001</td><td align="left" valign="top">&#x003C;.001</td><td align="left" valign="top">4.24</td></tr><tr><td align="left" valign="top" colspan="5">TER<sup><xref ref-type="table-fn" rid="table3fn4">d</xref></sup></td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Happy</td><td align="left" valign="top">23.96</td><td align="left" valign="top">&#x003C;.001</td><td align="left" valign="top">&#x003C;.001</td><td align="left" valign="top">8.66</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Sad</td><td align="left" valign="top">14.39</td><td align="left" valign="top">&#x003C;.001</td><td align="left" valign="top">.003</td><td align="left" valign="top">&#x2212;5.33</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Angry</td><td align="left" valign="top">1.82</td><td align="left" valign="top">.40</td><td align="left" valign="top">.52</td><td align="left" valign="top">&#x2212;0.89</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Fearful</td><td align="left" valign="top">6.52</td><td align="left" valign="top">.04</td><td align="left" valign="top">.09</td><td align="left" valign="top">&#x2212;2.91</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Surprise</td><td align="left" valign="top">7.71</td><td align="left" valign="top">.02</td><td align="left" valign="top">.06</td><td align="left" valign="top">&#x2212;4.13</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Disgust</td><td align="left" valign="top">1.15</td><td align="left" valign="top">.56</td><td align="left" valign="top">.63</td><td align="left" valign="top">&#x2212;0.15</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Trust</td><td align="left" valign="top">3.67</td><td align="left" valign="top">.16</td><td align="left" valign="top">.24</td><td align="left" valign="top">2.06</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Expectation</td><td align="left" valign="top">0.88</td><td align="left" valign="top">.64</td><td align="left" valign="top">.64</td><td align="left" valign="top">1.15</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Sentiment</td><td align="left" valign="top">4.27</td><td align="left" valign="top">.12</td><td align="left" valign="top">.21</td><td align="left" valign="top">1.58</td></tr></tbody></table><table-wrap-foot><fn id="table3fn1"><p><sup>a</sup>LR: likelihood ratio.</p></fn><fn id="table3fn2"><p><sup>b</sup>FDR: false discovery rate.</p></fn><fn id="table3fn3"><p><sup>c</sup>SER: speech-based emotion recognition.</p></fn><fn id="table3fn4"><p><sup>d</sup>TER: text-based emotion recognition.</p></fn></table-wrap-foot></table-wrap><p>For text-based emotion recognition, performance group remained significantly associated with happy and sad text-emotion proportions after FDR correction. Fearful and surprise text-emotion proportions showed nominal associations with performance group but did not remain significant after correction. Angry, disgust, trust, expectation, and sentiment were not significantly associated with performance group. The largest excellent vs poor differences were observed for happy text affect, which was 8.66 percentage points higher in excellent performances; sad text affect, which was 5.33 percentage points lower in excellent performances; and surprise text affect, which was 4.13 percentage points lower in excellent performances.</p><p>Together, these sensitivity analyses support the overall direction of the primary findings: higher-rated simulations showed higher positive, calm, or neutral affective patterns, whereas lower-rated simulations showed higher fearful or surprise-related affective patterns. However, the text-based findings were weaker and more emotion-specific than the speech-based findings, supporting the interpretation that speech- and text-based emotion recognition provide complementary but differently sensitive measures of affective communication.</p></sec><sec id="s3-4"><title>Affective Communication Across Simulation Phases</title><p>For this experiment, we compared the alignment (ie, agreement) between TER and SER across the simulation phase. For each phase, we allocated the 2 most prevalent emotion categories; if the highest-scoring emotion exceeded 50%, it was considered &#x201C;dominant.&#x201D; Phase-specific analyses revealed cross-modal concordance: negative SER and TER categories peaked during extreme deterioration and ACLS management, whereas calm and positive emotions predominated during the initial condition and recovery phases. <xref ref-type="table" rid="table4">Table 4</xref> summarizes the top 2 categories per modality and phase; differences are most pronounced between the excellent and poor performance groups during high stress phases.</p><table-wrap id="t4" position="float"><label>Table 4.</label><caption><p>Participants&#x2019; text-based emotion recognition (TER) and speech-based emotion recognition (SER) results with respect to the simulation phases. The percentage indicates the prevalence of emotions.</p></caption><table id="table4" frame="hsides" rules="groups"><thead><tr><td align="left" valign="bottom"/><td align="left" valign="bottom" colspan="2">Initial condition</td><td align="left" valign="bottom" colspan="2">Extreme deterioration</td><td align="left" valign="bottom" colspan="2">ACLS<sup><xref ref-type="table-fn" rid="table4fn1">a</xref></sup> management</td><td align="left" valign="bottom" colspan="2">Recovery</td></tr><tr><td align="left" valign="top"/><td align="left" valign="top">Top 2 vocal emotions (%)</td><td align="left" valign="top">Top 2 textual emotions (%)</td><td align="left" valign="top">Top 2 vocal emotions (%)</td><td align="left" valign="top">Top 2 textual emotions (%)</td><td align="left" valign="top">Top 2 vocal emotions (%)</td><td align="left" valign="top">Top 2 textual emotions (%)</td><td align="left" valign="top">Top 2 vocal emotions (%)</td><td align="left" valign="top">Top 2 textual emotions (%)</td></tr></thead><tbody><tr><td align="left" valign="top">Poor performance</td><td align="left" valign="top"><list list-type="bullet"><list-item><p><italic>Fearful (50.03)<sup><xref ref-type="table-fn" rid="table4fn2">b</xref></sup></italic></p></list-item><list-item><p>Angry (31.78)</p></list-item></list></td><td align="left" valign="top"><list list-type="bullet"><list-item><p>Fearful (36.70)</p></list-item><list-item><p>Sad (30.22)</p></list-item></list></td><td align="left" valign="top"><list list-type="bullet"><list-item><p><italic>Surprise (64.12)</italic></p></list-item><list-item><p>Fearful (18.95)</p></list-item></list></td><td align="left" valign="top"><list list-type="bullet"><list-item><p><italic>Surprise (59.81)</italic></p></list-item><list-item><p>Fearful (14.08)</p></list-item></list></td><td align="left" valign="top"><list list-type="bullet"><list-item><p>Angry (11.89)</p></list-item><list-item><p><italic>Fearful (64.57)</italic></p></list-item></list></td><td align="left" valign="top"><list list-type="bullet"><list-item><p><italic>Surprise (55.78)</italic></p></list-item><list-item><p>Fearful (23.08)</p></list-item></list></td><td align="left" valign="top"><list list-type="bullet"><list-item><p>Sad (30.73)</p></list-item><list-item><p>Angry (24.67)</p></list-item></list></td><td align="left" valign="top"><list list-type="bullet"><list-item><p><italic>Sad (50.89)</italic></p></list-item><list-item><p>Surprise (26.84)</p></list-item></list></td></tr><tr><td align="left" valign="top">Intermediate performance</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>Calm (40.32)</p></list-item><list-item><p>Neutral (16.70)</p></list-item></list></td><td align="left" valign="top"><list list-type="bullet"><list-item><p><italic>Happy (52.13)</italic></p></list-item><list-item><p>Trust (33.41)</p></list-item></list></td><td align="left" valign="top"><list list-type="bullet"><list-item><p><italic>Calm (53.25</italic>)</p></list-item><list-item><p>Sad (23.97)</p></list-item></list></td><td align="left" valign="top"><list list-type="bullet"><list-item><p><italic>Sad (50.22</italic>)</p></list-item><list-item><p>Sentiment (20.69)</p></list-item></list></td><td align="left" valign="top"><list list-type="bullet"><list-item><p>Calm (44.19)</p></list-item><list-item><p>Neutral (35.67)</p></list-item></list></td><td align="left" valign="top"><list list-type="bullet"><list-item><p>Surprise (34.89)</p></list-item><list-item><p>Expectation (32.22)</p></list-item></list></td><td align="left" valign="top"><list list-type="bullet"><list-item><p>Happy (37.21)</p></list-item><list-item><p>Calm (34.18)</p></list-item></list></td><td align="left" valign="top"><list list-type="bullet"><list-item><p><italic>Happy (53.29)</italic></p></list-item><list-item><p>Trust (30.62)</p></list-item></list></td></tr><tr><td align="left" valign="top">Excellent performance</td><td align="left" valign="top"><list list-type="bullet"><list-item><p><italic>Calm (57.10</italic>)</p></list-item><list-item><p>Happy (27.85)</p></list-item></list></td><td align="left" valign="top"><list list-type="bullet"><list-item><p><italic>Expectation (66.78</italic>)</p></list-item><list-item><p>Sentiment (11.67)</p></list-item></list></td><td align="left" valign="top"><list list-type="bullet"><list-item><p><italic>Surprise (53.06</italic>)</p></list-item><list-item><p>Calm (22.65)</p></list-item></list></td><td align="left" valign="top"><list list-type="bullet"><list-item><p><italic>Surprise (63.88)</italic></p></list-item><list-item><p>Expectation (9.23)</p></list-item></list></td><td align="left" valign="top"><list list-type="bullet"><list-item><p>Neutral (49.10)</p></list-item><list-item><p>Calm (37.66)</p></list-item></list></td><td align="left" valign="top"><list list-type="bullet"><list-item><p><italic>Expectation (50.77)</italic></p></list-item><list-item><p>Happy (31.05)</p></list-item></list></td><td align="left" valign="top"><list list-type="bullet"><list-item><p><italic>Neutral (53.28)</italic></p></list-item><list-item><p>Happy (26.85)</p></list-item></list></td><td align="left" valign="top"><list list-type="bullet"><list-item><p><italic>Happy (64.92)</italic></p></list-item><list-item><p>Fearful (14.78)</p></list-item></list></td></tr></tbody></table><table-wrap-foot><fn id="table4fn1"><p><sup>a</sup>ACLS: advanced cardiovascular life support.</p></fn><fn id="table4fn2"><p><sup>b</sup>Cases in which one of the emotions has a prevalence of over 50% (ie, &#x201C;dominant&#x201D;) are indicated in italics.</p></fn></table-wrap-foot></table-wrap><p>Across phases, 3 broad patterns emerged. First, differences between performance groups were most pronounced during the clinically demanding phases of extreme deterioration and ACLS management. Second, poor performance simulations showed a greater concentration of high-arousal negative emotion labels, particularly fear and surprise, during these phases. Third, excellent performance simulations more often showed calm, neutral, or positive affective labels during phases requiring coordinated action and recovery. Thus, the phase-based results suggest that affective communication differences were not uniform across the simulation but were most evident during moments of acute clinical demand.</p></sec></sec><sec id="s4" sec-type="discussion"><title>Discussion</title><sec id="s4-1"><title>Principal Findings</title><p>In this prospective, simulation-based study of senior anesthesia residents, we observed systematic differences in model-derived affective communication patterns across performance groups. In the primary omnibus analysis, SER profiles differed across performance groups (<italic>F</italic><sub>2,161</sub>=4.61; <italic>P</italic>=.02): residents in the poor performance group (scores &#x2264;2) expressed predominantly fear and surprise, whereas those in the excellent performance group (score=5) showed higher proportions of calm, happy, and neutral speech affect. TER profiles showed a directionally similar but nonsignificant omnibus effect (<italic>F</italic><sub>2,161</sub>=2.57; <italic>P</italic>=.08). Exploratory mixed-effects sensitivity analyses supported the overall direction of the speech-based findings. Performance group remained associated with happy, angry, fearful, surprise, calm, and neutral speech-emotion proportions after FDR correction. The largest excellent vs poor differences were observed for neutral speech affect (+4.24 percentage points), happy speech affect (+4.11 percentage points), and surprise speech affect (&#x2212;3.17 percentage points). For text-based emotion recognition, performance group remained associated with happy and sad text-emotion proportions after correction, whereas fearful and surprise text-emotion proportions showed nominal associations that did not remain significant after correction. These findings suggest that speech-derived affective markers were more consistently associated with expert-rated performance than text-derived markers, although both modalities showed patterns consistent with the primary results.</p></sec><sec id="s4-2"><title>Interpretation of Affective Communication Patterns</title><p>The present study does not establish that emotion recognition outputs improve the accuracy, reliability, or predictive validity of simulation-based competency assessment. The findings show associations between model-derived affective communication patterns and expert-rated simulation performance. Therefore, these markers should be interpreted as exploratory and complementary signals that may inform formative feedback and future research, not as standalone assessment tools or substitutes for expert evaluation.</p><p>Phase-stratified summaries indicated that differences were most pronounced during extreme deterioration and ACLS management, when residents were required to recognize deterioration, coordinate the team, and execute time-sensitive actions. In these phases, higher-rated residents maintained calmer/neutral speech affect and more positive/neutral text tone. These associations suggest that affective communication patterns&#x2014;how emotion is expressed and regulated&#x2014;vary across performance groups and across the simulation phase. We interpret these findings as associations rather than causal effects. Elevated expression of emotions such as fear and surprise may co-occur with uncertainty or cognitive overload, whereas calm and neutral affect may co-occur with clearer commands, closed-loop confirmation, and anticipatory planning.</p></sec><sec id="s4-3"><title>Modality-Specific Findings: Speech vs Text</title><p>The discrepancy between SER and TER findings warrants careful interpretation. In the primary omnibus analysis, SER showed a statistically significant association with performance group, whereas TER showed a directionally similar but nonsignificant trend. This suggests that speech-derived affective information may have been more robustly captured than text-derived affective information in the present dataset. One explanation is that SER captures paralinguistic features, including pitch, intensity, tempo, pauses, hesitation, and vocal arousal, which may be especially informative during high-stress simulation. TER, in contrast, depends on a multistep pipeline; thus, errors or biases introduced at any stage may attenuate the estimates.</p><p>Several modality-specific limitations may have contributed to the weaker primary TER finding. Residents&#x2019; speech during crisis management was often brief, imperative, protocol-driven, and interspersed with Hebrew-English clinical terminology. Such utterances may contain limited lexical emotional information, even when vocal delivery conveys strong affective cues. Automatic transcription errors, word-level timestamping errors, code-switching, medical abbreviations, and domain-specific terminology may also introduce noise. In addition, TER models may misclassify clinically appropriate urgency, direct commands, or technical language as negative affect, whereas SER models may misinterpret raised volume or rapid speech as distress even when these features reflect appropriate leadership. Therefore, neither modality should be treated as a direct measure of psychological state or competency.</p><p>The mixed-effects sensitivity analyses suggested that some TER categories, particularly happy and sad, varied across performance groups after adjustment for repeated simulations and scenario-level clustering. However, these analyses were exploratory and emotion-specific; therefore, they should be interpreted as hypothesis-generating. Overall, the findings suggest that SER and TER provide complementary but imperfect views of affective communication, with different sources of noise, bias, and context dependence.</p></sec><sec id="s4-4"><title>Educational Implications</title><p>Clinically, the phase specificity matters. The clearest separation between excellent and poor performances emerged precisely when the team, specifically the resident, transitioned from recognizing deterioration to taking decisive action. That observation generates practical hypotheses for training: debriefings might explicitly review affective communication during these windows; phase-targeted coaching could rehearse calm and neutral delivery of closed-loop commands; and formative feedback could pair technical checklists with brief, objective summaries of affective patterns.</p><p>From an educational perspective, the most appropriate near-term use of these metrics is formative rather than summative. SER and TER outputs should not be used to make independent pass-fail decisions, rank residents, or replace expert faculty assessment. Instead, they may be used as debriefing aids that help educators identify specific moments in a simulation where affective communication changed. For example, an educator could review a timeline showing increases in high-arousal vocal affect, pauses, or shifts toward calmer communication and use these moments to guide reflective discussion about the clinical context, communication clarity, use of closed-loop communication, and whether the emotional tone supported or interfered with team coordination.</p></sec><sec id="s4-5"><title>Limitations and Future Directions</title><p>Several important limitations should be considered when interpreting these findings. First, this was an exploratory observational study; therefore, the results cannot determine whether the observed affective patterns reflect competency-relevant communication, cognitive load, stress response, individual communication style, or other latent constructs. Future studies should combine automated emotion recognition with expert-coded communication behaviors, physiological stress measures, and multi-rater performance assessments to clarify these mechanisms. Second, the primary analyses were univariate and did not fully adjust for potential confounders, including sex, postgraduate year, scenario, site, and language-related factors. Although we added exploratory mixed-effects sensitivity analyses, these models were not prespecified and should be interpreted as robustness checks rather than definitive confirmatory analyses. Future studies should use preregistered analytic plans, fully powered hierarchical models, and richer covariate structures to account for repeated simulations, scenario-level clustering, hospital affiliation, and other potential confounders.</p><p>Third, performance ratings were primarily based on single-rater expert assessments. Therefore, these ratings should be viewed as expert-rated indicators of simulation performance rather than definitive reference standards for clinical competency. Future work should incorporate multiple independent raters, standardized rating procedures, rater training, and reliability analyses to strengthen the validity of the performance outcome.</p><p>Fourth, SER may be affected by background alarms, overlapping speech, microphone placement, speech separation errors, speaker-specific vocal characteristics, and cultural or language-specific differences in vocal affect. TER may be affected by transcription errors, segmentation errors, code-switching, short utterance length, clinical abbreviations, and limited lexical emotional content in protocol-driven commands. These sources of noise may differentially affect SER and TER, partly explaining why the primary SER findings were stronger than the primary TER findings.</p><p>Additionally, the models operated on Hebrew speech with Hebrew-English code-switching, which may have amplified speech-to-text and emotion classification errors. While we accounted for domain adaptation in medical discourse, a known challenge given specialized vocabulary and context-dependent meaning [<xref ref-type="bibr" rid="ref27">27</xref>], we did not specifically adapt the models to the emotional and communicative context of anesthesiology crisis simulation.</p><p>Finally, time-aligned analyses could further examine whether short surges in high arousal negative affect precede delayed or missed actions, and whether calm or neutral affect during extreme deterioration and ACLS management is associated with more timely execution after adjustment for resident-, site-, and scenario-level factors.</p></sec><sec id="s4-6"><title>Conclusions</title><p>In conclusion, this exploratory observational study suggests that automated speech- and text-based emotion recognition can capture phase-specific affective communication patterns during anesthesiology critical incident simulation training. Speech-derived emotion markers, and to a lesser extent, text-derived markers, were associated with expert-rated performance groups. These findings should not be interpreted as evidence that emotion recognition can independently assess competency or improve assessment accuracy. Rather, they support the feasibility of using affective communication markers as complementary tools for simulation debriefing, formative feedback, and future research on communication under crisis conditions. Further validation using multirater outcomes, adjusted statistical models, and larger multicenter datasets is needed before such measures can inform competency assessment.</p></sec></sec></body><back><ack><p>The authors thank the Technion Autonomous Systems Program for institutional support of this work.</p><p>The authors used ChatGPT (OpenAI) to assist with language editing and improvement of manuscript readability. All scientific content, analytic decisions, interpretations, and final manuscript revisions were reviewed and approved by the authors, who take full responsibility for the content of the manuscript.</p></ack><notes><sec><title>Funding</title><p>No specific funding was received for this study.</p></sec><sec><title>Data Availability</title><p>The data generated and analyzed during this study are not publicly available because they include identifiable or potentially sensitive simulation audio and video recordings. Deidentified derived data may be made available from the corresponding author on reasonable request, subject to institutional review board approval and data-sharing restrictions. To improve reproducibility, the analysis scripts used to process model outputs, compute emotion category proportions, perform statistical analysis, and generate figures are available from the corresponding author on reasonable request.</p></sec></notes><fn-group><fn fn-type="con"><p>Conceptualization: SG, IB, AR, SL</p><p>Data curation: SG</p><p>Formal analysis: SG</p><p>Funding acquisition: SL, AR</p><p>Investigation: FM</p><p>Methodology: SG, IB, SL</p><p>Project administration: AR, FM</p><p>Resources: SL, AR</p><p>Software: SG</p><p>Supervision: AR, SL</p><p>Validation: SG, SL</p><p>Visualization: SG</p><p>Writing &#x2013; original draft: SG</p><p>Writing &#x2013; review &#x0026; editing: SG, IB, FM, AR, SL</p><p>All authors reviewed and approved the final manuscript.</p></fn><fn fn-type="conflict"><p>AR served as a consultant and received research support from Medtronic. He gave seminars for MSD (Merck &#x0026; Co) and for Medtechnica (Massimo). All other authors have no conflicts of interest.</p></fn></fn-group><glossary><title>Abbreviations</title><def-list><def-item><term id="abb1">ACLS</term><def><p>advanced cardiovascular life support</p></def></def-item><def-item><term id="abb2">ATLS</term><def><p>advanced trauma life support</p></def></def-item><def-item><term id="abb3">EEND-SS</term><def><p>end-to-end neural diarization with speaker subspace</p></def></def-item><def-item><term id="abb4">FDR</term><def><p>false discovery rate</p></def></def-item><def-item><term id="abb5">GPU</term><def><p>graphics processing unit</p></def></def-item><def-item><term id="abb6">ICU</term><def><p>intensive care unit</p></def></def-item><def-item><term id="abb7">NLP</term><def><p>natural language processing</p></def></def-item><def-item><term id="abb8">OR</term><def><p>operating room</p></def></def-item><def-item><term id="abb9">SBA</term><def><p>simulation-based assessment</p></def></def-item><def-item><term id="abb10">SER</term><def><p>speech-based emotion recognition</p></def></def-item><def-item><term id="abb11">TER</term><def><p>text-based emotion recognition</p></def></def-item></def-list></glossary><ref-list><title>References</title><ref id="ref1"><label>1</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Reader</surname><given-names>TW</given-names> </name><name name-style="western"><surname>Flin</surname><given-names>R</given-names> </name><name name-style="western"><surname>Mearns</surname><given-names>K</given-names> </name><name name-style="western"><surname>Cuthbertson</surname><given-names>BH</given-names> </name></person-group><article-title>Developing a team performance framework for the intensive care unit</article-title><source>Crit Care Med</source><year>2009</year><month>05</month><volume>37</volume><issue>5</issue><fpage>1787</fpage><lpage>1793</lpage><pub-id pub-id-type="doi">10.1097/CCM.0b013e31819f0451</pub-id><pub-id pub-id-type="medline">19325474</pub-id></nlm-citation></ref><ref id="ref2"><label>2</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Marcum</surname><given-names>JA</given-names> </name></person-group><article-title>The role of emotions in clinical reasoning and decision making</article-title><source>J Med Philos</source><year>2013</year><month>10</month><volume>38</volume><issue>5</issue><fpage>501</fpage><lpage>519</lpage><pub-id pub-id-type="doi">10.1093/jmp/jht040</pub-id><pub-id pub-id-type="medline">23975905</pub-id></nlm-citation></ref><ref id="ref3"><label>3</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Eisenberg</surname><given-names>EM</given-names> </name><name name-style="western"><surname>Murphy</surname><given-names>AG</given-names> </name><name name-style="western"><surname>Sutcliffe</surname><given-names>K</given-names> </name><etal/></person-group><article-title>Communication in emergency medicine: implications for patient safety</article-title><source>Commun Monogr</source><year>2005</year><volume>72</volume><issue>4</issue><fpage>390</fpage><lpage>413</lpage><pub-id pub-id-type="doi">10.1080/03637750500322602</pub-id></nlm-citation></ref><ref id="ref4"><label>4</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Isbell</surname><given-names>LM</given-names> </name><name name-style="western"><surname>Boudreaux</surname><given-names>ED</given-names> </name><name name-style="western"><surname>Chimowitz</surname><given-names>H</given-names> </name><name name-style="western"><surname>Liu</surname><given-names>G</given-names> </name><name name-style="western"><surname>Cyr</surname><given-names>E</given-names> </name><name name-style="western"><surname>Kimball</surname><given-names>E</given-names> </name></person-group><article-title>What do emergency department physicians and nurses feel? a qualitative study of emotions, triggers, regulation strategies, and effects on patient care</article-title><source>BMJ Qual Saf</source><year>2020</year><month>10</month><volume>29</volume><issue>10</issue><fpage>1</fpage><lpage>2</lpage><pub-id pub-id-type="doi">10.1136/bmjqs-2019-010179</pub-id><pub-id pub-id-type="medline">31941799</pub-id></nlm-citation></ref><ref id="ref5"><label>5</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Gurman</surname><given-names>GM</given-names> </name><name name-style="western"><surname>Klein</surname><given-names>M</given-names> </name><name name-style="western"><surname>Weksler</surname><given-names>N</given-names> </name></person-group><article-title>Professional stress in anesthesiology: a review</article-title><source>J Clin Monit Comput</source><year>2012</year><month>08</month><volume>26</volume><issue>4</issue><fpage>329</fpage><lpage>335</lpage><pub-id pub-id-type="doi">10.1007/s10877-011-9328-7</pub-id><pub-id pub-id-type="medline">22180163</pub-id></nlm-citation></ref><ref id="ref6"><label>6</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Gazoni</surname><given-names>FM</given-names> </name><name name-style="western"><surname>Amato</surname><given-names>PE</given-names> </name><name name-style="western"><surname>Malik</surname><given-names>ZM</given-names> </name><name name-style="western"><surname>Durieux</surname><given-names>ME</given-names> </name></person-group><article-title>The impact of perioperative catastrophes on anesthesiologists: results of a national survey</article-title><source>Anesth Analg</source><year>2012</year><month>03</month><volume>114</volume><issue>3</issue><fpage>596</fpage><lpage>603</lpage><pub-id pub-id-type="doi">10.1213/ANE.0b013e318227524e</pub-id><pub-id pub-id-type="medline">21737706</pub-id></nlm-citation></ref><ref id="ref7"><label>7</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Howie</surname><given-names>EE</given-names> </name><name name-style="western"><surname>Ambler</surname><given-names>O</given-names> </name><name name-style="western"><surname>Gunn</surname><given-names>EGM</given-names> </name><etal/></person-group><article-title>Surgical sabermetrics: a scoping review of technology-enhanced assessment of nontechnical skills in the operating room</article-title><source>Ann Surg</source><year>2024</year><month>06</month><day>1</day><volume>279</volume><issue>6</issue><fpage>973</fpage><lpage>984</lpage><pub-id pub-id-type="doi">10.1097/SLA.0000000000006211</pub-id><pub-id pub-id-type="medline">38258573</pub-id></nlm-citation></ref><ref id="ref8"><label>8</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Hu</surname><given-names>YY</given-names> </name><name name-style="western"><surname>Arriaga</surname><given-names>AF</given-names> </name><name name-style="western"><surname>Peyre</surname><given-names>SE</given-names> </name><name name-style="western"><surname>Corso</surname><given-names>KA</given-names> </name><name name-style="western"><surname>Roth</surname><given-names>EM</given-names> </name><name name-style="western"><surname>Greenberg</surname><given-names>CC</given-names> </name></person-group><article-title>Deconstructing intraoperative communication failures</article-title><source>J Surg Res</source><year>2012</year><month>09</month><volume>177</volume><issue>1</issue><fpage>37</fpage><lpage>42</lpage><pub-id pub-id-type="doi">10.1016/j.jss.2012.04.029</pub-id><pub-id pub-id-type="medline">22591922</pub-id></nlm-citation></ref><ref id="ref9"><label>9</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Hall</surname><given-names>A</given-names> </name><name name-style="western"><surname>Kawai</surname><given-names>K</given-names> </name><name name-style="western"><surname>Graber</surname><given-names>K</given-names> </name><etal/></person-group><article-title>Acoustic analysis of surgeons&#x2019; voices to assess change in the stress response during surgical in situ simulation</article-title><source>BMJ Simul Technol Enhanc Learn</source><year>2021</year><month>04</month><day>13</day><volume>7</volume><issue>6</issue><fpage>471</fpage><lpage>477</lpage><pub-id pub-id-type="doi">10.1136/bmjstel-2020-000727</pub-id><pub-id pub-id-type="medline">35520977</pub-id></nlm-citation></ref><ref id="ref10"><label>10</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Ross</surname><given-names>AJ</given-names> </name><name name-style="western"><surname>Kodate</surname><given-names>N</given-names> </name><name name-style="western"><surname>Anderson</surname><given-names>JE</given-names> </name><name name-style="western"><surname>Thomas</surname><given-names>L</given-names> </name><name name-style="western"><surname>Jaye</surname><given-names>P</given-names> </name></person-group><article-title>Review of simulation studies in anaesthesia journals, 2001-2010: mapping and content analysis</article-title><source>Br J Anaesth</source><year>2012</year><month>07</month><volume>109</volume><issue>1</issue><fpage>99</fpage><lpage>109</lpage><pub-id pub-id-type="doi">10.1093/bja/aes184</pub-id><pub-id pub-id-type="medline">22696559</pub-id></nlm-citation></ref><ref id="ref11"><label>11</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Reynolds</surname><given-names>T</given-names> </name><name name-style="western"><surname>Kong</surname><given-names>ML</given-names> </name></person-group><article-title>Shifting the learning curve</article-title><source>BMJ</source><year>2010</year><month>12</month><day>2</day><volume>341</volume><fpage>c6260</fpage><pub-id pub-id-type="doi">10.1136/bmj.c6260</pub-id><pub-id pub-id-type="medline">21127119</pub-id></nlm-citation></ref><ref id="ref12"><label>12</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Grantcharov</surname><given-names>TP</given-names> </name><name name-style="western"><surname>Kristiansen</surname><given-names>VB</given-names> </name><name name-style="western"><surname>Bendix</surname><given-names>J</given-names> </name><name name-style="western"><surname>Bardram</surname><given-names>L</given-names> </name><name name-style="western"><surname>Rosenberg</surname><given-names>J</given-names> </name><name name-style="western"><surname>Funch-Jensen</surname><given-names>P</given-names> </name></person-group><article-title>Randomized clinical trial of virtual reality simulation for laparoscopic skills training</article-title><source>Br J Surg</source><year>2004</year><month>02</month><volume>91</volume><issue>2</issue><fpage>146</fpage><lpage>150</lpage><pub-id pub-id-type="doi">10.1002/bjs.4407</pub-id><pub-id pub-id-type="medline">14760660</pub-id></nlm-citation></ref><ref id="ref13"><label>13</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Cecilio-Fernandes</surname><given-names>D</given-names> </name><name name-style="western"><surname>Brand&#x00E3;o</surname><given-names>CFS</given-names> </name><name name-style="western"><surname>de Oliveira</surname><given-names>DLC</given-names> </name><name name-style="western"><surname>Fernandes</surname><given-names>G</given-names> </name><name name-style="western"><surname>Tio</surname><given-names>RA</given-names> </name></person-group><article-title>Additional simulation training: does it affect students&#x2019; knowledge acquisition and retention?</article-title><source>BMJ Simul Technol Enhanc Learn</source><year>2018</year><month>06</month><day>22</day><volume>5</volume><issue>3</issue><fpage>140</fpage><lpage>143</lpage><pub-id pub-id-type="doi">10.1136/bmjstel-2018-000312</pub-id><pub-id pub-id-type="medline">35514946</pub-id></nlm-citation></ref><ref id="ref14"><label>14</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Weile</surname><given-names>J</given-names> </name><name name-style="western"><surname>Nebsbjerg</surname><given-names>MA</given-names> </name><name name-style="western"><surname>Ovesen</surname><given-names>SH</given-names> </name><name name-style="western"><surname>Paltved</surname><given-names>C</given-names> </name><name name-style="western"><surname>Ingeman</surname><given-names>ML</given-names> </name></person-group><article-title>Simulation-based team training in time-critical clinical presentations in emergency medicine and critical care: a review of the literature</article-title><source>Adv Simul (Lond)</source><year>2021</year><month>01</month><day>20</day><volume>6</volume><issue>1</issue><fpage>3</fpage><pub-id pub-id-type="doi">10.1186/s41077-021-00154-4</pub-id><pub-id pub-id-type="medline">33472706</pub-id></nlm-citation></ref><ref id="ref15"><label>15</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Robertson</surname><given-names>JM</given-names> </name><name name-style="western"><surname>Dias</surname><given-names>RD</given-names> </name><name name-style="western"><surname>Yule</surname><given-names>S</given-names> </name><name name-style="western"><surname>Smink</surname><given-names>DS</given-names> </name></person-group><article-title>Operating room team training with simulation: a systematic review</article-title><source>J Laparoendosc Adv Surg Tech A</source><year>2017</year><month>05</month><volume>27</volume><issue>5</issue><fpage>475</fpage><lpage>480</lpage><pub-id pub-id-type="doi">10.1089/lap.2017.0043</pub-id><pub-id pub-id-type="medline">28294695</pub-id></nlm-citation></ref><ref id="ref16"><label>16</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Bienstock</surname><given-names>J</given-names> </name><name name-style="western"><surname>Heuer</surname><given-names>A</given-names> </name><name name-style="western"><surname>Zhang</surname><given-names>Y</given-names> </name></person-group><article-title>Simulation-based training and its use amongst practicing paramedics and emergency medical technicians: an evidence-based systematic review</article-title><source>Intl J Paramedicine</source><year>2023</year><month>01</month><day>9</day><volume>1</volume><fpage>12</fpage><lpage>28</lpage><pub-id pub-id-type="doi">10.56068/VWHV8080</pub-id></nlm-citation></ref><ref id="ref17"><label>17</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>McGaghie</surname><given-names>WC</given-names> </name><name name-style="western"><surname>Issenberg</surname><given-names>SB</given-names> </name><name name-style="western"><surname>Petrusa</surname><given-names>ER</given-names> </name><name name-style="western"><surname>Scalese</surname><given-names>RJ</given-names> </name></person-group><article-title>A critical review of simulation-based medical education research: 2003-2009</article-title><source>Med Educ</source><year>2010</year><month>01</month><volume>44</volume><issue>1</issue><fpage>50</fpage><lpage>63</lpage><pub-id pub-id-type="doi">10.1111/j.1365-2923.2009.03547.x</pub-id><pub-id pub-id-type="medline">20078756</pub-id></nlm-citation></ref><ref id="ref18"><label>18</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Ziv</surname><given-names>A</given-names> </name><name name-style="western"><surname>Rubin</surname><given-names>O</given-names> </name><name name-style="western"><surname>Sidi</surname><given-names>A</given-names> </name><name name-style="western"><surname>Berkenstadt</surname><given-names>H</given-names> </name></person-group><article-title>Credentialing and certifying with simulation</article-title><source>Anesthesiol Clin</source><year>2007</year><month>06</month><volume>25</volume><issue>2</issue><fpage>261</fpage><lpage>269</lpage><pub-id pub-id-type="doi">10.1016/j.anclin.2007.03.002</pub-id><pub-id pub-id-type="medline">17574189</pub-id></nlm-citation></ref><ref id="ref19"><label>19</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Seehusen</surname><given-names>DA</given-names> </name><name name-style="western"><surname>Kleinheksel</surname><given-names>AJ</given-names> </name><name name-style="western"><surname>Huang</surname><given-names>H</given-names> </name><name name-style="western"><surname>Harrison</surname><given-names>Z</given-names> </name><name name-style="western"><surname>Ledford</surname><given-names>CJW</given-names> </name></person-group><article-title>The power of one word to paint a halo or a horn: demonstrating the halo effect in learner handover and subsequent evaluation</article-title><source>Acad Med</source><year>2023</year><month>08</month><day>1</day><volume>98</volume><issue>8</issue><fpage>929</fpage><lpage>933</lpage><pub-id pub-id-type="doi">10.1097/ACM.0000000000005161</pub-id><pub-id pub-id-type="medline">36724305</pub-id></nlm-citation></ref><ref id="ref20"><label>20</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Humphrey-Murto</surname><given-names>S</given-names> </name><name name-style="western"><surname>Shaw</surname><given-names>T</given-names> </name><name name-style="western"><surname>Touchie</surname><given-names>C</given-names> </name><name name-style="western"><surname>Pugh</surname><given-names>D</given-names> </name><name name-style="western"><surname>Cowley</surname><given-names>L</given-names> </name><name name-style="western"><surname>Wood</surname><given-names>TJ</given-names> </name></person-group><article-title>Are raters influenced by prior information about a learner? a review of assimilation and contrast effects in assessment</article-title><source>Adv Health Sci Educ Theory Pract</source><year>2021</year><month>08</month><volume>26</volume><issue>3</issue><fpage>1133</fpage><lpage>1156</lpage><pub-id pub-id-type="doi">10.1007/s10459-021-10032-3</pub-id><pub-id pub-id-type="medline">33566199</pub-id></nlm-citation></ref><ref id="ref21"><label>21</label><nlm-citation citation-type="confproc"><person-group person-group-type="author"><name name-style="western"><surname>Adoma</surname><given-names>AF</given-names> </name><name name-style="western"><surname>Henry</surname><given-names>NM</given-names> </name><name name-style="western"><surname>Chen</surname><given-names>W</given-names> </name></person-group><article-title>Comparative analyses of BERT, RoBERTa, DistilBERT, and XLNet for text-based emotion recognition</article-title><access-date>2026-08-07</access-date><conf-name>2020 17th International Computer Conference on Wavelet Active Media Technology and Information Processing (ICCWAMTIP)</conf-name><conf-date>Dec 18-20, 2020</conf-date><comment><ext-link ext-link-type="uri" xlink:href="https://ieeexplore.ieee.org/document/9317379">https://ieeexplore.ieee.org/document/9317379</ext-link></comment><pub-id pub-id-type="doi">10.1109/ICCWAMTIP51612.2020.9317379</pub-id></nlm-citation></ref><ref id="ref22"><label>22</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Acheampong</surname><given-names>FA</given-names> </name><name name-style="western"><surname>Wenyu</surname><given-names>C</given-names> </name><name name-style="western"><surname>Nunoo&#x2010;Mensah</surname><given-names>H</given-names> </name></person-group><article-title>Text&#x2010;based emotion detection: advances, challenges, and opportunities</article-title><source>Eng Rep</source><year>2020</year><month>07</month><volume>2</volume><issue>7</issue><fpage>e12189</fpage><pub-id pub-id-type="doi">10.1002/eng2.12189</pub-id></nlm-citation></ref><ref id="ref23"><label>23</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Khalil</surname><given-names>RA</given-names> </name><name name-style="western"><surname>Jones</surname><given-names>E</given-names> </name><name name-style="western"><surname>Babar</surname><given-names>MI</given-names> </name><name name-style="western"><surname>Jan</surname><given-names>T</given-names> </name><name name-style="western"><surname>Zafar</surname><given-names>MH</given-names> </name><name name-style="western"><surname>Alhussain</surname><given-names>T</given-names> </name></person-group><article-title>Speech emotion recognition using deep learning techniques: a review</article-title><source>IEEE Access</source><year>2019</year><volume>7</volume><fpage>117327</fpage><lpage>117345</lpage><pub-id pub-id-type="doi">10.1109/ACCESS.2019.2936124</pub-id></nlm-citation></ref><ref id="ref24"><label>24</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Wani</surname><given-names>TM</given-names> </name><name name-style="western"><surname>Gunawan</surname><given-names>TS</given-names> </name><name name-style="western"><surname>Qadri</surname><given-names>SAA</given-names> </name><name name-style="western"><surname>Kartiwi</surname><given-names>M</given-names> </name><name name-style="western"><surname>Ambikairajah</surname><given-names>E</given-names> </name></person-group><article-title>A comprehensive review of speech emotion recognition systems</article-title><source>IEEE Access</source><year>2021</year><volume>9</volume><fpage>47795</fpage><lpage>47814</lpage><pub-id pub-id-type="doi">10.1109/ACCESS.2021.3068045</pub-id></nlm-citation></ref><ref id="ref25"><label>25</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Latif</surname><given-names>S</given-names> </name><name name-style="western"><surname>Qadir</surname><given-names>J</given-names> </name><name name-style="western"><surname>Qayyum</surname><given-names>A</given-names> </name><name name-style="western"><surname>Usama</surname><given-names>M</given-names> </name><name name-style="western"><surname>Younis</surname><given-names>S</given-names> </name></person-group><article-title>Speech technology for healthcare: opportunities, challenges, and state of the art</article-title><source>IEEE Rev Biomed Eng</source><year>2021</year><volume>14</volume><fpage>342</fpage><lpage>356</lpage><pub-id pub-id-type="doi">10.1109/RBME.2020.3006860</pub-id><pub-id pub-id-type="medline">32746367</pub-id></nlm-citation></ref><ref id="ref26"><label>26</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Birjali</surname><given-names>M</given-names> </name><name name-style="western"><surname>Kasri</surname><given-names>M</given-names> </name><name name-style="western"><surname>Beni-Hssane</surname><given-names>A</given-names> </name></person-group><article-title>A comprehensive survey on sentiment analysis: approaches, challenges and trends</article-title><source>Knowl Based Syst</source><year>2021</year><month>08</month><day>17</day><volume>226</volume><fpage>107134</fpage><pub-id pub-id-type="doi">10.1016/j.knosys.2021.107134</pub-id></nlm-citation></ref><ref id="ref27"><label>27</label><nlm-citation citation-type="book"><person-group person-group-type="author"><name name-style="western"><surname>Holderness</surname><given-names>E</given-names> </name><name name-style="western"><surname>Cawkwell</surname><given-names>P</given-names> </name><name name-style="western"><surname>Bolton</surname><given-names>K</given-names> </name><name name-style="western"><surname>Pustejovsky</surname><given-names>J</given-names> </name><name name-style="western"><surname>Hall</surname><given-names>MH</given-names> </name></person-group><article-title>Distinguishing clinical sentiment: the importance of domain adaptation in psychiatric patient health records</article-title><source>Proceedings of the 2nd Clinical Natural Language Processing Workshop</source><year>2019</year><access-date>2026-08-07</access-date><publisher-name>Association for Computational Linguistics</publisher-name><fpage>117</fpage><lpage>123</lpage><comment><ext-link ext-link-type="uri" xlink:href="https://aclanthology.org/W19-1915/">https://aclanthology.org/W19-1915/</ext-link></comment><pub-id pub-id-type="doi">10.18653/v1/W19-1915</pub-id></nlm-citation></ref><ref id="ref28"><label>28</label><nlm-citation citation-type="confproc"><person-group person-group-type="author"><name name-style="western"><surname>Ahmed</surname><given-names>T</given-names> </name><name name-style="western"><surname>Gopala Krishnan</surname><given-names>C</given-names> </name></person-group><article-title>A comprehensive study on emotion recognition in healthcare by applying machine learning and deep learning techniques</article-title><access-date>2026-08-07</access-date><conf-name>2024 International Conference on Knowledge Engineering and Communication Systems (ICKECS)</conf-name><conf-date>Apr 18-19, 2024</conf-date><comment><ext-link ext-link-type="uri" xlink:href="https://ieeexplore.ieee.org/document/10617024">https://ieeexplore.ieee.org/document/10617024</ext-link></comment><pub-id pub-id-type="doi">10.1109/ICKECS61492.2024.10617024</pub-id></nlm-citation></ref><ref id="ref29"><label>29</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Yamamoto</surname><given-names>S</given-names> </name><name name-style="western"><surname>Tanaka</surname><given-names>P</given-names> </name><name name-style="western"><surname>Madsen</surname><given-names>MV</given-names> </name><name name-style="western"><surname>Macario</surname><given-names>A</given-names> </name></person-group><article-title>Comparing anesthesiology residency training structure and requirements in seven different countries on three continents</article-title><source>Cureus</source><year>2017</year><month>02</month><day>26</day><volume>9</volume><issue>2</issue><fpage>e1060</fpage><pub-id pub-id-type="doi">10.7759/cureus.1060</pub-id><pub-id pub-id-type="medline">28367396</pub-id></nlm-citation></ref><ref id="ref30"><label>30</label><nlm-citation citation-type="confproc"><person-group person-group-type="author"><name name-style="western"><surname>Maiti</surname><given-names>S</given-names> </name><name name-style="western"><surname>Ueda</surname><given-names>Y</given-names> </name><name name-style="western"><surname>Watanabe</surname><given-names>S</given-names> </name><etal/></person-group><article-title>EEND-SS: joint end-to-end neural speaker diarization and speech separation for flexible number of speakers</article-title><access-date>2026-08-07</access-date><conf-name>2022 IEEE Spoken Language Technology Workshop (SLT)</conf-name><conf-date>Jan 9-12, 2023</conf-date><comment><ext-link ext-link-type="uri" xlink:href="https://ieeexplore.ieee.org/document/10022924">https://ieeexplore.ieee.org/document/10022924</ext-link></comment><pub-id pub-id-type="doi">10.1109/SLT54892.2023.10022924</pub-id></nlm-citation></ref><ref id="ref31"><label>31</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Ravanelli</surname><given-names>M</given-names> </name><name name-style="western"><surname>Parcollet</surname><given-names>T</given-names> </name><name name-style="western"><surname>Moumen</surname><given-names>A</given-names> </name><etal/></person-group><article-title>Open-source conversational AI with SpeechBrain 1.0</article-title><source>J Mach Learn Res</source><year>2024</year><access-date>2026-08-07</access-date><volume>25</volume><issue>333</issue><fpage>1</fpage><lpage>11</lpage><comment><ext-link ext-link-type="uri" xlink:href="https://www.jmlr.org/papers/v25/24-0991.html">https://www.jmlr.org/papers/v25/24-0991.html</ext-link></comment></nlm-citation></ref><ref id="ref32"><label>32</label><nlm-citation citation-type="book"><person-group person-group-type="author"><name name-style="western"><surname>Baevski</surname><given-names>A</given-names> </name><name name-style="western"><surname>Zhou</surname><given-names>H</given-names> </name><name name-style="western"><surname>Mohamed</surname><given-names>A</given-names> </name><name name-style="western"><surname>Auli</surname><given-names>M</given-names> </name></person-group><article-title>Wav2vec 2.0: a framework for self-supervised learning of speech representations</article-title><source>NIPS&#x2019;20: Proceedings of the 34th International Conference on Neural Information Processing Systems</source><year>2020</year><access-date>2026-08-07</access-date><publisher-name>Neural Information Processing Systems Foundation</publisher-name><fpage>12449</fpage><lpage>12460</lpage><comment><ext-link ext-link-type="uri" xlink:href="https://dl.acm.org/doi/abs/10.5555/3495724.3496768">https://dl.acm.org/doi/abs/10.5555/3495724.3496768</ext-link></comment></nlm-citation></ref><ref id="ref33"><label>33</label><nlm-citation citation-type="confproc"><person-group person-group-type="author"><name name-style="western"><surname>Pepino</surname><given-names>L</given-names> </name><name name-style="western"><surname>Riera</surname><given-names>P</given-names> </name><name name-style="western"><surname>Ferrer</surname><given-names>L</given-names> </name></person-group><article-title>Emotion recognition from speech using wav2vec 2.0 embeddings</article-title><year>2021</year><access-date>2026-08-07</access-date><conf-name>Interspeech 2021</conf-name><conf-date>Aug 30 to Sep 3, 2021</conf-date><conf-loc>Brno, Czechia</conf-loc><fpage>3400</fpage><lpage>3404</lpage><comment><ext-link ext-link-type="uri" xlink:href="https://www.isca-archive.org/interspeech_2021/pepino21_interspeech.html">https://www.isca-archive.org/interspeech_2021/pepino21_interspeech.html</ext-link></comment><pub-id pub-id-type="doi">10.21437/Interspeech.2021-703</pub-id></nlm-citation></ref><ref id="ref34"><label>34</label><nlm-citation citation-type="book"><person-group person-group-type="author"><name name-style="western"><surname>Radford</surname><given-names>A</given-names> </name><name name-style="western"><surname>Kim</surname><given-names>JW</given-names> </name><name name-style="western"><surname>Xu</surname><given-names>T</given-names> </name><name name-style="western"><surname>Brockman</surname><given-names>G</given-names> </name><name name-style="western"><surname>McLeavey</surname><given-names>C</given-names> </name><name name-style="western"><surname>Sutskever</surname><given-names>I</given-names> </name></person-group><article-title>Robust speech recognition via large-scale weak supervision</article-title><source>Proceedings of the 40th International Conference on Machine Learning</source><year>2023</year><access-date>2026-08-07</access-date><publisher-name>PLMR</publisher-name><comment><ext-link ext-link-type="uri" xlink:href="https://proceedings.mlr.press/v202/radford23a.html">https://proceedings.mlr.press/v202/radford23a.html</ext-link></comment></nlm-citation></ref><ref id="ref35"><label>35</label><nlm-citation citation-type="confproc"><person-group person-group-type="author"><name name-style="western"><surname>Bain</surname><given-names>M</given-names> </name><name name-style="western"><surname>Huh</surname><given-names>J</given-names> </name><name name-style="western"><surname>Han</surname><given-names>T</given-names> </name><name name-style="western"><surname>Zisserman</surname><given-names>A</given-names> </name></person-group><article-title>WhisperX: time-accurate speech transcription of long-form audio</article-title><year>2023</year><access-date>2026-08-07</access-date><conf-name>Interspeech 2023</conf-name><conf-date>Aug 20-24, 2023</conf-date><conf-loc>Dublin, Ireland</conf-loc><fpage>4489</fpage><lpage>4493</lpage><comment><ext-link ext-link-type="uri" xlink:href="https://www.isca-archive.org/interspeech_2023/bain23_interspeech.html">https://www.isca-archive.org/interspeech_2023/bain23_interspeech.html</ext-link></comment><pub-id pub-id-type="doi">10.21437/Interspeech.2023-78</pub-id></nlm-citation></ref><ref id="ref36"><label>36</label><nlm-citation citation-type="book"><person-group person-group-type="author"><name name-style="western"><surname>Seker</surname><given-names>A</given-names> </name><name name-style="western"><surname>Bandel</surname><given-names>E</given-names> </name><name name-style="western"><surname>Bareket</surname><given-names>D</given-names> </name><name name-style="western"><surname>Brusilovsky</surname><given-names>I</given-names> </name><name name-style="western"><surname>Greenfeld</surname><given-names>R</given-names> </name><name name-style="western"><surname>Tsarfaty</surname><given-names>R</given-names> </name></person-group><article-title>AlephBERT: language model pre-training and evaluation from sub-word to sentence level</article-title><source>Proceedings of the 60th Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers)</source><year>2022</year><access-date>2026-08-07</access-date><publisher-name>Association for Computational Linguistics</publisher-name><fpage>46</fpage><lpage>56</lpage><comment><ext-link ext-link-type="uri" xlink:href="https://aclanthology.org/2022.acl-long">https://aclanthology.org/2022.acl-long</ext-link></comment><pub-id pub-id-type="doi">10.18653/v1/2022.acl-long.4</pub-id></nlm-citation></ref><ref id="ref37"><label>37</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Chriqui</surname><given-names>A</given-names> </name><name name-style="western"><surname>Yahav</surname><given-names>I</given-names> </name></person-group><article-title>HeBERT and HebEMO: A Hebrew BERT model and a tool for polarity analysis and emotion recognition</article-title><source>INFORMS J Dat Sci</source><year>2022</year><month>04</month><volume>1</volume><issue>1</issue><fpage>81</fpage><lpage>95</lpage><pub-id pub-id-type="doi">10.1287/ijds.2022.0016</pub-id></nlm-citation></ref><ref id="ref38"><label>38</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Gershov</surname><given-names>S</given-names> </name><name name-style="western"><surname>Braunold</surname><given-names>D</given-names> </name><name name-style="western"><surname>Spektor</surname><given-names>R</given-names> </name><name name-style="western"><surname>Ioscovich</surname><given-names>A</given-names> </name><name name-style="western"><surname>Raz</surname><given-names>A</given-names> </name><name name-style="western"><surname>Laufer</surname><given-names>S</given-names> </name></person-group><article-title>Automating medical simulations</article-title><source>J Biomed Inform</source><year>2023</year><month>08</month><volume>144</volume><fpage>104446</fpage><pub-id pub-id-type="doi">10.1016/j.jbi.2023.104446</pub-id><pub-id pub-id-type="medline">37467836</pub-id></nlm-citation></ref><ref id="ref39"><label>39</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Luo</surname><given-names>Y</given-names> </name><name name-style="western"><surname>Mesgarani</surname><given-names>N</given-names> </name></person-group><article-title>Conv-TasNet: surpassing ideal time-frequency magnitude masking for speech separation</article-title><source>IEEE/ACM Trans Audio Speech Lang Process</source><year>2019</year><month>08</month><volume>27</volume><issue>8</issue><fpage>1256</fpage><lpage>1266</lpage><pub-id pub-id-type="doi">10.1109/TASLP.2019.2915167</pub-id><pub-id pub-id-type="medline">31485462</pub-id></nlm-citation></ref></ref-list></back></article>