<?xml version="1.0" encoding="UTF-8"?><!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.0 20040830//EN" "journalpublishing.dtd"><article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" dtd-version="2.0" xml:lang="en" article-type="research-article"><front><journal-meta><journal-id journal-id-type="nlm-ta">JMIR Med Educ</journal-id><journal-id journal-id-type="publisher-id">mededu</journal-id><journal-id journal-id-type="index">20</journal-id><journal-title>JMIR Medical Education</journal-title><abbrev-journal-title>JMIR Med Educ</abbrev-journal-title><issn pub-type="epub">2369-3762</issn><publisher><publisher-name>JMIR Publications</publisher-name><publisher-loc>Toronto, Canada</publisher-loc></publisher></journal-meta><article-meta><article-id pub-id-type="publisher-id">v12i1e97822</article-id><article-id pub-id-type="doi">10.2196/97822</article-id><article-categories><subj-group subj-group-type="heading"><subject>Original Paper</subject></subj-group></article-categories><title-group><article-title>Teaching Model Context Protocol, Retrieval-Augmented Generation, and AI Agents to a Multidisciplinary Hospital Workforce: Single-Group Pre-Post Survey Study</article-title></title-group><contrib-group><contrib contrib-type="author"><name name-style="western"><surname>Baek</surname><given-names>Gakyoung</given-names></name><degrees>BS</degrees><xref ref-type="aff" rid="aff1">1</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Lee</surname><given-names>Hyunna</given-names></name><degrees>PhD</degrees><xref ref-type="aff" rid="aff1">1</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Yang</surname><given-names>Dong Hyun</given-names></name><degrees>MD, PhD</degrees><xref ref-type="aff" rid="aff2">2</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Kang</surname><given-names>Minseo</given-names></name><degrees>BS</degrees><xref ref-type="aff" rid="aff1">1</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Lee</surname><given-names>Kun Hee</given-names></name><degrees>BS</degrees><xref ref-type="aff" rid="aff1">1</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Choi</surname><given-names>Minji</given-names></name><degrees>BA</degrees><xref ref-type="aff" rid="aff1">1</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Lee</surname><given-names>Yura</given-names></name><degrees>MD, PhD</degrees><xref ref-type="aff" rid="aff3">3</xref></contrib><contrib contrib-type="author" corresp="yes"><name name-style="western"><surname>Lee</surname><given-names>Kye Hwa</given-names></name><degrees>MD, PhD</degrees><xref ref-type="aff" rid="aff3">3</xref></contrib></contrib-group><aff id="aff1"><institution>Big Data Research Center, Asan Institute for Life Sciences, Asan Medical Center</institution><addr-line>Seoul</addr-line><country>Republic of Korea</country></aff><aff id="aff2"><institution>Department of Radiology and Research Institute of Radiology, University of Ulsan College of Medicine, Asan Medical Center</institution><addr-line>Seoul</addr-line><country>Republic of Korea</country></aff><aff id="aff3"><institution>Department of Information Medicine, Asan Medical Center, University of Ulsan College of Medicine</institution><addr-line>88 Olympic-ro 43-gil, Songpa-gu</addr-line><addr-line>Seoul</addr-line><country>Republic of Korea</country></aff><contrib-group><contrib contrib-type="editor"><name name-style="western"><surname>Stone</surname><given-names>Alicia</given-names></name></contrib></contrib-group><contrib-group><contrib contrib-type="reviewer"><name name-style="western"><surname>Yoon</surname><given-names>Dukyong</given-names></name></contrib><contrib contrib-type="reviewer"><name name-style="western"><surname>Krive</surname><given-names>Jacob</given-names></name></contrib></contrib-group><author-notes><corresp>Correspondence to Kye Hwa Lee, MD, PhD, Department of Information Medicine, Asan Medical Center, University of Ulsan College of Medicine, 88 Olympic-ro 43-gil, Songpa-gu, Seoul, 05505, Republic of Korea, 82 2-3010-5991; <email>eva@amc.seoul.kr</email></corresp></author-notes><pub-date pub-type="collection"><year>2026</year></pub-date><pub-date pub-type="epub"><day>11</day><month>9</month><year>2026</year></pub-date><volume>12</volume><elocation-id>e97822</elocation-id><history><date date-type="received"><day>13</day><month>04</month><year>2026</year></date><date date-type="rev-recd"><day>10</day><month>07</month><year>2026</year></date><date date-type="accepted"><day>10</day><month>07</month><year>2026</year></date></history><copyright-statement>&#x00A9; Gakyoung Baek, Hyunna Lee, Dong Hyun Yang, Minseo Kang, Kun Hee Lee, Minji Choi, Yura Lee, Kye Hwa Lee. Originally published in JMIR Medical Education (<ext-link ext-link-type="uri" xlink:href="https://mededu.jmir.org">https://mededu.jmir.org</ext-link>), 11.9.2026. </copyright-statement><copyright-year>2026</copyright-year><license license-type="open-access" xlink:href="https://creativecommons.org/licenses/by/4.0/"><p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (<ext-link ext-link-type="uri" xlink:href="https://creativecommons.org/licenses/by/4.0/">https://creativecommons.org/licenses/by/4.0/</ext-link>), which permits unrestricted use, distribution, and reproduction in any medium, provided the original work, first published in JMIR Medical Education, is properly cited. The complete bibliographic information, a link to the original publication on <ext-link ext-link-type="uri" xlink:href="https://mededu.jmir.org/">https://mededu.jmir.org/</ext-link>, as well as this copyright and license information must be included.</p></license><self-uri xlink:type="simple" xlink:href="https://mededu.jmir.org/2026/1/e97822"/><abstract><sec><title>Background</title><p>Hospitals worldwide need to upskill their workforce in advanced AI technologies; yet, published guidance on how to design and deliver such training, particularly in agent-level tools like retrieval-augmented generation (RAG) and the model context protocol (MCP), remains virtually absent.</p></sec><sec><title>Objective</title><p>To describe the design, implementation, and lessons learned from an 8-week, 56-hour intensive generative AI training program for a multidisciplinary hospital workforce, drawing on both quantitative outcome data and participants&#x2019; own reflections on their learning experience.</p></sec><sec sec-type="methods"><title>Methods</title><p>The program was delivered on-site at Asan Medical Center with simultaneous online broadcast to 2 regional affiliate hospitals. The curriculum was built around the premise that MCP and AI agents would become the foundation of health care AI use, allocating 37% (11.5/31 hours) of on-site instructional time to MCP, and 71% (22/31 hours) to hands-on practice. Participants progressed from foundational concepts through RAG and MCP to team-based capstone projects, supported by funded AI tool subscriptions, a dedicated internal cloud platform, and 3&#x2010;6 hours of weekly mentoring per team. A pre-post survey (pre: n=83; post: n=64) evaluated outcomes across Kirkpatrick levels 1&#x2010;3, complemented by thematic analysis of open-ended reflections on self-perceived growth.</p></sec><sec sec-type="results"><title>Results</title><p>The technologies that received the greatest curricular investment were associated with the largest self-efficacy differences (MCP: <italic>d</italic>=1.57; overall effect: <italic>r</italic>=.574), and participants most frequently cited MCP and RAG when describing how abstract concepts &#x201C;became concrete and actionable.&#x201D; Non-IT professionals, clinicians, health information managers, researchers, and administrative staff showed consistently larger gains than IT specialists; several reported coding for the first time through vibe coding, challenging the assumption that advanced AI training requires technical backgrounds. Despite significant overall gains, a knowledge-practice gap persisted: job-specific competency remained below the scale midpoint, though participants spontaneously reported generating workplace application ideas. Curriculum pacing was rated lowest despite high overall satisfaction (4.03/5), signaling that even 56 hours may progress too quickly for mixed-expertise cohorts. Capstone projects with dedicated mentoring received the highest satisfaction ratings; 11 of 12 teams presented functional prototypes, and one has since entered active pilot use in clinical departments ahead of planned hospital-wide deployment.</p></sec><sec sec-type="conclusions"><title>Conclusions</title><p>To our knowledge, this is the first program to teach 4 agent-level generative AI technologies, MCP, RAG, LangGraph orchestration, and AI agent design, to both IT and non-IT hospital staff. This program suggests that transforming a multidisciplinary hospital workforce into AI-capable professionals is achievable through intensive, hands-on training centered on agent-level technologies, and that capstone projects with dedicated mentoring can serve as a pathway from classroom learning toward institutional AI adoption. The knowledge-practice gap highlights the need for posttraining support structures to translate self-efficacy gains into sustained workplace practice.</p></sec></abstract><kwd-group><kwd>artificial intelligence</kwd><kwd>medical education</kwd><kwd>large language models</kwd><kwd>self-efficacy</kwd><kwd>health personnel</kwd><kwd>model context protocol</kwd><kwd>staff development</kwd><kwd>curriculum design</kwd></kwd-group></article-meta></front><body><sec id="s1" sec-type="intro"><title>Introduction</title><sec id="s1-1"><title>Background</title><p>Hospitals that want to harness AI for clinical documentation, decision support, and administrative automation [<xref ref-type="bibr" rid="ref1">1</xref>-<xref ref-type="bibr" rid="ref4">4</xref>] face a workforce problem: the technologies that make these applications possible, retrieval-augmented generation (RAG) [<xref ref-type="bibr" rid="ref5">5</xref>], the model context protocol (MCP) [<xref ref-type="bibr" rid="ref6">6</xref>], and orchestration frameworks such as LangGraph [<xref ref-type="bibr" rid="ref7">7</xref>], are evolving faster than the people who must build, adapt, and evaluate them [<xref ref-type="bibr" rid="ref8">8</xref>,<xref ref-type="bibr" rid="ref9">9</xref>]. The question is no longer whether to train hospital staff in advanced AI, but how.</p><p>Existing training efforts offer limited guidance. The 2024 Best Evidence Medical Education (BEME) scoping review by Gordon et al [<xref ref-type="bibr" rid="ref10">10</xref>] confirmed that published AI education programs remain predominantly foundational&#x2014;covering AI literacy, prompt engineering, and learners&#x2019; perceptions of generative AI tools such as ChatGPT [<xref ref-type="bibr" rid="ref11">11</xref>-<xref ref-type="bibr" rid="ref13">13</xref>]&#x2014;and target mainly medical students and residents [<xref ref-type="bibr" rid="ref14">14</xref>,<xref ref-type="bibr" rid="ref15">15</xref>]. Evidence on intensive, multiweek programs that teach advanced technologies to the broader hospital workforce&#x2014;IT specialists, health information managers (HIM), clinicians, researchers, and administrative staff&#x2014;is scarce [<xref ref-type="bibr" rid="ref16">16</xref>-<xref ref-type="bibr" rid="ref18">18</xref>]. No published study, to our knowledge, has reported the design and outcomes of training that covers RAG, MCP, AI agent design, and workflow orchestration together.</p></sec><sec id="s1-2"><title>Prior Work</title><p>Self-efficacy theory [<xref ref-type="bibr" rid="ref19">19</xref>] predicts that mastery experiences, particularly hands-on practice with novel technologies, are the strongest driver of sustained technology adoption [<xref ref-type="bibr" rid="ref20">20</xref>-<xref ref-type="bibr" rid="ref22">22</xref>]. Consistent with this, recent ChatGPT education studies report that structured training significantly raises perceived AI competency among health care professionals [<xref ref-type="bibr" rid="ref8">8</xref>,<xref ref-type="bibr" rid="ref12">12</xref>], and experiential learning approaches outperform didactic instruction for skill-based outcomes [<xref ref-type="bibr" rid="ref23">23</xref>]. However, most evaluations remain limited to satisfaction surveys (Kirkpatrick level 1 [<xref ref-type="bibr" rid="ref24">24</xref>,<xref ref-type="bibr" rid="ref25">25</xref>]), with few assessing knowledge gains (Level 2) or behavioral transfer intentions (Level 3) [<xref ref-type="bibr" rid="ref26">26</xref>,<xref ref-type="bibr" rid="ref27">27</xref>]. Our previous evaluation of a health informatics analyst program at the same institution demonstrated that intensive education can reshape job roles and skill profiles [<xref ref-type="bibr" rid="ref18">18</xref>]; this study extends that work to advanced generative AI technologies.</p></sec><sec id="s1-3"><title>Study Objectives</title><p>This study has two complementary aims: (1) to describe the design rationale, curriculum structure, and implementation experience of an 8-week, 56-hour intensive generative AI training program for a multidisciplinary hospital workforce at a tertiary academic medical center; and (2) to evaluate its outcomes using the Kirkpatrick framework across Levels 1&#x2010;3, providing empirical evidence that contextualizes the lessons learned. We addressed four research questions (RQs):</p><list list-type="order"><list-item><p>RQ1: Are posttraining AI knowledge self-efficacy and job-specific AI competency scores higher than pretraining scores? (Level 2; S4, S6)</p></list-item><list-item><p>RQ2: How satisfied are participants with the training program, and which curricular components are rated highest and lowest? (Level 1; S9)</p></list-item><list-item><p>RQ3: What are participants&#x2019; behavioral intentions to apply learned skills, and where do gaps between knowledge and application emerge? (Level 3 proxy; S11)</p></list-item><list-item><p>RQ4: Do between-group differences in AI knowledge self-efficacy vary by professional group, and what does this imply for participant selection?</p></list-item></list></sec></sec><sec id="s2" sec-type="methods"><title>Methods</title><sec id="s2-1"><title>Study Design</title><p>This was a single-group, pre-post survey study using an independent samples design conducted at a single tertiary academic medical center in Seoul, South Korea. The training program was administered from September to November 2025, with surveys distributed immediately before (pretraining) and after (posttraining) the program. Because survey responses were collected anonymously without personal identifiers, individual-level matching between pre- and posttraining responses was not possible; therefore, pre- and posttraining groups were treated as independent samples. This study adhered to the STROBE (Strengthening the Reporting of Observational Studies in Epidemiology) guidelines for cross-sectional studies [<xref ref-type="bibr" rid="ref28">28</xref>].</p></sec><sec id="s2-2"><title>Participants</title><p>Eligible participants were employees of Asan Medical Center, Ulsan University Hospital, or Gangneung Asan Hospital who enrolled in the advanced generative AI training program. The program was delivered on-site at Asan Medical Center with simultaneous online broadcast to the 2 regional affiliate hospitals, enabling multisite participation. Inclusion criteria were (1) current employment at one of the participating institutions, (2) voluntary enrollment in the program, and (3) age 18 years or older. Participants who attended fewer than 80% of the total training hours were excluded; all 83 enrollees met this threshold. Recruitment did not use an open call; instead, the program was promoted internally within candidate departments through departmental announcements and circulated notices, and interested employees enrolled voluntarily. Because the program&#x2019;s primary aim was to build AI-agent development capability, enrollment was prioritized in the order of IT workforce, health-information staff, clinicians, and researchers, with IT and related departments approached first.</p><p>All 83 enrollees completed the pretraining survey. Of these, 64 responded to the posttraining survey (response rate: 77.1%). A total of 147 survey responses were thus analyzed (pre: n=83; post: n=64). Demographic and occupational characteristics, professional group and years of work experience, were self-reported by participants as part of both the pre- and posttraining surveys, without collection of personal identifiers, and are detailed in the Results section and <xref ref-type="table" rid="table1">Table 1</xref>.</p><table-wrap id="t1" position="float"><label>Table 1.</label><caption><p>Participant demographics.</p></caption><table id="table1" frame="hsides" rules="groups"><thead><tr><td align="left" valign="bottom">Characteristics</td><td align="left" valign="bottom">Pretraining (n=83)</td><td align="left" valign="bottom">Posttraining (n=64)</td><td align="left" valign="bottom"><italic>P</italic> value</td></tr></thead><tbody><tr><td align="left" valign="top" colspan="3">Occupation, n (%)</td><td align="left" valign="top">.87<sup><xref ref-type="table-fn" rid="table1fn1">a</xref></sup></td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>IT specialists</td><td align="left" valign="top">48 (57.8)</td><td align="left" valign="top">33 (51.6)</td><td align="left" valign="top"/></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>HIM<sup><xref ref-type="table-fn" rid="table1fn2">b</xref></sup></td><td align="left" valign="top">10 (12)</td><td align="left" valign="top">9 (14.1)</td><td align="left" valign="top"/></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Clinicians</td><td align="left" valign="top">9 (10.8)</td><td align="left" valign="top">7 (10.9)</td><td align="left" valign="top"/></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Researchers</td><td align="left" valign="top">9 (10.8)</td><td align="left" valign="top">8 (12.5)</td><td align="left" valign="top"/></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Administrative staff</td><td align="left" valign="top">7 (8.4)</td><td align="left" valign="top">7 (10.9)</td><td align="left" valign="top"/></tr><tr><td align="left" valign="top" colspan="3">Work experience, n (%)</td><td align="left" valign="top">.85<sup><xref ref-type="table-fn" rid="table1fn1">a</xref></sup></td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>&#x003C;1 year</td><td align="left" valign="top">3 (3.6)</td><td align="left" valign="top">3 (4.7)</td><td align="left" valign="top"/></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>1&#x2010;5 years</td><td align="left" valign="top">13 (15.7)</td><td align="left" valign="top">12 (18.8)</td><td align="left" valign="top"/></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>6&#x2010;10 years</td><td align="left" valign="top">18 (21.7)</td><td align="left" valign="top">14 (21.9)</td><td align="left" valign="top"/></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>&#x2265;11 years</td><td align="left" valign="top">49 (59.0)</td><td align="left" valign="top">35 (54.7)</td><td align="left" valign="top"/></tr></tbody></table><table-wrap-foot><fn id="table1fn1"><p><sup>a</sup>chi-squared test.</p></fn><fn id="table1fn2"><p><sup>b</sup>HIM: health information managers.</p></fn></table-wrap-foot></table-wrap></sec><sec id="s2-3"><title>Training Program: Design Rationale and Curriculum Structure</title><sec id="s2-3-1"><title>Overview</title><p>The training program was conducted as part of the Medical AI Healthcare Professional Continuing Education Project, a national initiative organized by the Ministry of Health and Welfare and the Korea Human Resource Development Institute for Health &#x0026; Welfare (KOHI). This institutional framework provided the mandate and funding structure for delivering advanced AI training to hospital staff, and situates the present program within a broader national effort to build health care AI workforce capacity.</p></sec><sec id="s2-3-2"><title>Design Philosophy</title><p>Within this framework, the program was designed around a central premise: that MCP and AI agents would become the foundational infrastructure for health care AI use, much as databases became foundational for health informatics. Rather than teaching AI as a collection of isolated tools, the curriculum treated agent-level capabilities, connecting AI models to institutional data sources, orchestrating multistep workflows, and building domain-specific applications, as the core competency target. This premise motivated two key design decisions: (1) allocating the largest share of instructional time to MCP (11.5 h, 37% of instructional time), and (2) requiring all participants to build functional MCP-based prototypes through team capstone projects, on the reasoning that hands-on development experience is the fastest path to genuine understanding and use of AI agent architectures. Relatedly, the inaugural cohort was deliberately composed of staff positioned to enable hospital-wide AI adoption, IT and health-information personnel, related-department clinicians, and researchers, as the first phase of a staged institutional strategy rather than as an end in itself. This IT-first prioritization reflected an assumption that the interface between clinicians and IT professionals will expand substantially in the AI-agent era; building hospital-needs-driven agent-development capability first was expected to lay the foundation for later, clinician-needs-driven AI adoption.</p><p>The program prioritized experiential learning. Hands-on practice comprised 71% of on-site instructional time (22 of 31 h); the remaining 9 hours (29%) were didactic lectures. The subsequent 3-week team capstone added 25 hours of project-based work, so experiential learning dominated the 56-hour program overall (47 of 56 h, 84%). This ratio was informed by self-efficacy theory [<xref ref-type="bibr" rid="ref19">19</xref>], which predicts that mastery experiences&#x2014;direct, successful practice&#x2014;are the most potent source of self-efficacy, and by evidence that experiential learning produces stronger skill-based outcomes than didactic instruction alone [<xref ref-type="bibr" rid="ref23">23</xref>]. Because the cohort spanned a wide range of baseline technical skill&#x2014;from health-information staff who had never written code to experienced IT specialists&#x2014;the curriculum also deliberately minimized AI theory (such as the mathematics of deep learning or the internal mechanics of large language models [LLMs] and agents), emphasizing instead practical terminology, core concepts, the essential agentic methods participants needed, and concrete MCP use cases and development techniques. Within the hands-on sessions, participants worked in 2 parallel tracks matched to baseline skill&#x2014;an application track (assembling solutions with existing frameworks, such as a RAG pipeline in LangChain) and a development track (implementing components directly, such as embedding and similarity search)&#x2014;so that both first-time coders and experienced developers were appropriately challenged.</p><p>To ensure that all participants could engage in hands-on practice without technical barriers, the program provided substantial infrastructure support. Each participant received 3-month funded subscriptions to commercial AI development tools (Cursor IDE [Anysphere], Claude API, and GPT API), removing cost as an obstacle to experimentation. Additionally, the institution deployed a dedicated internal cloud environment where participants could design, test, and run MCP servers using local LLMs within the hospital&#x2019;s secure network, addressing the data governance constraints that typically impede AI development in health care settings.</p></sec><sec id="s2-3-3"><title>Curriculum Structure</title><p>The 8-week, 56-hour program consisted of 5 on-site instructional sessions (31 h) and a 3-week team-based capstone project (25 h; Table S1 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>). The curriculum followed a scaffolded progression designed to build competencies incrementally:</p><list list-type="order"><list-item><p>Week 1 &#x2014; foundations (7 h): foundation models and prompt engineering. Established shared vocabulary and baseline skills across all professional groups.</p></list-item><list-item><p>Week 2 &#x2014; RAG and LangGraph workflow design (7 h): RAG [<xref ref-type="bibr" rid="ref5">5</xref>] and medical document question-answering systems. Introduced the concept of grounding AI outputs in institutional knowledge bases. In the hands-on session, participants built a working RAG pipeline that answered questions over a corpus of institutional clinical documents. The afternoon introduced LangGraph for graph-based clinical-workflow design, modeling a care process (registration &#x2192; triage &#x2192; testing &#x2192; diagnosis &#x2192; prescription) as nodes and edges.</p></list-item><list-item><p>Week 3 &#x2014; MCP and vector databases (7 h): MCP [<xref ref-type="bibr" rid="ref6">6</xref>] fundamentals and vector database integration. As a hands-on exercise, participants built their first MCP server and connected it to a vector database of internal guidelines.</p></list-item><list-item><p>Week 4 &#x2014; advanced agent design (7 h): advanced MCP server development and LangGraph-based AI agent workflows [<xref ref-type="bibr" rid="ref7">7</xref>,<xref ref-type="bibr" rid="ref29">29</xref>]. As capstone preparation, participants designed a multistep, tool-calling clinical workflow agent in LangGraph.</p></list-item><list-item><p>Weeks 5&#x2010;7 &#x2014; team capstone projects (25 h): Self-directed team projects in which participants developed working AI prototypes for real health care challenges. Teams were composed of mixed professional backgrounds (IT specialists, clinicians, HIMs, researchers, and administrative staff) to leverage domain expertise alongside technical skills.</p></list-item><list-item><p>Week 8 &#x2014; integration (3 h): Advanced MCP host development and workflow integration; the program concluded with capstone presentations and an awards ceremony.</p></list-item></list><p>The technology-specific instructional hours, calculated as a percentage of the 31 on-site instructional hours, were distributed as follows: MCP (11.5 h, 37%), LangGraph workflow design (8.5 h, 27%), foundation models and prompt engineering (7 h, 23%), RAG (5 h, 16%), medical Question and Answer (Q&#x0026;A) systems (5 h, 16%), and AI agent design (4 h, 13%). Because some sessions covered more than one technology, their hours were counted in each relevant category, so the listed percentages sum to more than 100%. The heavy MCP allocation reflected the program&#x2019;s core thesis that agent-level infrastructure would be the most transferable and enduring competency.</p><p>A detailed session-level syllabus, including per-session learning objectives, hands-on exercises, and tools for all 8 weeks, is provided in <xref ref-type="supplementary-material" rid="app2">Multimedia Appendix 2</xref>. To support reproducibility while respecting institutional and funding constraints, we release this syllabus and the capstone evaluation rubric (<xref ref-type="supplementary-material" rid="app3">Multimedia Appendix 3</xref>); full scaffold code, proprietary MCP server templates, and internal datasets are not released owing to institutional intellectual-property policy and national-program KOHI constraints.</p></sec><sec id="s2-3-4"><title>Capstone Projects</title><p>All participants were organized into 12 teams, including one team each from Ulsan University Hospital and Gangneung Asan Hospital, each comprising a mix of professional backgrounds. To preserve familiar communication channels while combining expertise, each team was built around an intact IT subdepartment as its technical core, augmented with non-IT domain members (for example, an IT infrastructure team joined by 2 health-information staff and a clinician); this structure was applied consistently across all teams. Each team identified a real-world health care problem within their work context and developed a functional MCP-based AI prototype addressing it over 3 weeks. A single dedicated mentor, selected from the program instructors, provided 3&#x2010;6 hours of structured mentoring per team per week throughout the capstone period, ensuring consistent guidance across all projects. The capstone served dual purposes: providing extended development time for skill consolidation and generating tangible evidence of practical competency beyond self-report measures.</p><p>Capstone projects were evaluated at the closing ceremony (November 11, 2025) by a panel of 7 assessors using a standardized 100-point rubric comprising four weighted domains&#x2014;AI technical use (40 points; 4 items), user experience and user interface (20 points; 2 items), completeness and stability (20 points; 2 items), and clinical applicability (20 points; 2 items)&#x2014;with each item rated on a 1&#x2010;10 scale. To adjust for differences in scoring stringency across assessors, each assessor&#x2019;s raw scores were standardized (z-transformed) within assessor, converted to rank-based scores, and aggregated across the 7 assessors to determine final standings. The full rubric is provided in <xref ref-type="supplementary-material" rid="app3">Multimedia Appendix 3</xref>.</p></sec></sec><sec id="s2-4"><title>Outcome Measures</title><p>Survey items were developed by the research team specifically for this program, grounded in the competency domains of the training curriculum (MCP, RAG, LangGraph, AI agent design, and domain-specific applications) and informed by Compeau and Higgins [<xref ref-type="bibr" rid="ref20">20</xref>] computer self-efficacy scale, adapted to the generative AI domain. The instrument underwent internal review by the program instructors prior to deployment. Post hoc psychometric analyses (Cronbach &#x03B1;, item-total correlations, exploratory factor analysis [EFA]) are reported in Tables S5-S6 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>.</p><p>Survey instruments consisted of 100 pretraining items and 154 posttraining items organized into 11 sections. All attitudinal and self-efficacy items used a 5-point Likert scale (1=strongly disagree to 5=strongly agree). The measures were mapped to the Kirkpatrick 4-level training evaluation model [<xref ref-type="bibr" rid="ref24">24</xref>,<xref ref-type="bibr" rid="ref25">25</xref>] as follows:</p><p>Primary outcomes (pre-post comparison):</p><list list-type="order"><list-item><p>AI knowledge self-efficacy (S4; 8 items): self-assessed competency across 8 AI domains, basic AI understanding, prompt engineering, RAG, local LLM deployment, LangGraph workflow, MCP, medical Q&#x0026;A systems, and AI agent design. This corresponded to Kirkpatrick level 2 (learning).</p></list-item><list-item><p>Job-specific AI competency (S6; 7 items per professional group): self-assessed ability to apply AI tools to job-specific tasks, corresponding to Kirkpatrick level 2. Although each professional group received items tailored to their work context, all items assessed the same underlying construct (perceived ability to apply AI to professional tasks) on an identical response scale, enabling aggregation for overall pre-post comparisons.</p></list-item></list><p>Secondary outcomes (pre-post comparison):</p><list list-type="order"><list-item><p>AI attitudes (S2; 9 items): perceptions toward AI adoption, including apprehension, replacement concerns, and adaptation confidence.</p></list-item><list-item><p>AI daily usage (S3; 2 categorical items): frequency and duration of AI tool use in daily work.</p></list-item><list-item><p>Digital and programming competency (S5; 4 Likert items for digital competency, plus multiselect programming language proficiency): self-assessed digital skills and programming ability. Scale analysis was based on the 4 digital competency items.</p></list-item></list><p>Posttraining only outcomes:</p><list list-type="order"><list-item><p>Training satisfaction (S9; 18 items across 6 domains): curriculum design, content and practice, instructor quality, teaching methods, team project experience, and learning environment, corresponding to Kirkpatrick level 1 (reaction).</p></list-item><list-item><p>Perceived changes (S10; 10 items): self-reported changes in knowledge, skills, attitudes, and workplace application confidence following training.</p></list-item><list-item><p>Workplace application intention (S11; 6 core Likert items + supplementary topic-specific and categorical items): plans for applying learned skills to current work, peer sharing, and continued learning, corresponding to Kirkpatrick level 3 (behavioral intention). S11 measures behavioral intention rather than observed behavioral change; true level 3 assessment would require longitudinal follow-up, which was beyond the scope of this study. Scale analysis was based on the 6 core items.</p></list-item><list-item><p>Self-perceived growth (1 open-ended item): &#x201C;What aspect of change or growth did you most strongly feel through this 8-week program?&#x201D; All 64 responses were read in full by the research team. Recurring themes were identified inductively from the responses rather than applied from an a priori framework, and a set of 5 themes was agreed upon. Each response was then assigned to 1 or 2 of these themes (responses too brief to interpret were set aside as unclassifiable), and the number of responses per theme was tallied. The responses were originally written in Korean; representative quotations were translated into English, and the complete set of translated responses is provided in <xref ref-type="supplementary-material" rid="app4">Multimedia Appendix 4</xref>.</p></list-item></list></sec><sec id="s2-5"><title>Statistical Analysis</title><p>All analyses were performed using Python 3.x (pandas v2.0, SciPy v1.11, factor_analyzer). The anonymous design required treating pre- and posttraining groups as independent samples. Mann-Whitney <italic>U</italic> tests [<xref ref-type="bibr" rid="ref30">30</xref>] were used for pre-post comparisons, with the rank-biserial correlation <italic>r</italic> (= |Z| / &#x221A;N; Rosenthal [<xref ref-type="bibr" rid="ref31">31</xref>]) as the primary effect size, interpreted as small (&#x2265;.10), medium (&#x2265;.30), or large (&#x2265;.50) per Cohen [<xref ref-type="bibr" rid="ref32">32</xref>]. Cohen <italic>d</italic> was additionally reported to facilitate comparison with prior literature. The significance threshold was &#x03B1;=0.05 (2-tailed); posttraining-only measures (S9, S10, and S11) were summarized descriptively. Chi-squared tests assessed demographic comparability between the 2 independent samples.</p><p>Subgroup analyses by professional group were exploratory: Cohen <italic>d</italic> with 95% CIs was computed without formal hypothesis testing or multiple-comparisons correction, given small cell sizes (n=7&#x2010;9 for non-IT groups). Full inferential statistics are available in Table S2 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>.</p><p>Internal consistency was assessed using Cronbach &#x03B1; (acceptable&#x2265;0.70, good&#x2265;0.80 [<xref ref-type="bibr" rid="ref33">33</xref>]). EFA (principal axis factoring, direct oblimin rotation) and corrected item-total correlations provided post hoc construct validity evidence (Tables S5-S6 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>).</p><p>A worst-case sensitivity analysis assigned all 19 nonrespondents&#x2019; posttraining scores equal to their professional group&#x2019;s pretraining mean (Table S7 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>). A post hoc power analysis indicated the minimum detectable effect at 80% power was <italic>r</italic>=.24 (Table S4 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>).</p></sec><sec id="s2-6"><title>Ethical Considerations</title><p>The study protocol was approved by the Institutional Review Board of Asan Medical Center (IRB number 2025&#x2010;1070; approved August 28, 2025). Written informed consent was waived because data were collected anonymously through voluntary online surveys, participation posed no more than minimal risk, and the study could not practicably be conducted without the waiver. The first page of the survey informed participants of the study purpose, voluntary nature, and data handling procedures. Participants received no monetary compensation; program tuition was fully funded through the national continuing-education initiative administered by the KOHI.</p></sec></sec><sec id="s3" sec-type="results"><title>Results</title><sec id="s3-1"><title>Participant Characteristics</title><p>All 83 enrollees completed the program and responded to the pretraining survey; 64 responded to the posttraining survey (response rate: 77.1%; pre: n=83; post: n=64; <xref ref-type="fig" rid="figure1">Figure 1</xref>). Demographic characteristics of the pre- and posttraining groups are presented in <xref ref-type="table" rid="table1">Table 1</xref>. The occupational distribution was comparable between the 2 groups (IT specialists: 57.8%, 48/83 pre vs 51.6%, 33/64 post; HIM: 12%, 10/83, vs 14.1%, 9/64; clinicians: 10.8%, 9/83, vs 10.9%, 7/64; researchers: 10.8%, 9/83, vs 12.5%, 8/64; administrative staff: 8.4%, 7/83, vs 10.9%, 7/64); retention was lower among IT specialists (33/48, 68.8%) than among non-IT groups (78%, 7/9 to 100%, 7/7; see Limitations). The majority of participants in both groups had over 11 years of experience (59%, 49/83 pre vs 54.7%, 35/64 post). Chi-squared tests confirmed no statistically significant differences in occupational (<italic>&#x03C7;</italic>&#x00B2;<sub>4</sub>=1.23; <italic>P</italic>=.87) or experience-level (<italic>&#x03C7;</italic>&#x00B2;<sub>3</sub>=0.80; <italic>P</italic>=.85) distributions between the 2 independent samples.</p><fig position="float" id="figure1"><label>Figure 1.</label><caption><p>Participant flow diagram showing enrollment (n=83), completion of the 8-week training program (n=83), and response rates for the pretraining (n=83; 100%) and posttraining (n=64; 77.1%) surveys. No participants were excluded, as all enrollees met the &#x2265;80% attendance criterion. Because the surveys were anonymous, the pretraining and posttraining groups were analyzed as independent samples.</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="mededu_v12i1e97822_fig01.png"/></fig></sec><sec id="s3-2"><title>Primary Outcomes: AI Knowledge Self-Efficacy (S4)</title><p>Posttraining AI knowledge self-efficacy scores were significantly higher than pretraining scores (mean 3.40, SD 0.69 vs mean 2.34, SD 0.83; <italic>U</italic>=875.5, <italic>r</italic>=.574, <italic>d</italic>=1.37; <italic>P</italic>=3.50&#x00D7;10<sup>&#x2013;</sup>&#x00B9;&#x00B2;; <xref ref-type="table" rid="table2">Table 2</xref>, <xref ref-type="fig" rid="figure2">Figure 2</xref>). Internal consistency was excellent (&#x03B1;=0.909 pre, &#x03B1;=0.915 post; Tables S5-S6 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>). At the domain level, the technologies that received the greatest curricular investment and had the lowest baseline scores showed the largest between-group differences: MCP (<italic>d</italic>=1.57), medical Q&#x0026;A systems (<italic>d</italic>=1.39), and AI agent design (<italic>d</italic>=1.39), whereas foundational domains with higher baselines showed smaller differences (basic AI understanding: <italic>d</italic>=0.55; <xref ref-type="table" rid="table2">Table 2</xref>).</p><table-wrap id="t2" position="float"><label>Table 2.</label><caption><p>Pre-Post comparison of primary and secondary outcomes. r=rank-biserial correlation (Rosenthal formula); significance threshold: &#x03B1;=0.05.</p></caption><table id="table2" frame="hsides" rules="groups"><thead><tr><td align="left" valign="bottom">Measures</td><td align="left" valign="bottom">Pre mean (SD)</td><td align="left" valign="bottom">Post mean (SD)</td><td align="left" valign="bottom"><italic>U</italic></td><td align="left" valign="bottom"><italic>Z</italic></td><td align="left" valign="bottom"><italic>P</italic></td><td align="left" valign="bottom"><italic>r</italic></td><td align="left" valign="bottom">Cohen <italic>d</italic></td><td align="left" valign="bottom">Interpretation</td></tr></thead><tbody><tr><td align="left" valign="top">S4: AI knowledge self-efficacy</td><td align="left" valign="top">2.34 (0.83)</td><td align="left" valign="top">3.40 (0.69)</td><td align="left" valign="top">875.5</td><td align="left" valign="top">&#x2013;6.956</td><td align="left" valign="top">3.50&#x00D7;10<sup>&#x2013;</sup>&#x00B9;&#x00B2;</td><td align="left" valign="top">.574</td><td align="left" valign="top">1.37</td><td align="left" valign="top">Large</td></tr><tr><td align="left" valign="top">S6: Job-specific AI competency</td><td align="left" valign="top">1.95 (1.00)</td><td align="left" valign="top">2.77 (0.95)</td><td align="left" valign="top">1470</td><td align="left" valign="top">&#x2013;4.634</td><td align="left" valign="top">3.59&#x00D7;10<sup>&#x2013;</sup>&#x2076;</td><td align="left" valign="top">.382</td><td align="left" valign="top">0.85</td><td align="left" valign="top">Medium</td></tr><tr><td align="left" valign="top">S2: AI attitudes</td><td align="left" valign="top">3.79 (0.41)</td><td align="left" valign="top">3.77 (0.51)</td><td align="left" valign="top">2733</td><td align="left" valign="top">&#x2013;0.301</td><td align="left" valign="top">.764</td><td align="left" valign="top">.025</td><td align="left" valign="top">&#x2013;0.04</td><td align="left" valign="top">Negligible</td></tr><tr><td align="left" valign="top">S5: Digital competency</td><td align="left" valign="top">3.38 (0.89)</td><td align="left" valign="top">3.63 (0.82)</td><td align="left" valign="top">2228.5</td><td align="left" valign="top">&#x2013;1.675</td><td align="left" valign="top">.094</td><td align="left" valign="top">.138</td><td align="left" valign="top">0.29</td><td align="left" valign="top">Small</td></tr></tbody></table></table-wrap><fig position="float" id="figure2"><label>Figure 2.</label><caption><p>Comparison of self-assessed AI competency scores across eight skill domains before (n=83) and after (n=64) the training program, with Cohen <italic>d</italic> effect sizes and Mann-Whitney <italic>U</italic> test significance levels. MCP: model context protocol; LLM: large-language model; RAG: retrieval-augmented generation.</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="mededu_v12i1e97822_fig02.png"/></fig></sec><sec id="s3-3"><title>Subgroup Analysis by Professional Group</title><p>All 5 professional groups showed the largest between-group difference in MCP, and non-IT groups consistently showed larger differences than IT specialists (Table S2 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>). For MCP, clinicians showed the largest difference (<italic>d</italic>=2.88, 95% CI 1.42-4.33), followed by HIM (<italic>d</italic>=2.45, 95% CI 1.23-3.67), researchers (<italic>d</italic>=2.07, 95% CI 0.87-3.28), and administrative staff (<italic>d</italic>=1.83, 95% CI 0.55-3.11); IT specialists showed the smallest difference (<italic>d</italic>=1.36, 95% CI 0.87-1.85). These subgroup findings should be considered exploratory given small cell sizes (n=7&#x2010;9 for non-IT groups; Table S2 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>).</p></sec><sec id="s3-4"><title>Secondary Outcomes and the Knowledge-Practice Gap</title><p>Job-specific AI competency (S6) scores were significantly higher posttraining (mean 2.77, SD 0.95 vs mean 1.95, SD 1; <italic>P</italic>=3.59&#x00D7;10<sup>&#x2013;</sup>&#x2076;, <italic>r</italic>=.382, <italic>d</italic>=0.85; &#x03B1;=0.938-0.973 pre, &#x03B1;=0.877-0.981 post; <xref ref-type="table" rid="table2">Table 2</xref>). However, unlike general AI self-efficacy (S4), job-specific competency remained below the scale midpoint (2.77/5.0), indicating a knowledge-practice gap.</p><p>Three secondary pre-post comparisons did not reach significance (<xref ref-type="table" rid="table2">Table 2</xref>; Table S4 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>): AI attitudes (S2: <italic>r</italic>=.025; <italic>P</italic>=.76), with the mean of the 6 positively worded items (Q4-Q9) already at ceiling at baseline (M=4.29/5); the full 9-item S2 scale mean, with reverse coding applied to Q1-Q3, was 3.79 (<xref ref-type="table" rid="table2">Table 2</xref>); digital competency (S5: <italic>r</italic>=.138; <italic>P</italic>=.09); and AI daily usage (S3), where the proportion using 3+ hours daily increased modestly from 9.6% to 17.2%.</p></sec><sec id="s3-5"><title>Training Satisfaction (S9; Kirkpatrick Level 1)</title><p>Overall training satisfaction was 4.03/5.0 (SD 0.85), while the domain-averaged score was 3.48 (SD 0.63; &#x03B1;=0.939; <xref ref-type="table" rid="table3">Table 3</xref>, <xref ref-type="fig" rid="figure3">Figure 3</xref>). Satisfaction varied substantially across domains: team project experience was rated highest (mean 3.98, SD 0.84), followed by instructor quality (mean 3.91, SD 0.72), teaching methods (mean 3.58, SD 0.69), and content &#x0026; practice (mean 3.40, SD 0.67). Curriculum design was rated lowest (mean 2.95, SD 0.87), particularly learning pace (mean 2.91, SD 1.11). Item-level results for all posttraining scales (S9-S11) are reported in Table S3 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>.</p><table-wrap id="t3" position="float"><label>Table 3.</label><caption><p>Training satisfaction by domain (S9).</p></caption><table id="table3" frame="hsides" rules="groups"><thead><tr><td align="left" valign="bottom">Domain</td><td align="left" valign="bottom">Mean (SD)</td><td align="left" valign="bottom">Cronbach &#x03B1;</td></tr></thead><tbody><tr><td align="left" valign="top">Team project experience</td><td align="left" valign="top">3.98 (0.84)</td><td align="left" valign="top">0.789</td></tr><tr><td align="left" valign="top">Instructor quality</td><td align="left" valign="top">3.91 (0.72)</td><td align="left" valign="top">0.857</td></tr><tr><td align="left" valign="top">Teaching methods</td><td align="left" valign="top">3.58 (0.69)</td><td align="left" valign="top">0.833</td></tr><tr><td align="left" valign="top">Content and practice</td><td align="left" valign="top">3.40 (0.67)</td><td align="left" valign="top">0.845</td></tr><tr><td align="left" valign="top">Learning environment</td><td align="left" valign="top">3.07 (1.06)</td><td align="left" valign="top">0.912</td></tr><tr><td align="left" valign="top">Curriculum design</td><td align="left" valign="top">2.95 (0.87)</td><td align="left" valign="top">0.867</td></tr><tr><td align="left" valign="top">Overall (single item)</td><td align="left" valign="top">4.03 (0.85)</td><td align="left" valign="top">&#x2014;<sup><xref ref-type="table-fn" rid="table3fn1">a</xref></sup></td></tr><tr><td align="left" valign="top">Domain average</td><td align="left" valign="top">3.48 (0.63)</td><td align="left" valign="top">0.939</td></tr></tbody></table><table-wrap-foot><fn id="table3fn1"><p><sup>a</sup>Not available.</p></fn></table-wrap-foot></table-wrap><fig position="float" id="figure3"><label>Figure 3.</label><caption><p>Radar chart of program satisfaction ratings across seven evaluation dimensions (N=64; overall Cronbach &#x03B1;=0.939, 18 items), where the dashed orange line represents the overall mean (3.48) and the gray dashed circle indicates the scale midpoint (3).</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="mededu_v12i1e97822_fig03.png"/></fig></sec><sec id="s3-6"><title>Capstone Outcomes and Real-World Deployment</title><p>Of the 12 capstone teams, 11 presented functional AI prototypes spanning clinical, administrative, and IT applications (representative examples in <xref ref-type="table" rid="table4">Table 4</xref>; full list in Table S8 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>); one regional affiliate team was unable to present due to scheduling constraints. One project, a medical dictionary-based RAG/query chatbot system (&#x201C;WorksBot&#x201D;), received the Minister of Health and Welfare Award. One prototype, the hospital document knowledge retrieval RAG system, has since been developed into an institutional service (&#x201C;AI-Docs&#x201D;) and has been in active pilot use since May 2026 across multiple departments, including the Department of Nursing and the Big Data Research Center, with hospital-wide deployment planned within 2026. WorksBot remains a prototype and is not yet in operational service; a refined version is planned following the RAG pilot. Both efforts are being advanced by the Department of Digital Innovation and Support.</p><table-wrap id="t4" position="float"><label>Table 4.</label><caption><p>Representative capstone projects. Representative projects illustrating the range of clinical, operational, and IT applications; the complete list of all 12 capstone projects is provided in Table S8 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>.</p></caption><table id="table4" frame="hsides" rules="groups"><thead><tr><td align="left" valign="bottom">Projects</td><td align="left" valign="bottom">Domain</td><td align="left" valign="bottom">Core technologies</td><td align="left" valign="bottom">Status</td></tr></thead><tbody><tr><td align="left" valign="top">WorksBot (medical-dictionary RAG<sup><xref ref-type="table-fn" rid="table4fn1">a</xref></sup>/ or Q&#x0026;A)<sup><xref ref-type="table-fn" rid="table4fn2">b</xref></sup></td><td align="left" valign="top">Clinical knowledge access</td><td align="left" valign="top">MCP<sup><xref ref-type="table-fn" rid="table4fn3">c</xref></sup>, RAG, vector DB<sup><xref ref-type="table-fn" rid="table4fn4">d</xref></sup></td><td align="left" valign="top">Top-ranked; prototype</td></tr><tr><td align="left" valign="top">AI-Docs (document-retrieval RAG)</td><td align="left" valign="top">Hospital operations</td><td align="left" valign="top">MCP, RAG, vector DB</td><td align="left" valign="top">Active pilot; hospital-wide planned 2026</td></tr><tr><td align="left" valign="top">Natural-language clinical-data retrieval</td><td align="left" valign="top">Clinical</td><td align="left" valign="top">MCP, LangGraph</td><td align="left" valign="top">Prototype</td></tr><tr><td align="left" valign="top">Medical-record de-identification</td><td align="left" valign="top">Data governance or IT</td><td align="left" valign="top">MCP, local LLM<sup><xref ref-type="table-fn" rid="table4fn5">e</xref></sup></td><td align="left" valign="top">Prototype</td></tr><tr><td align="left" valign="top">AGS<sup><xref ref-type="table-fn" rid="table4fn6">f</xref></sup> operational-guideline Q&#x0026;A chatbot</td><td align="left" valign="top">Hospital operations</td><td align="left" valign="top">MCP, RAG, local LLM</td><td align="left" valign="top">Prototype</td></tr></tbody></table><table-wrap-foot><fn id="table4fn1"><p><sup>a</sup>RAG: retrieval-augmented generation. </p></fn><fn id="table4fn2"><p><sup>b</sup>Q&#x0026;A: question and answer.</p></fn><fn id="table4fn3"><p><sup>c</sup>MCP: model context protocol.</p></fn><fn id="table4fn4"><p><sup>d</sup>Vector DB: vector database.</p></fn><fn id="table4fn5"><p><sup>e</sup>LLM: large-language model.</p></fn><fn id="table4fn6"><p><sup>f</sup>AGS: Asan Global Standard, the institution's internal operational guideline system.</p></fn></table-wrap-foot></table-wrap></sec><sec id="s3-7"><title>Posttraining Outcomes (S10, S11; Kirkpatrick Levels 2-3)</title><p>Perceived changes (S10; &#x03B1;=0.928; <xref ref-type="fig" rid="figure4">Figure 4</xref>) showed high perceived knowledge improvement (mean 4.27, SD 0.62) and technology acceptance (mean 4.20, SD 0.62), but lower workplace application confidence (mean 3.50, SD 0.84), reinforcing the knowledge-practice gap observed in S6. Workplace application intention (S11; &#x03B1;=0.909; <xref ref-type="fig" rid="figure5">Figure 5</xref>) was high overall (mean 4.02, SD 0.70), with continued learning rated highest (mean 4.39, SD 0.75) and novel problem solving rated lowest (mean 3.78, SD 0.95).</p><fig position="float" id="figure4"><label>Figure 4.</label><caption><p>Self-perceived educational effectiveness across 10 items categorized into knowledge, skills, attitudes, and application domains (N=64; Cronbach &#x03B1;=0.928, 10 items), with an overall mean of 3.90 and error bars representing standard deviations.</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="mededu_v12i1e97822_fig04.png"/></fig><fig position="float" id="figure5"><label>Figure 5.</label><caption><p>Transfer of learning intentions across six items measured on a 5-point Likert scale (N=64; Cronbach &#x03B1;=0.909, 6 items), with the dashed line indicating the overall mean benchmark of 4.0. Error bars represent SD.</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="mededu_v12i1e97822_fig05.png"/></fig></sec><sec id="s3-8"><title>Participants&#x2019; Self-Reported Growth (Open-Ended Responses)</title><p>All 64 posttraining respondents answered the open-ended question: &#x201C;What aspect of change or growth did you most strongly feel through this 8-week program?&#x201D; Eight responses (12.5%) were too brief to classify; the remaining 56 responses were grouped by the research team into 5 recurring themes. Because 8 of these responses expressed 2 distinct themes (none was assigned to more than 2), this grouping yielded 64 theme-level assignments in total. The most prevalent theme was conceptual clarity (24 assignments, 37.5% of all 64 respondents), with participants reporting that previously abstract concepts, particularly MCP and RAG, became concrete and actionable. The second was confidence and reduced fear (14 assignments, 21.9%); participants described a shift from anxiety about AI to self-efficacy, with representative responses including &#x201C;the vague fear I had has been resolved to some extent&#x201D; and &#x201C;my first goal was confidence recovery, and I fully achieved it.&#x201D; Third, hands-on tool experience (12 assignments, 18.8%) was frequently cited, with participants specifically naming vibe coding, Cursor, and Claude Code as transformative experiences; notably, several nondevelopers described coding for the first time. Fourth, workplace application ideation (10 assignments, 15.6%) emerged, with participants reporting that they began identifying specific AI applications in their own work contexts. Finally, a smaller group reported a general attitude shift (4 assignments, 6.2%).</p></sec></sec><sec id="s4" sec-type="discussion"><title>Discussion</title><sec id="s4-1"><title>Principal Findings</title><p>This study describes the design, implementation, and outcomes of an 8-week intensive generative AI training program for a multidisciplinary hospital workforce. To our knowledge, this is the first study to report both the design and the pre-post evaluation of a hospital-workforce program integrating 4 advanced, agent-level generative AI technologies, MCP, RAG, LangGraph orchestration, and AI agent design, across both IT and non-IT roles, spanning IT specialists, clinicians, HIM, researchers, and administrative staff. This advances AI education research beyond its prevailing focus on AI literacy and prompt engineering [<xref ref-type="bibr" rid="ref10">10</xref>-<xref ref-type="bibr" rid="ref13">13</xref>] toward agent-level competencies, and, through the capstone-to-deployment pathway, links classroom learning to institutional adoption. Five principal lessons emerged from the convergence of quantitative survey data and participants&#x2019; open-ended reflections.</p></sec><sec id="s4-2"><title>Lesson 1: Center the Curriculum on Agent-Level Technologies&#x2014;The Unfamiliar Yields the Greatest Returns</title><p>The program&#x2019;s most consequential design decision was allocating 37% of instructional time to MCP, a technology virtually none of the participants had encountered. The quantitative data validated this bet: MCP showed the largest between-group self-efficacy difference across all professional groups (<italic>d</italic>=1.57; <xref ref-type="table" rid="table2">Table 2</xref>), and the 3 novel domains (MCP, medical Q&#x0026;A, and AI agent design) consistently outperformed foundational topics. The open-ended responses corroborated this pattern&#x2014;MCP and RAG were the most frequently named technologies when participants described their growth, and the dominant theme was that &#x201C;previously abstract concepts became concrete and actionable.&#x201D; This inverse relationship between baseline familiarity and learning yield is consistent with self-efficacy theory&#x2019;s prediction that mastery experiences are most impactful where prior competency is lowest [<xref ref-type="bibr" rid="ref19">19</xref>], and suggests that training programs should resist the temptation to stay at the prompt-engineering level.</p></sec><sec id="s4-3"><title>Lesson 2: Include Non-IT Professionals&#x2014;They Benefit Most and Bring Irreplaceable Domain Expertise</title><p>Non-IT professionals showed consistently larger between-group differences than IT specialists across most domains (Table S2 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>), particularly for MCP (clinicians: <italic>d</italic>=2.88; HIM: <italic>d</italic>=2.45). This finding challenges the assumption that advanced AI training should be restricted to technical staff. Qualitatively, nondevelopers described transformative experiences: several reported coding for the first time through vibe coding, and one described the capstone as an opportunity to &#x201C;bring MCP and RAG into hospital work and implement visible results through hands-on practice&#x2014;a fundamentally different experience from watching lecture slides.&#x201D; The mixed-expertise team structure, which received the highest satisfaction rating (mean 3.98, SD 0.84), provided the technical scaffolding non-IT professionals needed while leveraging their domain knowledge for health care-specific applications.</p></sec><sec id="s4-4"><title>Lesson 3: Expect a Knowledge-Practice Gap&#x2014;and Design for It</title><p>Despite significant self-efficacy gains (<italic>r</italic>=.574) and high perceived knowledge improvement (mean 4.27, SD 0.62), job-specific competency remained below the scale midpoint (2.77/5) and workplace application confidence was rated only 3.50&#x2014;a gap that converged across 3 independent measures (S6, S10, and S11). Yet, the open-ended responses revealed a more nuanced picture: 15.6% of respondents spontaneously reported generating specific workplace application ideas, and participants described shifts from passive understanding to active problem identification (&#x201C;I started thinking about improvements to the systems we currently operate&#x201D;). The Unified Theory of Acceptance and Use of Technology (UTAUT) framework [<xref ref-type="bibr" rid="ref22">22</xref>] suggests that training without corresponding organizational change management may produce motivated but unsupported learners [<xref ref-type="bibr" rid="ref34">34</xref>,<xref ref-type="bibr" rid="ref35">35</xref>]; health care settings compound this through restricted electronic medical record (EMR) environments, data governance requirements, and concerns about LLM hallucination. The gap is not a failure of training but a predictable transition challenge&#x2014;one that infrastructure support (funded tool subscriptions, secure internal cloud environments) and structured posttraining mentoring can help bridge [<xref ref-type="bibr" rid="ref36">36</xref>,<xref ref-type="bibr" rid="ref37">37</xref>].</p></sec><sec id="s4-5"><title>Lesson 4: Plan for Differentiated Pacing Across Heterogeneous Cohorts</title><p>Curriculum pacing was rated lowest among satisfaction domains (mean 2.91/5.0, SD 1.11), even as overall satisfaction was substantially higher (4.03/5.0). This signals that participants valued the program but struggled to keep pace&#x2014;a predictable tension when teaching rapidly evolving technologies to professionals ranging from experienced developers to clinical staff encountering code for the first time. The qualitative data illuminated both sides: some participants celebrated &#x201C;coding fearlessly for the first time,&#x201D; while the pacing score suggests others felt left behind. In practice, hands-on exercises were already split into application and development tracks matched to baseline skill; future iterations should additionally consider prerequisite assessment and supplementary catch-up sessions to accommodate varying baseline competencies.</p><p>Teaching a mixed-expertise cohort posed distinct challenges. IT and health-information staff were released from their regular duties to attend, whereas some clinicians and researchers used personal annual leave. At the outset, the wide range in baseline skill, from staff who had never coded to experienced IT specialists, was anticipated as a key challenge; 2 design choices mitigated it (Methods): the curriculum minimized AI theory in favor of practical concepts, agentic methods, and concrete MCP use cases, and each capstone team paired an intact IT-subdepartment core with non-IT domain members. Collaboration across backgrounds still required deliberate team-level scaffolding, and recruitment skewed toward IT staff, consistent with the program&#x2019;s IT-first prioritization and department-based, voluntary enrollment&#x2014;with comparatively fewer clinical participants.</p></sec><sec id="s4-6"><title>Lesson 5: Capstone Projects Bridge Instruction and Institutional Adoption</title><p>The capstone component, supported by 3&#x2010;6 hours of weekly mentoring per team from a dedicated instructor, culminated in functional prototypes from 11 of 12 teams. The qualitative data revealed why: participants described the capstone as where abstract knowledge &#x201C;became real&#x201D;&#x2014;the transition from understanding concepts to building something that works. More importantly, this bridge extended beyond the classroom: one capstone project has since advanced into institutional pilot use (AI-Docs; see Results). This progression from capstone prototype into pilot clinical use provides emerging Kirkpatrick level 3 evidence that training-generated projects can move into institutional adoption, suggesting the capstone model can serve as a pathway for AI innovation that outlasts the training period itself.</p><p>This program was conceived as the deliberate first phase of a staged, institution-wide strategy: we began with staff positioned to enable broad adoption (IT, health-information, related-department clinicians, and researchers), expecting their agent-level competency to seed hospital-wide AI deployment&#x2014;a trajectory already visible in the AI-Docs pilot. Subsequent phases are planned to extend structured training to frontline clinicians and other departments and to embed the curriculum as a recurring continuing professional development (CPD) track, a required or elective module, or a template for other hospitals through the national framework that funded it. So framed, this voluntary first cohort is a starting point, not an endpoint: audiences not yet aware of agent-level AI are reached through later, role-targeted phases rather than left behind.</p></sec><sec id="s4-7"><title>Comparison With Prior Work</title><p>The observed between-group difference in AI knowledge self-efficacy (<italic>r</italic>=.574) is larger than improvements typically reported in prior digital health education studies, which have largely relied on uncontrolled before-after designs and reported heterogeneous, generally smaller gains [<xref ref-type="bibr" rid="ref38">38</xref>]. The larger effects likely reflect the intensive, hands-on design (71% of on-site time practical; 84% experiential including the capstone), consistent with experiential learning evidence [<xref ref-type="bibr" rid="ref19">19</xref>,<xref ref-type="bibr" rid="ref23">23</xref>]. Recent ChatGPT education studies for health care professionals report similar directions but with shorter interventions and smaller effects [<xref ref-type="bibr" rid="ref11">11</xref>,<xref ref-type="bibr" rid="ref12">12</xref>]; our previous health informatics analyst education program at the same institution [<xref ref-type="bibr" rid="ref18">18</xref>] provides a direct comparator, and the present study extends that evidence to advanced generative AI technologies.</p><p>The absence of significant attitudinal change (S2; <italic>P</italic>=.76) aligns with ceiling effects in health care AI attitude studies [<xref ref-type="bibr" rid="ref13">13</xref>,<xref ref-type="bibr" rid="ref39">39</xref>], where voluntary enrollees already hold favorable views (4.29/5). Notably, however, the qualitative data captured attitudinal shifts that the Likert scale missed: 21.9% of respondents spontaneously described reduced fear or increased confidence as their most significant growth&#x2014;suggesting that self-efficacy gains may be a more sensitive indicator of training impact than attitudinal scales in already AI-positive populations. Future studies should recruit participants with more heterogeneous baseline attitudes [<xref ref-type="bibr" rid="ref40">40</xref>] and consider AI adoption readiness [<xref ref-type="bibr" rid="ref21">21</xref>] as an alternative construct.</p></sec><sec id="s4-8"><title>Limitations</title><p>This study has several important limitations that should be considered when interpreting the findings.</p><sec id="s4-8-1"><title>Study Design Constraints</title><p>The independent samples design, necessitated by anonymous data collection, prevents individual-level tracking, precluding within-person effect sizes; future studies should consider self-generated identification codes to enable paired analysis [<xref ref-type="bibr" rid="ref41">41</xref>]. The absence of a control group precludes causal attribution; improvements may partly reflect maturation or concurrent self-study. Additionally, the inverse relationship between baseline scores and improvement magnitude may partly reflect regression to the mean, particularly in this design where groups are not matched at the individual level. Subgroup comparisons by professional group also warrant caution beyond their exploratory, small-cell nature, because professional role was confounded with both baseline competency and career stage. The groups with the lowest baseline self-efficacy&#x2014;HIM (pre mean 1.74, SD 0.44) and clinicians (1.92)&#x2014;showed the largest gains, consistent with regression to the mean and ceiling effects, whereas researchers, despite being the most early-career group (none with &#x2265;11 years of experience), reported the highest baseline and posttraining scores. Professional roles were, moreover, self-classified. The larger gains observed for non-IT groups therefore cannot be attributed to professional background independent of baseline ability and experience.</p></sec><sec id="s4-8-2"><title>Attrition and Selection Bias</title><p>Of 83 enrollees, 19 (22.9%) did not respond to the posttraining survey, and these individuals may have differed systematically from respondents. A worst-case sensitivity analysis demonstrated that primary outcomes remain statistically significant even assuming zero training effect in all nonrespondents (Table S7 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>). Nevertheless, self-selection bias limits generalizability: voluntary enrollees at affiliated institutions with high baseline positive-attitude scores (4.29/5) may not represent the broader health care workforce. Importantly, posttraining nonresponse was not evenly distributed across professional groups: retention was lowest among IT specialists (33 of 48; 68.8%) and higher among non-IT groups (78% &#x2010;100%). Because the subgroup analyses contrast non-IT groups against IT specialists, this differential attrition may modestly inflate the apparent non-IT advantage, and the comparable cross-sectional occupational distributions should be interpreted with caution given limited statistical power.</p></sec><sec id="s4-8-3"><title>Measurement Limitations</title><p>All outcomes relied on self-report measures subject to social desirability bias [<xref ref-type="bibr" rid="ref42">42</xref>]. Self-efficacy assessments administered immediately after training may be inflated by posttraining euphoria, and the Dunning-Kruger effect is particularly relevant for novel technology domains where participants may lack sufficient expertise to calibrate their self-assessment accurately. Future studies should complement self-efficacy measures with objective performance assessments [<xref ref-type="bibr" rid="ref26">26</xref>]. The survey instrument was developed without formal content validity index (CVI) assessment or pilot testing, limiting construct validity evidence beyond the post hoc psychometric analyses reported in Tables S5-S6 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>; future iterations should establish CVI &#x2265;.80 via a priori expert panel review.</p></sec><sec id="s4-8-4"><title>Generalizability and Temporal Scope</title><p>Despite multisite delivery, participants were drawn primarily from a single tertiary academic medical center, limiting external validity. Short-term assessment does not capture whether gains persist; longitudinal follow-up at 3&#x2010;6 months is needed to assess durability and actual behavioral transfer [<xref ref-type="bibr" rid="ref43">43</xref>].</p></sec></sec><sec id="s4-9"><title>Conclusions</title><p>This study suggests that transforming a multidisciplinary hospital workforce into AI-capable professionals is achievable through intensive, hands-on training centered on agent-level technologies. The curriculum&#x2019;s investment in MCP, the technology anticipated to underpin future health care AI infrastructure, was associated with the largest self-efficacy gains, and non-IT professionals showed the largest between-group differences, challenging the assumption that advanced AI training should be restricted to technical staff. The team-based capstone model proved particularly valuable, with prototypes that progressed toward institutional adoption, providing early evidence that training programs can serve as a pathway for institutional AI innovation. However, the knowledge-practice gap identified across multiple measures underscores that training alone is insufficient; posttraining support structures, mentoring, communities of practice, and supervised implementation are essential to translate self-efficacy gains into sustained workplace practice. Future research should use longitudinal designs with matched cohorts to assess the durability of training effects and to determine whether capstone-to-deployment pathways can be replicated across institutions.</p></sec></sec></body><back><ack><p>The authors would like to express their sincere gratitude to the Asan Medical Center Big Data Research Center, the Department of Digital Innovation and Support, and the Asan Academic Institute for their support in operating the educational program, and to the Korea Human Resource Development Institute for Health &#x0026; Welfare for their support as the supervising organization of this educational project. The authors also thank all participants who took part in the educational program, including IT professionals, medical record administrators, health care professionals, researchers, and administrative or clerical staff.</p><p>During the preparation of this manuscript, the authors used Claude (Anthropic; Claude Opus model) to assist with English-language editing and phrasing, to support drafting of the manuscript text, and to help organize the references and formatting. The tool was not used to generate, analyze, or interpret the study data or results. After using this tool, the authors reviewed and edited all content as needed and take full responsibility for the content of the publication.</p></ack><notes><sec><title>Funding</title><p>This research was supported by (1) a grant of the Research-Centered Hospital Development R&#x0026;D Project through the Korea Health Industry Development Institute (KHIDI), funded by the Ministry of Health and Welfare, Republic of Korea (grant number: HR21C0198); and (2) a grant of the Korea Health Technology R&#x0026;D Project, funded by the Ministry of Health and Welfare, Republic of Korea (grant number: RS-2025&#x2010;02213531).</p></sec><sec><title>Data Availability</title><p>The datasets generated during this study are available from the corresponding author upon reasonable request.</p></sec></notes><fn-group><fn fn-type="con"><p>GKB, Kye Hwa Lee, and DHY contributed to the conceptualization. GKB and HNL contributed to data curation. HNL, YRL, and GKB contributed to the formal analysis. MSK, Kun Hee Lee, and MJC contributed to data acquisition. GKB and Kye Hwa Lee contributed to writing&#x2014;original draft. Kye Hwa Lee and GKB contributed to writing&#x2014;review and editing.</p></fn><fn fn-type="conflict"><p>None declared.</p></fn></fn-group><glossary><title>Abbreviations</title><def-list><def-item><term id="abb1">BEME</term><def><p>Best Evidence Medical Education</p></def></def-item><def-item><term id="abb2">CPD</term><def><p>continuing professional development</p></def></def-item><def-item><term id="abb3">CVI</term><def><p>content validity index</p></def></def-item><def-item><term id="abb4">EFA</term><def><p>exploratory factor analysis</p></def></def-item><def-item><term id="abb5">EMR</term><def><p>electronic medical record</p></def></def-item><def-item><term id="abb6">HIM</term><def><p>health information manager</p></def></def-item><def-item><term id="abb7">KOHI</term><def><p>Korea Human Resource Development Institute for Health &#x0026; Welfare</p></def></def-item><def-item><term id="abb8">LLM</term><def><p>large language model</p></def></def-item><def-item><term id="abb9">MCP</term><def><p>model context protocol</p></def></def-item><def-item><term id="abb10">Q&#x0026;A</term><def><p>question and answer</p></def></def-item><def-item><term id="abb11">RAG</term><def><p>retrieval-augmented generation</p></def></def-item><def-item><term id="abb12">RQ</term><def><p>research question</p></def></def-item><def-item><term id="abb13">STROBE</term><def><p>Strengthening the Reporting of Observational Studies in Epidemiology</p></def></def-item><def-item><term id="abb14">UTAUT</term><def><p>Unified Theory of Acceptance and Use of Technology</p></def></def-item></def-list></glossary><ref-list><title>References</title><ref id="ref1"><label>1</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Seo</surname><given-names>J</given-names> </name><name name-style="western"><surname>Choi</surname><given-names>D</given-names> </name><name name-style="western"><surname>Kim</surname><given-names>T</given-names> </name><etal/></person-group><article-title>Evaluation framework of large language models in medical documentation: development and usability study</article-title><source>J Med Internet Res</source><year>2024</year><month>11</month><day>20</day><volume>26</volume><fpage>e58329</fpage><pub-id pub-id-type="doi">10.2196/58329</pub-id><pub-id pub-id-type="medline">39566044</pub-id></nlm-citation></ref><ref id="ref2"><label>2</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Topol</surname><given-names>EJ</given-names> </name></person-group><article-title>High-performance medicine: the convergence of human and artificial intelligence</article-title><source>Nat Med</source><year>2019</year><month>01</month><volume>25</volume><issue>1</issue><fpage>44</fpage><lpage>56</lpage><pub-id pub-id-type="doi">10.1038/s41591-018-0300-7</pub-id><pub-id pub-id-type="medline">30617339</pub-id></nlm-citation></ref><ref id="ref3"><label>3</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Thirunavukarasu</surname><given-names>AJ</given-names> </name><name name-style="western"><surname>Ting</surname><given-names>DSJ</given-names> </name><name name-style="western"><surname>Elangovan</surname><given-names>K</given-names> </name><name name-style="western"><surname>Gutierrez</surname><given-names>L</given-names> </name><name name-style="western"><surname>Tan</surname><given-names>TF</given-names> </name><name name-style="western"><surname>Ting</surname><given-names>DSW</given-names> </name></person-group><article-title>Large language models in medicine</article-title><source>Nat Med</source><year>2023</year><month>08</month><volume>29</volume><issue>8</issue><fpage>1930</fpage><lpage>1940</lpage><pub-id pub-id-type="doi">10.1038/s41591-023-02448-8</pub-id><pub-id pub-id-type="medline">37460753</pub-id></nlm-citation></ref><ref id="ref4"><label>4</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>He</surname><given-names>J</given-names> </name><name name-style="western"><surname>Baxter</surname><given-names>SL</given-names> </name><name name-style="western"><surname>Xu</surname><given-names>J</given-names> </name><name name-style="western"><surname>Xu</surname><given-names>J</given-names> </name><name name-style="western"><surname>Zhou</surname><given-names>X</given-names> </name><name name-style="western"><surname>Zhang</surname><given-names>K</given-names> </name></person-group><article-title>The practical implementation of artificial intelligence technologies in medicine</article-title><source>Nat Med</source><year>2019</year><month>01</month><volume>25</volume><issue>1</issue><fpage>30</fpage><lpage>36</lpage><pub-id pub-id-type="doi">10.1038/s41591-018-0307-0</pub-id><pub-id pub-id-type="medline">30617336</pub-id></nlm-citation></ref><ref id="ref5"><label>5</label><nlm-citation citation-type="confproc"><person-group person-group-type="author"><name name-style="western"><surname>Lewis</surname><given-names>P</given-names> </name><name name-style="western"><surname>Perez</surname><given-names>E</given-names> </name><name name-style="western"><surname>Piktus</surname><given-names>A</given-names> </name><name name-style="western"><surname>Petroni</surname><given-names>F</given-names> </name><name name-style="western"><surname>Karpukhin</surname><given-names>V</given-names> </name><name name-style="western"><surname>Goyal</surname><given-names>N</given-names> </name><etal/></person-group><article-title>Retrieval-augmented generation for knowledge-intensive NLP tasks</article-title><year>2020</year><access-date>2026-07-21</access-date><conf-name>34th Conference on Neural Information Processing Systems (NeurIPS 2020)</conf-name><conf-date>Dec 6-12, 2020</conf-date><fpage>9459</fpage><lpage>9474</lpage><comment><ext-link ext-link-type="uri" xlink:href="https://proceedings.neurips.cc/paper/2020/file/6b493230205f780e1bc26945df7481e5-Paper.pdf">https://proceedings.neurips.cc/paper/2020/file/6b493230205f780e1bc26945df7481e5-Paper.pdf</ext-link></comment></nlm-citation></ref><ref id="ref6"><label>6</label><nlm-citation citation-type="web"><article-title>What is the model context protocol (MCP)?</article-title><source>Model Context Protocol</source><year>2024</year><access-date>2025-09-08</access-date><comment><ext-link ext-link-type="uri" xlink:href="https://modelcontextprotocol.io/">https://modelcontextprotocol.io/</ext-link></comment></nlm-citation></ref><ref id="ref7"><label>7</label><nlm-citation citation-type="web"><article-title>LangGraph overview</article-title><source>LangChain Docs</source><year>2024</year><access-date>2025-09-08</access-date><comment><ext-link ext-link-type="uri" xlink:href="https://docs.langchain.com/oss/python/langgraph/overview">https://docs.langchain.com/oss/python/langgraph/overview</ext-link></comment></nlm-citation></ref><ref id="ref8"><label>8</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Sallam</surname><given-names>M</given-names> </name></person-group><article-title>ChatGPT utility in healthcare education, research, and practice: systematic review on the promising perspectives and valid concerns</article-title><source>Healthcare (Basel)</source><year>2023</year><month>03</month><day>19</day><volume>11</volume><issue>6</issue><fpage>887</fpage><pub-id pub-id-type="doi">10.3390/healthcare11060887</pub-id><pub-id pub-id-type="medline">36981544</pub-id></nlm-citation></ref><ref id="ref9"><label>9</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Laupichler</surname><given-names>MC</given-names> </name><name name-style="western"><surname>Aster</surname><given-names>A</given-names> </name><name name-style="western"><surname>Schirch</surname><given-names>J</given-names> </name><name name-style="western"><surname>Raupach</surname><given-names>T</given-names> </name></person-group><article-title>Artificial intelligence literacy in higher and adult education: a scoping literature review</article-title><source>Comput Educ Artif Intell</source><year>2022</year><volume>3</volume><fpage>100101</fpage><pub-id pub-id-type="doi">10.1016/j.caeai.2022.100101</pub-id></nlm-citation></ref><ref id="ref10"><label>10</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Gordon</surname><given-names>M</given-names> </name><name name-style="western"><surname>Daniel</surname><given-names>M</given-names> </name><name name-style="western"><surname>Ajiboye</surname><given-names>A</given-names> </name><etal/></person-group><article-title>A scoping review of artificial intelligence in medical education: BEME Guide No. 84</article-title><source>Med Teach</source><year>2024</year><month>04</month><volume>46</volume><issue>4</issue><fpage>446</fpage><lpage>470</lpage><pub-id pub-id-type="doi">10.1080/0142159X.2024.2314198</pub-id><pub-id pub-id-type="medline">38423127</pub-id></nlm-citation></ref><ref id="ref11"><label>11</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Wu</surname><given-names>Y</given-names> </name><name name-style="western"><surname>Zheng</surname><given-names>Y</given-names> </name><name name-style="western"><surname>Feng</surname><given-names>B</given-names> </name><name name-style="western"><surname>Yang</surname><given-names>Y</given-names> </name><name name-style="western"><surname>Kang</surname><given-names>K</given-names> </name><name name-style="western"><surname>Zhao</surname><given-names>A</given-names> </name></person-group><article-title>Embracing ChatGPT for medical education: exploring its impact on doctors and medical students</article-title><source>JMIR Med Educ</source><year>2024</year><month>04</month><day>10</day><volume>10</volume><fpage>e52483</fpage><pub-id pub-id-type="doi">10.2196/52483</pub-id><pub-id pub-id-type="medline">38598263</pub-id></nlm-citation></ref><ref id="ref12"><label>12</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Hu</surname><given-names>JM</given-names> </name><name name-style="western"><surname>Liu</surname><given-names>FC</given-names> </name><name name-style="western"><surname>Chu</surname><given-names>CM</given-names> </name><name name-style="western"><surname>Chang</surname><given-names>YT</given-names> </name></person-group><article-title>Health care trainees&#x2019; and professionals&#x2019; perceptions of ChatGPT in improving medical knowledge training: rapid survey study</article-title><source>J Med Internet Res</source><year>2023</year><month>10</month><day>18</day><volume>25</volume><fpage>e49385</fpage><pub-id pub-id-type="doi">10.2196/49385</pub-id><pub-id pub-id-type="medline">37851495</pub-id></nlm-citation></ref><ref id="ref13"><label>13</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Tangadulrat</surname><given-names>P</given-names> </name><name name-style="western"><surname>Sono</surname><given-names>S</given-names> </name><name name-style="western"><surname>Tangtrakulwanich</surname><given-names>B</given-names> </name></person-group><article-title>Using ChatGPT for clinical practice and medical education: cross-sectional survey of medical students&#x2019; and physicians&#x2019; perceptions</article-title><source>JMIR Med Educ</source><year>2023</year><month>12</month><day>22</day><volume>9</volume><fpage>e50658</fpage><pub-id pub-id-type="doi">10.2196/50658</pub-id><pub-id pub-id-type="medline">38133908</pub-id></nlm-citation></ref><ref id="ref14"><label>14</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Sapci</surname><given-names>AH</given-names> </name><name name-style="western"><surname>Sapci</surname><given-names>HA</given-names> </name></person-group><article-title>Artificial intelligence education and tools for medical and health informatics students: systematic review</article-title><source>JMIR Med Educ</source><year>2020</year><month>06</month><day>30</day><volume>6</volume><issue>1</issue><fpage>e19285</fpage><pub-id pub-id-type="doi">10.2196/19285</pub-id><pub-id pub-id-type="medline">32602844</pub-id></nlm-citation></ref><ref id="ref15"><label>15</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Preiksaitis</surname><given-names>C</given-names> </name><name name-style="western"><surname>Rose</surname><given-names>C</given-names> </name></person-group><article-title>Opportunities, challenges, and future directions of generative artificial intelligence in medical education: scoping review</article-title><source>JMIR Med Educ</source><year>2023</year><month>10</month><day>20</day><volume>9</volume><fpage>e48785</fpage><pub-id pub-id-type="doi">10.2196/48785</pub-id><pub-id pub-id-type="medline">37862079</pub-id></nlm-citation></ref><ref id="ref16"><label>16</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Charow</surname><given-names>R</given-names> </name><name name-style="western"><surname>Jeyakumar</surname><given-names>T</given-names> </name><name name-style="western"><surname>Younus</surname><given-names>S</given-names> </name><etal/></person-group><article-title>Artificial intelligence education programs for health care professionals: scoping review</article-title><source>JMIR Med Educ</source><year>2021</year><month>12</month><day>13</day><volume>7</volume><issue>4</issue><fpage>e31043</fpage><pub-id pub-id-type="doi">10.2196/31043</pub-id><pub-id pub-id-type="medline">34898458</pub-id></nlm-citation></ref><ref id="ref17"><label>17</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Lee</surname><given-names>YM</given-names> </name><name name-style="western"><surname>Kim</surname><given-names>S</given-names> </name><name name-style="western"><surname>Lee</surname><given-names>YH</given-names> </name><etal/></person-group><article-title>Defining medical AI competencies for medical school graduates: outcomes of a delphi survey and medical student/educator questionnaire of South Korean medical schools</article-title><source>Acad Med</source><year>2024</year><month>05</month><day>1</day><volume>99</volume><issue>5</issue><fpage>524</fpage><lpage>533</lpage><pub-id pub-id-type="doi">10.1097/ACM.0000000000005618</pub-id><pub-id pub-id-type="medline">38207056</pub-id></nlm-citation></ref><ref id="ref18"><label>18</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Lee</surname><given-names>KH</given-names> </name><name name-style="western"><surname>Lee</surname><given-names>JH</given-names> </name><name name-style="western"><surname>Lee</surname><given-names>Y</given-names> </name><etal/></person-group><article-title>Impact of health informatics analyst education on job role, career transition, and skill development: survey study</article-title><source>JMIR Med Educ</source><year>2024</year><month>09</month><day>25</day><volume>10</volume><fpage>e54427</fpage><pub-id pub-id-type="doi">10.2196/54427</pub-id><pub-id pub-id-type="medline">39320368</pub-id></nlm-citation></ref><ref id="ref19"><label>19</label><nlm-citation citation-type="book"><person-group person-group-type="author"><name name-style="western"><surname>Bandura</surname><given-names>A</given-names> </name></person-group><source>Self-Efficacy: The Exercise of Control</source><year>1997</year><publisher-name>W.H. Freeman</publisher-name><pub-id pub-id-type="other">9780716728504</pub-id></nlm-citation></ref><ref id="ref20"><label>20</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Compeau</surname><given-names>DR</given-names> </name><name name-style="western"><surname>Higgins</surname><given-names>CA</given-names> </name></person-group><article-title>Computer self-efficacy: development of a measure and initial test</article-title><source>MIS Q</source><year>1995</year><month>06</month><day>1</day><volume>19</volume><issue>2</issue><fpage>189</fpage><lpage>211</lpage><pub-id pub-id-type="doi">10.2307/249688</pub-id></nlm-citation></ref><ref id="ref21"><label>21</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Davis</surname><given-names>FD</given-names> </name></person-group><article-title>Perceived usefulness, perceived ease of use, and user acceptance of information technology</article-title><source>MIS Q</source><year>1989</year><month>09</month><day>1</day><volume>13</volume><issue>3</issue><fpage>319</fpage><lpage>340</lpage><pub-id pub-id-type="doi">10.2307/249008</pub-id></nlm-citation></ref><ref id="ref22"><label>22</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Venkatesh</surname><given-names>V</given-names> </name><name name-style="western"><surname>Morris</surname><given-names>MG</given-names> </name><name name-style="western"><surname>Davis</surname><given-names>GB</given-names> </name><name name-style="western"><surname>Davis</surname><given-names>FD</given-names> </name></person-group><article-title>User acceptance of information technology: toward a unified view1</article-title><source>MIS Q</source><year>2003</year><month>09</month><day>1</day><volume>27</volume><issue>3</issue><fpage>425</fpage><lpage>478</lpage><pub-id pub-id-type="doi">10.2307/30036540</pub-id></nlm-citation></ref><ref id="ref23"><label>23</label><nlm-citation citation-type="book"><person-group person-group-type="author"><name name-style="western"><surname>Kolb</surname><given-names>DA</given-names> </name></person-group><source>Experiential Learning: Experience as the Source of Learning and Development</source><year>2015</year><edition>2</edition><publisher-name>Pearson Education</publisher-name><pub-id pub-id-type="other">9780133892406</pub-id></nlm-citation></ref><ref id="ref24"><label>24</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Kirkpatrick</surname><given-names>DL</given-names> </name></person-group><article-title>Evaluating training programs: evidence vs. proof</article-title><source>Train Dev J</source><year>1977</year><access-date>2026-08-14</access-date><volume>31</volume><issue>11</issue><fpage>9</fpage><lpage>12</lpage><comment><ext-link ext-link-type="uri" xlink:href="https://eric.ed.gov/?id=EJ169223">https://eric.ed.gov/?id=EJ169223</ext-link></comment></nlm-citation></ref><ref id="ref25"><label>25</label><nlm-citation citation-type="book"><person-group person-group-type="author"><name name-style="western"><surname>Kirkpatrick</surname><given-names>JD</given-names> </name><name name-style="western"><surname>Kirkpatrick</surname><given-names>WK</given-names> </name></person-group><source>Kirkpatrick&#x2019;s Four Levels of Training Evaluation</source><year>2016</year><publisher-name>ATD Press</publisher-name><pub-id pub-id-type="other">9781607280088</pub-id></nlm-citation></ref><ref id="ref26"><label>26</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Yardley</surname><given-names>S</given-names> </name><name name-style="western"><surname>Dornan</surname><given-names>T</given-names> </name></person-group><article-title>Kirkpatrick&#x2019;s levels and education &#x201C;evidence&#x201D;</article-title><source>Med Educ</source><year>2012</year><month>01</month><volume>46</volume><issue>1</issue><fpage>97</fpage><lpage>106</lpage><pub-id pub-id-type="doi">10.1111/j.1365-2923.2011.04076.x</pub-id><pub-id pub-id-type="medline">22150201</pub-id></nlm-citation></ref><ref id="ref27"><label>27</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Alhassan</surname><given-names>AI</given-names> </name></person-group><article-title>Implementing faculty development programs in medical education utilizing Kirkpatrick&#x2019;s model</article-title><source>Adv Med Educ Pract</source><year>2022</year><volume>13</volume><fpage>945</fpage><lpage>954</lpage><pub-id pub-id-type="doi">10.2147/AMEP.S372652</pub-id><pub-id pub-id-type="medline">36039186</pub-id></nlm-citation></ref><ref id="ref28"><label>28</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>von Elm</surname><given-names>E</given-names> </name><name name-style="western"><surname>Altman</surname><given-names>DG</given-names> </name><name name-style="western"><surname>Egger</surname><given-names>M</given-names> </name><etal/></person-group><article-title>The Strengthening the Reporting of Observational Studies in Epidemiology (STROBE) statement: guidelines for reporting observational studies</article-title><source>Lancet</source><year>2007</year><month>10</month><day>20</day><volume>370</volume><issue>9596</issue><fpage>1453</fpage><lpage>1457</lpage><pub-id pub-id-type="doi">10.1016/S0140-6736(07)61602-X</pub-id><pub-id pub-id-type="medline">18064739</pub-id></nlm-citation></ref><ref id="ref29"><label>29</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Wang</surname><given-names>L</given-names> </name><name name-style="western"><surname>Ma</surname><given-names>C</given-names> </name><name name-style="western"><surname>Feng</surname><given-names>X</given-names> </name><etal/></person-group><article-title>A survey on large language model based autonomous agents</article-title><source>Front Comput Sci</source><year>2024</year><month>12</month><volume>18</volume><issue>6</issue><fpage>186345</fpage><pub-id pub-id-type="doi">10.1007/s11704-024-40231-1</pub-id></nlm-citation></ref><ref id="ref30"><label>30</label><nlm-citation citation-type="book"><person-group person-group-type="author"><name name-style="western"><surname>Hollander</surname><given-names>M</given-names> </name><name name-style="western"><surname>Wolfe</surname><given-names>DA</given-names> </name><name name-style="western"><surname>Chicken</surname><given-names>E</given-names> </name></person-group><source>Nonparametric Statistical Methods</source><year>2014</year><edition>3</edition><publisher-name>John Wiley &#x0026; Sons</publisher-name><pub-id pub-id-type="doi">10.1002/9781119196037</pub-id></nlm-citation></ref><ref id="ref31"><label>31</label><nlm-citation citation-type="book"><person-group person-group-type="author"><name name-style="western"><surname>Rosenthal</surname><given-names>R</given-names> </name></person-group><source>Meta-Analytic Procedures for Social Research</source><year>1991</year><publisher-name>Sage</publisher-name><pub-id pub-id-type="other">9780803942462</pub-id></nlm-citation></ref><ref id="ref32"><label>32</label><nlm-citation citation-type="book"><person-group person-group-type="author"><name name-style="western"><surname>Cohen</surname><given-names>J</given-names> </name></person-group><source>Statistical Power Analysis for the Behavioral Sciences</source><year>1988</year><edition>2</edition><publisher-name>Lawrence Erlbaum Associates</publisher-name><pub-id pub-id-type="other">9780805802832</pub-id></nlm-citation></ref><ref id="ref33"><label>33</label><nlm-citation citation-type="book"><person-group person-group-type="author"><name name-style="western"><surname>Nunnally</surname><given-names>JC</given-names> </name><name name-style="western"><surname>Bernstein</surname><given-names>IH</given-names> </name></person-group><source>Psychometric Theory</source><year>1994</year><edition>3</edition><publisher-name>McGraw-Hill</publisher-name><pub-id pub-id-type="other">9780070478497</pub-id></nlm-citation></ref><ref id="ref34"><label>34</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Hassan</surname><given-names>M</given-names> </name><name name-style="western"><surname>Kushniruk</surname><given-names>A</given-names> </name><name name-style="western"><surname>Borycki</surname><given-names>E</given-names> </name></person-group><article-title>Barriers to and facilitators of artificial intelligence adoption in health care: scoping review</article-title><source>JMIR Hum Factors</source><year>2024</year><month>08</month><day>29</day><volume>11</volume><fpage>e48633</fpage><pub-id pub-id-type="doi">10.2196/48633</pub-id><pub-id pub-id-type="medline">39207831</pub-id></nlm-citation></ref><ref id="ref35"><label>35</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Greenhalgh</surname><given-names>T</given-names> </name><name name-style="western"><surname>Wherton</surname><given-names>J</given-names> </name><name name-style="western"><surname>Papoutsi</surname><given-names>C</given-names> </name><etal/></person-group><article-title>Beyond adoption: a new framework for theorizing and evaluating nonadoption, abandonment, and challenges to the scale-up, spread, and sustainability of health and care technologies</article-title><source>J Med Internet Res</source><year>2017</year><month>11</month><day>1</day><volume>19</volume><issue>11</issue><fpage>e367</fpage><pub-id pub-id-type="doi">10.2196/jmir.8775</pub-id><pub-id pub-id-type="medline">29092808</pub-id></nlm-citation></ref><ref id="ref36"><label>36</label><nlm-citation citation-type="book"><person-group person-group-type="author"><name name-style="western"><surname>Pfeffer</surname><given-names>J</given-names> </name><name name-style="western"><surname>Sutton</surname><given-names>RI</given-names> </name></person-group><source>The Knowing-Doing Gap: How Smart Companies Turn Knowledge into Action</source><year>2000</year><publisher-name>Harvard Business School Press</publisher-name><pub-id pub-id-type="doi">10.1108/scm.2001.6.3.142.1</pub-id></nlm-citation></ref><ref id="ref37"><label>37</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Grimshaw</surname><given-names>JM</given-names> </name><name name-style="western"><surname>Eccles</surname><given-names>MP</given-names> </name><name name-style="western"><surname>Lavis</surname><given-names>JN</given-names> </name><name name-style="western"><surname>Hill</surname><given-names>SJ</given-names> </name><name name-style="western"><surname>Squires</surname><given-names>JE</given-names> </name></person-group><article-title>Knowledge translation of research findings</article-title><source>Implement Sci</source><year>2012</year><month>05</month><day>31</day><volume>7</volume><fpage>50</fpage><pub-id pub-id-type="doi">10.1186/1748-5908-7-50</pub-id><pub-id pub-id-type="medline">22651257</pub-id></nlm-citation></ref><ref id="ref38"><label>38</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Tudor Car</surname><given-names>L</given-names> </name><name name-style="western"><surname>Kyaw</surname><given-names>BM</given-names> </name><name name-style="western"><surname>Nannan Panday</surname><given-names>RS</given-names> </name><etal/></person-group><article-title>Digital health training programs for medical students: scoping review</article-title><source>JMIR Med Educ</source><year>2021</year><month>07</month><day>21</day><volume>7</volume><issue>3</issue><fpage>e28275</fpage><pub-id pub-id-type="doi">10.2196/28275</pub-id><pub-id pub-id-type="medline">34287206</pub-id></nlm-citation></ref><ref id="ref39"><label>39</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Chen</surname><given-names>M</given-names> </name><name name-style="western"><surname>Zhang</surname><given-names>B</given-names> </name><name name-style="western"><surname>Cai</surname><given-names>Z</given-names> </name><etal/></person-group><article-title>Acceptance of clinical artificial intelligence among physicians and medical students: a systematic review with cross-sectional survey</article-title><source>Front Med (Lausanne)</source><year>2022</year><volume>9</volume><fpage>990604</fpage><pub-id pub-id-type="doi">10.3389/fmed.2022.990604</pub-id></nlm-citation></ref><ref id="ref40"><label>40</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Wartman</surname><given-names>SA</given-names> </name><name name-style="western"><surname>Combs</surname><given-names>CD</given-names> </name></person-group><article-title>Reimagining medical education in the age of AI</article-title><source>AMA J Ethics</source><year>2019</year><month>02</month><day>1</day><volume>21</volume><issue>2</issue><fpage>E146</fpage><lpage>152</lpage><pub-id pub-id-type="doi">10.1001/amajethics.2019.146</pub-id><pub-id pub-id-type="medline">30794124</pub-id></nlm-citation></ref><ref id="ref41"><label>41</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Yurek</surname><given-names>LA</given-names> </name><name name-style="western"><surname>Vasey</surname><given-names>J</given-names> </name><name name-style="western"><surname>Sullivan Havens</surname><given-names>D</given-names> </name></person-group><article-title>The use of self-generated identification codes in longitudinal research</article-title><source>Eval Rev</source><year>2008</year><month>10</month><volume>32</volume><issue>5</issue><fpage>435</fpage><lpage>452</lpage><pub-id pub-id-type="doi">10.1177/0193841X08316676</pub-id><pub-id pub-id-type="medline">18477737</pub-id></nlm-citation></ref><ref id="ref42"><label>42</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Podsakoff</surname><given-names>PM</given-names> </name><name name-style="western"><surname>MacKenzie</surname><given-names>SB</given-names> </name><name name-style="western"><surname>Lee</surname><given-names>JY</given-names> </name><name name-style="western"><surname>Podsakoff</surname><given-names>NP</given-names> </name></person-group><article-title>Common method biases in behavioral research: a critical review of the literature and recommended remedies</article-title><source>J Appl Psychol</source><year>2003</year><month>10</month><volume>88</volume><issue>5</issue><fpage>879</fpage><lpage>903</lpage><pub-id pub-id-type="doi">10.1037/0021-9010.88.5.879</pub-id><pub-id pub-id-type="medline">14516251</pub-id></nlm-citation></ref><ref id="ref43"><label>43</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Salas</surname><given-names>E</given-names> </name><name name-style="western"><surname>Tannenbaum</surname><given-names>SI</given-names> </name><name name-style="western"><surname>Kraiger</surname><given-names>K</given-names> </name><name name-style="western"><surname>Smith-Jentsch</surname><given-names>KA</given-names> </name></person-group><article-title>The science of training and development in organizations: what matters in practice</article-title><source>Psychol Sci Public Interest</source><year>2012</year><month>06</month><volume>13</volume><issue>2</issue><fpage>74</fpage><lpage>101</lpage><pub-id pub-id-type="doi">10.1177/1529100612436661</pub-id><pub-id pub-id-type="medline">26173283</pub-id></nlm-citation></ref></ref-list><app-group><supplementary-material id="app1"><label>Multimedia Appendix 1</label><p>Tables S1-S8 (S1: training curriculum by week; S2: S4 full domain statistics by subgroup with 95% CIs; S3: posttraining item-level results; S4: post-hoc power analysis; S5: exploratory factor analysis results; S6: item-total correlation analysis; S7: worst-case nonresponse sensitivity analysis; S8: capstone project descriptions and outcomes).</p><media xlink:href="mededu_v12i1e97822_app1.docx" xlink:title="DOCX File, 50 KB"/></supplementary-material><supplementary-material id="app2"><label>Multimedia Appendix 2</label><p>Detailed session-level training syllabus (per-session learning objectives, hands-on exercises, and tools).</p><media xlink:href="mededu_v12i1e97822_app2.docx" xlink:title="DOCX File, 28 KB"/></supplementary-material><supplementary-material id="app3"><label>Multimedia Appendix 3</label><p>Capstone project evaluation rubric (standardized 100-point rubric used by the assessment panel).</p><media xlink:href="mededu_v12i1e97822_app3.docx" xlink:title="DOCX File, 26 KB"/></supplementary-material><supplementary-material id="app4"><label>Multimedia Appendix 4</label><p>Open-ended survey responses (English translations of all 64 free-text responses to the posttraining self-perceived growth question).</p><media xlink:href="mededu_v12i1e97822_app4.docx" xlink:title="DOCX File, 29 KB"/></supplementary-material><supplementary-material id="app5"><label>Checklist 1</label><p>STROBE checklist for cross-sectional studies.</p><media xlink:href="mededu_v12i1e97822_app5.docx" xlink:title="DOCX File, 22 KB"/></supplementary-material></app-group></back></article>