<?xml version="1.0" encoding="UTF-8"?><!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.0 20040830//EN" "journalpublishing.dtd"><article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" dtd-version="2.0" xml:lang="en" article-type="research-article"><front><journal-meta><journal-id journal-id-type="nlm-ta">JMIR Med Educ</journal-id><journal-id journal-id-type="publisher-id">mededu</journal-id><journal-id journal-id-type="index">20</journal-id><journal-title>JMIR Medical Education</journal-title><abbrev-journal-title>JMIR Med Educ</abbrev-journal-title><issn pub-type="epub">2369-3762</issn><publisher><publisher-name>JMIR Publications</publisher-name><publisher-loc>Toronto, Canada</publisher-loc></publisher></journal-meta><article-meta><article-id pub-id-type="publisher-id">v12i1e99520</article-id><article-id pub-id-type="doi">10.2196/99520</article-id><article-categories><subj-group subj-group-type="heading"><subject>Viewpoint</subject></subj-group></article-categories><title-group><article-title>AI-Mediated Assessment of Continuing Medical Education: The Case-based Learning Intelligence Credit System (CLICS) Framework</article-title></title-group><contrib-group><contrib contrib-type="author" corresp="yes"><name name-style="western"><surname>Wattanasirichaigoon</surname><given-names>Somkiat</given-names></name><degrees>MD</degrees><xref ref-type="aff" rid="aff1">1</xref><xref ref-type="aff" rid="aff2">2</xref><xref ref-type="aff" rid="aff3">3</xref></contrib></contrib-group><aff id="aff1"><institution>Center for Continuing Medical Education, Medical Council of Thailand</institution><addr-line>The Royal Golden Jubilee Building, 2nd Floor Tiwanon Road Talat Khwan, Mueang Nonthaburi District</addr-line><addr-line>Nonthaburi</addr-line><country>Thailand</country></aff><aff id="aff2"><institution>Dhurakij Pundit University</institution><addr-line>Bangkok</addr-line><country>Thailand</country></aff><aff id="aff3"><institution>Division of Information Technology Management, Faculty of Engineering, Mahidol University</institution><addr-line>Nakhon Pathom</addr-line><country>Thailand</country></aff><contrib-group><contrib contrib-type="editor"><name name-style="western"><surname>Cardoso</surname><given-names>Taiane de Azevedo</given-names></name></contrib></contrib-group><contrib-group><contrib contrib-type="reviewer"><name name-style="western"><surname>McMahon</surname><given-names>Graham</given-names></name></contrib><contrib contrib-type="reviewer"><name name-style="western"><surname>Cervero</surname><given-names>Ronald</given-names></name></contrib><contrib contrib-type="reviewer"><name name-style="western"><surname>Khan</surname><given-names>Uzma</given-names></name></contrib></contrib-group><author-notes><corresp>Correspondence to Somkiat Wattanasirichaigoon, MD, Center for Continuing Medical Education, Medical Council of Thailand, The Royal Golden Jubilee Building, 2nd Floor Tiwanon Road Talat Khwan, Mueang Nonthaburi District, Nonthaburi, 11000, Thailand, 66 812891665; <email>somkiat.wat@mahidol.ac.th</email></corresp></author-notes><pub-date pub-type="collection"><year>2026</year></pub-date><pub-date pub-type="epub"><day>30</day><month>9</month><year>2026</year></pub-date><volume>12</volume><elocation-id>e99520</elocation-id><history><date date-type="received"><day>26</day><month>04</month><year>2026</year></date><date date-type="rev-recd"><day>29</day><month>08</month><year>2026</year></date><date date-type="accepted"><day>31</day><month>08</month><year>2026</year></date></history><copyright-statement>&#x00A9; Somkiat Wattanasirichaigoon. Originally published in JMIR Medical Education (<ext-link ext-link-type="uri" xlink:href="https://mededu.jmir.org">https://mededu.jmir.org</ext-link>), 30.9.2026. </copyright-statement><copyright-year>2026</copyright-year><license license-type="open-access" xlink:href="https://creativecommons.org/licenses/by/4.0/"><p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (<ext-link ext-link-type="uri" xlink:href="https://creativecommons.org/licenses/by/4.0/">https://creativecommons.org/licenses/by/4.0/</ext-link>), which permits unrestricted use, distribution, and reproduction in any medium, provided the original work, first published in JMIR Medical Education, is properly cited. The complete bibliographic information, a link to the original publication on <ext-link ext-link-type="uri" xlink:href="https://mededu.jmir.org/">https://mededu.jmir.org/</ext-link>, as well as this copyright and license information must be included.</p></license><self-uri xlink:type="simple" xlink:href="https://mededu.jmir.org/2026/1/e99520"/><abstract><p>Continuing medical education (CME) and continuing professional development (CPD) systems have traditionally relied on time-based credit allocation, using participation duration as a proxy for professional learning. Although administratively simple and scalable, this model does not reliably demonstrate whether physicians have engaged in meaningful learning, improved clinical reasoning, or critically appraised evidence. The emergence of generative AI creates an opportunity to rethink how physician learning is documented, assessed, and credited. This viewpoint proposes the Case-based Learning Intelligence Credit System (CLICS), a conceptual framework for translating AI-mediated clinical learning interactions into auditable evidence of reasoning-related engagement that could support CME/CPD credit. CLICS introduces the professional learning episode (PLE) as the basic unit of creditable learning: a coherent AI-mediated interaction demonstrating a clinically meaningful problem, reasoning development through iterative inquiry, contextual or evidentiary integration, and reflective synthesis. PLEs are evaluated using the proposed Practice Intelligence Score-7 (PIS-7) rubric, subject to human calibration and oversight; the rubric assesses observable reasoning behavior within the episode rather than the AI&#x2019;s answer, and qualifying PLEs may be translated into CME/CPD credit through threshold-based, human-auditable conversion rules. CLICS is not intended to replace traditional CME but to extend it as an optional, evidence-generating pathway for personalized, practice-embedded professional development. Its implementation requires iterative validation, stratified human audit, privacy-by-design architecture, antigaming controls, bias monitoring, and professional oversight. If validated, CLICS may enable identification of domain-specific areas for improvement and support personalized, adaptive learning pathways.</p></abstract><kwd-group><kwd>continuing medical education</kwd><kwd>continuing professional development</kwd><kwd>generative AI</kwd><kwd>clinical reasoning</kwd><kwd>medical education</kwd><kwd>physician learning</kwd><kwd>large language models</kwd><kwd>learning analytics</kwd><kwd>AI assessment</kwd></kwd-group></article-meta></front><body><sec id="s1" sec-type="intro"><title>Introduction</title><p>Continuing medical education (CME) and continuing professional development (CPD) systems commonly translate attendance, completion, or duration into professional learning credit [<xref ref-type="bibr" rid="ref1">1</xref>-<xref ref-type="bibr" rid="ref3">3</xref>]. Although practical and scalable, these participation metrics leave the quality of reasoning and evidence appraisal during learning largely unobserved.</p><p>Conversational generative AI creates a new opportunity, because physicians increasingly use it to explore clinical questions, compare alternatives, and reflect on decisions through iterative interactions embedded in clinical practice [<xref ref-type="bibr" rid="ref4">4</xref>-<xref ref-type="bibr" rid="ref6">6</xref>].</p><p>This viewpoint proposes the Case-based Learning Intelligence Credit System (CLICS), a framework for translating AI-mediated learning traces into auditable CME/CPD evidence. CLICS makes 3 conceptual shifts: from time-based toward evidence-informed CME; from AI as a learning assistant alone toward AI-assisted evaluation of learning traces, subject to human calibration and oversight; and from isolated learning activities toward practice-embedded professional learning infrastructure.</p><p>With CLICS, we do not propose that AI independently certify physician competence or that token count, conversation length, or prompt complexity determine credit. Instead, we distinguish superficial AI use from qualifying learning interactions through the professional learning episode (PLE), defined below. <xref ref-type="table" rid="table1">Table 1</xref> contrasts conventional CME/CPD with CLICS to show how the proposed framework may complement established systems by recognizing reasoning-centered learning that participation metrics do not capture.</p><table-wrap id="t1" position="float"><label>Table 1.</label><caption><p>A conceptual comparison between traditional continuing medical education (CME)/continuing professional development (CPD) and the Case-based Learning Intelligence Credit System (CLICS) framework, showing shifts from participation to reasoning-related evidence.</p></caption><table id="table1" frame="hsides" rules="groups"><thead><tr><td align="left" valign="bottom">Dimension</td><td align="left" valign="bottom">Traditional CME/CPD (established paradigm)</td><td align="left" valign="bottom">CLICS<break/>(proposed paradigm)</td><td align="left" valign="bottom">Interpretive implication</td></tr></thead><tbody><tr><td align="left" valign="top">Foundational assumption</td><td align="left" valign="top">Learning is approximated by time spent in accredited activities</td><td align="left" valign="top">Learning is inferred from observable, reasoning-related engagement within PLEs<sup><xref ref-type="table-fn" rid="table1fn1">a</xref></sup></td><td align="left" valign="top">Shifts the proxy of learning from duration to observable evidence</td></tr><tr><td align="left" valign="top">Unit of analysis</td><td align="left" valign="top">Educational event (eg, lecture, module, workshop)</td><td align="left" valign="top">PLE (AI-mediated reasoning interaction)</td><td align="left" valign="top">Reframes learning as episodic reasoning rather than attendance-based</td></tr><tr><td align="left" valign="top">Measurement target</td><td align="left" valign="top">Participation and completion</td><td align="left" valign="top">Reasoning-related interaction behavior, inquiry, and reflection</td><td align="left" valign="top">Moves from exposure metrics to reasoning-behavior metrics</td></tr><tr><td align="left" valign="top">Nature of evidence</td><td align="left" valign="top">Indirect (attendance as proxy)</td><td align="left" valign="top">Direct evidence of interaction behavior; inferential evidence of reasoning</td><td align="left" valign="top">Enables partial observation of reasoning-related behaviors</td></tr><tr><td align="left" valign="top">Assessment approach</td><td align="left" valign="top">Often absent or knowledge-based (eg, quizzes)</td><td align="left" valign="top">Multidomain assessment of reasoning-related behavior (PIS-7<sup><xref ref-type="table-fn" rid="table1fn2">b</xref></sup>)</td><td align="left" valign="top">Introduces a structured framework for evaluating reasoning behavior</td></tr><tr><td align="left" valign="top">Temporal context of learning</td><td align="left" valign="top">Scheduled and externally organized</td><td align="left" valign="top">Embedded within real-time clinical problem solving</td><td align="left" valign="top">Aligns CME with practice-based learning moments</td></tr><tr><td align="left" valign="top">Degree of personalization</td><td align="left" valign="top">Limited and content-driven</td><td align="left" valign="top">Potentially high; physician-driven and context-specific</td><td align="left" valign="top">Supports individualized learning trajectories</td></tr><tr><td align="left" valign="top">Feedback structure</td><td align="left" valign="top">Minimal or delayed</td><td align="left" valign="top">Immediate, domain-specific (eg, reasoning depth, critical appraisal)</td><td align="left" valign="top">Enables formative, actionable feedback loops</td></tr><tr><td align="left" valign="top">Auditability</td><td align="left" valign="top">Administrative verification (attendance logs)</td><td align="left" valign="top">Transcript-based audit with human calibration</td><td align="left" valign="top">Expands audit from compliance to reasoning-related evidence</td></tr><tr><td align="left" valign="top">Vulnerability to superficial completion</td><td align="left" valign="top">Possible (passive attendance)</td><td align="left" valign="top">Designed to be mitigated by PLE criteria and reflective synthesis requirements</td><td align="left" valign="top">Introduces structural resistance to low-effort credit accumulation</td></tr><tr><td align="left" valign="top">Governance model</td><td align="left" valign="top">Established regulatory frameworks</td><td align="left" valign="top">Emerging governance (audit, bias control, privacy, AI oversight)</td><td align="left" valign="top">Necessitates new regulatory and ethical architectures</td></tr><tr><td align="left" valign="top">Position within CME ecosystem</td><td align="left" valign="top">Foundational and established</td><td align="left" valign="top">Complementary and exploratory</td><td align="left" valign="top">Supports coexistence rather than replacement</td></tr></tbody></table><table-wrap-foot><fn id="table1fn1"><p><sup>a</sup>PLE: professional learning episode.</p></fn><fn id="table1fn2"><p><sup>b</sup>PIS-7: Practice Intelligence Score-7.</p></fn></table-wrap-foot></table-wrap></sec><sec id="s2"><title>Limitations of Time-Based CME</title><p>Traditional CME systems have played an important role in supporting lifelong learning, and systematic reviews show that CME can influence physician performance and, in some cases, patient outcomes, particularly when learning is interactive and linked to clinical practice [<xref ref-type="bibr" rid="ref1">1</xref>,<xref ref-type="bibr" rid="ref2">2</xref>]. However, CME effectiveness is highly variable and depends on how learning is designed, delivered, and assessed [<xref ref-type="bibr" rid="ref2">2</xref>,<xref ref-type="bibr" rid="ref7">7</xref>]. Furthermore, time-based credit allocation has 3 fundamental limitations. First, it relies on participation as a proxy for learning: attendance can be verified administratively, but it does not reflect whether meaningful cognitive engagement has occurred; assessment systems should be grounded in meaningful interpretation of performance rather than exposure time alone [<xref ref-type="bibr" rid="ref3">3</xref>,<xref ref-type="bibr" rid="ref8">8</xref>,<xref ref-type="bibr" rid="ref9">9</xref>]. Second, many CME activities are weakly aligned with real-world clinical reasoning&#x2014;delivered in structured, preplanned formats that do not reflect the complexity and contextual variability of practice&#x2014;even though reasoning, defined as the process of collecting information, generating hypotheses, and evaluating evidence, develops through iterative problem-solving and exposure to uncertainty rather than passive information acquisition [<xref ref-type="bibr" rid="ref10">10</xref>-<xref ref-type="bibr" rid="ref15">15</xref>]. Third, time-based systems offer limited personalization. Physicians differ in specialty, experience, and practice environment yet typically receive standardized content, even though tailored, context-specific feedback is known to enhance retention and reasoning performance [<xref ref-type="bibr" rid="ref3">3</xref>,<xref ref-type="bibr" rid="ref7">7</xref>,<xref ref-type="bibr" rid="ref15">15</xref>,<xref ref-type="bibr" rid="ref16">16</xref>].</p><p>CLICS addresses these limitations by shifting the focus from learning exposure to learning evidence&#x2014;asking not how long a physician participated but what observable evidence of reasoning-related engagement is present&#x2014;consistent with performance-informed assessment principles that emphasize observable evidence of learning rather than participation time alone [<xref ref-type="bibr" rid="ref8">8</xref>,<xref ref-type="bibr" rid="ref17">17</xref>].</p></sec><sec id="s3"><title>The CLICS Framework</title><p>CLICS integrates AI-mediated interaction, learning-trace capture, the Practice Intelligence Score-7 (PIS-7), human calibration, and credit conversion within a coordinated CME/CPD infrastructure. Although the underlying CLICS architecture is designed to be applicable across regulated professions requiring CPD, this viewpoint focuses specifically on physician CME, reflecting its current development for physician CME and planned phase 0 feasibility and reliability/agreement evaluation. Profession-specific adaptation, calibration, and governance would be required before extension beyond this scope.</p><p>At the foundation of CLICS is the AI-mediated learning trace: the sequence of physician-AI exchanges generated during a problem-centered inquiry [<xref ref-type="bibr" rid="ref4">4</xref>-<xref ref-type="bibr" rid="ref6">6</xref>,<xref ref-type="bibr" rid="ref18">18</xref>-<xref ref-type="bibr" rid="ref20">20</xref>]. Such traces provide observable, though indirect, evidence of reasoning behavior and must be interpreted cautiously because linguistic output does not fully represent underlying cognition [<xref ref-type="bibr" rid="ref10">10</xref>,<xref ref-type="bibr" rid="ref14">14</xref>].</p><p>CLICS defines the PLE as its fundamental unit of creditable learning: a coherent, AI-mediated interaction demonstrating (1) a clinically meaningful problem, (2) reasoning development through iterative inquiry, (3) integration of contextual or evidentiary information, and (4) reflective synthesis. Interactions limited to factual lookup, repetitive prompting, or passive copying of AI output do not qualify, a distinction essential to prevent superficial AI use from being misread as meaningful learning.</p><p>PLEs are evaluated by an AI-based evaluator that is conceptually distinct from the AI system generating educational content. Its purpose is to assess how the physician engages with the interaction rather than to judge the correctness of the AI&#x2019;s answer; the PIS-7 rubric used for this assessment is described below [<xref ref-type="bibr" rid="ref8">8</xref>,<xref ref-type="bibr" rid="ref9">9</xref>]. Moreover, only qualifying PLEs are eligible for proposed CME/CPD credit. <xref ref-type="fig" rid="figure1">Figure 1</xref> summarizes the closed-loop workflow from inquiry to evaluation, feedback, and subsequent learning.</p><fig position="float" id="figure1"><label>Figure 1.</label><caption><p>Closed-loop learning model of the Case-based Learning Intelligence Credit System (CLICS). AI-mediated clinical inquiries generate professional learning episodes (PLEs), which are evaluated using the proposed Practice Intelligence Score-7 (PIS-7) framework.</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="mededu_v12i1e99520_fig01.png"/></fig></sec><sec id="s4"><title>PLEs and PIS-7 Scoring</title><p>PLE components are grounded in the established perspectives that meaningful medical learning involves problem-solving, reasoning, contextualization, and reflection rather than passive exposure [<xref ref-type="bibr" rid="ref13">13</xref>-<xref ref-type="bibr" rid="ref15">15</xref>]. Problem framing supports hypothesis generation [<xref ref-type="bibr" rid="ref10">10</xref>,<xref ref-type="bibr" rid="ref11">11</xref>]; iterative reasoning supports hypothesis testing and uncertainty management [<xref ref-type="bibr" rid="ref11">11</xref>,<xref ref-type="bibr" rid="ref12">12</xref>]; contextual integration incorporates patient, evidence, and system constraints [<xref ref-type="bibr" rid="ref12">12</xref>,<xref ref-type="bibr" rid="ref17">17</xref>]; and reflective synthesis supports deeper learning and knowledge transfer [<xref ref-type="bibr" rid="ref14">14</xref>,<xref ref-type="bibr" rid="ref15">15</xref>,<xref ref-type="bibr" rid="ref21">21</xref>]. All 4 components must be present for an interaction to qualify as a PLE.</p><p>To operationalize the evaluation of PLEs, CLICS uses the PIS-7 framework, a multidomain rubric designed to assess physician reasoning behavior within a PLE. The 7 domains and their conceptual definitions are summarized in <xref ref-type="table" rid="table2">Table 2</xref>. Each PIS-7 domain is scored from 0 to 5, and the scores are converted into a weighted composite ranging from 0 to 100:</p><disp-formula id="equWL1"><mml:math id="eqn1"><mml:mstyle displaystyle="true" scriptlevel="0"><mml:mrow><mml:mstyle displaystyle="true" scriptlevel="0"><mml:mstyle displaystyle="true" scriptlevel="0"><mml:mtext>PIS-7</mml:mtext><mml:mo>=</mml:mo><mml:munderover><mml:mo movablelimits="false">&#x2211;</mml:mo><mml:mrow><mml:mi>i</mml:mi><mml:mo>=</mml:mo><mml:mn>1</mml:mn></mml:mrow><mml:mrow><mml:mn>7</mml:mn></mml:mrow></mml:munderover><mml:mrow><mml:mo>(</mml:mo><mml:mrow><mml:mfrac><mml:msub><mml:mi>D</mml:mi><mml:mrow><mml:mi>i</mml:mi></mml:mrow></mml:msub><mml:mn>5</mml:mn></mml:mfrac><mml:mo>&#x00D7;</mml:mo><mml:msub><mml:mi>W</mml:mi><mml:mrow><mml:mi>i</mml:mi></mml:mrow></mml:msub></mml:mrow><mml:mo>)</mml:mo></mml:mrow></mml:mstyle></mml:mstyle></mml:mrow></mml:mstyle></mml:math></disp-formula><p>where D&#x1D62; is the domain score (0&#x2010;5), W&#x1D62; is the domain weight expressed in percentage points, and &#x03A3;W&#x1D62;=100. For example, a score of 4/5 in a 20% domain contributes 16 points.</p><table-wrap id="t2" position="float"><label>Table 2.</label><caption><p>Practice Intelligence Score-7 (PIS-7) domains and proposed weighting for evaluating the quality of professional learning episodes.</p></caption><table id="table2" frame="hsides" rules="groups"><thead><tr><td align="left" valign="bottom">Domain</td><td align="left" valign="bottom">Construct assessed</td><td align="left" valign="bottom">Indicators of high-quality engagement</td><td align="left" valign="bottom">Weight (%)</td></tr></thead><tbody><tr><td align="left" valign="top">Professional question intelligence</td><td align="left" valign="top">Ability to frame clinically meaningful, context-aware questions</td><td align="left" valign="top">Defines relevant clinical problems with appropriate clinical context</td><td align="left" valign="top">15</td></tr><tr><td align="left" valign="top">Reasoning depth</td><td align="left" valign="top">Depth of clinical reasoning and uncertainty management</td><td align="left" valign="top">Demonstrates differential reasoning, prioritization, and risk-benefit analysis</td><td align="left" valign="top">20</td></tr><tr><td align="left" valign="top">Contextual integration</td><td align="left" valign="top">Integration of patient, evidence, and system factors</td><td align="left" valign="top">Incorporates patient factors, guidelines, and local practice constraints</td><td align="left" valign="top">15</td></tr><tr><td align="left" valign="top">Iterative inquiry</td><td align="left" valign="top">Refinement of reasoning through follow-up questions</td><td align="left" valign="top">Uses iterative questioning to clarify uncertainty and explore alternatives</td><td align="left" valign="top">15</td></tr><tr><td align="left" valign="top">Critical appraisal</td><td align="left" valign="top">Evaluation of AI-generated information</td><td align="left" valign="top">Requests evidence, challenges unsupported claims, and recognizes limitations</td><td align="left" valign="top">15</td></tr><tr><td align="left" valign="top">Reflective synthesis</td><td align="left" valign="top">Consolidation of learning into actionable insight</td><td align="left" valign="top">Summarizes learning and articulates implications for clinical practice</td><td align="left" valign="top">10</td></tr><tr><td align="left" valign="top">AI literacy and prompt skill</td><td align="left" valign="top">Responsible and effective use of AI tools</td><td align="left" valign="top">Provides context, verifies outputs, and maintains clinical responsibility</td><td align="left" valign="top">10</td></tr><tr><td align="left" valign="top" colspan="3">Total</td><td align="left" valign="top">100</td></tr></tbody></table></table-wrap><p>The weighting scheme assigns the highest weight to reasoning depth (20%), with the other 6 domains weighted at 10% to 15% each. This author-derived asymmetry reflects the conceptual centrality of diagnostic and therapeutic reasoning to physician decision-making [<xref ref-type="bibr" rid="ref10">10</xref>-<xref ref-type="bibr" rid="ref12">12</xref>]. The weights have not undergone formal consensus-based elicitation (eg, Delphi) and remain provisional pending empirical validation. After ethics approval and consent, a planned calibration workshop will use synthetic anchor PLEs to train 20 to 25 expert raters on the PIS-7 behavioral anchors before blinded scoring of study PLEs. Human-AI agreement and human interrater reliability will then be assessed using prespecified intraclass correlation coefficient (ICC) analyses to refine the scoring rubric and evaluator configuration rather than to rederive the domain weights.</p><p>The composite PIS-7 score summarizes observable reasoning behavior within a PLE rather than physician competence. Outputs should include domain-level explanations, transcript-based justification, and confidence indicators to support transparent human interpretation [<xref ref-type="bibr" rid="ref8">8</xref>,<xref ref-type="bibr" rid="ref9">9</xref>]. The PIS-7 is not intended to rank physicians or assign labels of intelligence; its formative purpose is to identify episode-specific areas for improvement and support personalized learning [<xref ref-type="bibr" rid="ref9">9</xref>,<xref ref-type="bibr" rid="ref15">15</xref>].</p><p>A preliminary credit conversion model may classify PLEs into noncreditable (&#x003C;40), basic (40&#x2013;59), meaningful (60&#x2013;74), advanced (75&#x2013;89), and very high (&#x2265;90) engagement tiers, each associated with a proposed base credit unit of 0, 0.25, 0.5, 0.75, and 1, respectively&#x2014;consistent with the 0.25&#x2010;1 base credit unit range specified in the associated patent claims&#x2014;with final CPD credit equal to this unit multiplied by a jurisdiction-specific scale factor. Tier labels refer only to engagement within a PLE. These thresholds are conceptual pending empirical calibration; <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref> provides 3 illustrative, hypothetical worked examples spanning a nonqualifying interaction and lower- and higher-scoring qualifying PLEs, with domain-level scores, rationale, and proposed credit equivalence. <xref ref-type="table" rid="table3">Table 3</xref> summarizes the conceptual workflow from interaction screening to feedback.</p><table-wrap id="t3" position="float"><label>Table 3.</label><caption><p>Conceptual screening, scoring, and credit workflow in the Case-based Learning Intelligence Credit System (CLICS).</p></caption><table id="table3" frame="hsides" rules="groups"><thead><tr><td align="left" valign="bottom">Stage</td><td align="left" valign="bottom">Operational question or action</td><td align="left" valign="bottom">Output</td></tr></thead><tbody><tr><td align="left" valign="top">PLE<sup><xref ref-type="table-fn" rid="table3fn1">a</xref></sup> eligibility gating</td><td align="left" valign="top">Does the interaction demonstrate a clinically meaningful problem, iterative reasoning, contextual or evidentiary integration, and reflective synthesis? Single factual lookups, repetitive prompting, or passive copying do not qualify.</td><td align="left" valign="top">Qualifying PLE, or no PIS-7<sup><xref ref-type="table-fn" rid="table3fn2">b</xref></sup> score and no credit</td></tr><tr><td align="left" valign="top">PIS-7 scoring</td><td align="left" valign="top">Score 7 reasoning-related domains from 0 to 5 using the proposed weights.</td><td align="left" valign="top">Domain profile and weighted 0 to 100 composite</td></tr><tr><td align="left" valign="top">Blinded human reference and calibration</td><td align="left" valign="top">For phase 0, use a primary set of at least 150 eligible PLEs, with &#x2265;3 independent expert ratings per PLE and AI scores hidden until human scoring is final; if &#x003E;150 PLEs are available, select 150 by score-stratified random sampling.</td><td align="left" valign="top">Human reference score, human-AI agreement, interrater reliability, and poststudy refinement of scoring anchors</td></tr><tr><td align="left" valign="top">Conceptual credit conversion</td><td align="left" valign="top">Map the composite score to an author-derived engagement tier and base credit unit only as a proposed future conversion model pending empirical validation and regulatory approval.</td><td align="left" valign="top">Illustrative base credit unit (0&#x2010;1) &#x00D7; jurisdiction-specific scale factor; no actual continuing medical education/continuing professional development credit in phase 0</td></tr><tr><td align="left" valign="top">Feedback and next cycle</td><td align="left" valign="top">Use domain-level results to identify episode-specific areas for improvement.</td><td align="left" valign="top">Personalized feedback and subsequent PLEs within the closed learning loop</td></tr></tbody></table><table-wrap-foot><fn id="table3fn1"><p><sup>a</sup>PLE: professional learning episode. </p></fn><fn id="table3fn2"><p><sup>b</sup>PIS-7: Practice Intelligence Score-7.</p></fn></table-wrap-foot></table-wrap></sec><sec id="s5"><title>Validation Roadmap</title><p>A framework that translates AI-mediated learning interactions into CME/CPD credit requires systematic validation of content validity, scoring reliability, construct validity, and implementation validity. Content validity concerns whether the PLE construct and PIS-7 domains adequately represent the intended reasoning-related learning behaviors [<xref ref-type="bibr" rid="ref8">8</xref>,<xref ref-type="bibr" rid="ref9">9</xref>]; scoring reliability requires comparison of AI-generated scores with human expert ratings using measures such as the ICC, weighted kappa, and Bland-Altman analysis [<xref ref-type="bibr" rid="ref22">22</xref>-<xref ref-type="bibr" rid="ref24">24</xref>]. A recent comparison of ChatGPT (OpenAI) and faculty scoring of formative medical education assessments reported 67% overall exact agreement [<xref ref-type="bibr" rid="ref25">25</xref>], reinforcing the need for human calibration. Construct validity must establish that higher scores reflect deeper reasoning rather than verbosity [<xref ref-type="bibr" rid="ref10">10</xref>,<xref ref-type="bibr" rid="ref11">11</xref>], while implementation validity concerns feasibility, scalability, and auditability.</p><p>In the phase 0 pilot study, we plan to enroll up to 75 licensed physicians, each contributing up to 3 PLEs (maximum 225 PLEs), together with an expert-rater pool of 20 to 25 physicians or medical educators. The primary reliability analysis requires a minimum evaluable set of 150 eligible PLEs, each independently scored by at least 3 blinded expert raters; the mean of the assigned human ratings (minimum 3) will serve as the human reference score. Phase 0 is designed to evaluate workflow feasibility and preliminary PIS-7 scoring reliability/agreement; study scores will not be used to award actual CME/CPD credit, affect licensure, or make decisions about physician competence.</p><p>If more than 150 eligible PLEs are available, the primary set of 150 will be selected by score-stratified random sampling across low-, middle-, and high-score ranges using locked AI scores before human review; if fewer than 150 are available, feasibility and agreement will be reported as exploratory without claiming validation. Primary human-AI agreement will compare the locked AI composite with the mean human reference using a prespecified absolute agreement ICC with 95% CI, while human interrater reliability will be reported separately; weighted kappa and Bland-Altman analysis will provide secondary agreement checks.</p><p>Scoring discrepancies and rubric-clarity feedback will inform subsequent refinement, but the primary evaluator configuration will remain locked throughout the primary analysis. If a critical safety or security change becomes necessary, affected scoring will be paused and the change will be documented under formal change control, with pre- and postchange data separated. These safeguards will preserve the distinction between developmental calibration and post hoc adjustment to observed results.</p></sec><sec id="s6"><title>Governance, Ethics, and Risk Control</title><p>CLICS requires governance for transparency, accountability, fairness, data protection, and human oversight. It should function as an educational assessment system rather than a surveillance mechanism: physicians should know what data are collected and how scores are generated, while expert educators and CME regulatory bodies define thresholds, review audits, and adjudicate disputes [<xref ref-type="bibr" rid="ref8">8</xref>,<xref ref-type="bibr" rid="ref9">9</xref>,<xref ref-type="bibr" rid="ref26">26</xref>-<xref ref-type="bibr" rid="ref29">29</xref>].</p><p>The phase 0 pilot study protocol was submitted to the Central Research Ethics Committee (CREC Thailand), under the supervision of the Foundation for Human Research Promotion in Thailand, on June 24, 2026 (case number CREC(S)24.06.69_01). The CREC reviewed the protocol on August 17, 2026, and requested revision for approval; the revised phase 0 protocol (version 4.1) was resubmitted on August 26, 2026, and remains under review at the time of this revision. The study was prospectively registered with the Thai Clinical Trials Registry on August 1, 2026 (TCTR20260801001), following registry submission on July 23, 2026. No participant or expert-rater scoring data collection has begun, and no such data will be collected before CREC approval is obtained.</p><p>For phase 0, research PLEs will use hypothetical or general clinical-learning problems without intentional collection of identifiable patient data, medical records, images, or patient outcomes. Research transcripts will be processed within CCME Version 8.0, the digital infrastructure/platform developed and operated by the Center for Continuing Medical Education (CCME), Medical Council of Thailand, and hosted on Thailand&#x2019;s Government Data Center and Cloud Service using a prespecified, locally deployed, open-weight, primary PIS-7 evaluator; transcripts will not be transmitted to external commercial AI APIs or used to train or fine-tune external models. Before the first participant PLE, the evaluator model artifact/build, prompt and rubric versions, inference configuration, and relevant run-time environment will be locked and versioned to support reproducibility and auditability.</p><p>AI-mediated transcripts may contain patient-adjacent information and indirect identifiers, creating reidentification risk. CLICS should therefore incorporate data minimization, pseudonymization, secure storage, role-based access, and audit logging [<xref ref-type="bibr" rid="ref27">27</xref>,<xref ref-type="bibr" rid="ref28">28</xref>], with General Data Protection Regulation (GDPR)&#x2013;comparable safeguards; under the GDPR, pseudonymized data that remain attributable to an identifiable person are still personal data [<xref ref-type="bibr" rid="ref30">30</xref>]. Participation should initially be voluntary and restricted to educational use. AI-based evaluators may also be influenced by linguistic style independent of reasoning quality [<xref ref-type="bibr" rid="ref26">26</xref>]. Generative AI may produce plausible but incorrect information or sycophantically agree with flawed premises [<xref ref-type="bibr" rid="ref31">31</xref>-<xref ref-type="bibr" rid="ref35">35</xref>]. Platform-level safeguards are designed to encourage challenges to unsupported premises, explicit uncertainty, and evidence traceability; these measures aim to mitigate, not eliminate, hallucinations and confirmation bias and require empirical validation.</p><p>The framework should also include safeguards against gaming, including autonomous or delegated AI agents generating a PLE without the physician&#x2019;s direct engagement. In any future operational implementation in which qualifying PLEs are converted to CME/CPD credit, reauthentication at the point of credit attribution could strengthen identity assurance. Reauthentication confirms the identity of the individual claiming credit but cannot, by itself, establish that every preceding reasoning step was personally generated by that individual. Additional safeguards may therefore include session-continuity and behavioral-consistency checks, similarity and anomalous-pattern analyses, and credit caps. The reliable technical detection of agent-delegated interactions has not yet been validated; this remains an unresolved governance challenge rather than a solved capability.</p></sec><sec id="s7"><title>Discussion and Policy Implications</title><p>CLICS may encourage a shift from participation-based metrics toward observable evidence of reasoning-related engagement, aligning with broader movements toward performance-informed and evidence-oriented assessment while remaining distinct from formal competency assessment [<xref ref-type="bibr" rid="ref8">8</xref>,<xref ref-type="bibr" rid="ref17">17</xref>]. This framing enables domain-specific feedback and individualized learning pathways [<xref ref-type="bibr" rid="ref14">14</xref>,<xref ref-type="bibr" rid="ref15">15</xref>,<xref ref-type="bibr" rid="ref21">21</xref>], while the AI literacy domain reflects the growing integration of AI into clinical workflows [<xref ref-type="bibr" rid="ref4">4</xref>-<xref ref-type="bibr" rid="ref6">6</xref>,<xref ref-type="bibr" rid="ref18">18</xref>-<xref ref-type="bibr" rid="ref20">20</xref>].</p><p>CLICS differs in 3 specific architectural respects from simpler structured reflective CME approaches, such as asking a physician, &#x201C;What did you learn from this interaction?&#x201D; First, a standardized Socratic rescue mode, implemented through a response integrity middleware pipeline, can activate when, after 4 physician-AI exchanges, the interaction remains nonqualifying or off-topic. Rather than simply supplying an answer, it returns targeted, nonleading Socratic questions to redirect the interaction toward a qualifying learning process. In phase 0, this mechanism is a standardized learning scaffold rather than an experimental arm, and its activation is analyzed only as an exploratory process metric. Second, once an interaction qualifies as a PLE, the evidence-based response protocol is designed to require traceable literature citations and explicit evidence-level or uncertainty markers for factual claims, providing an auditable evidentiary check that unstructured reflection does not offer. Third, the PIS-7 scores 7 distinct reasoning domains rather than a single undifferentiated judgment, enabling longitudinal, domain-level tracking of reasoning-related behavior and cohort-level analysis&#x2014;a level of granularity that free-text reflection is not designed to provide.</p><p>These structural distinctions provide plausible mechanisms by which CLICS may capture signals unavailable from simpler reflective approaches. This remains a theoretical argument, not a demonstrated finding: CLICS has not yet been shown to be empirically superior to structured reflective CME, and direct comparative evaluation is identified as a priority for future research.</p><p>Because AI-mediated learning tools have developed largely outside traditional educational oversight structures, CME/CPD accreditors and regulators are increasingly relevant to quality assurance for AI-mediated learning. This is consistent with guidance from the Accreditation Council for Continuing Medical Education (ACCME), including its January 2026 guidance on responsible AI use in accredited continuing education, its April 2026 alert clarifying that responsibility for learner-facing AI-generated content remains with the accredited provider rather than shifting to a vendor, and the recent literature on governance of generative AI in medical education [<xref ref-type="bibr" rid="ref36">36</xref>-<xref ref-type="bibr" rid="ref38">38</xref>]. The response integrity middleware pipeline and evidence-based response protocol function as procedural and evidence-traceability safeguards within a single interaction&#x2014;enforcing citation and evidence-level marking&#x2014;but do not by themselves constitute the content-validation infrastructure described in ACCME guidance, which includes predeployment validation against defined clinical scenarios, a defined clinical oversight structure with authority to intervene, ongoing monitoring and revalidation, and vendor accountability sufficient for provider oversight. Commercial bias that is subtle and consistent with evidence is difficult to detect at the level of a single response and is distinct from physician reasoning quality, which is what the PIS-7 measures. Phase 0 intentionally prioritizes reproducibility and research data protection by using a single locked, locally deployed, open-weight, primary evaluator rather than randomizing across external providers. In future operational implementations that incorporate multiple AI providers or knowledge pathways, aggregated PIS-7 and evidence-marker data could support the investigation of systematic models or vendor-specific patterns as 1 supplementary input to independent content validation and accreditor oversight.</p><p>A potential strength of CLICS is its alignment with practice-embedded learning: physicians often learn while addressing real problems involving uncertainty [<xref ref-type="bibr" rid="ref10">10</xref>,<xref ref-type="bibr" rid="ref13">13</xref>]. CLICS is also designed to mitigate epistemic homogenization by structuring AI-mediated learning around inquiry rather than answer conformity. Beyond the phase 0 evaluation setting, its modular architecture could support multiple models or knowledge pathways to expose learners to alternative interpretations and competing hypotheses; phase 0 intentionally uses a single, locked primary evaluator to maximize reproducibility. In the pilot, the design seeks to support epistemic pluralism through open-ended problems, Socratic counterquestioning, requests for disconfirming evidence, and PIS-7 domains&#x2014;including iterative inquiry, critical appraisal, contextual integration, and reflective synthesis&#x2014;which reward reasoned challenge and reflection rather than agreement with a predetermined answer. Repeated PLEs may also provide longitudinal signals relevant to metacognition, including recognition of uncertainty, consideration of alternatives, and self-monitoring of reasoning. These mechanisms aim to support evidence-constrained epistemic pluralism, although their effectiveness in preventing homogenization or improving metacognition remains to be empirically validated.</p><p>CLICS has important limitations. Learning traces remain a partial, indirect view of cognition [<xref ref-type="bibr" rid="ref10">10</xref>,<xref ref-type="bibr" rid="ref11">11</xref>], and implementation feasibility has not yet been established. The residual evaluator and AI-content risks described above require empirical testing. These limitations support phased validation and cautious interpretation.</p><p>From a policy perspective, CLICS is best understood as a complementary pathway that could eventually allow physicians to obtain a portion of CME credit through validated PLEs while continuing established educational activities. Earlier efforts to move beyond time-based participation include the American Nurses Credentialing Center outcome-based continuing education model and competency-based continuing professional development frameworks described in Canada [<xref ref-type="bibr" rid="ref39">39</xref>,<xref ref-type="bibr" rid="ref40">40</xref>]. Experience with competency-based assessment in Canada has also highlighted the administrative burden that can accompany frequent assessment and documentation [<xref ref-type="bibr" rid="ref41">41</xref>]. We do not assume that the use of CLICS will remove such implementation challenges by default. Its plausible points of difference are structural: credit derives directly from a machine-readable transcript rather than a separately authored document, scoring criteria are operationalized as explicit behavioral anchors, and the architecture integrates with AI-mediated learning that physicians already use. Furthermore, because it can engage with novel, nonroutine problems as they arise, CLICS is designed to support CPD at the point of practice in real time rather than only through retrospective reflection on resolved cases. Whether automated capture sufficiently reduces administrative burden and assessor variability to improve sustained adoption remains an empirical question for phase 0 and subsequent phases, not a conclusion asserted here in advance of evidence.</p></sec><sec id="s8" sec-type="conclusions"><title>Conclusions</title><p>We propose CLICS, a complementary pathway for recognizing AI-mediated, practice-embedded learning through the PLE and PIS-7. The PIS-7 provides an episode-specific summary of observable reasoning behavior, not a definitive measure of physician competence or clinical performance. If validated, CLICS may support adaptive, evidence-informed professional development by linking feedback and iterative learning to auditable, reasoning-related evidence. Its value will depend on empirical validation, human oversight, and governance addressing content integrity, privacy, fairness, and auditability.</p></sec></body><back><ack><p>The author acknowledges the global community of medical educators, clinicians, and regulators working to define responsible approaches to AI in medical education and continuing professional development.</p><p>The author used generative AI (ChatGPT, OpenAI) for text generation, editing/language polishing and proofreading, and literature search and summarization during preparation of the original manuscript. During preparation of the revision, the author used both ChatGPT (OpenAI) and Claude (Anthropic) to assist with the analysis of editor and reviewer comments, literature identification and summarization, drafting and restructuring of selected passages, condensation of text, and language refinement. The underlying conceptual architecture of both the Case-based Learning Intelligence Credit System (CLICS) and Practice Intelligence Score-7 (PIS-7) was developed by the author. The author independently reviewed and verified all AI-assisted content, references, interpretations, and final wording and retains full responsibility for the manuscript and its conclusions. Generative AI was not used to conduct participant-level data analysis, generate study code, or determine the research design or selection of research methods.</p></ack><notes><sec><title>Funding</title><p>Preparation of this manuscript received no external funding. The phase 0 pilot study described herein is funded by the Center for Continuing Medical Education, the Medical Council of Thailand.</p></sec><sec><title>Data Availability</title><p>No participant-level or expert-rater scoring datasets were generated or analyzed for this viewpoint. The article presents a conceptual framework and proposed validation strategy. Any future sharing of deidentified transcripts, scoring rubrics, evaluator prompts, analytic code, or aggregate datasets will remain subject to ethics, privacy, and institutional governance requirements.</p></sec></notes><fn-group><fn fn-type="conflict"><p>The corresponding author is the named inventor on a Thai patent application and an international patent application covering the Case-based Learning Intelligence Credit System (CLICS) architecture, both filed May 27, 2026, and on associated Thai trademark applications for CLICS and the Practice Intelligence Score-7 (PIS-7), filed May 25, 2026. Background intellectual property is held by MD Products and Services Co Ltd, a company with which the corresponding author is affiliated. The corresponding author also serves as director of the Center for Continuing Medical Education (CCME), the Medical Council of Thailand, which is funding the planned phase 0 pilot study. The CCME executive committee approved proceeding with the phase 0 pilot study under the director&#x2019;s leadership before any operational implementation of CLICS; the corresponding author was not present and did not participate in the meeting at which this decision was made. These relationships are disclosed for transparency and will be kept current in subsequent empirical, pilot, or implementation studies arising from this work.</p></fn></fn-group><glossary><title>Abbreviations</title><def-list><def-item><term id="abb1">ACCME</term><def><p>Accreditation Council for Continuing Medical Education</p></def></def-item><def-item><term id="abb2">CCME</term><def><p>Center for Continuing Medical Education</p></def></def-item><def-item><term id="abb3">CLICS</term><def><p>Case-based Learning Intelligence Credit System</p></def></def-item><def-item><term id="abb4">CME</term><def><p>continuing medical education</p></def></def-item><def-item><term id="abb5">CPD</term><def><p>continuing professional development</p></def></def-item><def-item><term id="abb6">ICC</term><def><p>intraclass correlation coefficient</p></def></def-item><def-item><term id="abb7">PIS-7</term><def><p>Practice Intelligence Score-7</p></def></def-item><def-item><term id="abb8">PLE</term><def><p>professional learning episode</p></def></def-item></def-list></glossary><ref-list><title>References</title><ref id="ref1"><label>1</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Davis</surname><given-names>DA</given-names> </name><name name-style="western"><surname>Thomson</surname><given-names>MA</given-names> </name><name name-style="western"><surname>Oxman</surname><given-names>AD</given-names> </name><name name-style="western"><surname>Haynes</surname><given-names>RB</given-names> </name></person-group><article-title>Changing physician performance. A systematic review of the effect of continuing medical education strategies</article-title><source>JAMA</source><year>1995</year><month>09</month><day>6</day><volume>274</volume><issue>9</issue><fpage>700</fpage><lpage>705</lpage><pub-id pub-id-type="doi">10.1001/jama.274.9.700</pub-id><pub-id pub-id-type="medline">7650822</pub-id></nlm-citation></ref><ref id="ref2"><label>2</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Cervero</surname><given-names>RM</given-names> </name><name name-style="western"><surname>Gaines</surname><given-names>JK</given-names> </name></person-group><article-title>The impact of CME on physician performance and patient health outcomes: an updated synthesis of systematic reviews</article-title><source>J Contin Educ Health Prof</source><year>2015</year><volume>35</volume><issue>2</issue><fpage>131</fpage><lpage>138</lpage><pub-id pub-id-type="doi">10.1002/chp.21290</pub-id><pub-id pub-id-type="medline">26115113</pub-id></nlm-citation></ref><ref id="ref3"><label>3</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Moore</surname><given-names>DE</given-names>  <suffix>Jr</suffix></name><name name-style="western"><surname>Green</surname><given-names>JS</given-names> </name><name name-style="western"><surname>Gallis</surname><given-names>HA</given-names> </name></person-group><article-title>Achieving desired results and improved outcomes: integrating planning and assessment throughout learning activities</article-title><source>J Contin Educ Health Prof</source><year>2009</year><volume>29</volume><issue>1</issue><fpage>1</fpage><lpage>15</lpage><pub-id pub-id-type="doi">10.1002/chp.20001</pub-id><pub-id pub-id-type="medline">19288562</pub-id></nlm-citation></ref><ref id="ref4"><label>4</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Wartman</surname><given-names>SA</given-names> </name><name name-style="western"><surname>Combs</surname><given-names>CD</given-names> </name></person-group><article-title>Medical education must move from the information age to the age of artificial intelligence</article-title><source>Acad Med</source><year>2018</year><month>08</month><volume>93</volume><issue>8</issue><fpage>1107</fpage><lpage>1109</lpage><pub-id pub-id-type="doi">10.1097/ACM.0000000000002044</pub-id><pub-id pub-id-type="medline">29095704</pub-id></nlm-citation></ref><ref id="ref5"><label>5</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Masters</surname><given-names>K</given-names> </name></person-group><article-title>Artificial intelligence in medical education</article-title><source>Med Teach</source><year>2019</year><month>09</month><volume>41</volume><issue>9</issue><fpage>976</fpage><lpage>980</lpage><pub-id pub-id-type="doi">10.1080/0142159X.2019.1595557</pub-id><pub-id pub-id-type="medline">31007106</pub-id></nlm-citation></ref><ref id="ref6"><label>6</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Chan</surname><given-names>KS</given-names> </name><name name-style="western"><surname>Zary</surname><given-names>N</given-names> </name></person-group><article-title>Applications and challenges of implementing artificial intelligence in medical education: integrative review</article-title><source>JMIR Med Educ</source><year>2019</year><month>06</month><day>15</day><volume>5</volume><issue>1</issue><fpage>e13930</fpage><pub-id pub-id-type="doi">10.2196/13930</pub-id><pub-id pub-id-type="medline">31199295</pub-id></nlm-citation></ref><ref id="ref7"><label>7</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Cook</surname><given-names>DA</given-names> </name><name name-style="western"><surname>Levinson</surname><given-names>AJ</given-names> </name><name name-style="western"><surname>Garside</surname><given-names>S</given-names> </name><name name-style="western"><surname>Dupras</surname><given-names>DM</given-names> </name><name name-style="western"><surname>Erwin</surname><given-names>PJ</given-names> </name><name name-style="western"><surname>Montori</surname><given-names>VM</given-names> </name></person-group><article-title>Internet-based learning in the health professions: a meta-analysis</article-title><source>JAMA</source><year>2008</year><month>09</month><day>10</day><volume>300</volume><issue>10</issue><fpage>1181</fpage><lpage>1196</lpage><pub-id pub-id-type="doi">10.1001/jama.300.10.1181</pub-id><pub-id pub-id-type="medline">18780847</pub-id></nlm-citation></ref><ref id="ref8"><label>8</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Cook</surname><given-names>DA</given-names> </name><name name-style="western"><surname>Hatala</surname><given-names>R</given-names> </name></person-group><article-title>Validation of educational assessments: a primer for simulation and beyond</article-title><source>Adv Simul (Lond)</source><year>2016</year><volume>1</volume><fpage>31</fpage><pub-id pub-id-type="doi">10.1186/s41077-016-0033-y</pub-id><pub-id pub-id-type="medline">29450000</pub-id></nlm-citation></ref><ref id="ref9"><label>9</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Downing</surname><given-names>SM</given-names> </name></person-group><article-title>Validity: on meaningful interpretation of assessment data</article-title><source>Med Educ</source><year>2003</year><month>09</month><volume>37</volume><issue>9</issue><fpage>830</fpage><lpage>837</lpage><pub-id pub-id-type="doi">10.1046/j.1365-2923.2003.01594.x</pub-id><pub-id pub-id-type="medline">14506816</pub-id></nlm-citation></ref><ref id="ref10"><label>10</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Eva</surname><given-names>KW</given-names> </name></person-group><article-title>What every teacher needs to know about clinical reasoning</article-title><source>Med Educ</source><year>2005</year><month>01</month><volume>39</volume><issue>1</issue><fpage>98</fpage><lpage>106</lpage><pub-id pub-id-type="doi">10.1111/j.1365-2929.2004.01972.x</pub-id><pub-id pub-id-type="medline">15612906</pub-id></nlm-citation></ref><ref id="ref11"><label>11</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Norman</surname><given-names>G</given-names> </name></person-group><article-title>Research in clinical reasoning: past history and current trends</article-title><source>Med Educ</source><year>2005</year><month>04</month><volume>39</volume><issue>4</issue><fpage>418</fpage><lpage>427</lpage><pub-id pub-id-type="doi">10.1111/j.1365-2929.2005.02127.x</pub-id><pub-id pub-id-type="medline">15813765</pub-id></nlm-citation></ref><ref id="ref12"><label>12</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Gruppen</surname><given-names>LD</given-names> </name></person-group><article-title>Clinical reasoning: defining it, teaching it, assessing it, studying it</article-title><source>West J Emerg Med</source><year>2017</year><month>01</month><volume>18</volume><issue>1</issue><fpage>4</fpage><lpage>7</lpage><pub-id pub-id-type="doi">10.5811/westjem.2016.11.33191</pub-id><pub-id pub-id-type="medline">28115999</pub-id></nlm-citation></ref><ref id="ref13"><label>13</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Bowen</surname><given-names>JL</given-names> </name></person-group><article-title>Educational strategies to promote clinical diagnostic reasoning</article-title><source>N Engl J Med</source><year>2006</year><month>11</month><day>23</day><volume>355</volume><issue>21</issue><fpage>2217</fpage><lpage>2225</lpage><pub-id pub-id-type="doi">10.1056/NEJMra054782</pub-id><pub-id pub-id-type="medline">17124019</pub-id></nlm-citation></ref><ref id="ref14"><label>14</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Mamede</surname><given-names>S</given-names> </name><name name-style="western"><surname>Schmidt</surname><given-names>HG</given-names> </name></person-group><article-title>The structure of reflective practice in medicine</article-title><source>Med Educ</source><year>2004</year><month>12</month><volume>38</volume><issue>12</issue><fpage>1302</fpage><lpage>1308</lpage><pub-id pub-id-type="doi">10.1111/j.1365-2929.2004.01917.x</pub-id><pub-id pub-id-type="medline">15566542</pub-id></nlm-citation></ref><ref id="ref15"><label>15</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Sandars</surname><given-names>J</given-names> </name></person-group><article-title>The use of reflection in medical education: AMEE Guide No. 44</article-title><source>Med Teach</source><year>2009</year><month>08</month><volume>31</volume><issue>8</issue><fpage>685</fpage><lpage>695</lpage><pub-id pub-id-type="doi">10.1080/01421590903050374</pub-id><pub-id pub-id-type="medline">19811204</pub-id></nlm-citation></ref><ref id="ref16"><label>16</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Ericsson</surname><given-names>KA</given-names> </name></person-group><article-title>Deliberate practice and acquisition of expert performance: a general overview</article-title><source>Acad Emerg Med</source><year>2008</year><month>11</month><volume>15</volume><issue>11</issue><fpage>988</fpage><lpage>994</lpage><pub-id pub-id-type="doi">10.1111/j.1553-2712.2008.00227.x</pub-id><pub-id pub-id-type="medline">18778378</pub-id></nlm-citation></ref><ref id="ref17"><label>17</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Epstein</surname><given-names>RM</given-names> </name><name name-style="western"><surname>Hundert</surname><given-names>EM</given-names> </name></person-group><article-title>Defining and assessing professional competence</article-title><source>JAMA</source><year>2002</year><month>01</month><day>9</day><volume>287</volume><issue>2</issue><fpage>226</fpage><lpage>235</lpage><pub-id pub-id-type="doi">10.1001/jama.287.2.226</pub-id><pub-id pub-id-type="medline">11779266</pub-id></nlm-citation></ref><ref id="ref18"><label>18</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Topol</surname><given-names>EJ</given-names> </name></person-group><article-title>High-performance medicine: the convergence of human and artificial intelligence</article-title><source>Nat Med</source><year>2019</year><month>01</month><volume>25</volume><issue>1</issue><fpage>44</fpage><lpage>56</lpage><pub-id pub-id-type="doi">10.1038/s41591-018-0300-7</pub-id><pub-id pub-id-type="medline">30617339</pub-id></nlm-citation></ref><ref id="ref19"><label>19</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Rajpurkar</surname><given-names>P</given-names> </name><name name-style="western"><surname>Chen</surname><given-names>E</given-names> </name><name name-style="western"><surname>Banerjee</surname><given-names>O</given-names> </name><name name-style="western"><surname>Topol</surname><given-names>EJ</given-names> </name></person-group><article-title>AI in health and medicine</article-title><source>Nat Med</source><year>2022</year><month>01</month><volume>28</volume><issue>1</issue><fpage>31</fpage><lpage>38</lpage><pub-id pub-id-type="doi">10.1038/s41591-021-01614-0</pub-id><pub-id pub-id-type="medline">35058619</pub-id></nlm-citation></ref><ref id="ref20"><label>20</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Lin</surname><given-names>Y</given-names> </name><name name-style="western"><surname>Luo</surname><given-names>Z</given-names> </name><name name-style="western"><surname>Ye</surname><given-names>Z</given-names> </name><etal/></person-group><article-title>Applications, challenges, and prospects of generative artificial intelligence empowering medical education: scoping review</article-title><source>JMIR Med Educ</source><year>2025</year><month>10</month><day>23</day><volume>11</volume><fpage>e71125</fpage><pub-id pub-id-type="doi">10.2196/71125</pub-id><pub-id pub-id-type="medline">41128430</pub-id></nlm-citation></ref><ref id="ref21"><label>21</label><nlm-citation citation-type="book"><person-group person-group-type="author"><name name-style="western"><surname>Sch&#x00F6;n</surname><given-names>DA</given-names> </name></person-group><source>The Reflective Practitioner: How Professionals Think in Action</source><year>1983</year><publisher-name>Basic Books</publisher-name><pub-id pub-id-type="other">9780465068784</pub-id></nlm-citation></ref><ref id="ref22"><label>22</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Koo</surname><given-names>TK</given-names> </name><name name-style="western"><surname>Li</surname><given-names>MY</given-names> </name></person-group><article-title>A guideline of selecting and reporting intraclass correlation coefficients for reliability research</article-title><source>J Chiropr Med</source><year>2016</year><month>06</month><volume>15</volume><issue>2</issue><fpage>155</fpage><lpage>163</lpage><pub-id pub-id-type="doi">10.1016/j.jcm.2016.02.012</pub-id><pub-id pub-id-type="medline">27330520</pub-id></nlm-citation></ref><ref id="ref23"><label>23</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>McHugh</surname><given-names>ML</given-names> </name></person-group><article-title>Interrater reliability: the kappa statistic</article-title><source>Biochem Med (Zagreb)</source><year>2012</year><volume>22</volume><issue>3</issue><fpage>276</fpage><lpage>282</lpage><pub-id pub-id-type="medline">23092060</pub-id></nlm-citation></ref><ref id="ref24"><label>24</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Bland</surname><given-names>JM</given-names> </name><name name-style="western"><surname>Altman</surname><given-names>DG</given-names> </name></person-group><article-title>Statistical methods for assessing agreement between two methods of clinical measurement</article-title><source>Lancet</source><year>1986</year><month>02</month><day>8</day><volume>1</volume><issue>8476</issue><fpage>307</fpage><lpage>310</lpage><pub-id pub-id-type="medline">2868172</pub-id></nlm-citation></ref><ref id="ref25"><label>25</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Sreedhar</surname><given-names>R</given-names> </name><name name-style="western"><surname>Chang</surname><given-names>L</given-names> </name><name name-style="western"><surname>Gangopadhyaya</surname><given-names>A</given-names> </name><etal/></person-group><article-title>Comparing scoring consistency of large language models with faculty for formative assessments in medical education</article-title><source>J Gen Intern Med</source><year>2025</year><month>01</month><volume>40</volume><issue>1</issue><fpage>127</fpage><lpage>134</lpage><pub-id pub-id-type="doi">10.1007/s11606-024-09050-9</pub-id><pub-id pub-id-type="medline">39402411</pub-id></nlm-citation></ref><ref id="ref26"><label>26</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Char</surname><given-names>DS</given-names> </name><name name-style="western"><surname>Shah</surname><given-names>NH</given-names> </name><name name-style="western"><surname>Magnus</surname><given-names>D</given-names> </name></person-group><article-title>Implementing machine learning in health care - addressing ethical challenges</article-title><source>N Engl J Med</source><year>2018</year><month>03</month><day>15</day><volume>378</volume><issue>11</issue><fpage>981</fpage><lpage>983</lpage><pub-id pub-id-type="doi">10.1056/NEJMp1714229</pub-id><pub-id pub-id-type="medline">29539284</pub-id></nlm-citation></ref><ref id="ref27"><label>27</label><nlm-citation citation-type="web"><article-title>Ethics and governance of artificial intelligence for health: WHO guidance</article-title><source>World Health Organization</source><year>2021</year><access-date>2026-08-23</access-date><publisher-name>World Health Organization</publisher-name><comment><ext-link ext-link-type="uri" xlink:href="https://www.who.int/publications/i/item/9789240029200">https://www.who.int/publications/i/item/9789240029200</ext-link></comment></nlm-citation></ref><ref id="ref28"><label>28</label><nlm-citation citation-type="web"><article-title>Regulatory considerations on artificial intelligence for health</article-title><source>World Health Organization</source><year>2023</year><access-date>2026-08-23</access-date><publisher-name>World Health Organization</publisher-name><comment><ext-link ext-link-type="uri" xlink:href="https://www.who.int/publications/i/item/9789240078871">https://www.who.int/publications/i/item/9789240078871</ext-link></comment></nlm-citation></ref><ref id="ref29"><label>29</label><nlm-citation citation-type="web"><article-title>Recommendation on the Ethics of Artificial Intelligence</article-title><source>United Nations Educational, Scientific and Cultural Organization</source><year>2021</year><access-date>2026-08-23</access-date><publisher-name>UNESCO</publisher-name><comment><ext-link ext-link-type="uri" xlink:href="https://www.unesco.org/en/legal-affairs/recommendation-ethics-artificial-intelligence">https://www.unesco.org/en/legal-affairs/recommendation-ethics-artificial-intelligence</ext-link></comment></nlm-citation></ref><ref id="ref30"><label>30</label><nlm-citation citation-type="web"><article-title>Regulation (EU) 2016/679 of the European Parliament and of the Council of 27 April 2016 on the protection of natural persons with regard to the processing of personal data and on the free movement of such data, and repealing Directive 95/46/EC (General Data Protection Regulation) (Text with EEA relevance)</article-title><source>EUR-Lex</source><year>2016</year><access-date>2026-08-23</access-date><comment><ext-link ext-link-type="uri" xlink:href="https://eur-lex.europa.eu/eli/reg/2016/679/oj/">https://eur-lex.europa.eu/eli/reg/2016/679/oj/</ext-link></comment></nlm-citation></ref><ref id="ref31"><label>31</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Sallam</surname><given-names>M</given-names> </name></person-group><article-title>ChatGPT utility in healthcare education, research, and practice: systematic review on the promising perspectives and valid concerns</article-title><source>Healthcare (Basel)</source><year>2023</year><month>03</month><day>19</day><volume>11</volume><issue>6</issue><fpage>887</fpage><pub-id pub-id-type="doi">10.3390/healthcare11060887</pub-id><pub-id pub-id-type="medline">36981544</pub-id></nlm-citation></ref><ref id="ref32"><label>32</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Kung</surname><given-names>TH</given-names> </name><name name-style="western"><surname>Cheatham</surname><given-names>M</given-names> </name><name name-style="western"><surname>Medenilla</surname><given-names>A</given-names> </name><etal/></person-group><article-title>Performance of ChatGPT on USMLE: potential for AI-assisted medical education using large language models</article-title><source>PLOS Digit Health</source><year>2023</year><month>02</month><volume>2</volume><issue>2</issue><fpage>e0000198</fpage><pub-id pub-id-type="doi">10.1371/journal.pdig.0000198</pub-id><pub-id pub-id-type="medline">36812645</pub-id></nlm-citation></ref><ref id="ref33"><label>33</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Thirunavukarasu</surname><given-names>AJ</given-names> </name><name name-style="western"><surname>Ting</surname><given-names>DSJ</given-names> </name><name name-style="western"><surname>Elangovan</surname><given-names>K</given-names> </name><name name-style="western"><surname>Gutierrez</surname><given-names>L</given-names> </name><name name-style="western"><surname>Tan</surname><given-names>TF</given-names> </name><name name-style="western"><surname>Ting</surname><given-names>DSW</given-names> </name></person-group><article-title>Large language models in medicine</article-title><source>Nat Med</source><year>2023</year><month>08</month><volume>29</volume><issue>8</issue><fpage>1930</fpage><lpage>1940</lpage><pub-id pub-id-type="doi">10.1038/s41591-023-02448-8</pub-id><pub-id pub-id-type="medline">37460753</pub-id></nlm-citation></ref><ref id="ref34"><label>34</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Kanjee</surname><given-names>Z</given-names> </name><name name-style="western"><surname>Crowe</surname><given-names>B</given-names> </name><name name-style="western"><surname>Rodman</surname><given-names>A</given-names> </name></person-group><article-title>Accuracy of a generative artificial intelligence model in a complex diagnostic challenge</article-title><source>JAMA</source><year>2023</year><month>07</month><day>3</day><volume>330</volume><issue>1</issue><fpage>78</fpage><lpage>80</lpage><pub-id pub-id-type="doi">10.1001/jama.2023.8288</pub-id><pub-id pub-id-type="medline">37318797</pub-id></nlm-citation></ref><ref id="ref35"><label>35</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Chen</surname><given-names>S</given-names> </name><name name-style="western"><surname>Gao</surname><given-names>M</given-names> </name><name name-style="western"><surname>Sasse</surname><given-names>K</given-names> </name><etal/></person-group><article-title>When helpfulness backfires: LLMs and the risk of false medical information due to sycophantic behavior</article-title><source>NPJ Digit Med</source><year>2025</year><month>10</month><day>17</day><volume>8</volume><issue>1</issue><fpage>605</fpage><pub-id pub-id-type="doi">10.1038/s41746-025-02008-z</pub-id><pub-id pub-id-type="medline">41107408</pub-id></nlm-citation></ref><ref id="ref36"><label>36</label><nlm-citation citation-type="web"><article-title>Guidance on the responsible use of artificial intelligence (AI) in accredited continuing education (CE)</article-title><source>Accreditation Council for Continuing Medical Education</source><year>2026</year><access-date>2026-08-23</access-date><publisher-name>ACCME</publisher-name><comment><ext-link ext-link-type="uri" xlink:href="https://accme.org/wp-content/uploads/2026/01/1098_20260130_Guidance_on_Artificial_Intelligence_in_Accredited_CE_ACCME.pdf">https://accme.org/wp-content/uploads/2026/01/1098_20260130_Guidance_on_Artificial_Intelligence_in_Accredited_CE_ACCME.pdf</ext-link></comment></nlm-citation></ref><ref id="ref37"><label>37</label><nlm-citation citation-type="web"><article-title>Urgent alert on the use of AI in accredited CE</article-title><source>Accreditation Council for Continuing Medical Education</source><year>2026</year><access-date>2026-08-23</access-date><publisher-name>ACCME</publisher-name><comment><ext-link ext-link-type="uri" xlink:href="https://accme.org/news/urgent-alert-on-the-use-of-ai-in-accredited-ce/">https://accme.org/news/urgent-alert-on-the-use-of-ai-in-accredited-ce/</ext-link></comment></nlm-citation></ref><ref id="ref38"><label>38</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Tran</surname><given-names>M</given-names> </name><name name-style="western"><surname>Balasooriya</surname><given-names>C</given-names> </name><name name-style="western"><surname>Jonnagaddala</surname><given-names>J</given-names> </name><etal/></person-group><article-title>Situating governance and regulatory concerns for generative artificial intelligence and large language models in medical education</article-title><source>NPJ Digit Med</source><year>2025</year><month>05</month><day>27</day><volume>8</volume><issue>1</issue><fpage>315</fpage><pub-id pub-id-type="doi">10.1038/s41746-025-01721-z</pub-id><pub-id pub-id-type="medline">40425695</pub-id></nlm-citation></ref><ref id="ref39"><label>39</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Graebe</surname><given-names>J</given-names> </name></person-group><article-title>Continuing professional development: utilizing competency-based education and the American Nurses Credentialing Center outcome-based continuing education model&#x00A9;</article-title><source>J Contin Educ Nurs</source><year>2019</year><month>03</month><day>1</day><volume>50</volume><issue>3</issue><fpage>100</fpage><lpage>102</lpage><pub-id pub-id-type="doi">10.3928/00220124-20190218-02</pub-id><pub-id pub-id-type="medline">30835317</pub-id></nlm-citation></ref><ref id="ref40"><label>40</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Campbell</surname><given-names>C</given-names> </name><name name-style="western"><surname>Silver</surname><given-names>I</given-names> </name><name name-style="western"><surname>Sherbino</surname><given-names>J</given-names> </name><name name-style="western"><surname>Cate</surname><given-names>OT</given-names> </name><name name-style="western"><surname>Holmboe</surname><given-names>ES</given-names> </name></person-group><article-title>Competency-based continuing professional development</article-title><source>Med Teach</source><year>2010</year><volume>32</volume><issue>8</issue><fpage>657</fpage><lpage>662</lpage><pub-id pub-id-type="doi">10.3109/0142159X.2010.500708</pub-id><pub-id pub-id-type="medline">20662577</pub-id></nlm-citation></ref><ref id="ref41"><label>41</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Cheung</surname><given-names>K</given-names> </name><name name-style="western"><surname>Rogoza</surname><given-names>C</given-names> </name><name name-style="western"><surname>Chung</surname><given-names>AD</given-names> </name><name name-style="western"><surname>Kwan</surname><given-names>BYM</given-names> </name></person-group><article-title>Analyzing the administrative burden of competency based medical education</article-title><source>Can Assoc Radiol J</source><year>2022</year><month>05</month><volume>73</volume><issue>2</issue><fpage>299</fpage><lpage>304</lpage><pub-id pub-id-type="doi">10.1177/08465371211038963</pub-id><pub-id pub-id-type="medline">34449283</pub-id></nlm-citation></ref></ref-list><app-group><supplementary-material id="app1"><label>Multimedia Appendix 1</label><p>Illustrative worked examples of Practice Intelligence Score-7 (PIS-7) scoring and continuing medical education (CME) credit conversion.</p><media xlink:href="mededu_v12i1e99520_app1.pdf" xlink:title="PDF File, 164 KB"/></supplementary-material></app-group></back></article>