<?xml version="1.0" encoding="UTF-8"?><!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.0 20040830//EN" "journalpublishing.dtd"><article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" dtd-version="2.0" xml:lang="en" article-type="research-article"><front><journal-meta><journal-id journal-id-type="nlm-ta">JMIR Med Educ</journal-id><journal-id journal-id-type="publisher-id">mededu</journal-id><journal-id journal-id-type="index">20</journal-id><journal-title>JMIR Medical Education</journal-title><abbrev-journal-title>JMIR Med Educ</abbrev-journal-title><issn pub-type="epub">2369-3762</issn><publisher><publisher-name>JMIR Publications</publisher-name><publisher-loc>Toronto, Canada</publisher-loc></publisher></journal-meta><article-meta><article-id pub-id-type="publisher-id">v12i1e95904</article-id><article-id pub-id-type="doi">10.2196/95904</article-id><article-categories><subj-group subj-group-type="heading"><subject>Original Paper</subject></subj-group></article-categories><title-group><article-title>Refining the Immersive Technology Evaluation Measure for Virtual Reality Clinical Skills Among Native Arabic-Speaking Medical Students: Qualitative Study</article-title></title-group><contrib-group><contrib contrib-type="author" equal-contrib="yes"><name name-style="western"><surname>Alblooshi</surname><given-names>Afaf Sulaiman</given-names></name><degrees>MD, PhD</degrees><xref ref-type="aff" rid="aff1">1</xref><xref ref-type="fn" rid="equal-contrib1">*</xref></contrib><contrib contrib-type="author" equal-contrib="yes"><name name-style="western"><surname>Almarzooqi</surname><given-names>Falah Mohammed</given-names></name><degrees>MBchB, CCFP</degrees><xref ref-type="aff" rid="aff2">2</xref><xref ref-type="fn" rid="equal-contrib1">*</xref></contrib><contrib contrib-type="author" corresp="yes"><name name-style="western"><surname>Almansoori</surname><given-names>Taleb M</given-names></name><degrees>MBBS, FRCPC</degrees><xref ref-type="aff" rid="aff3">3</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Punn</surname><given-names>Anmol</given-names></name><degrees>MD, MPH</degrees><xref ref-type="aff" rid="aff1">1</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Alameen</surname><given-names>Marwa Gaffar</given-names></name><degrees>MD, MScHPE</degrees><xref ref-type="aff" rid="aff1">1</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Kieu</surname><given-names>Alexander</given-names></name><degrees>MD</degrees><xref ref-type="aff" rid="aff2">2</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>AlRadini</surname><given-names>Faten Abdullah</given-names></name><degrees>MD</degrees><xref ref-type="aff" rid="aff4">4</xref></contrib></contrib-group><aff id="aff1"><institution>Department of Medical Education, College of Medicine and Health Sciences, United Arab Emirates University</institution><addr-line>Al Ain</addr-line><addr-line>Abu Dhabi</addr-line><country>United Arab Emirates</country></aff><aff id="aff2"><institution>Department of Family Medicine, College of Medicine and Health Sciences, United Arab Emirates University</institution><addr-line>Al Ain</addr-line><addr-line>Abu Dhabi</addr-line><country>United Arab Emirates</country></aff><aff id="aff3"><institution>Department of Radiology, College of Medicine and Health Sciences, United Arab Emirates University</institution><addr-line>Sheikh Khalifa Street</addr-line><addr-line>Al Ain</addr-line><addr-line>Abu Dhabi</addr-line><country>United Arab Emirates</country></aff><aff id="aff4"><institution>Department of Family and Community Medicine, College of Medicine, Princess Nourah bint Abdulrahman University</institution><addr-line>Riyadh</addr-line><addr-line>Riyadh Region</addr-line><country>Saudi Arabia</country></aff><contrib-group><contrib contrib-type="editor"><name name-style="western"><surname>Stone</surname><given-names>Alicia</given-names></name></contrib></contrib-group><contrib-group><contrib contrib-type="reviewer"><name name-style="western"><surname>Clark</surname><given-names>Adrian</given-names></name></contrib><contrib contrib-type="reviewer"><name name-style="western"><surname>Yang</surname><given-names>Yue</given-names></name></contrib></contrib-group><author-notes><corresp>Correspondence to Taleb M Almansoori, MBBS, FRCPC, Department of Radiology, College of Medicine and Health Sciences, United Arab Emirates University, Sheikh Khalifa Street, Al Ain, Abu Dhabi, United Arab Emirates, 971 37137558; <email>taleb.almansoor@uaeu.ac.ae</email></corresp><fn fn-type="equal" id="equal-contrib1"><label>*</label><p>these authors contributed equally</p></fn></author-notes><pub-date pub-type="collection"><year>2026</year></pub-date><pub-date pub-type="epub"><day>28</day><month>9</month><year>2026</year></pub-date><volume>12</volume><elocation-id>e95904</elocation-id><history><date date-type="received"><day>23</day><month>03</month><year>2026</year></date><date date-type="rev-recd"><day>26</day><month>08</month><year>2026</year></date><date date-type="accepted"><day>26</day><month>08</month><year>2026</year></date></history><copyright-statement>&#x00A9; Afaf Sulaiman Alblooshi, Falah Mohammed Almarzooqi, Taleb Mohamed Almansoori, Anmol Punn, Marwa Gaffar Alameen, Alexander Kieu, Faten Abdullah AlRadini. Originally published in JMIR Medical Education (<ext-link ext-link-type="uri" xlink:href="https://mededu.jmir.org">https://mededu.jmir.org</ext-link>), 28.9.2026. </copyright-statement><copyright-year>2026</copyright-year><license license-type="open-access" xlink:href="https://creativecommons.org/licenses/by/4.0/"><p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (<ext-link ext-link-type="uri" xlink:href="https://creativecommons.org/licenses/by/4.0/">https://creativecommons.org/licenses/by/4.0/</ext-link>), which permits unrestricted use, distribution, and reproduction in any medium, provided the original work, first published in JMIR Medical Education, is properly cited. The complete bibliographic information, a link to the original publication on <ext-link ext-link-type="uri" xlink:href="https://mededu.jmir.org/">https://mededu.jmir.org/</ext-link>, as well as this copyright and license information must be included.</p></license><self-uri xlink:type="simple" xlink:href="https://mededu.jmir.org/2026/1/e95904"/><abstract><sec><title>Background</title><p>Extended reality technologies, including virtual reality (VR), augmented reality, and mixed reality, are increasingly used in medical education to create immersive and interactive learning environments. As these modalities expand, validated instruments are needed to measure learners&#x2019; experiences accurately. The Immersive Technology Evaluation Measure (ITEM) is a multidomain questionnaire assessing immersion, motivation, cognitive load, usability, and debriefing. Although cognitive interviewing informed its original development, less is known about how ITEM questions function when used in a different linguistic and educational context.</p></sec><sec><title>Objective</title><p>This study aimed to evaluate the clarity, comprehensibility, and response processes of ITEM among native Arabic-speaking medical students enrolled in an English-medium medical program. We also sought to identify linguistic, referential, and contextual sources of comprehension difficulty and use these findings to inform proposed item-level refinements.</p></sec><sec sec-type="methods"><title>Methods</title><p>We conducted a qualitative cognitive interviewing study with 10 third-year and fourth-year medical students at United Arab Emirates University following VR-based clinical skills activities. Using concurrent think-aloud and verbal probing techniques, participants explained how they interpreted each ITEM question and arrived at their responses. Interview transcripts, recordings, and interviewer notes were independently reviewed using a descriptive, item-focused approach. Participant feedback was examined for recurring comprehension problems, and proposed item-level decisions were reviewed through research team consensus. Item-level saturation was reached after 8 interviews and confirmed with 2 additional interviews.</p></sec><sec sec-type="results"><title>Results</title><p>Participants identified comprehension or response-process problems in 47.5% (19/40) of ITEM questions. Difficulties occurred across all 5 domains and were most frequent in immersion (6/9, 66.7%) and usability (6/10, 60%), followed by debriefing (4/5, 80%), motivation (2/10, 20%), and cognitive load (1/6, 16.7%). Common difficulties involved nonspecific references such as &#x201C;activity&#x201D; and &#x201C;technology,&#x201D; unfamiliar or abstract terminology, negative wording, and unclear temporal or contextual framing. For example, &#x201C;concern&#x201D; was interpreted by some participants as worry rather than focus, and &#x201C;mentally demanding&#x201D; was interpreted in relation to mental health rather than cognitive effort. Cognitive interviewing informed different item-level decisions rather than a uniform revision: 11 problematic questions received proposed wording or contextual revisions, 2 received presentation-only modifications, and 6 were retained unchanged where clarification could introduce a meaning not clearly established in the original item. Six additional questions received limited contextual or terminology-standardizing changes for consistency. No items were removed, and the original 5-domain structure was retained.</p></sec><sec sec-type="conclusions"><title>Conclusions</title><p>Cognitive interviewing identified response-process difficulties that were not apparent from the questionnaire wording alone and provided a systematic basis for determining when clarification was warranted and when the original wording should be preserved. The findings extend response-process evidence for ITEM and illustrate the value of examining established educational measures in new linguistic and educational settings. The proposed refinements provide a foundation for further cognitive testing and psychometric evaluation across immersive learning contexts.</p></sec></abstract><kwd-group><kwd>interview</kwd><kwd>extended reality</kwd><kwd>virtual reality</kwd><kwd>immersive technology</kwd><kwd>interactive learning environments</kwd><kwd>medical education</kwd></kwd-group></article-meta></front><body><sec id="s1" sec-type="intro"><title>Introduction</title><p>Extended reality (XR) technologies, including virtual reality (VR), augmented reality (AR), and mixed reality (MR), are increasingly used in medical education to create immersive and interactive learning environments [<xref ref-type="bibr" rid="ref1">1</xref>-<xref ref-type="bibr" rid="ref4">4</xref>]. These modalities allow learners to engage with simulated clinical scenarios that support decision-making, procedural skills practice, and teamwork in a safe and engaging learning environment [<xref ref-type="bibr" rid="ref5">5</xref>,<xref ref-type="bibr" rid="ref6">6</xref>]. As immersive learning platforms expand in medical education, there is a parallel need for validated instruments that can reliably assess learners&#x2019; experiences and outcomes [<xref ref-type="bibr" rid="ref7">7</xref>].</p><p>Several established questionnaires assess individual components of immersive learning; however, these tools are often applied independently and may fail to capture the broader, multidimensional nature of immersive learning experiences. For example, usability and technology acceptance are commonly assessed using the System Usability Scale (SUS) [<xref ref-type="bibr" rid="ref8">8</xref>] and the technology acceptance model [<xref ref-type="bibr" rid="ref9">9</xref>], presence and immersion using the iGroup Presence Questionnaire [<xref ref-type="bibr" rid="ref10">10</xref>], perceived cognitive workload using the NASA Task Load Index (NASA-TLX) [<xref ref-type="bibr" rid="ref11">11</xref>], and learner motivation using the intrinsic motivation inventory [<xref ref-type="bibr" rid="ref12">12</xref>]. This limitation is particularly relevant in XR-based education, where learners&#x2019; perceptions may be shaped simultaneously by the level of immersion, system usability, motivation, cognitive demands, and opportunities for structured reflection. A comprehensive measure integrating complementary constructs can therefore support a more complete evaluation of learner experience in immersive technology-enhanced learning (TEL).</p><p>To address the need for a multidomain evaluation instrument, Jacobs et al [<xref ref-type="bibr" rid="ref7">7</xref>] developed the Immersive Technology Evaluation Measure (ITEM), a 40-item self-report questionnaire designed to assess key dimensions of the learner&#x2019;s experience in immersive environments. ITEM integrates 5 established measures and frameworks: the Adapted Immersion Experience Questionnaire (AIEQ) [<xref ref-type="bibr" rid="ref13">13</xref>,<xref ref-type="bibr" rid="ref14">14</xref>], abridged intrinsic motivation inventory (AIMI) [<xref ref-type="bibr" rid="ref14">14</xref>,<xref ref-type="bibr" rid="ref15">15</xref>], NASA-TLX [<xref ref-type="bibr" rid="ref11">11</xref>], SUS [<xref ref-type="bibr" rid="ref8">8</xref>], and PEARLS (Prompts for Engaging and Reflective Learning in Simulation) debriefing tool [<xref ref-type="bibr" rid="ref16">16</xref>]. Together, these domains assess immersion, intrinsic motivation, cognitive load, system usability, and debriefing. Cognitive interviewing was included in the original development of ITEM to assess participants&#x2019; understanding of its questions [<xref ref-type="bibr" rid="ref7">7</xref>]. However, cognitive evaluation during initial instrument development does not necessarily establish that questions will be interpreted similarly when an instrument is subsequently used in different linguistic, cultural, or educational contexts.</p><p>Questionnaires are widely used in medical education research to evaluate learners&#x2019; perceptions, attitudes, and experiences [<xref ref-type="bibr" rid="ref17">17</xref>-<xref ref-type="bibr" rid="ref20">20</xref>]. Notably, even survey items that appear straightforward to researchers can be interpreted differently by respondents, potentially introducing misunderstanding and response error [<xref ref-type="bibr" rid="ref21">21</xref>-<xref ref-type="bibr" rid="ref23">23</xref>]. Cognitive interviewing is a qualitative questionnaire-evaluation approach that examines how respondents understand questions and formulate responses [<xref ref-type="bibr" rid="ref22">22</xref>-<xref ref-type="bibr" rid="ref24">24</xref>]. The survey response process has commonly been conceptualized in terms of comprehension, retrieval of relevant information, judgment, and response formulation [<xref ref-type="bibr" rid="ref24">24</xref>]. Cognitive interviewing uses techniques such as think-aloud (TA) and verbal probing (VP) to explore these processes and identify problems involving terminology, ambiguous wording, question complexity, and contextual interpretation [<xref ref-type="bibr" rid="ref22">22</xref>,<xref ref-type="bibr" rid="ref25">25</xref>-<xref ref-type="bibr" rid="ref28">28</xref>]. It has also been used to examine whether questionnaire items function similarly across cultural and linguistic contexts [<xref ref-type="bibr" rid="ref28">28</xref>-<xref ref-type="bibr" rid="ref30">30</xref>].</p><p>This consideration is particularly relevant when respondents complete questionnaires in a language that is not their first language. Previous research has demonstrated that language proficiency and linguistic characteristics of survey questions, including complex syntax, ambiguous wording, and less familiar terminology, warrant consideration when questionnaires are administered across linguistic groups [<xref ref-type="bibr" rid="ref31">31</xref>,<xref ref-type="bibr" rid="ref32">32</xref>]. Such issues are relevant to medical education in the United Arab Emirates, where medical training may be delivered in English within a predominantly Arabic-speaking clinical and social environment [<xref ref-type="bibr" rid="ref33">33</xref>]. Research involving Arabic-speaking university students has also demonstrated differences in knowledge and attitude scores when the same health-related questionnaire was administered in Arabic and English [<xref ref-type="bibr" rid="ref34">34</xref>]. These findings do not imply inadequate English proficiency among learners studying in English-medium programs; rather, they highlight the potential value of examining how specific questionnaire terms and expressions are interpreted when the survey language differs from respondents&#x2019; first language.</p><p>Although cognitive interviewing was incorporated into the original development of ITEM [<xref ref-type="bibr" rid="ref7">7</xref>], limited evidence is available regarding how its established English-language questions are interpreted by learners whose first language is not English when the instrument is administered without translation. Accordingly, this study aimed to evaluate the clarity, comprehensibility, and response processes of ITEM among native Arabic-speaking medical students enrolled in an English-medium medical program in the United Arab Emirates after VR-based clinical skills activities. Using cognitive interviewing with TA and VP techniques, we examined how participants interpreted individual ITEM questions and identified linguistic, referential, and contextual sources of comprehension difficulty. The findings were used to inform proposed linguistic or contextual clarifications where appropriate, while retaining the original wording when a modification could not be made confidently without potentially altering the intended meaning.</p></sec><sec id="s2" sec-type="methods"><title>Methods</title><sec id="s2-1"><title>Study Design</title><p>This qualitative study used cognitive interviewing to evaluate the clarity, comprehensibility, and response processes of ITEM among native Arabic-speaking medical students enrolled in an English-medium medical program. We selected cognitive interviewing as an evidence-based questionnaire evaluation method to identify potential sources of misunderstanding and response error prior to broader administration of the instrument. Reporting was guided by the Cognitive Interviewing Reporting Framework (CIRF) [<xref ref-type="bibr" rid="ref28">28</xref>] (<xref ref-type="supplementary-material" rid="app5">Checklist 1</xref>).</p></sec><sec id="s2-2"><title>Setting</title><p>We conducted the study in September 2025 at the Simulation Center (iSTAR), College of Medicine and Health Sciences, United Arab Emirates University (UAEU), in the United Arab Emirates. The undergraduate medical curriculum is delivered in English, while Arabic is the first language of the enrolled United Arab Emirates&#x2013;national student population. The UAEU medical program is a 6-year curriculum with premedical, preclinical, and clinical phases; we recruited participants from years 3 and 4 because these preclinical years represent a transition into structured clinical skills training supported by simulation-based education.</p></sec><sec id="s2-3"><title>Ethical Considerations</title><p>The Social Sciences Ethics Committee at the UAEU approved the study (ERSC_2025_8044). Participation was voluntary. Participants received no financial compensation or other incentives for participation. We obtained written informed consent electronically prior to participation using REDCap, a secure, web-based data capture platform designed to support research data collection and management [<xref ref-type="bibr" rid="ref35">35</xref>,<xref ref-type="bibr" rid="ref36">36</xref>]. Participants consented to audiovisual recording for transcription and analysis and were informed that they could decline to answer any question or withdraw at any time without penalty. We stored data securely with access restricted to the research team. We removed identifying information during transcript processing and analyzed anonymized transcripts to protect participant privacy and confidentiality.</p></sec><sec id="s2-4"><title>Instrument</title><p>ITEM is a 40-item multidomain questionnaire designed to capture user experience in immersive TEL environments [<xref ref-type="bibr" rid="ref7">7</xref>] (<xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>). Participants completed the original English-language ITEM wording during the cognitive interviews. The questionnaire integrates 5 domains derived from established frameworks and validated measures: AIEQ [<xref ref-type="bibr" rid="ref13">13</xref>,<xref ref-type="bibr" rid="ref14">14</xref>] measures immersion and system fidelity, AIMI [<xref ref-type="bibr" rid="ref14">14</xref>,<xref ref-type="bibr" rid="ref15">15</xref>] assesses learner motivation and perceived educational value, NASA-TLX [<xref ref-type="bibr" rid="ref11">11</xref>] evaluates cognitive and physical load associated with task performance, SUS [<xref ref-type="bibr" rid="ref8">8</xref>] assesses confidence and perceived accessibility of the technology, and PEARLS [<xref ref-type="bibr" rid="ref16">16</xref>] explores reflective learning and debriefing processes.</p><p>Proposed wording decisions arising from the cognitive interviews are presented separately and were not administered in a second round of cognitive testing.</p></sec><sec id="s2-5"><title>Participants and Recruitment</title><p>Participants were undergraduate medical students in years 3 and 4 of the UAEU undergraduate medical program who had completed the preliminary VR clinical skills activity. Twenty students participated in the preliminary activity and were informed verbally about the cognitive interview study; all were subsequently invited by email. Ten students volunteered to participate. Interviews were intended to continue until no new item-level comprehension issues emerged rather than to achieve a predetermined fixed sample size. Saturation was identified after 8 interviews through ongoing review and team discussion, and 2 additional interviews were conducted to confirm that no new relevant item-level issues emerged. All 10 participants were United Arab Emirates nationals and native Arabic speakers. English-language proficiency and prior language of schooling were not formally assessed or collected.</p></sec><sec id="s2-6"><title>VR Experience (Context for Questionnaire Evaluation)</title><p>Prior to the cognitive interviews, participants completed a curriculum-aligned VR clinical skills session. Participants also completed a presession demographic and XR-background questionnaire in REDCap (<xref ref-type="supplementary-material" rid="app2">Multimedia Appendix 2</xref>). Facilitators first provided a technical orientation explaining the VR equipment, headset controls, and activity sequence. All participants used the same VR platform and followed the same instructional process, although the clinical module differed according to year-level curricular content. Year 3 students were randomly allocated to 1 of 2 modules (vital signs, n=3; and intravenous cannula insertion, n=1), and year 4 students completed a nasogastric tube insertion module (n=6). Each student completed the guided procedure individually. After the VR activity, participants provided verbal reflections organized broadly around strengths, weaknesses, opportunities, and threats. This activity was intended to elicit general feedback on usability, engagement, and instructional value and did not constitute a formal structured simulation debrief or use the PEARLS framework. The study did not compare cognitive interview findings across VR modules. <xref ref-type="fig" rid="figure1">Figure 1</xref> provides an overview of the study process.</p><fig position="float" id="figure1"><label>Figure 1.</label><caption><p>Workflow of the structured virtual reality (VR) clinical skills session and subsequent cognitive interview evaluation of the Immersive Technology Evaluation Measure (ITEM). IV: intravenous; SWOT: strengths, weaknesses, opportunities, and threats.</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="mededu_v12i1e95904_fig01.png"/></fig></sec><sec id="s2-7"><title>Cognitive Interview Procedure</title><p>Following the VR learning sessions, year 3 (n=4) and year 4 (n=6) students participated in cognitive interviews. Two faculty members (FMA and AK) conducted the interviews online via Microsoft Teams (version 25240.1603.3956.3103). Both interviewers were experienced in medical education and familiar with the use of cognitive interviewing procedures. Neither interviewer participated in the students&#x2019; VR teaching sessions or assessment. With participant permission, we recorded interviews and transcribed them verbatim. Interviews were scheduled for 30 minutes and lasted 20 to 45 minutes depending on participant engagement. Participants were asked to keep their cameras on to support observation of nonverbal cues such as facial expressions or hesitation. Interviewers used a combination of concurrent and retrospective cognitive interviewing techniques [<xref ref-type="bibr" rid="ref22">22</xref>]: TA (concurrent) [<xref ref-type="bibr" rid="ref26">26</xref>], whereby participants were prompted to think aloud while reading each survey item aloud, verbalizing their thoughts, interpretations, and reasoning as they answered the questions; VP (concurrent) [<xref ref-type="bibr" rid="ref27">27</xref>], whereby interviewers used a combination of prescripted and spontaneous probes to clarify interpretation, meaning, and response selection; and retrospective debriefing [<xref ref-type="bibr" rid="ref22">22</xref>], whereby, after completing the questionnaire, participants were asked to reflect on overall clarity, confusing terms, and any culturally unfamiliar wording.</p><p>Interviewers used a combination of prescripted probes and spontaneous follow-up probes when participants hesitated, requested clarification, or provided an interpretation requiring further exploration.</p></sec><sec id="s2-8"><title>Data Management and Item-Level Analysis</title><p>We transcribed recordings using Otter AI (version 3.93.0; Otter.ai, Inc) and manually checked each transcript against the recording for accuracy. Identifying information was removed prior to analysis. Two researchers (AP and MGA) independently reviewed all transcripts, recordings, and interviewer notes. Participant comments were organized in Microsoft Excel by ITEM question, together with observations of hesitation, requests for clarification, and alternative interpretations.</p><p>The analysis was descriptive and item-focused rather than a formal thematic analysis. For each ITEM question, the researchers first summarized the specific comprehension or response-process difficulties raised by participants. These item-level observations were then compared across questions to identify recurring types of difficulty. Recurring problems were organized into descriptive categories comprising ambiguous referents, unfamiliar or abstract terminology, negative wording, temporal ambiguity, and contextual relevance. These categories were used to organize and summarize the item-level findings; they were not intended to represent a formal qualitative coding framework or thematic analysis. An item was considered problematic when participants expressed confusion, requested clarification, hesitated because of uncertainty about meaning, or provided differing interpretations that indicated potential ambiguity. Item-level findings, participant-reported problems, modification decisions, and the rationale for retaining or revising each question were documented in an item-level audit (<xref ref-type="supplementary-material" rid="app3">Multimedia Appendix 3</xref>).</p><p>AP and MGA compared their item-level assessments and assignment of descriptive problem categories after independent review and did not identify substantive disagreements. The wider research team subsequently reviewed the participant comments and proposed item-level decisions. Any differences during team discussion were resolved through consensus. No specialized qualitative analysis software or interrater agreement statistic was used.</p><p>Proposed wording changes were kept deliberately conservative. Where participant feedback indicated a comprehension problem that could be clarified without intentionally changing the item meaning, the team considered a brief linguistic or contextual clarification. Where the intended scope of an item could not be established confidently from the original wording or source measure, the original wording was retained rather than imposing a new interpretation. <xref ref-type="supplementary-material" rid="app4">Multimedia Appendix 4</xref> presents the proposed wording and presentation modifications. A native English-speaking language expert from the United States reviewed the proposed wording for linguistic clarity. The proposed wording was not subjected to a second round of cognitive interviews in this study.</p></sec><sec id="s2-9"><title>Rigor and Reflexivity</title><p>To support consistency and transparency, 2 researchers independently reviewed the interview transcripts, recordings, and interviewer notes, and item-level findings were subsequently discussed with the wider research team. Neither cognitive interviewer participated in the students&#x2019; VR teaching sessions or assessment. Because several members of the research team were involved in immersive TEL implementation, the team considered how their familiarity with VR and ITEM could influence interpretation of participant comments. Proposed item-level decisions were therefore discussed collectively and grounded in participants&#x2019; statements, observed comprehension difficulties, and interviewer notes, with any differences in interpretation resolved through consensus.</p></sec></sec><sec id="s3" sec-type="results"><title>Results</title><sec id="s3-1"><title>Participants</title><p>Ten undergraduate medical students participated in the cognitive interviews (n=4 in year 3 and n=6 in year 4). The mean age was 20 (SD 0.5; range 20&#x2010;21) years for year 3 students and 21 (SD 0.9; range 20&#x2010;22) years for year 4 students. Nine participants were female and 1 was male. All participants were United Arab Emirates nationals and native Arabic speakers enrolled in an English-medium medical program. All participants completed a VR-based clinical skills session prior to the interviews.</p></sec><sec id="s3-2"><title>Cognitive Interview Duration and Item-Level Saturation</title><p>Interviews were scheduled for 30 minutes and lasted 20 to 45 minutes, with a mean duration of approximately 33 minutes. Item-level saturation was reached after 8 interviews, with no new comprehension problems identified. Two additional interviews were conducted to confirm saturation, and no new relevant item-level issues emerged.</p></sec><sec id="s3-3"><title>Overview of Item-Level Findings</title><p>Across the 40 questions of ITEM, participants identified 47.5% (19/40) as presenting a comprehension or response-process problem. Problems occurred across all 5 domains: immersion (6/9, 66.7%), usability (6/10, 60%), debriefing (4/5, 80%), motivation (2/10, 20%), and cognitive load (1/6, 16.7%). The most common difficulties involved ambiguous or nonspecific referents, unfamiliar or abstract terminology, negative wording, temporal ambiguity, and contextual relevance.</p><p>Of the 19 participant-identified problematic questions, 11 (57.9%) received proposed wording or contextual revisions, 2 (10.5%) received presentation-only modifications without alteration of the original wording, and 6 (31.6%) were retained in their original wording. An additional 6 questions that were not independently identified as problematic received limited contextual or terminology-standardizing changes for consistency. The complete item-level findings and decision rationale are provided in <xref ref-type="supplementary-material" rid="app3">Multimedia Appendix 3</xref>. Illustrative examples of revisions made following participant feedback are presented in <xref ref-type="table" rid="table1">Table 1</xref>.</p><table-wrap id="t1" position="float"><label>Table 1.</label><caption><p>Illustrative cognitive interview findings and proposed item-level decisions for the Immersive Technology Evaluation Measure (ITEM)<sup><xref ref-type="table-fn" rid="table1fn1">a</xref></sup>.</p></caption><table id="table1" frame="hsides" rules="groups"><thead><tr><td align="left" valign="top">Original question (Q)</td><td align="left" valign="top">Problem identified</td><td align="left" valign="top">Representative participant comment</td><td align="left" valign="top">Item-level decision</td></tr></thead><tbody><tr><td align="left" valign="top" colspan="4">Section 1: immersion</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Q1: I was interested in seeing how the activity would progress</td><td align="left" valign="top">Ambiguous reference to &#x201C;activity&#x201D;</td><td align="left" valign="top">&#x201C;Which activity are they talking about here?&#x201D;<break/>&#x201C;What activity do you mean? So, the VR in general, like the whole study, just my experience with the VR? So, I think specifying the activity would be good.&#x201D;</td><td align="left" valign="top">Proposed contextual clarification: &#x201C;I was interested in seeing how the activity (simulation session using VR) would progress.&#x201D;</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Q6: At the time, the activity was my only concern</td><td align="left" valign="top">&#x201C;Concern&#x201D; interpreted as worry rather than focus</td><td align="left" valign="top">&#x201C;Why concern? It was not a concern. It was something nice.&#x201D;</td><td align="left" valign="top">Proposed plain-language clarification: &#x201C;At the time, the activity was my only concern (focus).&#x201D;</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Q7: I wanted to learn more about the outcome following the activity</td><td align="left" valign="top">Unclear &#x201C;outcome&#x201D; referent.</td><td align="left" valign="top">&#x201C;Outcome of what?&#x201D;<break/>&#x201C;Like, would it be applied or introduced in our university? Is that what they mean?&#x201D;</td><td align="left" valign="top">Original wording retained. Participants interpreted &#x201C;outcome&#x201D; differently, and the intended referent could not be established confidently without imposing a new interpretation.</td></tr><tr><td align="left" valign="top" colspan="4">Section 2: motivation</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Q13: This activity did not hold my attention at all</td><td align="left" valign="top">Negative wording and response directionality</td><td align="left" valign="top">&#x201C;What do they mean by this question?&#x201D;<break/>&#x201C;Okay? No, I was like immersed in it, but I am still confused actually with the DID NOT HOLD.&#x201D;</td><td align="left" valign="top">No wording revision. &#x201C;DID NOT&#x201D; was capitalized as a typographic emphasis to draw attention to the negative construction: &#x201C;This activity DID NOT hold my attention at all.&#x201D;</td></tr><tr><td align="left" valign="top" colspan="4">Section 3: cognitive load</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Q20: How mentally demanding was the task?</td><td align="left" valign="top">Misinterpretation of the phrase &#x201C;mentally demanding&#x201D;</td><td align="left" valign="top">&#x201C;How much I had focused into the task, or did I understand it wrong?&#x201D;<break/>&#x201C;When I read mentally demanding, I thought about mental health and not in a thinking way.&#x201D;</td><td align="left" valign="top">Proposed plain-language clarification: &#x201C;How mentally demanding (needs a lot of focus) was the task?&#x201D;</td></tr><tr><td align="left" valign="top" colspan="4">Section 4: usability</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Q27: I found the technology unnecessarily complex</td><td align="left" valign="top">Difficulty interpreting &#x201C;unnecessarily complex&#x201D;</td><td align="left" valign="top">&#x201C;So, it is not complex, or it is complex?&#x201D;<break/>&#x201C;This word. Do they mean I found the technology complex?&#x201D;</td><td align="left" valign="top">Proposed clarification and contextual specification: &#x201C;I found the VR technology unnecessarily complex (more complicated than necessary).&#x201D;</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Q29: I think that I would need the support of a technical person to be able to use this technology</td><td align="left" valign="top">Unclear temporal scope (initial vs ongoing support)</td><td align="left" valign="top">&#x201C;Do you mean every time I use it, or only the first time?&#x201D;<break/>&#x201C;I think the first time, I think the first time we need someone, but after that, I don&#x2019;t think we need someone. So, I don&#x2019;t know what to put my answer.&#x201D;</td><td align="left" valign="top">Original wording retained. The source item does not specify a temporal frame; therefore, no clarification was introduced.</td></tr></tbody></table><table-wrap-foot><fn id="table1fn1"><p><sup>a</sup>Examples are illustrative; complete item-level findings and decision rationales for all 40 ITEM questions are provided in <xref ref-type="supplementary-material" rid="app3">Multimedia Appendix 3</xref>. Proposed wording and presentation modifications were not cognitively retested in this study.</p></fn></table-wrap-foot></table-wrap></sec><sec id="s3-4"><title>Patterns of Comprehension Difficulty</title><p>A recurring issue across domains was the use of nonspecific references such as &#x201C;activity&#x201D; and &#x201C;technology.&#x201D; Participants variously interpreted these as referring to the VR software, the clinical procedure, or the overall simulation session. Where appropriate, proposed contextual clarifications specified the VR activity or VR technology to reduce uncertainty regarding the referent.</p><p>Participants also reported difficulty with abstract or unfamiliar terms. For example, some were uncertain about the meaning of &#x201C;immersion,&#x201D; and one participant interpreted &#x201C;mentally demanding&#x201D; as referring to mental health rather than the amount of cognitive effort required. Negatively worded statements also created response uncertainty. In response, typographic emphasis was proposed for &#x201C;DID NOT&#x201D; in question 13 without altering the wording of the item.</p><p>Not all identified problems resulted in wording changes. For question 7, participants differed in their interpretation of &#x201C;outcome,&#x201D; and for question 29, they questioned whether technical support referred to initial or ongoing use. Because these meanings could not be established confidently from the original wording, the original questions were retained rather than introducing a clarification that could alter their intended scope.</p></sec><sec id="s3-5"><title>Debriefing Domain</title><p>Participants identified comprehension difficulties in 4 of the 5 debriefing-domain questions. However, participants had not undergone a formal structured postactivity debrief; the brief strengths, weaknesses, opportunities, and threats&#x2013;style reflection following the VR activity was not based on the PEARLS framework. Consequently, these difficulties could not be attributed confidently to questionnaire wording alone and may have reflected limited contextual applicability. The original wording of all debriefing-domain questions was therefore retained.</p></sec><sec id="s3-6"><title>General Feedback</title><p>Across interviews, participants described some questions as broad, formal, or unfamiliar and indicated that brief contextual cues or plain-language explanations could support comprehension of selected terms. Most identified issues concerned wording or contextual interpretation rather than the overall structure of ITEM. No participant feedback indicated a need to remove an ITEM domain or change the questionnaire&#x2019;s overall multidomain structure.</p></sec></sec><sec id="s4" sec-type="discussion"><title>Discussion</title><sec id="s4-1"><title>Principal Findings</title><p>This study used cognitive interviewing to evaluate the clarity, comprehensibility, and response processes of the ITEM among native Arabic-speaking medical students enrolled in an English-medium medical program. Participants identified comprehension or response-process difficulties across all 5 ITEM domains, including ambiguous referents, unfamiliar or abstract terminology, negative wording, temporal ambiguity, and contextual relevance.</p><p>The findings demonstrated that response-process difficulties can emerge when an established English-language instrument is used in a different linguistic and educational context, even when learners study in English. Importantly, cognitive interviewing informed not only where clarification might be useful but also where retaining the original wording was more appropriate because the intended meaning could not be established confidently. The resulting item-level decisions therefore included proposed linguistic clarifications, contextual specifications, presentation changes, and deliberate retention of original wording. These findings provide response-process evidence for ITEM in this context but do not demonstrate that the proposed revisions themselves improved interpretability.</p></sec><sec id="s4-2"><title>Comparison With Prior Work</title><p>Our findings are consistent with cognitive interviewing literature showing that respondents may interpret apparently straightforward questionnaire items differently according to wording, context, and linguistic experience [<xref ref-type="bibr" rid="ref22">22</xref>,<xref ref-type="bibr" rid="ref23">23</xref>,<xref ref-type="bibr" rid="ref28">28</xref>,<xref ref-type="bibr" rid="ref29">29</xref>]. Cross-cultural cognitive interviewing research similarly emphasizes that equivalent interpretation should not be assumed solely because respondents can complete an instrument in the language in which it was developed [<xref ref-type="bibr" rid="ref29">29</xref>,<xref ref-type="bibr" rid="ref30">30</xref>].</p><p>Research on survey response across linguistic groups also supports consideration of language proficiency and linguistic complexity when questionnaires are administered across languages [<xref ref-type="bibr" rid="ref31">31</xref>,<xref ref-type="bibr" rid="ref32">32</xref>]. In an Arabic-speaking context, differences have been reported when the same health-related questionnaire was administered in Arabic and English [<xref ref-type="bibr" rid="ref34">34</xref>]. Although these studies differ from the present study in design and population, they support examining response processes when questionnaire language differs from respondents&#x2019; first language. The present study extends this work by examining an established English-language immersive technology measure without translating it, reflecting how ITEM would be used in an English-medium medical education setting in the United Arab Emirates [<xref ref-type="bibr" rid="ref33">33</xref>].</p></sec><sec id="s4-3"><title>Domain-Specific and Contextual Findings</title><p>Comprehension problems occurred across all ITEM domains but were most frequent in the immersion and usability domains. Several of these questions contained abstract terminology or broad referents such as &#x201C;activity,&#x201D; &#x201C;technology,&#x201D; and &#x201C;outcome,&#x201D; which participants interpreted in different ways. Where clarification could be introduced without intentionally changing the item scope, brief contextual or plain-language explanations were proposed.</p><p>Findings from the debriefing domain required separate interpretation. Participants identified difficulties with 4 of the 5 debriefing questions; however, they had not undergone a formal structured postactivity debrief. The brief strengths, weaknesses, opportunities, and threats&#x2013;style reflection was not a PEARLS-based debriefing process. Because PEARLS was developed specifically for structured simulation debriefing [<xref ref-type="bibr" rid="ref7">7</xref>,<xref ref-type="bibr" rid="ref16">16</xref>], these difficulties may have reflected limited contextual applicability rather than wording alone. The original wording of all debriefing questions was therefore retained.</p><p>The cognitive interview findings also highlighted the importance of avoiding overcorrection. For question 7, the meaning of &#x201C;outcome&#x201D; remained uncertain, and for question 29, the original item did not specify whether technical support referred to initial or ongoing use. Rather than introducing a meaning not clearly established in the source items, we retained the original wording. This conservative approach recognizes that cognitive interviewing can identify ambiguity without necessarily providing sufficient evidence to redefine an item [<xref ref-type="bibr" rid="ref28">28</xref>,<xref ref-type="bibr" rid="ref29">29</xref>], and even apparently minor questionnaire adaptations should be considered carefully because they may affect item functioning [<xref ref-type="bibr" rid="ref37">37</xref>].</p></sec><sec id="s4-4"><title>Implications for Medical Education and Instrument Use</title><p>These findings support cognitive interviewing as a practical method for examining response-process evidence when established educational measures are applied beyond their original development context [<xref ref-type="bibr" rid="ref22">22</xref>,<xref ref-type="bibr" rid="ref23">23</xref>,<xref ref-type="bibr" rid="ref28">28</xref>,<xref ref-type="bibr" rid="ref29">29</xref>]. For ITEM, this study provides an item-level account of where participants encountered uncertainty and documents the rationale for proposed linguistic, contextual, or presentation modifications. Because the proposed wording was not cognitively retested, these modifications should be considered provisional rather than validated improvements.</p><p>The addition of &#x201C;VR&#x201D; to selected proposed clarifications does not make ITEM a VR-only instrument. ITEM was developed for immersive TEL and was originally evaluated in both VR and AR contexts [<xref ref-type="bibr" rid="ref7">7</xref>]. In this study, &#x201C;VR&#x201D; was added only where participants found the generic term &#x201C;technology&#x201D; insufficiently specific. In other immersive settings, the original technology-neutral wording could be retained or the relevant modality, such as AR or MR, could be specified where needed.</p><p>More broadly, studying medicine in English does not necessarily mean that all English-language questionnaire expressions will be interpreted uniformly. Response-process evaluation may therefore complement psychometric testing when established instruments are used in different linguistic or educational settings.</p></sec><sec id="s4-5"><title>Limitations</title><p>This study has several limitations. First, the proposed wording was not cognitively retested; therefore, the study identifies response-process difficulties and proposes potential refinements but cannot establish that the revisions resolved the identified comprehension problems. Second, English-language proficiency and prior language of schooling were not formally assessed, so the findings cannot determine whether the observed comprehension difficulties varied according to individual language proficiency or educational language background. Third, only 1 male student participated, which limited our ability to examine whether interpretation differed by sex. The sample was also drawn from a single institution. These characteristics constrain subgroup comparisons and broader generalizability; however, the purpose of cognitive interviewing was to identify and characterize item-level comprehension problems rather than estimate their prevalence in the wider student population. Item-level saturation was reached after 8 interviews and confirmed with 2 additional interviews, supporting the adequacy of the sample for the study&#x2019;s primary cognitive interviewing objective. Fourth, participants did not undergo a formal structured debrief; therefore, difficulties with the PEARLS-derived items may reflect limited contextual applicability rather than questionnaire wording alone. Finally, interviews were conducted online via Microsoft Teams. In some cases, screen sharing disabled participants&#x2019; cameras, a common constraint when using tablets or iPads, which limited observation of nonverbal cues such as facial expressions and hesitations. Online interviewing may also reduce opportunities for rapport building compared with in-person interviews, which could have affected the depth of spontaneous verbalizations during the TA process.</p></sec><sec id="s4-6"><title>Future Directions</title><p>Future research should cognitively retest the proposed item-level modifications to determine whether they resolve the identified problems without introducing new interpretations. Subsequent studies should evaluate reliability, construct validity, and measurement performance in larger and more diverse samples, including learners with different linguistic backgrounds and levels of English proficiency. Further investigations across VR, AR, and MR settings and in activities involving formal structured debriefing would also help determine when modality-specific contextual wording is appropriate.</p></sec><sec id="s4-7"><title>Conclusions</title><p>Cognitive interviewing identified linguistic, referential, and contextual response-process difficulties in several ITEM questions among native Arabic-speaking medical students enrolled in an English-medium medical program. The findings informed proposed item-level clarifications while also identifying questions for which retaining the original wording was more appropriate because further modification could impose an interpretation not clearly established in the source item.</p><p>These findings provide additional response-process evidence for ITEM in a different linguistic and educational context but do not establish that the proposed wording changes improve interpretability. Further cognitive testing and subsequent psychometric evaluation are needed before the proposed refinements can be considered confirmed.</p></sec></sec></body><back><ack><p>The authors gratefully acknowledge Dr Chris Jacobs for granting permission to use the Immersive Technology Evaluation Measure (ITEM) in this study.</p><p>The authors acknowledge support from the Princess Nourah bint Abdulrahman University Researchers Supporting Project (PNURSP2026R290), Princess Nourah bint Abdulrahman University, Riyadh, Saudi Arabia.</p><p>The authors also acknowledge MediSim VR for providing access to the VR software used in this study and Dr Pradeesh Sathyan, Senior Consultant for Medical Technology, Arrownex LLC, for technical and content-related support.</p><p>The authors acknowledge Ms Jane Koester (College of Medicine and Health Sciences, United Arab Emirates University [UAEU]) for support with linguistic clarification and contextual refinement of selected questionnaire items to enhance clarity for the local context. The authors acknowledge the contributions of the medical students Latifa Mohammed Alderei, Aryam Muhsen Albreiki, Raghad Salem Alharthi, and Shaikha Ahmed Alzaabi (College of Medicine and Health Sciences, UAEU) for assistance with participant recruitment. The authors are especially appreciative of all the medical students who volunteered their time to take part in this study.</p><p>The authors acknowledge the use of OpenAI&#x2019;s ChatGPT (GPT 5.2) to support language editing and improve clarity and readability of the manuscript. All content was critically reviewed, verified, and approved by the authors. No generative AI tools were used to generate, analyze, or interpret study data.</p></ack><notes><sec><title>Funding</title><p>This research was supported by a research grant from the United Arab Emirates University (grant 12M225).</p></sec><sec><title>Data Availability</title><p>The qualitative interview transcripts are not publicly available because of the sensitive nature of the interview data and the potential risk of participant identification. Deidentified excerpts supporting the findings are included in the manuscript, and complete item-level findings and decision rationales are provided in <xref ref-type="supplementary-material" rid="app4">Multimedia Appendix 4</xref>. Additional information may be made available by the corresponding author (TMA) upon reasonable request and subject to institutional ethical approval.</p></sec></notes><fn-group><fn fn-type="con"><p>Conceptualization: ASA</p><p>Data curation: FMA, AK</p><p>Formal analysis: AP, MGA</p><p>Investigation: FMA, AK</p><p>Methodology: ASA, FAA, TMA</p><p>Validation: ASA, FMA, AP, MGA</p><p>Writing&#x2014;original draft: ASA, AP</p><p>Writing&#x2014;review and editing: ASA, AP, FMA, MGA, AK, TMA, FAA</p><p>All authors reviewed and approved the final manuscript and agree to be accountable for all aspects of the work.</p></fn><fn fn-type="conflict"><p>None declared.</p></fn></fn-group><glossary><title>Abbreviations</title><def-list><def-item><term id="abb1">AIEQ</term><def><p>Adapted Immersion Experience Questionnaire</p></def></def-item><def-item><term id="abb2">AIMI</term><def><p>abridged intrinsic motivation inventory</p></def></def-item><def-item><term id="abb3">AR</term><def><p>augmented reality</p></def></def-item><def-item><term id="abb4">ITEM</term><def><p>Immersive Technology Evaluation Measure</p></def></def-item><def-item><term id="abb5">MR</term><def><p>mixed reality</p></def></def-item><def-item><term id="abb6">NASA-TLX</term><def><p>NASA Task Load Index</p></def></def-item><def-item><term id="abb7">PEARLS</term><def><p>Prompts for Engaging and Reflective Learning in Simulation</p></def></def-item><def-item><term id="abb8">SUS</term><def><p>System Usability Scale</p></def></def-item><def-item><term id="abb9">TA</term><def><p>think-aloud</p></def></def-item><def-item><term id="abb10">TEL</term><def><p>technology-enhanced learning</p></def></def-item><def-item><term id="abb11">UAEU</term><def><p>United Arab Emirates University</p></def></def-item><def-item><term id="abb12">VP</term><def><p>verbal probing</p></def></def-item><def-item><term id="abb13">VR</term><def><p>virtual reality</p></def></def-item><def-item><term id="abb14">XR</term><def><p>extended reality</p></def></def-item></def-list></glossary><ref-list><title>References</title><ref id="ref1"><label>1</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Mergen</surname><given-names>M</given-names> </name><name name-style="western"><surname>Graf</surname><given-names>N</given-names> </name><name name-style="western"><surname>Meyerheim</surname><given-names>M</given-names> </name></person-group><article-title>Reviewing the current state of virtual reality integration in medical education - a scoping review</article-title><source>BMC Med Educ</source><year>2024</year><month>07</month><day>23</day><volume>24</volume><issue>1</issue><fpage>788</fpage><pub-id pub-id-type="doi">10.1186/s12909-024-05777-5</pub-id><pub-id pub-id-type="medline">39044186</pub-id></nlm-citation></ref><ref id="ref2"><label>2</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Tene</surname><given-names>T</given-names> </name><name name-style="western"><surname>Vique L&#x00F3;pez</surname><given-names>DF</given-names> </name><name name-style="western"><surname>Valverde Aguirre</surname><given-names>PE</given-names> </name><name name-style="western"><surname>Orna Puente</surname><given-names>LM</given-names> </name><name name-style="western"><surname>Vacacela Gomez</surname><given-names>C</given-names> </name></person-group><article-title>Virtual reality and augmented reality in medical education: an umbrella review</article-title><source>Front Digit Health</source><year>2024</year><volume>6</volume><fpage>1365345</fpage><pub-id pub-id-type="doi">10.3389/fdgth.2024.1365345</pub-id><pub-id pub-id-type="medline">38550715</pub-id></nlm-citation></ref><ref id="ref3"><label>3</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Sung</surname><given-names>H</given-names> </name><name name-style="western"><surname>Kim</surname><given-names>M</given-names> </name><name name-style="western"><surname>Park</surname><given-names>J</given-names> </name><name name-style="western"><surname>Shin</surname><given-names>N</given-names> </name><name name-style="western"><surname>Han</surname><given-names>Y</given-names> </name></person-group><article-title>Effectiveness of virtual reality in healthcare education: systematic review and meta-analysis</article-title><source>Sustainability</source><year>2024</year><volume>16</volume><issue>19</issue><fpage>8520</fpage><pub-id pub-id-type="doi">10.3390/su16198520</pub-id></nlm-citation></ref><ref id="ref4"><label>4</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Rojas-S&#x00E1;nchez</surname><given-names>MA</given-names> </name><name name-style="western"><surname>Palos-S&#x00E1;nchez</surname><given-names>PR</given-names> </name><name name-style="western"><surname>Folgado-Fern&#x00E1;ndez</surname><given-names>JA</given-names> </name></person-group><article-title>Systematic literature review and bibliometric analysis on virtual reality and education</article-title><source>Educ Inf Technol</source><year>2023</year><month>01</month><volume>28</volume><issue>1</issue><fpage>155</fpage><lpage>192</lpage><pub-id pub-id-type="doi">10.1007/s10639-022-11167-5</pub-id></nlm-citation></ref><ref id="ref5"><label>5</label><nlm-citation citation-type="book"><person-group person-group-type="author"><name name-style="western"><surname>Forrest</surname><given-names>K</given-names> </name><name name-style="western"><surname>McKimm</surname><given-names>J</given-names> </name></person-group><source>Healthcare Simulation at a Glance</source><year>2019</year><edition>1</edition><publisher-name>John Wiley &#x0026; Sons</publisher-name><pub-id pub-id-type="other">9781118871843</pub-id></nlm-citation></ref><ref id="ref6"><label>6</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Ryan</surname><given-names>GV</given-names> </name><name name-style="western"><surname>Callaghan</surname><given-names>S</given-names> </name><name name-style="western"><surname>Rafferty</surname><given-names>A</given-names> </name><name name-style="western"><surname>Higgins</surname><given-names>MF</given-names> </name><name name-style="western"><surname>Mangina</surname><given-names>E</given-names> </name><name name-style="western"><surname>McAuliffe</surname><given-names>F</given-names> </name></person-group><article-title>Learning outcomes of immersive technologies in health care student education: systematic review of the literature</article-title><source>J Med Internet Res</source><year>2022</year><month>02</month><day>1</day><volume>24</volume><issue>2</issue><fpage>e30082</fpage><pub-id pub-id-type="doi">10.2196/30082</pub-id><pub-id pub-id-type="medline">35103607</pub-id></nlm-citation></ref><ref id="ref7"><label>7</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Jacobs</surname><given-names>C</given-names> </name><name name-style="western"><surname>Wheeler</surname><given-names>J</given-names> </name><name name-style="western"><surname>Williams</surname><given-names>M</given-names> </name><name name-style="western"><surname>Joiner</surname><given-names>R</given-names> </name></person-group><article-title>Cognitive interviewing as a method to inform questionnaire design and validity - Immersive Technology Evaluation Measure (ITEM) for healthcare education</article-title><source>Comput Educ X Real</source><year>2023</year><volume>2</volume><fpage>100027</fpage><pub-id pub-id-type="doi">10.1016/j.cexr.2023.100027</pub-id></nlm-citation></ref><ref id="ref8"><label>8</label><nlm-citation citation-type="book"><person-group person-group-type="author"><name name-style="western"><surname>Brooke</surname><given-names>J</given-names> </name></person-group><person-group person-group-type="editor"><name name-style="western"><surname>Jordan</surname><given-names>PW</given-names> </name><name name-style="western"><surname>Thomas</surname><given-names>B</given-names> </name><name name-style="western"><surname>McClelland</surname><given-names>IL</given-names> </name><name name-style="western"><surname>Weerdmeester</surname><given-names>B</given-names> </name></person-group><article-title>SUS: a 'quick and dirty' usability scale</article-title><source>Usability Evaluation in Industry</source><year>1996</year><publisher-name>CRC Press</publisher-name><pub-id pub-id-type="doi">10.1201/9781498710411</pub-id></nlm-citation></ref><ref id="ref9"><label>9</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Davis</surname><given-names>FD</given-names> </name></person-group><article-title>Perceived usefulness, perceived ease of use, and user acceptance of information technology</article-title><source>MIS Q</source><year>1989</year><volume>13</volume><issue>3</issue><fpage>319</fpage><lpage>340</lpage><pub-id pub-id-type="doi">10.2307/249008</pub-id></nlm-citation></ref><ref id="ref10"><label>10</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Schubert</surname><given-names>T</given-names> </name><name name-style="western"><surname>Friedmann</surname><given-names>F</given-names> </name><name name-style="western"><surname>Regenbrecht</surname><given-names>H</given-names> </name></person-group><article-title>The experience of presence: factor analytic insights</article-title><source>Presence Teleoperators Virtual Environ</source><year>2001</year><month>06</month><volume>10</volume><issue>3</issue><fpage>266</fpage><lpage>281</lpage><pub-id pub-id-type="doi">10.1162/105474601300343603</pub-id></nlm-citation></ref><ref id="ref11"><label>11</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Hart</surname><given-names>SG</given-names> </name></person-group><article-title>NASA-task load index (NASA-TLX); 20 years later</article-title><source>Proc Hum Factors Ergon Soc Annu Meet</source><year>2006</year><month>10</month><volume>50</volume><issue>9</issue><fpage>904</fpage><lpage>908</lpage><pub-id pub-id-type="doi">10.1177/154193120605000909</pub-id></nlm-citation></ref><ref id="ref12"><label>12</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>McAuley</surname><given-names>E</given-names> </name><name name-style="western"><surname>Duncan</surname><given-names>T</given-names> </name><name name-style="western"><surname>Tammen</surname><given-names>VV</given-names> </name></person-group><article-title>Psychometric properties of the intrinsic motivation inventory in a competitive sport setting: a confirmatory factor analysis</article-title><source>Res Q Exerc Sport</source><year>1989</year><month>03</month><volume>60</volume><issue>1</issue><fpage>48</fpage><lpage>58</lpage><pub-id pub-id-type="doi">10.1080/02701367.1989.10607413</pub-id><pub-id pub-id-type="medline">2489825</pub-id></nlm-citation></ref><ref id="ref13"><label>13</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Jennett</surname><given-names>C</given-names> </name><name name-style="western"><surname>Cox</surname><given-names>AL</given-names> </name><name name-style="western"><surname>Cairns</surname><given-names>P</given-names> </name><etal/></person-group><article-title>Measuring and defining the experience of immersion in games</article-title><source>Int J Hum Comput Stud</source><year>2008</year><month>09</month><volume>66</volume><issue>9</issue><fpage>641</fpage><lpage>661</lpage><pub-id pub-id-type="doi">10.1016/j.ijhcs.2008.04.004</pub-id></nlm-citation></ref><ref id="ref14"><label>14</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Jacobs</surname><given-names>C</given-names> </name><name name-style="western"><surname>Rigby</surname><given-names>JM</given-names> </name></person-group><article-title>Developing measures of immersion and motivation for learning technologies in healthcare simulation: a pilot study</article-title><source>J Adv Med Educ Prof</source><year>2022</year><month>07</month><volume>10</volume><issue>3</issue><fpage>163</fpage><lpage>171</lpage><pub-id pub-id-type="doi">10.30476/JAMP.2022.95226.1632</pub-id><pub-id pub-id-type="medline">35910517</pub-id></nlm-citation></ref><ref id="ref15"><label>15</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Ryan</surname><given-names>RM</given-names> </name></person-group><article-title>Control and information in the intrapersonal sphere: an extension of cognitive evaluation theory</article-title><source>J Pers Soc Psychol</source><year>1982</year><volume>43</volume><issue>3</issue><fpage>450</fpage><lpage>461</lpage><pub-id pub-id-type="doi">10.1037/0022-3514.43.3.450</pub-id></nlm-citation></ref><ref id="ref16"><label>16</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Eppich</surname><given-names>W</given-names> </name><name name-style="western"><surname>Cheng</surname><given-names>A</given-names> </name></person-group><article-title>Promoting Excellence and Reflective Learning in Simulation (PEARLS): development and rationale for a blended approach to health care simulation debriefing</article-title><source>Simul Healthc</source><year>2015</year><month>04</month><volume>10</volume><issue>2</issue><fpage>106</fpage><lpage>115</lpage><pub-id pub-id-type="doi">10.1097/SIH.0000000000000072</pub-id><pub-id pub-id-type="medline">25710312</pub-id></nlm-citation></ref><ref id="ref17"><label>17</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Roff</surname><given-names>S</given-names> </name><name name-style="western"><surname>McAleer</surname><given-names>S</given-names> </name><name name-style="western"><surname>Harden</surname><given-names>RM</given-names> </name><etal/></person-group><article-title>Development and validation of the Dundee Ready Education Environment Measure (DREEM)</article-title><source>Med Teach</source><year>1997</year><volume>19</volume><issue>4</issue><fpage>295</fpage><lpage>299</lpage><pub-id pub-id-type="doi">10.3109/01421599709034208</pub-id></nlm-citation></ref><ref id="ref18"><label>18</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Marshall</surname><given-names>RE</given-names> </name></person-group><article-title>Measuring the medical school learning environment</article-title><source>J Med Educ</source><year>1978</year><month>02</month><volume>53</volume><issue>2</issue><fpage>98</fpage><lpage>104</lpage><pub-id pub-id-type="doi">10.1097/00001888-197802000-00003</pub-id><pub-id pub-id-type="medline">633337</pub-id></nlm-citation></ref><ref id="ref19"><label>19</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Wilson</surname><given-names>KL</given-names> </name><name name-style="western"><surname>Lizzio</surname><given-names>A</given-names> </name><name name-style="western"><surname>Ramsden</surname><given-names>P</given-names> </name></person-group><article-title>The development, validation and application of the Course Experience Questionnaire</article-title><source>Stud High Educ</source><year>1997</year><volume>22</volume><issue>1</issue><fpage>33</fpage><lpage>53</lpage><pub-id pub-id-type="doi">10.1080/03075079712331381121</pub-id></nlm-citation></ref><ref id="ref20"><label>20</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>AlHaqwi</surname><given-names>AI</given-names> </name><name name-style="western"><surname>Kuntze</surname><given-names>J</given-names> </name><name name-style="western"><surname>van der Molen</surname><given-names>HT</given-names> </name></person-group><article-title>Development of the Clinical Learning Evaluation Questionnaire for undergraduate clinical education: factor structure, validity, and reliability study</article-title><source>BMC Med Educ</source><year>2014</year><month>03</month><day>4</day><volume>14</volume><fpage>44</fpage><pub-id pub-id-type="doi">10.1186/1472-6920-14-44</pub-id><pub-id pub-id-type="medline">24592913</pub-id></nlm-citation></ref><ref id="ref21"><label>21</label><nlm-citation citation-type="book"><person-group person-group-type="author"><name name-style="western"><surname>Tourangeau</surname><given-names>R</given-names> </name><name name-style="western"><surname>Rips</surname><given-names>LJ</given-names> </name><name name-style="western"><surname>Rasinski</surname><given-names>K</given-names> </name></person-group><source>The Psychology of Survey Response</source><year>2000</year><publisher-name>Cambridge University Press</publisher-name><pub-id pub-id-type="doi">10.1017/CBO9780511819322</pub-id></nlm-citation></ref><ref id="ref22"><label>22</label><nlm-citation citation-type="book"><person-group person-group-type="author"><name name-style="western"><surname>Willis</surname><given-names>GB</given-names> </name></person-group><source>Cognitive Interviewing: A Tool for Improving Questionnaire Design</source><year>2004</year><publisher-name>SAGE Publications</publisher-name><pub-id pub-id-type="other">9780761928034</pub-id></nlm-citation></ref><ref id="ref23"><label>23</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Willis</surname><given-names>GB</given-names> </name><name name-style="western"><surname>Artino</surname><given-names>AR Jr</given-names> </name></person-group><article-title>What do our respondents think we&#x2019;re asking? Using cognitive interviewing to improve medical education surveys</article-title><source>J Grad Med Educ</source><year>2013</year><month>09</month><volume>5</volume><issue>3</issue><fpage>353</fpage><lpage>356</lpage><pub-id pub-id-type="doi">10.4300/JGME-D-13-00154.1</pub-id><pub-id pub-id-type="medline">24404294</pub-id></nlm-citation></ref><ref id="ref24"><label>24</label><nlm-citation citation-type="book"><person-group person-group-type="author"><name name-style="western"><surname>Tourangeau</surname><given-names>R</given-names> </name></person-group><person-group person-group-type="editor"><name name-style="western"><surname>Jabine</surname><given-names>TB</given-names> </name><name name-style="western"><surname>Straf</surname><given-names>ML</given-names> </name><name name-style="western"><surname>Tanur</surname><given-names>JM</given-names> </name><name name-style="western"><surname>Tourangeau</surname><given-names>R</given-names> </name></person-group><article-title>Cognitive sciences and survey methods</article-title><source>Cognitive Aspects of Survey Methodology: Building a Bridge between Disciplines</source><year>1984</year><publisher-name>National Academies Press</publisher-name><fpage>73</fpage><lpage>100</lpage><pub-id pub-id-type="other">9780309077842</pub-id></nlm-citation></ref><ref id="ref25"><label>25</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Beatty</surname><given-names>PC</given-names> </name><name name-style="western"><surname>Willis</surname><given-names>GB</given-names> </name></person-group><article-title>Research synthesis: the practice of cognitive interviewing</article-title><source>Public Opin Q</source><year>2007</year><volume>71</volume><issue>2</issue><fpage>287</fpage><lpage>311</lpage><pub-id pub-id-type="doi">10.1093/poq/nfm006</pub-id></nlm-citation></ref><ref id="ref26"><label>26</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Ericsson</surname><given-names>KA</given-names> </name><name name-style="western"><surname>Simon</surname><given-names>HA</given-names> </name></person-group><article-title>Verbal reports as data</article-title><source>Psychol Rev</source><year>1980</year><volume>87</volume><issue>3</issue><fpage>215</fpage><lpage>251</lpage><pub-id pub-id-type="doi">10.1037/0033-295X.87.3.215</pub-id></nlm-citation></ref><ref id="ref27"><label>27</label><nlm-citation citation-type="book"><person-group person-group-type="author"><name name-style="western"><surname>Forsyth</surname><given-names>BH</given-names> </name><name name-style="western"><surname>Lessler</surname><given-names>JT</given-names> </name></person-group><person-group person-group-type="editor"><name name-style="western"><surname>Biemer</surname><given-names>PP</given-names> </name><name name-style="western"><surname>Groves</surname><given-names>RM</given-names> </name><name name-style="western"><surname>Lyberg</surname><given-names>LE</given-names> </name><name name-style="western"><surname>Mathiowetz</surname><given-names>NA</given-names> </name><name name-style="western"><surname>Sudman</surname><given-names>S</given-names> </name></person-group><article-title>Cognitive laboratory methods: a taxonomy</article-title><source>Measurement Errors in Surveys</source><year>2004</year><publisher-name>John Wiley &#x0026; Sons</publisher-name><fpage>393</fpage><lpage>418</lpage><pub-id pub-id-type="doi">10.1002/9781118150382</pub-id></nlm-citation></ref><ref id="ref28"><label>28</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Boeije</surname><given-names>H</given-names> </name><name name-style="western"><surname>Willis</surname><given-names>G</given-names> </name></person-group><article-title>The Cognitive Interviewing Reporting Framework (CIRF): towards the harmonization of cognitive testing reports</article-title><source>Methodology</source><year>2013</year><volume>9</volume><issue>3</issue><fpage>87</fpage><lpage>95</lpage><pub-id pub-id-type="doi">10.1027/1614-2241/a000075</pub-id></nlm-citation></ref><ref id="ref29"><label>29</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Willis</surname><given-names>GB</given-names> </name><name name-style="western"><surname>Miller</surname><given-names>K</given-names> </name></person-group><article-title>Cross-cultural cognitive interviewing: seeking comparability and enhancing understanding</article-title><source>Field Methods</source><year>2011</year><volume>23</volume><issue>4</issue><fpage>331</fpage><lpage>341</lpage><pub-id pub-id-type="doi">10.1177/1525822X11416092</pub-id></nlm-citation></ref><ref id="ref30"><label>30</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Schildmann</surname><given-names>EK</given-names> </name><name name-style="western"><surname>Groeneveld</surname><given-names>EI</given-names> </name><name name-style="western"><surname>Denzel</surname><given-names>J</given-names> </name><etal/></person-group><article-title>Discovering the hidden benefits of cognitive interviewing in two languages: the first phase of a validation study of the Integrated Palliative care Outcome Scale</article-title><source>Palliat Med</source><year>2016</year><month>06</month><volume>30</volume><issue>6</issue><fpage>599</fpage><lpage>610</lpage><pub-id pub-id-type="doi">10.1177/0269216315608348</pub-id><pub-id pub-id-type="medline">26415736</pub-id></nlm-citation></ref><ref id="ref31"><label>31</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Park</surname><given-names>H</given-names> </name><name name-style="western"><surname>Sha</surname><given-names>MM</given-names> </name><name name-style="western"><surname>Willis</surname><given-names>G</given-names> </name></person-group><article-title>Influence of English-language proficiency on the cognitive processing of survey questions</article-title><source>Field Methods</source><year>2016</year><volume>28</volume><issue>4</issue><fpage>415</fpage><lpage>430</lpage><pub-id pub-id-type="doi">10.1177/1525822X16630262</pub-id></nlm-citation></ref><ref id="ref32"><label>32</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Wenz</surname><given-names>A</given-names> </name><name name-style="western"><surname>Al Baghal</surname><given-names>T</given-names> </name><name name-style="western"><surname>Gaia</surname><given-names>A</given-names> </name></person-group><article-title>Language proficiency among respondents: implications for data quality in a longitudinal face-to-face survey</article-title><source>J Surv Stat Methodol</source><year>2021</year><volume>9</volume><issue>1</issue><fpage>73</fpage><lpage>93</lpage><pub-id pub-id-type="doi">10.1093/jssam/smz045</pub-id></nlm-citation></ref><ref id="ref33"><label>33</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Ismaiel</surname><given-names>S</given-names> </name><name name-style="western"><surname>AlGhafari</surname><given-names>D</given-names> </name><name name-style="western"><surname>Ibrahim</surname><given-names>H</given-names> </name></person-group><article-title>Promoting physician-patient language concordance in undergraduate medical education: a peer assisted learning approach</article-title><source>BMC Med Educ</source><year>2023</year><month>01</month><day>3</day><volume>23</volume><issue>1</issue><fpage>1</fpage><pub-id pub-id-type="doi">10.1186/s12909-022-03986-4</pub-id><pub-id pub-id-type="medline">36593450</pub-id></nlm-citation></ref><ref id="ref34"><label>34</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Gazzaz</surname><given-names>ZJ</given-names> </name><name name-style="western"><surname>Baig</surname><given-names>M</given-names> </name><name name-style="western"><surname>Albarakati</surname><given-names>M</given-names> </name><name name-style="western"><surname>Alfalig</surname><given-names>HA</given-names> </name><name name-style="western"><surname>Jameel</surname><given-names>T</given-names> </name></person-group><article-title>Language barriers in understanding healthcare information: Arabic-speaking students&#x2019; comprehension of diabetic questionnaires in Arabic and English languages</article-title><source>Cureus</source><year>2023</year><month>10</month><volume>15</volume><issue>10</issue><fpage>e46777</fpage><pub-id pub-id-type="doi">10.7759/cureus.46777</pub-id><pub-id pub-id-type="medline">37954810</pub-id></nlm-citation></ref><ref id="ref35"><label>35</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Harris</surname><given-names>PA</given-names> </name><name name-style="western"><surname>Taylor</surname><given-names>R</given-names> </name><name name-style="western"><surname>Thielke</surname><given-names>R</given-names> </name><name name-style="western"><surname>Payne</surname><given-names>J</given-names> </name><name name-style="western"><surname>Gonzalez</surname><given-names>N</given-names> </name><name name-style="western"><surname>Conde</surname><given-names>JG</given-names> </name></person-group><article-title>Research electronic data capture (REDCap)--a metadata-driven methodology and workflow process for providing translational research informatics support</article-title><source>J Biomed Inform</source><year>2009</year><month>04</month><volume>42</volume><issue>2</issue><fpage>377</fpage><lpage>381</lpage><pub-id pub-id-type="doi">10.1016/j.jbi.2008.08.010</pub-id><pub-id pub-id-type="medline">18929686</pub-id></nlm-citation></ref><ref id="ref36"><label>36</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Harris</surname><given-names>PA</given-names> </name><name name-style="western"><surname>Taylor</surname><given-names>R</given-names> </name><name name-style="western"><surname>Minor</surname><given-names>BL</given-names> </name><etal/></person-group><article-title>The REDCap consortium: building an international community of software platform partners</article-title><source>J Biomed Inform</source><year>2019</year><month>07</month><volume>95</volume><fpage>103208</fpage><pub-id pub-id-type="doi">10.1016/j.jbi.2019.103208</pub-id><pub-id pub-id-type="medline">31078660</pub-id></nlm-citation></ref><ref id="ref37"><label>37</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Sousa</surname><given-names>VE</given-names> </name><name name-style="western"><surname>Matson</surname><given-names>J</given-names> </name><name name-style="western"><surname>Dunn Lopez</surname><given-names>K</given-names> </name></person-group><article-title>Questionnaire adapting: little changes mean a lot</article-title><source>West J Nurs Res</source><year>2017</year><month>09</month><volume>39</volume><issue>9</issue><fpage>1289</fpage><lpage>1300</lpage><pub-id pub-id-type="doi">10.1177/0193945916678212</pub-id><pub-id pub-id-type="medline">28322671</pub-id></nlm-citation></ref></ref-list><app-group><supplementary-material id="app1"><label>Multimedia Appendix 1</label><p>Immersive Technology Evaluation Measure questionnaire administered during cognitive interviews (original wording).</p><media xlink:href="mededu_v12i1e95904_app1.docx" xlink:title="DOCX File, 26 KB"/></supplementary-material><supplementary-material id="app2"><label>Multimedia Appendix 2</label><p>Presession participant questionnaire capturing demographic and background information.</p><media xlink:href="mededu_v12i1e95904_app2.docx" xlink:title="DOCX File, 18 KB"/></supplementary-material><supplementary-material id="app3"><label>Multimedia Appendix 3</label><p>Immersive Technology Evaluation Measure item-level cognitive interview audit.</p><media xlink:href="mededu_v12i1e95904_app3.xlsx" xlink:title="XLSX File, 24 KB"/></supplementary-material><supplementary-material id="app4"><label>Multimedia Appendix 4</label><p>Proposed modified Immersive Technology Evaluation Measure following cognitive interview review.</p><media xlink:href="mededu_v12i1e95904_app4.docx" xlink:title="DOCX File, 27 KB"/></supplementary-material><supplementary-material id="app5"><label>Checklist 1</label><p>Cognitive Interviewing Reporting Framework (CIRF) checklist.</p><media xlink:href="mededu_v12i1e95904_app5.docx" xlink:title="DOCX File, 21 KB"/></supplementary-material></app-group></back></article>