<?xml version="1.0" encoding="UTF-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.0 20040830//EN" "http://dtd.nlm.nih.gov/publishing/2.0/journalpublishing.dtd">
<article xmlns:xlink="http://www.w3.org/1999/xlink" article-type="research-article" dtd-version="2.0">
  <front>
    <journal-meta>
      <journal-id journal-id-type="publisher-id">JMIR Human Factors</journal-id>
      <journal-id journal-id-type="nlm-ta">JMIR Hum Factors</journal-id>
      <journal-title>JMIR Human Factors</journal-title>
      <issn pub-type="epub">2292-9495</issn>
      <publisher>
        <publisher-name>JMIR Publications</publisher-name>
        <publisher-loc>Toronto, Canada</publisher-loc>
      </publisher>
    </journal-meta>
    <article-meta>
      <article-id pub-id-type="publisher-id">v13i1e97649</article-id>
      <article-id pub-id-type="pmid">42849041</article-id>
      <article-id pub-id-type="doi">10.2196/97649</article-id>
      <article-categories>
        <subj-group subj-group-type="heading">
          <subject>Original Paper</subject>
        </subj-group>
        <subj-group subj-group-type="article-type">
          <subject>Original Paper</subject>
        </subj-group>
      </article-categories>
      <title-group>
        <article-title>Clinicians’ Trust in AI-Based Clinical Recommendations Across Controlled Primary Care–Style Clinical Vignettes: Nurse-Dominant Experimental Pilot Study</article-title>
      </title-group>
      <contrib-group>
        <contrib contrib-type="editor">
          <name>
            <surname>Law</surname>
            <given-names>Stephanie</given-names>
          </name>
        </contrib>
      </contrib-group>
      <contrib-group>
        <contrib contrib-type="reviewer">
          <name>
            <surname>Saremi</surname>
            <given-names>Mostaan</given-names>
          </name>
        </contrib>
        <contrib contrib-type="reviewer">
          <name>
            <surname>Harada</surname>
            <given-names>Yukinori</given-names>
          </name>
        </contrib>
      </contrib-group>
      <contrib-group>
        <contrib id="contrib1" contrib-type="author" corresp="yes" equal-contrib="yes">
          <name name-style="western">
            <surname>Choudhury</surname>
            <given-names>Avishek</given-names>
          </name>
          <degrees>MS, PhD</degrees>
          <xref rid="aff1" ref-type="aff">1</xref>
          <address>
            <institution>Industrial and Management Systems Engineering</institution>
            <institution>Benjamin M Statler College of Engineering and Mineral Resources</institution>
            <institution>West Virginia University</institution>
            <addr-line>321 Engineering Sciences Bldg</addr-line>
            <addr-line>1306 Evansdale Drive</addr-line>
            <addr-line>Morgantown, WV, 26506</addr-line>
            <country>United States</country>
            <phone>1 3042939431</phone>
            <email>avishek.choudhury@mail.wvu.edu</email>
          </address>
          <ext-link ext-link-type="orcid">https://orcid.org/0000-0002-5342-0709</ext-link>
        </contrib>
        <contrib id="contrib2" contrib-type="author" equal-contrib="yes">
          <name name-style="western">
            <surname>Shahsavar</surname>
            <given-names>Yeganeh</given-names>
          </name>
          <degrees>PhD</degrees>
          <xref rid="aff2" ref-type="aff">2</xref>
          <ext-link ext-link-type="orcid">https://orcid.org/0000-0003-3422-7257</ext-link>
        </contrib>
        <contrib id="contrib3" contrib-type="author">
          <name name-style="western">
            <surname>Gurses</surname>
            <given-names>Ayse P</given-names>
          </name>
          <degrees>MS, MPH, PhD</degrees>
          <xref rid="aff2" ref-type="aff">2</xref>
          <xref rid="aff3" ref-type="aff">3</xref>
          <ext-link ext-link-type="orcid">https://orcid.org/0000-0001-7422-6852</ext-link>
        </contrib>
      </contrib-group>
      <aff id="aff1">
        <label>1</label>
        <institution>Industrial and Management Systems Engineering</institution>
        <institution>Benjamin M Statler College of Engineering and Mineral Resources</institution>
        <institution>West Virginia University</institution>
        <addr-line>Morgantown, WV</addr-line>
        <country>United States</country>
      </aff>
      <aff id="aff2">
        <label>2</label>
        <institution>Armstrong Institute Center for Health Care Human Factors</institution>
        <institution>Johns Hopkins Medicine</institution>
        <institution>Johns Hopkins University</institution>
        <addr-line>Baltimore, MD</addr-line>
        <country>United States</country>
      </aff>
      <aff id="aff3">
        <label>3</label>
        <institution>Department of Health Policy and Management</institution>
        <institution>Johns Hopkins Bloomberg School of Public Health</institution>
        <institution>Johns Hopkins University</institution>
        <addr-line>Baltimore, MD</addr-line>
        <country>United States</country>
      </aff>
      <author-notes>
        <corresp>Corresponding Author: Avishek Choudhury <email>avishek.choudhury@mail.wvu.edu</email></corresp>
      </author-notes>
      <pub-date pub-type="collection">
        <year>2026</year>
      </pub-date>
      <pub-date pub-type="epub">
        <day>8</day>
        <month>10</month>
        <year>2026</year>
      </pub-date>
      <volume>13</volume>
      <elocation-id>e97649</elocation-id>
      <history>
        <date date-type="received">
          <day>8</day>
          <month>4</month>
          <year>2026</year>
        </date>
        <date date-type="rev-request">
          <day>23</day>
          <month>6</month>
          <year>2026</year>
        </date>
        <date date-type="rev-recd">
          <day>11</day>
          <month>9</month>
          <year>2026</year>
        </date>
        <date date-type="accepted">
          <day>11</day>
          <month>9</month>
          <year>2026</year>
        </date>
      </history>
      <copyright-statement>©Avishek Choudhury, Yeganeh Shahsavar, Ayse P Gurses. Originally published in JMIR Human Factors (https://humanfactors.jmir.org), 08.10.2026.</copyright-statement>
      <copyright-year>2026</copyright-year>
      <license license-type="open-access" xlink:href="https://creativecommons.org/licenses/by/4.0/">
        <p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (https://creativecommons.org/licenses/by/4.0/), which permits unrestricted use, distribution, and reproduction in any medium, provided the original work, first published in JMIR Human Factors, is properly cited. The complete bibliographic information, a link to the original publication on https://humanfactors.jmir.org, as well as this copyright and license information must be included.</p>
      </license>
      <self-uri xlink:href="https://humanfactors.jmir.org/2026/1/e97649" xlink:type="simple"/>
      <abstract>
        <sec sec-type="background">
          <title>Background</title>
          <p>Safe integration of AI-enabled clinical decision support requires understanding whether users’ trust and behavioral reliance are appropriately calibrated to recommendation quality.</p>
        </sec>
        <sec sec-type="objective">
          <title>Objective</title>
          <p>This pilot study examined a 3-item vignette-level trust construct and accept or reject behavior across sequential primary care–style clinical vignettes.</p>
        </sec>
        <sec sec-type="methods">
          <title>Methods</title>
          <p>A total of 68 health care professionals, including 59 (86.8%) registered nurses and 9 (13.2%) physicians, each evaluated 21 clinical vignettes from 1 of 2 series. Each series contained 10 correct and 11 intentionally incorrect AI recommendations. Participants accepted or rejected each recommendation and responded to 3 questionnaire items measuring trust, perceived transparency, and likelihood of acting on the recommendation before receiving correctness feedback and points. The unrounded item mean formed the trust composite. In the behavioral analysis, 50 unrecorded accept or reject responses were classified as rejections. A pooled cross-classified linear mixed-effects model was constructed to assess associations between the trust composite, recommendation correctness, vignette position, and participant and vignette characteristics. It included random intercepts for participant and vignette item, with case series included as a nuisance adjustment.</p>
        </sec>
        <sec sec-type="results">
          <title>Results</title>
          <p>Participants accepted 374 of 748 (50%) incorrect recommendations and rejected 112 of 680 (16.47%) correct recommendations. Incorrect recommendations received lower prefeedback trust than correct recommendations (β=−0.772, 95% CI −1.036 to −0.507; standardized effect=−0.419). Trust declined modestly across vignette positions (β=−0.027, 95% CI −0.048 to −0.005; standardized effect=−0.088), with a decline in case series A but not case series B. Baseline intention to use AI was positively associated with trust (β=0.391, 95% CI 0.110-0.672; standardized effect=0.243), whereas higher perceived diagnostic difficulty was negatively associated with trust (β=−0.474, 95% CI −0.648 to −0.301; standardized effect=−0.257). In the exploratory lagged model, trust in the preceding recommendation was associated with trust in the current recommendation (β=0.266, 95% CI 0.210-0.322; standardized effect=0.264).</p>
        </sec>
        <sec sec-type="conclusions">
          <title>Conclusions</title>
          <p>Self-reported trust differentiated correct from incorrect recommendations in aggregate, while acceptance and rejection were not fully aligned with recommendation correctness. These preliminary findings identify a potential evaluation problem: lower trust in incorrect advice does not necessarily imply that users will reject it. The findings do not establish real-world clinical effects.</p>
        </sec>
      </abstract>
      <kwd-group>
        <kwd>artificial intelligence</kwd>
        <kwd>clinical decision support systems</kwd>
        <kwd>trust composite</kwd>
        <kwd>human-AI interaction</kwd>
        <kwd>trust calibration</kwd>
        <kwd>behavioral reliance</kwd>
        <kwd>registered nurses</kwd>
        <kwd>clinical vignettes</kwd>
        <kwd>perceived difficulty</kwd>
        <kwd>repeated interaction</kwd>
      </kwd-group>
    </article-meta>
  </front>
  <body>
    <sec sec-type="introduction">
      <title>Introduction</title>
      <p>AI-enabled clinical decision support systems (CDSS) have the potential to improve diagnostic and treatment workflows across health care settings by augmenting clinicians’ productivity, particularly in environments characterized by workload and uncertainty [<xref ref-type="bibr" rid="ref1">1</xref>-<xref ref-type="bibr" rid="ref3">3</xref>]. Yet, the adoption of AI at the bedside has been slower than expected [<xref ref-type="bibr" rid="ref4">4</xref>]. One major reason for this slow uptake is clinicians’ trust in AI systems [<xref ref-type="bibr" rid="ref5">5</xref>]. Clinicians must feel confident that AI recommendations are reliable and safe before they are willing to incorporate them into their clinical workflow. However, integrating AI into health care without fully understanding how clinicians interact with these systems and how their trust develops over time can create risks for patient safety [<xref ref-type="bibr" rid="ref6">6</xref>,<xref ref-type="bibr" rid="ref7">7</xref>].</p>
      <p>AI performance may fluctuate depending on the clinical scenario, data quality, or system limitations [<xref ref-type="bibr" rid="ref8">8</xref>]. If clinicians overtrust AI recommendations, they may accept incorrect advice, while undertrust may lead them to ignore useful guidance. Both situations can negatively affect patient care. Therefore, to safely and effectively integrate AI into clinical practice, it is essential to understand how clinicians develop, adjust, and act on trust in AI systems, as well as the factors that influence this trust [<xref ref-type="bibr" rid="ref6">6</xref>,<xref ref-type="bibr" rid="ref9">9</xref>]. In many ways, understanding these trust dynamics is just as important as improving the technical performance of AI algorithms themselves [<xref ref-type="bibr" rid="ref10">10</xref>].</p>
      <p>In clinical environments, which are high-paced and high-stakes work systems, maintaining appropriate trust in AI systems is critical. Health care workers (nurses and physicians) often need to make rapid decisions while managing multiple sources of information, and in such settings, clinicians must quickly determine when to accept, question, or override AI recommendations. While they commonly adopt a trust-but-verify approach when working with residents or colleagues, AI systems differ in that they lack shared accountability and interactive clarification. Trust in AI can be confounded with a user’s self-confidence [<xref ref-type="bibr" rid="ref11">11</xref>,<xref ref-type="bibr" rid="ref12">12</xref>]. If a user lacks self-confidence, they may report lower trust in the AI not because the AI is unreliable, but because they doubt their own capacity to use it safely. Distinguishing self-confidence from trust is necessary to identify the right remedies (improve model explanations vs improve user training). Additionally, positive past performance may gradually increase trust in AI recommendations [<xref ref-type="bibr" rid="ref13">13</xref>]. At the same time, as trust grows, users may become less likely to critically scrutinize AI outputs, allowing erroneous recommendations to pass through the decision-making process without sufficient evaluation. Conversely, negative past experiences with AI may make users more skeptical of the system, reducing their willingness to rely on AI recommendations even when they are accurate [<xref ref-type="bibr" rid="ref5">5</xref>,<xref ref-type="bibr" rid="ref7">7</xref>,<xref ref-type="bibr" rid="ref14">14</xref>-<xref ref-type="bibr" rid="ref16">16</xref>]. These dynamics highlight the importance of understanding how nurses’ and physicians’ trust evolve through repeated interactions with AI systems. Although AI decision support tools often perform with high accuracy, infrequent errors can hinder trust in the technology.</p>
      <p>Real-world clinical work is demanding, and trust in automation is known to vary with workload, uncertainty, and perceived task difficulty [<xref ref-type="bibr" rid="ref17">17</xref>-<xref ref-type="bibr" rid="ref19">19</xref>]. A growing literature has examined global attitudes toward AI in health care, yet far fewer studies have explored how trust in AI can change in a clinical setting [<xref ref-type="bibr" rid="ref10">10</xref>,<xref ref-type="bibr" rid="ref20">20</xref>]. Existing work has largely focused on isolated psychological or demographic predictors such as prior trust, confidence, or experience without simultaneously examining how these individual characteristics interact with contextual pressures such as time constraints and perceived diagnostic complexity [<xref ref-type="bibr" rid="ref21">21</xref>,<xref ref-type="bibr" rid="ref22">22</xref>].</p>
      <p>Building on trust calibration and human-automation reliance frameworks, we conceptualized trust as a dynamic evaluation that should ideally correspond to the perceived trustworthiness and observed reliability of the AI system. Foundational trust-in-automation theory argues that trust guides reliance when users cannot fully inspect or understand an automated system, and that safe human-automation interaction requires “appropriate reliance” rather than either blind acceptance or blanket rejection of automation [<xref ref-type="bibr" rid="ref23">23</xref>,<xref ref-type="bibr" rid="ref24">24</xref>]. Automation-reliance research further distinguishes subjective trust from observable patterns of use, misuse, disuse, and overreliance, emphasizing that users may behave in ways that do not perfectly match their stated trust. This distinction is especially important in clinical decision support, where automation bias can occur when clinicians overrely on CDSS recommendations and reduce independent information seeking or verification [<xref ref-type="bibr" rid="ref25">25</xref>]. Recent health care AI studies and reviews similarly emphasize that clinicians’ trust in AI-CDSS is shaped by clinical reliability, transparency, validation, usability, prior experience, and professional expertise, and that trust should be examined separately from behavioral reliance on AI advice [<xref ref-type="bibr" rid="ref16">16</xref>,<xref ref-type="bibr" rid="ref26">26</xref>,<xref ref-type="bibr" rid="ref27">27</xref>]. Therefore, appropriate reliance occurs when clinicians accept useful or correct AI recommendations and question or reject recommendations that are unsafe, incorrect, or insufficiently supported; overreliance and underreliance represent mismatches between AI reliability and clinician behavior. In repeated AI-assisted decisions, trust may be influenced by dispositional attitudes toward AI, the visible plausibility of each recommendation, feedback from prior clinical vignettes, and task conditions such as difficulty and time constraints [<xref ref-type="bibr" rid="ref13">13</xref>]. Therefore, this pilot study distinguished reported trust-related evaluation from behavioral acceptance of AI recommendations.</p>
      <p>This experimental pilot study had 3 aims. The primary aim was to characterize how the vignette-level trust composite changed across 21 sequential AI-assisted clinical vignettes and how it was associated with recommendation correctness, vignette position, baseline AI attitudes, participant characteristics, and perceived diagnostic difficulty. The second aim was to describe whether behavioral reliance, measured by accepting or rejecting recommendations, aligned with programmed correctness. The exploratory aim was to examine the association between trust in consecutive recommendations.</p>
    </sec>
    <sec sec-type="methods">
      <title>Methods</title>
      <sec>
        <title>Study Design</title>
        <p>This online experimental pilot study used repeated observations across 2 fixed series of 21 clinical vignettes. For each clinical vignette, participants reviewed a fictional patient presentation and a large language model (LLM)–generated diagnosis and recommendation, indicated whether they accepted or rejected the recommendation, and rated their trust. The same response process was repeated sequentially for all 21 clinical vignettes in the assigned series. No real patient encounters or patient outcomes were evaluated.</p>
      </sec>
      <sec>
        <title>Clinical Vignettes</title>
        <p>Before generating any study materials, the authors provided Current Medical Diagnosis and Treatment 2010 [<xref ref-type="bibr" rid="ref28">28</xref>] to ChatGPT (GPT-4; OpenAI) as reference material. ChatGPT was then used to draft short, fictional clinical vignettes and their associated diagnoses and recommendations. The clinical vignettes concerned common adult presentations intended to be understandable without subspecialty knowledge. They were not drawn from real patient records. The generation tasks involved creating the patient descriptions and associated reference diagnoses and recommendations, with intentionally incorrect participant-facing alternatives included where required by the experimental design.</p>
        <p>The authors reviewed the generated materials before expert physician review. They compared the initial diagnoses with the textbook and checked the clinical vignettes and recommendations for readability, completeness, clinical coherence, and internal consistency. A physician subsequently conducted the expert physician review, assessing the plausibility of each presentation, the compatibility of the proposed diagnosis with the clinical information, and the consistency of each participant-facing recommendation with its intended correct or incorrect status. She finalized the reference diagnoses and approved the materials for use in the experiment. No corrections to the initial reference diagnoses were required during either the authors’ textbook-based review or the physician’s review.</p>
        <p>The absence of diagnosis-level corrections should be distinguished from the deliberately incorrect recommendations used in the experiment. Reference diagnoses established what was considered correct for each clinical vignette; participant-facing recommendations could intentionally differ from those diagnoses to implement the planned correctness sequence. These incorrect recommendations were experimental stimuli, not diagnostic errors left unaddressed during review.</p>
        <p><xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref> describes the production and review procedure and provides prompt templates for reference-guided vignette generation, diagnosis and recommendation generation, and intentionally incorrect alternatives.</p>
      </sec>
      <sec>
        <title>Participant Recruitment</title>
        <p>Participants were recruited between August and September 2025 through Centiment, a paid online audience-panel service [<xref ref-type="bibr" rid="ref29">29</xref>]. The survey was distributed to the provider’s health care worker panel. Eligibility was based on self-reported professional designation as a registered nurse (RN) or physician and reported clinical experience. RNs were included because they are important users of digital clinical information and decision support and routinely contribute to assessment, monitoring, triage, escalation, and care coordination [<xref ref-type="bibr" rid="ref21">21</xref>,<xref ref-type="bibr" rid="ref22">22</xref>]. The clinical-vignette task examined evaluation of AI recommendations and reliance behavior; it was not intended to test participants’ independent diagnostic or prescribing competence.</p>
        <p>No independent credential, licensure, employer, or National Provider Identifier verification was performed. Centiment applies proprietary quality procedures, including duplicate and automated-response detection, device or browser fingerprinting, ReCAPTCHA, and fraud scoring. These procedures were performed by the panel provider; the research team did not receive IP addresses or device identifiers. The number invited, prescreened, or excluded by the provider and the numerical fraud-score threshold were not available to the research team.</p>
      </sec>
      <sec>
        <title>Study Procedures</title>
        <p>The 42 clinical vignettes were allocated to two 21-vignette series of broadly comparable scope and expected difficulty, without formal matching on independent prestudy difficulty ratings. Each series included 10 correct and 11 intentionally incorrect recommendations, giving a programmed accuracy rate of 47.62%. Both began with 4 correct recommendations followed by 2 incorrect recommendations; later correctness positions differed. <xref rid="figure1" ref-type="fig">Figure 1</xref> shows the series-specific sequences. Check marks (✓) indicate recommendations classified as correct; cross marks (x) indicate intentionally incorrect recommendations. The two case series shared the same status through Trial 8 but differed at selected later trial positions. Accept or reject decisions and ratings were submitted before correctness feedback, and points were displayed for each trial. Participants rated the perceived difficulty of every clinical vignette so that variation in subjective difficulty could be described and included in the analysis.</p>
        <fig id="figure1" position="float">
          <label>Figure 1</label>
          <caption>
            <p>Actual programmed sequence of AI recommendation correctness across the 21 analyzed cases.</p>
          </caption>
          <graphic xlink:href="humanfactors_v13i1e97649_fig1.png" alt-version="no" mimetype="image" position="float" xlink:type="simple"/>
        </fig>
        <p>Participants completed the study in one Qualtrics session and were randomly assigned to case series A or B. The session included electronic consent, task instructions, baseline questions, a demonstration vignette, and the 21 analyzed clinical vignettes. A programmed attention-check screen appeared between the fifth and sixth analyzed clinical vignettes. The demonstration vignette and attention-check screen were excluded from the analyses. Case series A had no fixed response window; case series B used a 45-second response window.</p>
        <p>On each clinical-vignette screen, participants reviewed the patient presentation and AI recommendation, indicated whether they accepted or rejected the recommendation, and responded to 3 trust questionnaire items and a perceived-difficulty question. The accept or reject item preceded the trust questions to prioritize eliciting behavioral reliance before explicitly asking participants to reflect on trust. All responses appeared on the same screen and were submitted before correctness feedback or points were displayed.</p>
        <p>After submission, participants received feedback about the programmed correctness of the AI recommendation. They received +10 points for accepting a correct recommendation or rejecting an incorrect recommendation and −10 points for the opposite decisions. Feedback and scores from preceding clinical vignettes were intentionally available to inform responses to subsequent recommendations. In case series B, the page advanced when the 45-second response window expired; the analytic handling of unrecorded decisions is specified below under the “Analysis” section.</p>
      </sec>
      <sec>
        <title>Outcome Measurement</title>
        <p>The primary outcome was vignette-level trust in the AI recommendation. Trust was operationalized through 3 investigator-developed questionnaire items addressing perceived trust, perceived transparency, and likelihood of acting on the recommendation. These items were specified as indicators of the study’s trust construct rather than as 3 separate primary outcomes; their unrounded mean was the trust composite. The conceptual framing drew on trust-in-automation literature [<xref ref-type="bibr" rid="ref23">23</xref>,<xref ref-type="bibr" rid="ref24">24</xref>]. The observed accept or reject response was a separate behavioral-reliance outcome. Baseline measures are summarized in <xref ref-type="table" rid="table1">Table 1</xref>, and the clinical-vignette measures are given in <xref ref-type="table" rid="table2">Table 2</xref>.</p>
        <table-wrap position="float" id="table1">
          <label>Table 1</label>
          <caption>
            <p>Baseline participant measures.</p>
          </caption>
          <table width="1000" cellpadding="5" cellspacing="0" border="1" rules="groups" frame="hsides">
            <col width="190"/>
            <col width="450"/>
            <col width="360"/>
            <thead>
              <tr valign="top">
                <td>Construct</td>
                <td>Survey items</td>
                <td>Scale or scoring</td>
              </tr>
            </thead>
            <tbody>
              <tr valign="top">
                <td>Professional designation</td>
                <td>What is your professional or academic background?</td>
                <td>
                  <list list-type="bullet">
                    <list-item>
                      <p>Registered nurse</p>
                    </list-item>
                    <list-item>
                      <p>Physician</p>
                    </list-item>
                  </list>
                </td>
              </tr>
              <tr valign="top">
                <td>Clinical experience</td>
                <td>How many years of experience do you have in your current field or line of work?</td>
                <td>
                  <list list-type="bullet">
                    <list-item>
                      <p>&#60;1 year</p>
                    </list-item>
                    <list-item>
                      <p>1-2 years</p>
                    </list-item>
                    <list-item>
                      <p>3-5 years</p>
                    </list-item>
                    <list-item>
                      <p>6-10 years</p>
                    </list-item>
                    <list-item>
                      <p>&#62;10 years</p>
                    </list-item>
                  </list>
                </td>
              </tr>
              <tr valign="top">
                <td>Prior AI use</td>
                <td>In the last 6 months, have you ever used AI in medical practice?</td>
                <td>
                  <list list-type="bullet">
                    <list-item>
                      <p>Yes</p>
                    </list-item>
                    <list-item>
                      <p>No</p>
                    </list-item>
                  </list>
                </td>
              </tr>
              <tr valign="top">
                <td>Baseline trust in AI (3-item composite)</td>
                <td>General trust in AI, perceived transparency of AI recommendations, and likelihood of acting on an AI recommendation.</td>
                <td>
                  <list list-type="bullet">
                    <list-item>
                      <p>Each item: 1=none at all to 7=completely</p>
                    </list-item>
                    <list-item>
                      <p>Composite=unrounded mean of available items (minimum 2 of 3)</p>
                    </list-item>
                  </list>
                </td>
              </tr>
              <tr valign="top">
                <td>Intent to use AI</td>
                <td>I would like to use AI in my clinical practice.</td>
                <td>
                  <list list-type="bullet">
                    <list-item>
                      <p>1=strongly disagree to 5=strongly agree</p>
                    </list-item>
                  </list>
                </td>
              </tr>
              <tr valign="top">
                <td>Self-confidence</td>
                <td>Nine retained items: I handle new situations with relative comfort and ease; I feel positive and energized about life; I keep trying even after others have given up; If I work hard to solve a problem, I will find the answer; I achieve the goals I set for myself; people give me positive feedback on my work and achievements; when I overcome an obstacle, I think about the lessons I have learned; I believe that if I work hard, I will achieve my goals; I have contact with people with similar skills and experience whom I consider successful.</td>
                <td>
                  <list list-type="bullet">
                    <list-item>
                      <p>Each item: 1=not at all to 5=very often</p>
                    </list-item>
                    <list-item>
                      <p>Mean of 9 retained items</p>
                    </list-item>
                  </list>
                </td>
              </tr>
            </tbody>
          </table>
        </table-wrap>
        <table-wrap position="float" id="table2">
          <label>Table 2</label>
          <caption>
            <p>Clinical-vignette measures completed before correctness feedback.</p>
          </caption>
          <table width="1000" cellpadding="5" cellspacing="0" border="1" rules="groups" frame="hsides">
            <col width="220"/>
            <col width="560"/>
            <col width="220"/>
            <thead>
              <tr valign="top">
                <td>Construct</td>
                <td>Survey item</td>
                <td>Scale</td>
              </tr>
            </thead>
            <tbody>
              <tr valign="top">
                <td>Perceived diagnostic difficulty</td>
                <td>
                  <list list-type="bullet">
                    <list-item>
                      <p>How difficult do you think it would be to diagnose this patient case based on the information provided without AI assistance?</p>
                    </list-item>
                  </list>
                </td>
                <td>
                  <list list-type="bullet">
                    <list-item>
                      <p>1=very easy to 5=very difficult</p>
                    </list-item>
                  </list>
                </td>
              </tr>
              <tr valign="top">
                <td>Vignette-level trust in the AI recommendation (3-item composite)</td>
                <td>
                  <list list-type="bullet">
                    <list-item>
                      <p>How much do you trust the AI-based recommendation?</p>
                    </list-item>
                    <list-item>
                      <p>How transparent do you find the AI-based recommendation?</p>
                    </list-item>
                    <list-item>
                      <p>How likely are you to act upon the AI-based recommendation?</p>
                    </list-item>
                  </list>
                </td>
                <td>
                  <list list-type="bullet">
                    <list-item>
                      <p>Each item: 1=none at all to 7=completely</p>
                    </list-item>
                    <list-item>
                      <p>Composite=unrounded mean of the 3 items</p>
                    </list-item>
                  </list>
                </td>
              </tr>
            </tbody>
          </table>
        </table-wrap>
        <p>Separate 1-factor confirmatory factor analysis (CFA) summaries based on polychoric correlations were aligned to the 21 analyzed clinical-vignette positions. Average variance extracted (AVE) ranged from 0.76 to 0.92, coefficient omega from 0.90 to 0.97, and Cronbach α from 0.87 to 0.96. The 3 items showed strong internal consistency; these results do not make the self-reported composite equivalent to observed acceptance or rejection. Vignette-specific factor loadings and reliability estimates are reported in <xref ref-type="supplementary-material" rid="app2">Multimedia Appendix 2</xref>.</p>
        <p>Baseline trust was calculated as the unrounded mean of the corresponding 3 baseline items, using at least 2 available responses, as specified in <xref ref-type="table" rid="table1">Table 1</xref>. Cronbach α was 0.938 among participants with complete baseline items. Self-confidence was the unrounded mean of the 9 retained 1-5 items in <xref ref-type="table" rid="table1">Table 1</xref>, informed by self-efficacy theory [<xref ref-type="bibr" rid="ref30">30</xref>]. The supplied self-confidence CFA reported AVE=0.65, coefficient Ω=0.94, and Cronbach α=0.91 (<xref ref-type="supplementary-material" rid="app2">Multimedia Appendix 2</xref>).</p>
      </sec>
      <sec>
        <title>Data Collection</title>
        <p>Before the clinical-vignette sequence, participants reported their professional designation, clinical experience, prior AI use, baseline trust, intention to use AI, and self-confidence. During each clinical vignette, the accept or reject decision was elicited as a binary choice, the three trust items as 1-7 ratings, and perceived diagnostic difficulty as a 1-5 rating. Responses were linked by anonymous study record, case series, and clinical-vignette position for repeated-measures analysis.</p>
        <p>The first 20 eligible participants formed a pilot test to assess survey access, clarity of instructions, screen functioning, and whether participants could understand and answer the questions. They met the same eligibility criteria and completed the same consent process, measures, clinical-vignette sequence, and feedback procedure as the remaining participants. No technical or comprehension problems were identified, and no changes were made; these responses were retained in the final analytic sample.</p>
      </sec>
      <sec>
        <title>Analysis</title>
        <p>The hypotheses concerned the trust composite: H1a predicted a difference between correct and intentionally incorrect recommendations after accounting for vignette position and participant and vignette variability; H1b predicted a positive association between trust in consecutive recommendations; H2 predicted higher trust among participants with stronger baseline intention to use AI; and H3 predicted lower trust for clinical vignettes perceived as more difficult.</p>
        <p>The demonstration vignette and attention-check screen were excluded. All 3 trust items and perceived-difficulty responses were complete for the 68 analyzed participants across 21 clinical vignettes, yielding 1428 participant-vignette observations. The data were organized in long format. The 5-point difficulty rating was categorized as lower difficulty (very easy or easy) or higher difficulty (moderate, difficult, or very difficult); a sensitivity model retained the original 1-5 score. Perceived difficulty was treated as a concurrent subjective task-context predictor rather than as a preexposure causal confounder.</p>
        <p>Recommendation correctness was coded separately for each case series using the programmed status and cross-checked against Qualtrics task-score fields. Clinical-vignette position was centered at position 11. Correctness remained partly linked to vignette content, position, and the fixed reliability sequence; coefficients involving correctness were therefore interpreted as adjusted associations rather than causal effects.</p>
        <p>The primary analysis pooled all 68 participants. A parsimonious cross-classified linear mixed-effects model assessed associations with the trust composite and included random intercepts for participant and vignette item, with 42 distinct vignette items. Fixed effects were recommendation correctness, centered clinical-vignette position, case series, professional designation, baseline trust, intention to use AI, self-confidence, clinical experience (≤10 vs &#62;10 years), and perceived difficulty. Continuous participant-level predictors were mean-centered. Case series was a nuisance adjustment because vignette content and response window varied together; between-series differences were not attributed to response timing.</p>
        <p>The primary model was fitted separately within each case series, and a pooled vignette position × case series model examined differences in sequence slopes. A secondary interaction model examined whether the incorrect-versus-correct trust difference varied with vignette position, baseline trust, intention to use AI, self-confidence, experience, or perceived difficulty. These interactions were excluded from the parsimonious primary model and interpreted as hypothesis-generating. An exploratory lagged model used positions 2-21 to relate trust in the current recommendation to trust in the immediately preceding recommendation while adjusting for current correctness, vignette position, case series, difficulty, and participant variables. Models with and without preceding-vignette trust were compared by maximum likelihood; the lagged coefficient was interpreted as a temporal association rather than evidence of a causal mechanism.</p>
        <p>The study and analyses were not preregistered. Models were estimated by restricted maximum likelihood in Python 3.13 (Python Software Foundation) using <italic>statsmodels</italic> (version 0.14.6). Maximum likelihood was used for the lagged-model comparison. All reported models converged. Estimates were summarized using coefficients, SEs, 2-sided Wald tests, and 95% CIs. Standardized effects came from parallel models with z-standardized outcomes and continuous predictors, retaining the original categorical coding. Marginal <italic>R</italic>² described variance explained by fixed effects; conditional <italic>R</italic>² included fixed and random effects. Trust-composite end point frequencies were also examined.</p>
        <p>Behavioral reliance was summarized by acceptance and rejection relative to programmed recommendation correctness. A total of 50 case series B observations had no recorded accept or reject response. For the primary descriptive analysis, these observations were assigned to the rejection category under the assumption that participants who trusted a recommendation sufficiently to act on it would record acceptance within the 45-second response window. This rule retained all presented clinical vignettes and treated unrecorded decisions as nonacceptance; it did not establish that participants actively rejected or distrusted the recommendation. A descriptive sensitivity summary excluded unrecorded decisions instead of recoding them; the resulting counts and denominators are reported in <xref ref-type="supplementary-material" rid="app3">Multimedia Appendix 3</xref>.</p>
      </sec>
      <sec>
        <title>Ethical Considerations</title>
        <p>The West Virginia University (WVU) Institutional Review Board (IRB) approved the study under the WVU Flexibility Review Model (protocol number 2502117081). All participants provided electronic informed consent before beginning the survey. The study did not request names, contact information, IP addresses, licensure identifiers, employer information, or other direct identifiers; the analytic dataset contained anonymous responses. Study files were stored in password-protected locations accessible only to authorized research personnel. Participants were recruited and compensated through Centiment according to the provider’s panel procedures; the exact participant payment was not available to the research team.</p>
        <p>This was an online clinical-vignette experiment without real patient encounters or patient outcomes. Applicable CONSORT-AI (Consolidated Standards of Reporting Trials–Artificial Intelligence) reporting principles informed descriptions of eligibility, allocation, participant flow, recommendation presentation, human-AI interaction, and analysis [<xref ref-type="bibr" rid="ref31">31</xref>].</p>
      </sec>
    </sec>
    <sec sec-type="results">
      <title>Results</title>
      <sec>
        <title>Participant Flow and Characteristics</title>
        <p>A total of 113 individuals participated in the survey. Sixty-eight met the eligibility and completeness criteria and were included in the analysis. The panel provider did not supply the number invited or prescreened, the number excluded through provider-level fraud or quality procedures, or sufficient information to distinguish ineligible from incomplete records among the remaining 45 participants. The final sample included 59/68 (86.76%) RNs and 9/68 (13.24%) physicians, with 34 participants assigned to each case series. All 68 analyzed participants provided complete responses to the 3 trust questionnaire items and the perceived-difficulty item across 21 clinical vignettes.</p>
        <p>Baseline descriptive statistics, reported as means and SDs, are presented in <xref ref-type="table" rid="table3">Table 3</xref>. The mean baseline trust composite was 4.37 (SD 1.64); mean intention to use AI was 3.35 (SD 1.16); and mean self-confidence was 4.02 (SD 0.70).</p>
        <table-wrap position="float" id="table3">
          <label>Table 3</label>
          <caption>
            <p>Participant characteristics and baseline measures<sup>a</sup>.</p>
          </caption>
          <table width="1000" cellpadding="5" cellspacing="0" border="1" rules="groups" frame="hsides">
            <col width="30"/>
            <col width="330"/>
            <col width="230"/>
            <col width="230"/>
            <col width="180"/>
            <thead>
              <tr valign="top">
                <td colspan="2">Characteristic</td>
                <td>Case series A (n=34)</td>
                <td>Case series B (n=34)</td>
                <td>Total (n=68)</td>
              </tr>
            </thead>
            <tbody>
              <tr valign="top">
                <td colspan="5">Professional designation, n (%)</td>
              </tr>
              <tr valign="top">
                <td>
                  <break/>
                </td>
                <td>Registered nurse</td>
                <td>31 (91.18)</td>
                <td>28 (82.35)</td>
                <td>59 (86.76)</td>
              </tr>
              <tr valign="top">
                <td>
                  <break/>
                </td>
                <td>Physician</td>
                <td>3 (8.82)</td>
                <td>6 (17.65)</td>
                <td>9 (13.24)</td>
              </tr>
              <tr valign="top">
                <td colspan="5">Experience (years), n (%)</td>
              </tr>
              <tr valign="top">
                <td>
                  <break/>
                </td>
                <td>&#60;1</td>
                <td>2 (5.88)</td>
                <td>0 (0.00)</td>
                <td>2 (2.94)</td>
              </tr>
              <tr valign="top">
                <td>
                  <break/>
                </td>
                <td>1-2</td>
                <td>1 (2.94)</td>
                <td>3 (8.82)</td>
                <td>4 (5.88)</td>
              </tr>
              <tr valign="top">
                <td>
                  <break/>
                </td>
                <td>3-5</td>
                <td>8 (23.53)</td>
                <td>2 (5.88)</td>
                <td>10 (14.71)</td>
              </tr>
              <tr valign="top">
                <td>
                  <break/>
                </td>
                <td>6-10</td>
                <td>10 (29.41)</td>
                <td>4 (11.76)</td>
                <td>14 (20.59)</td>
              </tr>
              <tr valign="top">
                <td>
                  <break/>
                </td>
                <td>&#62;10</td>
                <td>13 (38.24)</td>
                <td>25 (73.53)</td>
                <td>38 (55.88)</td>
              </tr>
              <tr valign="top">
                <td colspan="5">Prior AI use, n (%)</td>
              </tr>
              <tr valign="top">
                <td>
                  <break/>
                </td>
                <td>Yes</td>
                <td>16 (47.06)</td>
                <td>16 (47.06)</td>
                <td>32 (47.06)</td>
              </tr>
              <tr valign="top">
                <td>
                  <break/>
                </td>
                <td>No</td>
                <td>18 (52.94)</td>
                <td>18 (52.94)</td>
                <td>36 (52.94)</td>
              </tr>
              <tr valign="top">
                <td colspan="2">Baseline trust composite, mean (SD)</td>
                <td>4.28 (1.66)</td>
                <td>4.46 (1.64)</td>
                <td>4.37 (1.64)</td>
              </tr>
              <tr valign="top">
                <td colspan="2">Intent to use AI, mean (SD)</td>
                <td>3.41 (1.08)</td>
                <td>3.29 (1.24)</td>
                <td>3.35 (1.16)</td>
              </tr>
              <tr valign="top">
                <td colspan="2">Self-confidence, mean (SD)</td>
                <td>3.95 (0.66)</td>
                <td>4.09 (0.74)</td>
                <td>4.02 (0.70)</td>
              </tr>
            </tbody>
          </table>
          <table-wrap-foot>
            <fn id="table3fn1">
              <p><sup>a</sup>The baseline trust composite is the unrounded mean of general trust, perceived transparency, and likelihood-to-act items. Self-confidence is the unrounded mean of 9 retained items.</p>
            </fn>
          </table-wrap-foot>
        </table-wrap>
      </sec>
      <sec>
        <title>Perceived Diagnostic Difficulty</title>
        <p>Of the 714 difficulty ratings in each series, 342 (47.90%) were categorized as higher difficulty in case series A and 355 (49.72%) in case series B. The position-specific distributions are shown in <xref ref-type="supplementary-material" rid="app3">Multimedia Appendix 3</xref>.</p>
      </sec>
      <sec>
        <title>Trust Across Clinical Vignettes</title>
        <p>The trust composite varied across the 21 clinical vignettes in both series (Figure S2 in <xref ref-type="supplementary-material" rid="app3">Multimedia Appendix 3</xref>). Case series A increased from a mean of 5.04 at position 1 to 5.38 at position 4, decreased to 3.56 at the first incorrect recommendation at position 5, and ended at 3.21. Case series B began at 4.80, decreased to 3.54 at position 5, and ended at 4.32. Grouped means by case series, difficulty, and correctness are shown in <xref ref-type="supplementary-material" rid="app3">Multimedia Appendix 3</xref>.</p>
        <p>Across 1428 trust composites, 123 (8.61%) were at the minimum of 1 and 177 (12.39%) were at the maximum of 7; 300 scores (21.01%) occurred at either end point. Scores of 2 or lower accounted for 17.72%, and scores of 6 or higher accounted for 27.03%.</p>
      </sec>
      <sec>
        <title>Behavioral Reliance</title>
        <p>Across both series, 374 of 748 (50%) incorrect recommendations were accepted, and 112 of 680 (16.47%) correct recommendations were in the rejection category. Each series contributed 714 decision instances: 374 with incorrect and 340 with correct recommendations. Incorrect recommendations were accepted 185 times in case series A (49.47% of incorrect recommendations; 25.91% of all series A observations) and 189 times in case series B (50.53%; 26.47% of all series B observations). Correct recommendations were in the rejection category 52 times in case series A (15.29% of correct recommendations; 7.28% of all series A observations) and 60 times in case series B (17.65%; 8.40%).</p>
        <p>At clinical vignette 5, the first incorrect recommendation in both series, 16 participants in case series A accepted and 18 rejected the recommendation; in case series B, 12 accepted and 22 responses were classified in the rejection category. Acceptance was above zero for every incorrect recommendation. When both series returned to a correct recommendation at position 17, 25 participants in case series A and 28 in case series B accepted it. Position-specific counts are shown in <xref ref-type="supplementary-material" rid="app3">Multimedia Appendix 3</xref>.</p>
        <p>Of the 50 unrecorded case series B decisions, 26 occurred for correct and 24 for incorrect recommendations. When only recorded responses were counted, incorrect recommendations were accepted in 374 of 724 (51.66%) observations, and correct recommendations were explicitly rejected in 86 of 654 (13.15%) observations. In case series B alone, explicit rejection of correct recommendations was 34 of 314 (10.83%) recorded responses, compared with 60 of 340 (17.65%) under rejection-category coding. Complete denominators for both summaries are provided in <xref ref-type="supplementary-material" rid="app3">Multimedia Appendix 3</xref>.</p>
      </sec>
      <sec>
        <title>Primary Mixed-Effects Model</title>
        <p><xref ref-type="table" rid="table4">Table 4</xref> presents the pooled parsimonious cross-classified model. Incorrect AI recommendations were associated with lower prefeedback trust than correct recommendations (β=−0.772, 95% CI −1.036 to −0.507; standardized effect=−0.419). Trust also declined modestly across clinical vignettes (β=−0.027, 95% CI −0.048 to −0.005; standardized effect=−0.088). Intention to use AI was positively associated with trust (β=0.391, 95% CI 0.110-0.672; standardized effect=0.243), whereas higher perceived difficulty was negatively associated with trust (β=−0.474, 95% CI −0.648 to −0.301; standardized effect=−0.257). CIs for the remaining fixed effects included zero. The sensitivity model retaining the original 1-5 difficulty score yielded the same substantive pattern; each 1-point increase in difficulty was associated with a 0.404-point lower trust composite (95% CI −0.495 to −0.312).</p>
        <table-wrap position="float" id="table4">
          <label>Table 4</label>
          <caption>
            <p>Parsimonious cross-classified mixed-effects model predicting the vignette-level trust composite<sup>a</sup>.</p>
          </caption>
          <table width="1000" cellpadding="5" cellspacing="0" border="1" rules="groups" frame="hsides">
            <col width="330"/>
            <col width="180"/>
            <col width="220"/>
            <col width="100"/>
            <col width="170"/>
            <thead>
              <tr valign="top">
                <td>Predictor</td>
                <td>β (SE)</td>
                <td>95% CI</td>
                <td>
                  <italic>P</italic>
                </td>
                <td>Standardized effect</td>
              </tr>
            </thead>
            <tbody>
              <tr valign="top">
                <td>Intercept</td>
                <td>4.795 (0.249)</td>
                <td>4.307 to 5.283</td>
                <td>&#60;.001</td>
                <td>—<sup>b</sup></td>
              </tr>
              <tr valign="top">
                <td>Vignette position (centered)</td>
                <td>−0.027 (0.011)</td>
                <td>−0.048 to −0.005</td>
                <td>.02</td>
                <td>−0.088</td>
              </tr>
              <tr valign="top">
                <td>AI recommendation incorrect</td>
                <td>−0.772 (0.135)</td>
                <td>−1.036 to −0.507</td>
                <td>&#60;.001</td>
                <td>−0.419</td>
              </tr>
              <tr valign="top">
                <td>Case series B vs A (nuisance adjustment)</td>
                <td>0.174 (0.300)</td>
                <td>−0.414 to 0.762</td>
                <td>.56</td>
                <td>0.094</td>
              </tr>
              <tr valign="top">
                <td>Professional designation (physician)</td>
                <td>0.243 (0.401)</td>
                <td>−0.542 to 1.029</td>
                <td>.54</td>
                <td>0.132</td>
              </tr>
              <tr valign="top">
                <td>Baseline trust composite</td>
                <td>0.072 (0.095)</td>
                <td>−0.114 to 0.259</td>
                <td>.45</td>
                <td>0.064</td>
              </tr>
              <tr valign="top">
                <td>Intent to use AI</td>
                <td>0.391 (0.143)</td>
                <td>0.110 to 0.672</td>
                <td>.006</td>
                <td>0.243</td>
              </tr>
              <tr valign="top">
                <td>Self-confidence</td>
                <td>0.047 (0.198)</td>
                <td>−0.341 to 0.435</td>
                <td>.81</td>
                <td>0.018</td>
              </tr>
              <tr valign="top">
                <td>Experience &#62;10 years</td>
                <td>0.007 (0.294)</td>
                <td>−0.569 to 0.583</td>
                <td>.98</td>
                <td>0.004</td>
              </tr>
              <tr valign="top">
                <td>Perceived diagnostic difficulty: higher</td>
                <td>−0.474 (0.089)</td>
                <td>−0.648 to −0.301</td>
                <td>&#60;.001</td>
                <td>−0.257</td>
              </tr>
            </tbody>
          </table>
          <table-wrap-foot>
            <fn id="table4fn1">
              <p><sup>a</sup>N=1428 observations from 68 participants and 42 vignette items. Reference categories: correct recommendation, case series A, registered nurse, ≤10 years of experience, and lower perceived difficulty. The case series was included only as a nuisance adjustment. Marginal <italic>R</italic>²=0.176; conditional <italic>R</italic>²=0.500. Standardized effects were obtained from a parallel model in which the outcome and continuous predictors were z-standardized; categorical predictors retained their original coding.</p>
            </fn>
            <fn id="table4fn2">
              <p><sup>b</sup>Not applicable.</p>
            </fn>
          </table-wrap-foot>
        </table-wrap>
      </sec>
      <sec>
        <title>Sensitivity Analyses by Case Series</title>
        <p>Sensitivity analyses showed that lower trust for incorrect recommendations was evident in both case series: case series A, β=−0.850 (95% CI −1.197 to −0.502), and case series B, β=−0.676 (95% CI −0.996 to −0.356). The longitudinal pattern differed across vignette sets. Trust declined with vignette position in case series A (β=−0.056, 95% CI −0.085 to −0.027) but not in case series B (β=0.004, 95% CI −0.022 to 0.030). In the pooled vignette position × case series model, the case series B slope was 0.064 points per clinical vignette less negative than the case series A slope (95% CI 0.027-0.102). Higher perceived difficulty was negatively associated with trust in both series, whereas intention to use AI was positively associated with trust only in case series B (<xref ref-type="table" rid="table5">Table 5</xref>).</p>
        <table-wrap position="float" id="table5">
          <label>Table 5</label>
          <caption>
            <p>Parsimonious models fitted separately within each case series<sup>a</sup>.</p>
          </caption>
          <table width="1000" cellpadding="5" cellspacing="0" border="1" rules="groups" frame="hsides">
            <col width="200"/>
            <col width="130"/>
            <col width="150"/>
            <col width="150"/>
            <col width="0"/>
            <col width="140"/>
            <col width="140"/>
            <col width="90"/>
            <thead>
              <tr valign="top">
                <td>Predictor</td>
                <td colspan="4">Series A</td>
                <td colspan="3">Series B</td>
              </tr>
              <tr valign="top">
                <td>
                  <break/>
                </td>
                <td>β (SE)</td>
                <td>95% CI</td>
                <td><italic>P</italic> value</td>
                <td colspan="2">β (SE)</td>
                <td>95% CI</td>
                <td><italic>P</italic> value</td>
              </tr>
            </thead>
            <tbody>
              <tr valign="top">
                <td>Intercept</td>
                <td>4.918 (0.269)</td>
                <td>4.390 to 5.446</td>
                <td>&#60;.001</td>
                <td colspan="2">4.846 (0.458)</td>
                <td>3.948 to 5.744</td>
                <td><break/><break/><break/>&#60;.001</td>
              </tr>
              <tr valign="top">
                <td>Vignette position (centered)</td>
                <td>−0.056 (0.015)</td>
                <td>−0.085 to −0.027</td>
                <td><break/><break/><break/>&#60;.001</td>
                <td colspan="2">0.004 (0.013)</td>
                <td>−0.022 to 0.030</td>
                <td><break/><break/><break/>.77</td>
              </tr>
              <tr valign="top">
                <td>AI recommendation incorrect</td>
                <td>−0.850 (0.177)</td>
                <td>−1.197 to −0.502</td>
                <td><break/><break/><break/>&#60;.001</td>
                <td colspan="2">−0.676 (0.163)</td>
                <td>−0.996 to 0.356</td>
                <td>&#60;.001</td>
              </tr>
              <tr valign="top">
                <td>Professional designation: physician</td>
                <td>1.144 (0.642)</td>
                <td>−0.113 to 2.402</td>
                <td>.07</td>
                <td colspan="2">−0.354 (0.564)</td>
                <td>−1.459 to 0.752</td>
                <td><break/><break/><break/>.53</td>
              </tr>
              <tr valign="top">
                <td>Baseline trust composite</td>
                <td>−0.060 (0.131)</td>
                <td>−0.317 to 0.198</td>
                <td>.65</td>
                <td colspan="2">0.134 (0.128)</td>
                <td>−0.118 to 0.385</td>
                <td>.30</td>
              </tr>
              <tr valign="top">
                <td>Intent to use AI</td>
                <td>0.236 (0.203)</td>
                <td>−0.163 to 0.634</td>
                <td>.25</td>
                <td colspan="2">0.619 (0.211)</td>
                <td>0.205 to 1.033</td>
                <td>.003</td>
              </tr>
              <tr valign="top">
                <td>Self-confidence</td>
                <td>−0.211 (0.273)</td>
                <td>−0.745 to 0.324</td>
                <td>.44</td>
                <td colspan="2">0.045 (0.313)</td>
                <td>−0.568 to 0.657</td>
                <td>.89</td>
              </tr>
              <tr valign="top">
                <td>Experience &#62;10 years</td>
                <td>−0.305 (0.367)</td>
                <td>−1.024 to 0.414</td>
                <td>.41</td>
                <td colspan="2">0.173 (0.486)</td>
                <td>−0.780 to 1.126</td>
                <td>.72</td>
              </tr>
              <tr valign="top">
                <td>Perceived diagnostic difficulty: higher</td>
                <td>−0.604 (0.144)</td>
                <td>−0.887 to −0.321</td>
                <td>&#60;.001</td>
                <td colspan="2">−0.344 (0.103)</td>
                <td>−0.546 to −0.141</td>
                <td>&#60;.001</td>
              </tr>
            </tbody>
          </table>
          <table-wrap-foot>
            <fn id="table5fn1">
              <p><sup>a</sup>Each model included 34 participants, 714 observations, 21 vignette items, and random intercepts for participant and vignette item. Marginal or conditional <italic>R</italic>²=0.178/0.438 for case series A and 0.323/0.601 for case series B. The pooled vignette position × case series model is summarized in the text.</p>
            </fn>
          </table-wrap-foot>
        </table-wrap>
      </sec>
      <sec>
        <title>Exploratory Interaction Model</title>
        <p>The exploratory model suggested that sensitivity to incorrect recommendations varied with several characteristics. The incorrect-versus-correct difference was more negative among participants with more than 10 years of experience (interaction β=−0.703, 95% CI −0.998 to −0.407) and at higher self-confidence (interaction β=−0.360 per scale point, 95% CI −0.571 to −0.150), but less negative on clinical vignettes categorized as higher difficulty (interaction β=0.327, 95% CI 0.045-0.609). Interactions with vignette position, baseline trust, and intention to use AI had CIs that included zero. Because this interaction-heavy model was secondary and participant-level power was limited, these findings are considered hypothesis-generating. Full model estimates are reported in <xref ref-type="supplementary-material" rid="app3">Multimedia Appendix 3</xref>.</p>
      </sec>
      <sec>
        <title>Exploratory Lagged Model</title>
        <p>In the exploratory lagged model, the preceding clinical vignette’s trust composite was positively associated with current-vignette trust (β=0.266, 95% CI 0.210-0.322; standardized effect=0.264). Incorrect current-vignette recommendations remained associated with lower trust (β=−0.753, 95% CI −1.037 to −0.469; standardized effect=−0.405), intention to use AI remained positively associated with trust, and higher perceived difficulty remained negatively associated with trust. Vignette position and the remaining participant-level variables had CIs that included zero. Adding preceding-vignette trust improved model fit relative to the otherwise identical model without the lagged term. Full coefficients and model-fit statistics are reported in <xref ref-type="supplementary-material" rid="app3">Multimedia Appendix 3</xref>.</p>
      </sec>
      <sec>
        <title>Hypothesis Summary</title>
        <p>H1a was supported because incorrect recommendations received lower vignette-level trust than correct recommendations. H1b was supported statistically because the preceding clinical vignette’s trust composite was positively associated with current-vignette trust, although this result is interpreted as a temporal association rather than evidence of a causal carryover mechanism. H2 was supported because baseline intention to use AI was positively associated with vignette-level trust. H3 was supported because higher perceived diagnostic difficulty was negatively associated with vignette-level trust.</p>
      </sec>
    </sec>
    <sec sec-type="discussion">
      <title>Discussion</title>
      <sec>
        <title>Principal Findings and Contributions</title>
        <p>Participants assigned lower trust to incorrect than to correct AI recommendations, yet accepted half of the incorrect recommendations and rejected approximately 1 in 6 correct recommendations under the stated behavioral coding rule. This aggregate pattern identifies a potential gap between self-reported trust and reliance calibrated to recommendation correctness; it does not establish that individual participants recognized errors before accepting them. Trust in consecutive recommendations was positively associated; baseline intention to use AI was associated with higher trust, and greater perceived difficulty with lower trust. The pooled decline across vignette positions was evident in case series A but not case series B.</p>
        <p>Together, these results suggest that clinical AI evaluation should focus on 2 related but distinct questions: whether users can judge the quality of individual AI recommendations and whether those judgments lead to appropriate reliance. The findings should nevertheless be interpreted as preliminary because they were obtained from a small, nurse-dominant sample completing simulated primary care–style tasks.</p>
      </sec>
      <sec>
        <title>Dynamic Trust and Recommendation Correctness</title>
        <p>Participants gave substantially lower trust ratings to incorrect recommendations than to correct recommendations. Because the trust questions were completed before current-vignette feedback was shown, this difference suggests that participants were, on average, sensitive to features of the vignette and the apparent plausibility of the AI recommendation. The average difference was therefore present before explicit feedback about that recommendation. This result is consistent with recent repeated-interaction studies. Kahr et al [<xref ref-type="bibr" rid="ref32">32</xref>] found that people reported greater trust and reliance when AI advice came from a more accurate model, while Yang et al [<xref ref-type="bibr" rid="ref33">33</xref>] showed that users adjust trust from one interaction to the next and react particularly strongly to automation failures.</p>
        <p>The pattern of trust across the 21 clinical vignettes was less consistent. Trust declined in case series A but remained comparatively stable in case series B. This qualification is important because it means that the pooled decline should not be interpreted as a universal process in which trust steadily erodes whenever users encounter AI errors. Instead, the course of trust appears to depend on which cases users see, how convincing the recommendations appear, and where correct and incorrect recommendations occur in the sequence. This result is consistent with Rittenberg et al [<xref ref-type="bibr" rid="ref34">34</xref>], who found that trust did not always rise and fall in direct proportion to automation reliability. In their experiments, trust sometimes declined even when system reliability improved, and trust was difficult to rebuild after users had experienced a poorly performing system. Kahr et al [<xref ref-type="bibr" rid="ref32">32</xref>], by contrast, observed increasing trust when participants repeatedly interacted with high-accuracy AI advice. Considered together, these studies and the present findings indicate that there may not be one standard trust trajectory. Trust depends on reliability, but it is also shaped by starting impressions, the order and type of errors, task characteristics, and the user’s ability to evaluate each recommendation. Because recommendation correctness, case content, and vignette position were not independently randomized in this study, the lower trust assigned to incorrect recommendations cannot be attributed solely to correctness. Some of the difference may also reflect the particular cases used or their position in the sequence. The main conclusion is therefore that participants showed aggregate sensitivity to recommendation quality, while the way trust changed over time remained dependent on the specific case series.</p>
      </sec>
      <sec>
        <title>Trust in Consecutive Recommendations</title>
        <p>Trust in the immediately preceding clinical vignette was positively associated with trust in the current clinical vignette. In everyday terms, participants appeared to carry part of their recent impression of the AI into the next case rather than evaluating every recommendation in complete isolation. This makes sense in repeated AI use: after seeing several useful recommendations, a person may approach the next recommendation more favorably, whereas recent errors may create greater caution. This finding is consistent with evidence from Kahr et al [<xref ref-type="bibr" rid="ref32">32</xref>], Yang et al [<xref ref-type="bibr" rid="ref33">33</xref>], and Rittenberg et al [<xref ref-type="bibr" rid="ref34">34</xref>] that early reliability experiences can influence later trust and its recovery. This study extends this work to repeated clinical-vignette tasks. However, the lagged finding should not be described as proof of trust inertia. Consecutive ratings from the same participant will often resemble one another potentially because individuals have relatively stable response styles and general attitudes. The most defensible interpretation is that trust showed temporal continuity where the previous rating provided useful information about the next rating, but the study cannot determine whether this occurred because of memory of previous AI performance, accumulated feedback, a stable personal response pattern, or another cognitive process.</p>
        <p>The practical implication remains important. Trust in AI should be studied as a developing process rather than as a one-time opinion.</p>
      </sec>
      <sec>
        <title>Reported Trust and Behavioral Reliance</title>
        <p>One of the most important contributions of this study is the separate examination of reported trust and actual accept or reject behavior. Lower average trust for incorrect recommendations coexisted with frequent acceptance of those recommendations. This finding supports prior concerns about automation bias in clinical decision support. Gaube et al [<xref ref-type="bibr" rid="ref35">35</xref>] found that inaccurate advice reduced physicians’ diagnostic accuracy and that substantial subgroups of both radiologists and physicians with less radiology expertise repeatedly followed incorrect advice. This study identifies a related concern in a different task: lower average trust for incorrect advice did not coincide with uniformly low acceptance of that advice. A recommendation may be accepted because it appears plausible, reduces the effort required to reach a decision, or is difficult to verify independently.</p>
        <p>At the same time, the present behavioral pattern differs from findings reported by Küper et al [<xref ref-type="bibr" rid="ref27">27</xref>]. In their 2025 study of 223 dermatologists, participants relied more strongly on correct than incorrect AI advice, but overall reliance on AI was low and self-reliance was high. In this study, by contrast, half of the incorrect recommendations were accepted. This difference may reflect several features: Küper et al [<xref ref-type="bibr" rid="ref27">27</xref>] studied specialists performing dermatology image-classification tasks, whereas this study included a nurse-dominant sample evaluating short primary care–style vignettes.</p>
        <p>The binary accept or reject decision used in this study provides a clear behavioral indicator, but actual clinical reliance may be more complicated. Sivaraman et al [<xref ref-type="bibr" rid="ref9">9</xref>] found that intensive care unit (ICU) clinicians did not always simply accept or reject AI treatment advice; they often accepted some parts, rejected others, or delayed action, a process described as negotiation. Thus, the current findings clearly demonstrate misalignment between reference correctness and behavior, but future studies should allow more nuanced responses that better reflect how health care professionals use recommendations in practice.</p>
        <p>Together, the aggregate findings support measuring trust and behavioral reliance separately rather than assuming that lower self-reported trust ensures rejection of incorrect advice. The present analyses did not directly test a within-participant causal pathway from trust to the accept or reject decision.</p>
      </sec>
      <sec>
        <title>Baseline AI Attitudes and Perceived Difficulty</title>
        <p>Participants who entered the study with a stronger intention to use AI generally reported higher trust across the individual cases. This finding is consistent with recent reviews showing that perceived usefulness, openness to technology, previous experience, and expectations about AI influence health care professionals’ acceptance of AI systems [<xref ref-type="bibr" rid="ref5">5</xref>,<xref ref-type="bibr" rid="ref6">6</xref>,<xref ref-type="bibr" rid="ref16">16</xref>]. This study extends that literature by showing that a general willingness to use AI was related not only to overall acceptance, but also to evaluations of specific recommendations during repeated interactions. However, the baseline trust composite itself was not clearly associated with vignette-level trust after the other variables were considered. This suggests that general trust and intention to use AI may not represent the same thing. A participant may express general trust in AI but remain unwilling to use it because of workflow, accountability, or professional concerns. Conversely, someone may be willing to use AI as a helpful tool without assuming that every recommendation is trustworthy. This distinction should be studied more directly in future work.</p>
        <p>Participants also reported lower trust in cases they perceived as more difficult. This finding suggests that difficult or ambiguous cases may make participants less certain not only about their own judgment, but also about the quality of the AI recommendation. It complements the finding reported by Gaube et al [<xref ref-type="bibr" rid="ref35">35</xref>] that the effect of inaccurate clinical advice varied across individual cases and was especially concerning when cases involved difficult-to-recognize diagnostic features. Because perceived difficulty was measured during the same clinical vignette, this study cannot conclude that difficulty caused lower trust. A case may have seemed difficult because the vignette was ambiguous, because the AI recommendation appeared inconsistent, because the participant lacked relevant knowledge, or because several of these factors occurred together. The appropriate takeaway is that participants trusted AI recommendations less on cases they personally experienced as more difficult.</p>
      </sec>
      <sec>
        <title>Limitations and Future Research</title>
        <p>Several limitations constrain interpretation. First, the study used short primary care–style clinical vignettes rather than real encounters. The cases were drafted with LLM assistance, anchored to a 2010 textbook, and reviewed by one physician, but were not independently adjudicated by multiple practicing primary care clinicians. They may not reflect current guidelines, contemporary diagnostic pathways, or the complexity of clinical practice.</p>
        <p>Second, the sample was small and nurse-dominant. Registered nurses represented 86.76% (59/68) of participants, whereas physicians represented 13.24% (9/68). Although nurses are important users of clinical information and decision support, they do not ordinarily make all primary care diagnoses or prescribe treatments independently. The mismatch between the simulated task and usual RN scope may have affected perceived difficulty, trust, and reliance, limiting ecological validity and generalization to either physician diagnosis or routine nursing practice.</p>
        <p>Third, professional designation was self-reported and was not independently verified through licensure, employer, or National Provider Identifier records. The third-party panel did not provide complete counts for invitations, prescreening, or provider-level quality exclusions. Fourth, the participant-level sample limited power for professional-group comparisons and the exploratory interaction tests; their estimates should be interpreted as hypothesis-generating.</p>
        <p>Fifth, the two case series differed in vignette content, later correctness positions, and response window. These features were not independently varied, so between-series differences cannot be attributed to response timing or case characteristics. The differing clinical vignette slopes demonstrate that the pooled longitudinal pattern was sensitive to vignette set. Sixth, correctness followed fixed series-specific sequences rather than being randomized or counterbalanced, leaving correctness partly linked to vignette content, vignette position, and reliability phase.</p>
        <p>The handling of unrecorded decisions is an additional limitation. Classifying the 50 unrecorded case series B responses as rejection represented nonacceptance, not a verified active decision. It may overstate deliberate rejection and understate acceptance among participants who provided an explicit response. This matters particularly for apparent underreliance: 26 unrecorded responses to correct recommendations contributed to the rejection category. The recorded-response sensitivity summary illustrates this dependence on coding but cannot establish what the unrecorded decisions would have been. Future studies should retain timeout or nonresponse as a separate behavioral outcome rather than assume that it is equivalent to rejection.</p>
      </sec>
      <sec>
        <title>Conclusions</title>
        <p>This nurse-dominant pilot study illustrates why clinical AI evaluation should examine both self-reported trust and observed reliance across clinical vignettes. Incorrect recommendations received lower trust on average, but half were accepted; approximately 1 in 6 correct recommendations fell into the rejection category under the stated coding rule. These findings identify a potential evaluation problem rather than a demonstrated patient-safety effect. Larger studies using independently validated, role-aligned clinical vignettes, counterbalanced experimental factors, and separate recording of active rejection and nonresponse are needed to determine how recommendation quality, feedback, and difficulty shape appropriate reliance.</p>
      </sec>
    </sec>
  </body>
  <back>
    <app-group>
      <supplementary-material id="app1">
        <label>Multimedia Appendix 1</label>
        <p>Clinical vignettes, generation procedure, and prompts.</p>
        <media xlink:href="humanfactors_v13i1e97649_app1.docx" xlink:title="DOCX File , 42 KB"/>
      </supplementary-material>
      <supplementary-material id="app2">
        <label>Multimedia Appendix 2</label>
        <p>Measurement structure and reliability.</p>
        <media xlink:href="humanfactors_v13i1e97649_app2.docx" xlink:title="DOCX File , 40 KB"/>
      </supplementary-material>
      <supplementary-material id="app3">
        <label>Multimedia Appendix 3</label>
        <p>Supplementary Analyses and Results.</p>
        <media xlink:href="humanfactors_v13i1e97649_app3.docx" xlink:title="DOCX File , 536 KB"/>
      </supplementary-material>
    </app-group>
    <glossary>
      <title>Abbreviations</title>
      <def-list>
        <def-item>
          <term id="abb1">AVE</term>
          <def>
            <p>average variance extracted</p>
          </def>
        </def-item>
        <def-item>
          <term id="abb2">CDSS</term>
          <def>
            <p>clinical decision support systems</p>
          </def>
        </def-item>
        <def-item>
          <term id="abb3">CFA</term>
          <def>
            <p>confirmatory factor analysis</p>
          </def>
        </def-item>
        <def-item>
          <term id="abb4">CONSORT-AI</term>
          <def>
            <p>Consolidated Standards of Reporting Trials–Artificial Intelligence</p>
          </def>
        </def-item>
        <def-item>
          <term id="abb5">ICU</term>
          <def>
            <p>intensive care unit</p>
          </def>
        </def-item>
        <def-item>
          <term id="abb6">IRB</term>
          <def>
            <p>Institutional Review Board</p>
          </def>
        </def-item>
        <def-item>
          <term id="abb7">LLM</term>
          <def>
            <p>large language model</p>
          </def>
        </def-item>
        <def-item>
          <term id="abb8">RN</term>
          <def>
            <p>registered nurse</p>
          </def>
        </def-item>
        <def-item>
          <term id="abb9">WVU</term>
          <def>
            <p>West Virginia University</p>
          </def>
        </def-item>
      </def-list>
    </glossary>
    <ack>
      <p>The authors gratefully acknowledge Dr Zaira Chaudhry for her expert physician review and finalization of the clinical vignettes, reference diagnoses, and AI recommendations following the authors’ textbook-based checks. Her review assessed clinical coherence, plausibility, and consistency with each recommendation’s intended correctness status before deployment. The authors declare the use of generative AI (GenAI) in the research and writing process. According to the GAIDeT (Generative AI Delegation Taxonomy; 2025), the following tasks were delegated to GenAI tools under full human supervision: literature search and systematization, code generation, reproducibility testing, proofreading and editing, and summarizing text. The GenAI tools used were Apple AI and ChatGPT 5.6. Responsibility for the final manuscript lies entirely with the authors. GenAI tools are not listed as authors and do not bear responsibility for the final outcomes. The declaration was submitted by all authors.</p>
    </ack>
    <notes>
      <title>Funding</title>
      <p>The authors declared no financial support was received for this work.</p>
    </notes>
    <fn-group>
      <fn fn-type="con">
        <p>AC was responsible for conceptualization, survey development, figure development, data collection, experimental design, and manuscript drafting. YS contributed to data analysis, survey development, figure development, data collection, experimental design, and manuscript drafting. APG contributed to critical manuscript revision. All authors reviewed and approved the final manuscript.</p>
      </fn>
      <fn fn-type="conflict">
        <p>None declared.</p>
      </fn>
    </fn-group>
    <ref-list>
      <ref id="ref1">
        <label>1</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Chustecki</surname>
              <given-names>M</given-names>
            </name>
          </person-group>
          <article-title>Benefits and risks of AI in health care: narrative review</article-title>
          <source>Interact J Med Res</source>
          <year>2024</year>
          <volume>13</volume>
          <fpage>e53616</fpage>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://www.i-jmr.org/2024//e53616/"/>
          </comment>
          <pub-id pub-id-type="doi">10.2196/53616</pub-id>
          <pub-id pub-id-type="medline">39556817</pub-id>
          <pub-id pub-id-type="pii">v13i1e53616</pub-id>
          <pub-id pub-id-type="pmcid">PMC11612599</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref2">
        <label>2</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Jiang</surname>
              <given-names>F</given-names>
            </name>
            <name name-style="western">
              <surname>Jiang</surname>
              <given-names>Y</given-names>
            </name>
            <name name-style="western">
              <surname>Zhi</surname>
              <given-names>H</given-names>
            </name>
            <name name-style="western">
              <surname>Dong</surname>
              <given-names>Y</given-names>
            </name>
            <name name-style="western">
              <surname>Li</surname>
              <given-names>H</given-names>
            </name>
            <name name-style="western">
              <surname>Ma</surname>
              <given-names>S</given-names>
            </name>
            <name name-style="western">
              <surname>Wang</surname>
              <given-names>Y</given-names>
            </name>
            <name name-style="western">
              <surname>Dong</surname>
              <given-names>Q</given-names>
            </name>
            <name name-style="western">
              <surname>Shen</surname>
              <given-names>H</given-names>
            </name>
            <name name-style="western">
              <surname>Wang</surname>
              <given-names>Y</given-names>
            </name>
          </person-group>
          <article-title>Artificial intelligence in healthcare: past, present and future</article-title>
          <source>Stroke Vasc Neurol</source>
          <year>2017</year>
          <volume>2</volume>
          <issue>4</issue>
          <fpage>230</fpage>
          <lpage>243</lpage>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://svn.bmj.com/lookup/pmidlookup?view=long&#38;pmid=29507784"/>
          </comment>
          <pub-id pub-id-type="doi">10.1136/svn-2017-000101</pub-id>
          <pub-id pub-id-type="medline">29507784</pub-id>
          <pub-id pub-id-type="pii">svn-2017-000101</pub-id>
          <pub-id pub-id-type="pmcid">PMC5829945</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref3">
        <label>3</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Taylor</surname>
              <given-names>RA</given-names>
            </name>
            <name name-style="western">
              <surname>Sangal</surname>
              <given-names>RB</given-names>
            </name>
            <name name-style="western">
              <surname>Smith</surname>
              <given-names>ME</given-names>
            </name>
            <name name-style="western">
              <surname>Haimovich</surname>
              <given-names>AD</given-names>
            </name>
            <name name-style="western">
              <surname>Rodman</surname>
              <given-names>A</given-names>
            </name>
            <name name-style="western">
              <surname>Iscoe</surname>
              <given-names>MS</given-names>
            </name>
            <name name-style="western">
              <surname>Pavuluri</surname>
              <given-names>SK</given-names>
            </name>
            <name name-style="western">
              <surname>Rose</surname>
              <given-names>C</given-names>
            </name>
            <name name-style="western">
              <surname>Janke</surname>
              <given-names>AT</given-names>
            </name>
            <name name-style="western">
              <surname>Wright</surname>
              <given-names>DS</given-names>
            </name>
            <name name-style="western">
              <surname>Socrates</surname>
              <given-names>V</given-names>
            </name>
            <name name-style="western">
              <surname>Declan</surname>
              <given-names>A</given-names>
            </name>
          </person-group>
          <article-title>Leveraging artificial intelligence to reduce diagnostic errors in emergency medicine: challenges, opportunities, and future directions</article-title>
          <source>Acad Emerg Med</source>
          <year>2025</year>
          <volume>32</volume>
          <issue>3</issue>
          <fpage>327</fpage>
          <lpage>339</lpage>
          <pub-id pub-id-type="doi">10.1111/acem.15066</pub-id>
          <pub-id pub-id-type="medline">39676165</pub-id>
          <pub-id pub-id-type="pmcid">PMC11921089</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref4">
        <label>4</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Fackler</surname>
              <given-names>J</given-names>
            </name>
            <name name-style="western">
              <surname>Ghobadi</surname>
              <given-names>K</given-names>
            </name>
            <name name-style="western">
              <surname>Gurses</surname>
              <given-names>A</given-names>
            </name>
          </person-group>
          <article-title>Algorithms at the bedside: moving past development and validation</article-title>
          <source>Pediatr Crit Care Med</source>
          <year>2024</year>
          <volume>25</volume>
          <issue>3</issue>
          <fpage>276</fpage>
          <lpage>278</lpage>
          <pub-id pub-id-type="doi">10.1097/PCC.0000000000003437</pub-id>
          <pub-id pub-id-type="medline">38451799</pub-id>
          <pub-id pub-id-type="pii">00130478-202403000-00012</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref5">
        <label>5</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Lambert</surname>
              <given-names>SI</given-names>
            </name>
            <name name-style="western">
              <surname>Madi</surname>
              <given-names>M</given-names>
            </name>
            <name name-style="western">
              <surname>Sopka</surname>
              <given-names>S</given-names>
            </name>
            <name name-style="western">
              <surname>Lenes</surname>
              <given-names>A</given-names>
            </name>
            <name name-style="western">
              <surname>Stange</surname>
              <given-names>H</given-names>
            </name>
            <name name-style="western">
              <surname>Buszello</surname>
              <given-names>C</given-names>
            </name>
            <name name-style="western">
              <surname>Stephan</surname>
              <given-names>A</given-names>
            </name>
          </person-group>
          <article-title>An integrative review on the acceptance of artificial intelligence among healthcare professionals in hospitals</article-title>
          <source>NPJ Digit Med</source>
          <year>2023</year>
          <volume>6</volume>
          <issue>1</issue>
          <fpage>111</fpage>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://doi.org/10.1038/s41746-023-00852-5"/>
          </comment>
          <pub-id pub-id-type="doi">10.1038/s41746-023-00852-5</pub-id>
          <pub-id pub-id-type="medline">37301946</pub-id>
          <pub-id pub-id-type="pii">10.1038/s41746-023-00852-5</pub-id>
          <pub-id pub-id-type="pmcid">PMC10257646</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref6">
        <label>6</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Steerling</surname>
              <given-names>E</given-names>
            </name>
            <name name-style="western">
              <surname>Siira</surname>
              <given-names>E</given-names>
            </name>
            <name name-style="western">
              <surname>Nilsen</surname>
              <given-names>P</given-names>
            </name>
            <name name-style="western">
              <surname>Svedberg</surname>
              <given-names>P</given-names>
            </name>
            <name name-style="western">
              <surname>Nygren</surname>
              <given-names>J</given-names>
            </name>
          </person-group>
          <article-title>Implementing AI in healthcare-the relevance of trust: a scoping review</article-title>
          <source>Front Health Serv</source>
          <year>2023</year>
          <volume>3</volume>
          <fpage>1211150</fpage>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://europepmc.org/abstract/MED/37693234"/>
          </comment>
          <pub-id pub-id-type="doi">10.3389/frhs.2023.1211150</pub-id>
          <pub-id pub-id-type="medline">37693234</pub-id>
          <pub-id pub-id-type="pmcid">PMC10484529</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref7">
        <label>7</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Scott</surname>
              <given-names>IA</given-names>
            </name>
            <name name-style="western">
              <surname>van der Vegt</surname>
              <given-names>A</given-names>
            </name>
            <name name-style="western">
              <surname>Lane</surname>
              <given-names>P</given-names>
            </name>
            <name name-style="western">
              <surname>McPhail</surname>
              <given-names>S</given-names>
            </name>
            <name name-style="western">
              <surname>Magrabi</surname>
              <given-names>F</given-names>
            </name>
          </person-group>
          <article-title>Achieving large-scale clinician adoption of AI-enabled decision support</article-title>
          <source>BMJ Health Care Inform</source>
          <year>2024</year>
          <volume>31</volume>
          <issue>1</issue>
          <fpage>e100971</fpage>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://informatics.bmj.com/lookup/pmidlookup?view=long&#38;pmid=38816209"/>
          </comment>
          <pub-id pub-id-type="doi">10.1136/bmjhci-2023-100971</pub-id>
          <pub-id pub-id-type="medline">38816209</pub-id>
          <pub-id pub-id-type="pii">bmjhci-2023-100971</pub-id>
          <pub-id pub-id-type="pmcid">PMC11141172</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref8">
        <label>8</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Choudhury</surname>
              <given-names>A</given-names>
            </name>
            <name name-style="western">
              <surname>Chaudhry</surname>
              <given-names>Z</given-names>
            </name>
          </person-group>
          <article-title>Large language models and user trust: consequence of self-referential learning loop and the deskilling of health care professionals</article-title>
          <source>J Med Internet Res</source>
          <year>2024</year>
          <volume>26</volume>
          <fpage>e56764</fpage>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://www.jmir.org/2024//e56764/"/>
          </comment>
          <pub-id pub-id-type="doi">10.2196/56764</pub-id>
          <pub-id pub-id-type="medline">38662419</pub-id>
          <pub-id pub-id-type="pii">v26i1e56764</pub-id>
          <pub-id pub-id-type="pmcid">PMC11082730</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref9">
        <label>9</label>
        <nlm-citation citation-type="confproc">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Sivaraman</surname>
              <given-names>V</given-names>
            </name>
            <name name-style="western">
              <surname>Bukowski</surname>
              <given-names>LA</given-names>
            </name>
            <name name-style="western">
              <surname>Levin</surname>
              <given-names>J</given-names>
            </name>
            <name name-style="western">
              <surname>Kahn</surname>
              <given-names>JM</given-names>
            </name>
            <name name-style="western">
              <surname>Perer</surname>
              <given-names>A</given-names>
            </name>
          </person-group>
          <article-title>Ignore, trust, or negotiate: understanding clinician acceptance of AI-based treatment recommendations in health care</article-title>
          <year>2023</year>
          <conf-name>CHI '23: Proceedings of the 2023 CHI Conference on Human Factors in Computing Systems</conf-name>
          <conf-date>2023 April 23 - 28</conf-date>
          <conf-loc>Hamburg, Germany</conf-loc>
          <fpage>1</fpage>
          <lpage>18</lpage>
          <pub-id pub-id-type="doi">10.1145/3544548.3581075</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref10">
        <label>10</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Choudhury</surname>
              <given-names>A</given-names>
            </name>
            <name name-style="western">
              <surname>Shamszare</surname>
              <given-names>H</given-names>
            </name>
          </person-group>
          <article-title>Human factors influencing trust in healthcare artificial intelligence: systematic literature review</article-title>
          <source>IISE Trans Occup Ergon Hum Factors</source>
          <year>2026</year>
          <volume>14</volume>
          <issue>2</issue>
          <fpage>116</fpage>
          <lpage>131</lpage>
          <pub-id pub-id-type="doi">10.1080/24725838.2026.2625655</pub-id>
          <pub-id pub-id-type="medline">41674135</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref11">
        <label>11</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Wong</surname>
              <given-names>KKL</given-names>
            </name>
            <name name-style="western">
              <surname>Han</surname>
              <given-names>Y</given-names>
            </name>
            <name name-style="western">
              <surname>Cai</surname>
              <given-names>Y</given-names>
            </name>
            <name name-style="western">
              <surname>Ouyang</surname>
              <given-names>W</given-names>
            </name>
            <name name-style="western">
              <surname>Du</surname>
              <given-names>H</given-names>
            </name>
            <name name-style="western">
              <surname>Liu</surname>
              <given-names>C</given-names>
            </name>
          </person-group>
          <article-title>From trust in automation to trust in AI in healthcare: a 30-year longitudinal review and an interdisciplinary framework</article-title>
          <source>Bioengineering (Basel)</source>
          <year>2025</year>
          <volume>12</volume>
          <issue>10</issue>
          <fpage>1070</fpage>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://www.mdpi.com/resolver?pii=bioengineering12101070"/>
          </comment>
          <pub-id pub-id-type="doi">10.3390/bioengineering12101070</pub-id>
          <pub-id pub-id-type="medline">41155069</pub-id>
          <pub-id pub-id-type="pii">bioengineering12101070</pub-id>
          <pub-id pub-id-type="pmcid">PMC12562135</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref12">
        <label>12</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Li</surname>
              <given-names>X</given-names>
            </name>
            <name name-style="western">
              <surname>Zong</surname>
              <given-names>Q</given-names>
            </name>
            <name name-style="western">
              <surname>Cheng</surname>
              <given-names>M</given-names>
            </name>
          </person-group>
          <article-title>The impact of medical explainable artificial intelligence on nurses' innovation behaviour: a structural equation modelling approach</article-title>
          <source>J Nurs Manag</source>
          <year>2024</year>
          <volume>2024</volume>
          <fpage>8885760</fpage>
          <pub-id pub-id-type="doi">10.1155/2024/8885760</pub-id>
          <pub-id pub-id-type="medline">40224836</pub-id>
          <pub-id pub-id-type="pmcid">PMC11918505</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref13">
        <label>13</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Shamszare</surname>
              <given-names>H</given-names>
            </name>
            <name name-style="western">
              <surname>Chaudhry</surname>
              <given-names>Z</given-names>
            </name>
            <name name-style="western">
              <surname>Berenji</surname>
              <given-names>M</given-names>
            </name>
            <name name-style="western">
              <surname>Choudhury</surname>
              <given-names>A</given-names>
            </name>
            <name name-style="western">
              <surname>CHOUDHURY</surname>
              <given-names>A</given-names>
            </name>
            <name name-style="western">
              <surname>CHOUDHURY</surname>
              <given-names>A</given-names>
            </name>
            <name name-style="western">
              <surname>CHOUDHURY</surname>
              <given-names>A</given-names>
            </name>
          </person-group>
          <article-title>Conceptualizing clinicians’ trust in artificial intelligence as a function of their expertise, workload, patient outcome, diagnosis difficulty, and AI accuracy: a systems thinking approach</article-title>
          <source>IEEE Access</source>
          <year>2025</year>
          <volume>13</volume>
          <fpage>119601</fpage>
          <lpage>119618</lpage>
          <pub-id pub-id-type="doi">10.1109/access.2025.3586555</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref14">
        <label>14</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Goddard</surname>
              <given-names>K</given-names>
            </name>
            <name name-style="western">
              <surname>Roudsari</surname>
              <given-names>A</given-names>
            </name>
            <name name-style="western">
              <surname>Wyatt</surname>
              <given-names>JC</given-names>
            </name>
          </person-group>
          <article-title>Automation bias: a systematic review of frequency, effect mediators, and mitigators</article-title>
          <source>J Am Med Inform Assoc</source>
          <year>2012</year>
          <volume>19</volume>
          <issue>1</issue>
          <fpage>121</fpage>
          <lpage>127</lpage>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://europepmc.org/abstract/MED/21685142"/>
          </comment>
          <pub-id pub-id-type="doi">10.1136/amiajnl-2011-000089</pub-id>
          <pub-id pub-id-type="medline">21685142</pub-id>
          <pub-id pub-id-type="pii">amiajnl-2011-000089</pub-id>
          <pub-id pub-id-type="pmcid">PMC3240751</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref15">
        <label>15</label>
        <nlm-citation citation-type="confproc">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Vereschak</surname>
              <given-names>O</given-names>
            </name>
            <name name-style="western">
              <surname>Alizadeh</surname>
              <given-names>F</given-names>
            </name>
            <name name-style="western">
              <surname>Bailly</surname>
              <given-names>G</given-names>
            </name>
            <name name-style="western">
              <surname>Caramiaux</surname>
              <given-names>B</given-names>
            </name>
          </person-group>
          <article-title>Trust in ai-assisted decision making: perspectives from those behind the system and those for whom the decision is made</article-title>
          <year>2024</year>
          <conf-name>CHI '24: Proceedings of the 2024 CHI Conference on Human Factors in Computing Systems</conf-name>
          <conf-date>2024 May 11 - 16</conf-date>
          <conf-loc>Honolulu, HI, USA</conf-loc>
          <fpage>1</fpage>
          <lpage>14</lpage>
          <pub-id pub-id-type="doi">10.1145/3613904.3642018</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref16">
        <label>16</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Tun</surname>
              <given-names>HM</given-names>
            </name>
            <name name-style="western">
              <surname>Rahman</surname>
              <given-names>HA</given-names>
            </name>
            <name name-style="western">
              <surname>Naing</surname>
              <given-names>L</given-names>
            </name>
            <name name-style="western">
              <surname>Malik</surname>
              <given-names>OA</given-names>
            </name>
          </person-group>
          <article-title>Trust in artificial intelligence-based clinical decision support systems among health care workers: systematic review</article-title>
          <source>J Med Internet Res</source>
          <year>2025</year>
          <volume>27</volume>
          <fpage>e69678</fpage>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://www.jmir.org/2025//e69678/"/>
          </comment>
          <pub-id pub-id-type="doi">10.2196/69678</pub-id>
          <pub-id pub-id-type="medline">40772775</pub-id>
          <pub-id pub-id-type="pii">v27i1e69678</pub-id>
          <pub-id pub-id-type="pmcid">PMC12440830</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref17">
        <label>17</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Sato</surname>
              <given-names>T</given-names>
            </name>
            <name name-style="western">
              <surname>Inman</surname>
              <given-names>J</given-names>
            </name>
            <name name-style="western">
              <surname>Politowicz</surname>
              <given-names>MS</given-names>
            </name>
            <name name-style="western">
              <surname>Chancey</surname>
              <given-names>ET</given-names>
            </name>
            <name name-style="western">
              <surname>Yamani</surname>
              <given-names>Y</given-names>
            </name>
          </person-group>
          <article-title>A meta-analytic approach to investigating the relationship between human-automation trust and attention allocation</article-title>
          <source>Proc Hum Factors Ergon Soc Annu Meet</source>
          <year>2023</year>
          <volume>67</volume>
          <issue>1</issue>
          <fpage>959</fpage>
          <lpage>964</lpage>
          <pub-id pub-id-type="doi">10.1177/21695067231194333</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref18">
        <label>18</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Kiani</surname>
              <given-names>A</given-names>
            </name>
            <name name-style="western">
              <surname>Uyumazturk</surname>
              <given-names>B</given-names>
            </name>
            <name name-style="western">
              <surname>Rajpurkar</surname>
              <given-names>P</given-names>
            </name>
            <name name-style="western">
              <surname>Wang</surname>
              <given-names>A</given-names>
            </name>
            <name name-style="western">
              <surname>Gao</surname>
              <given-names>R</given-names>
            </name>
            <name name-style="western">
              <surname>Jones</surname>
              <given-names>E</given-names>
            </name>
            <name name-style="western">
              <surname>Yu</surname>
              <given-names>Y</given-names>
            </name>
            <name name-style="western">
              <surname>Langlotz</surname>
              <given-names>CP</given-names>
            </name>
            <name name-style="western">
              <surname>Ball</surname>
              <given-names>RL</given-names>
            </name>
            <name name-style="western">
              <surname>Montine</surname>
              <given-names>TJ</given-names>
            </name>
            <name name-style="western">
              <surname>Martin</surname>
              <given-names>BA</given-names>
            </name>
            <name name-style="western">
              <surname>Berry</surname>
              <given-names>GJ</given-names>
            </name>
            <name name-style="western">
              <surname>Ozawa</surname>
              <given-names>MG</given-names>
            </name>
            <name name-style="western">
              <surname>Hazard</surname>
              <given-names>FK</given-names>
            </name>
            <name name-style="western">
              <surname>Brown</surname>
              <given-names>RA</given-names>
            </name>
            <name name-style="western">
              <surname>Chen</surname>
              <given-names>SB</given-names>
            </name>
            <name name-style="western">
              <surname>Wood</surname>
              <given-names>M</given-names>
            </name>
            <name name-style="western">
              <surname>Allard</surname>
              <given-names>LS</given-names>
            </name>
            <name name-style="western">
              <surname>Ylagan</surname>
              <given-names>L</given-names>
            </name>
            <name name-style="western">
              <surname>Ng</surname>
              <given-names>AY</given-names>
            </name>
            <name name-style="western">
              <surname>Shen</surname>
              <given-names>J</given-names>
            </name>
          </person-group>
          <article-title>Impact of a deep learning assistant on the histopathologic classification of liver cancer</article-title>
          <source>NPJ Digit Med</source>
          <year>2020</year>
          <volume>3</volume>
          <fpage>23</fpage>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://doi.org/10.1038/s41746-020-0232-8"/>
          </comment>
          <pub-id pub-id-type="doi">10.1038/s41746-020-0232-8</pub-id>
          <pub-id pub-id-type="medline">32140566</pub-id>
          <pub-id pub-id-type="pii">232</pub-id>
          <pub-id pub-id-type="pmcid">PMC7044422</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref19">
        <label>19</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Tschandl</surname>
              <given-names>P</given-names>
            </name>
            <name name-style="western">
              <surname>Rinner</surname>
              <given-names>C</given-names>
            </name>
            <name name-style="western">
              <surname>Apalla</surname>
              <given-names>Z</given-names>
            </name>
            <name name-style="western">
              <surname>Argenziano</surname>
              <given-names>G</given-names>
            </name>
            <name name-style="western">
              <surname>Codella</surname>
              <given-names>N</given-names>
            </name>
            <name name-style="western">
              <surname>Halpern</surname>
              <given-names>A</given-names>
            </name>
            <name name-style="western">
              <surname>Janda</surname>
              <given-names>M</given-names>
            </name>
            <name name-style="western">
              <surname>Lallas</surname>
              <given-names>A</given-names>
            </name>
            <name name-style="western">
              <surname>Longo</surname>
              <given-names>C</given-names>
            </name>
            <name name-style="western">
              <surname>Malvehy</surname>
              <given-names>J</given-names>
            </name>
            <name name-style="western">
              <surname>Paoli</surname>
              <given-names>J</given-names>
            </name>
            <name name-style="western">
              <surname>Puig</surname>
              <given-names>S</given-names>
            </name>
            <name name-style="western">
              <surname>Rosendahl</surname>
              <given-names>C</given-names>
            </name>
            <name name-style="western">
              <surname>Soyer</surname>
              <given-names>HP</given-names>
            </name>
            <name name-style="western">
              <surname>Zalaudek</surname>
              <given-names>I</given-names>
            </name>
            <name name-style="western">
              <surname>Kittler</surname>
              <given-names>H</given-names>
            </name>
          </person-group>
          <article-title>Human-computer collaboration for skin cancer recognition</article-title>
          <source>Nat Med</source>
          <year>2020</year>
          <volume>26</volume>
          <issue>8</issue>
          <fpage>1229</fpage>
          <lpage>1234</lpage>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://hdl.handle.net/11368/2968193"/>
          </comment>
          <pub-id pub-id-type="doi">10.1038/s41591-020-0942-0</pub-id>
          <pub-id pub-id-type="medline">32572267</pub-id>
          <pub-id pub-id-type="pii">10.1038/s41591-020-0942-0</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref20">
        <label>20</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Chaudhry</surname>
              <given-names>ZS</given-names>
            </name>
            <name name-style="western">
              <surname>Choudhury</surname>
              <given-names>A</given-names>
            </name>
          </person-group>
          <article-title>Clinical applications of artificial intelligence in occupational health: a systematic literature review</article-title>
          <source>J Occup Environ Med</source>
          <year>2024</year>
          <volume>66</volume>
          <issue>12</issue>
          <fpage>943</fpage>
          <lpage>955</lpage>
          <pub-id pub-id-type="doi">10.1097/JOM.0000000000003212</pub-id>
          <pub-id pub-id-type="medline">39190393</pub-id>
          <pub-id pub-id-type="pii">00043764-202412000-00001</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref21">
        <label>21</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Choudhury</surname>
              <given-names>A</given-names>
            </name>
          </person-group>
          <article-title>Toward an ecologically valid conceptual framework for the use of artificial intelligence in clinical settings: need for systems thinking, accountability, decision-making, trust, and patient safety considerations in safeguarding the technology and clinicians</article-title>
          <source>JMIR Hum Factors</source>
          <year>2022</year>
          <volume>9</volume>
          <issue>2</issue>
          <fpage>e35421</fpage>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://humanfactors.jmir.org/2022/2/e35421/"/>
          </comment>
          <pub-id pub-id-type="doi">10.2196/35421</pub-id>
          <pub-id pub-id-type="medline">35727615</pub-id>
          <pub-id pub-id-type="pii">v9i2e35421</pub-id>
          <pub-id pub-id-type="pmcid">PMC9257623</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref22">
        <label>22</label>
        <nlm-citation citation-type="book">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Browne</surname>
              <given-names>JT</given-names>
            </name>
            <name name-style="western">
              <surname>Bakker</surname>
              <given-names>S</given-names>
            </name>
            <name name-style="western">
              <surname>Yu</surname>
              <given-names>B</given-names>
            </name>
            <name name-style="western">
              <surname>Lloyd</surname>
              <given-names>P</given-names>
            </name>
            <name name-style="western">
              <surname>Ben</surname>
              <given-names>AS</given-names>
            </name>
          </person-group>
          <article-title>Trust in clinical AI: expanding the unit of analysis</article-title>
          <source>HHAI2022: Augmenting Human Intellect</source>
          <year>2022</year>
          <publisher-loc>Amsterdam, the Netherlands</publisher-loc>
          <publisher-name>IOS Press</publisher-name>
          <fpage>96</fpage>
          <lpage>113</lpage>
        </nlm-citation>
      </ref>
      <ref id="ref23">
        <label>23</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Lee</surname>
              <given-names>JD</given-names>
            </name>
            <name name-style="western">
              <surname>See</surname>
              <given-names>KA</given-names>
            </name>
          </person-group>
          <article-title>Trust in automation: designing for appropriate reliance</article-title>
          <source>Hum Factors</source>
          <year>2004</year>
          <volume>46</volume>
          <issue>1</issue>
          <fpage>50</fpage>
          <lpage>80</lpage>
          <pub-id pub-id-type="doi">10.1518/hfes.46.1.50_30392</pub-id>
          <pub-id pub-id-type="medline">15151155</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref24">
        <label>24</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Hoff</surname>
              <given-names>KA</given-names>
            </name>
            <name name-style="western">
              <surname>Bashir</surname>
              <given-names>M</given-names>
            </name>
          </person-group>
          <article-title>Trust in automation: integrating empirical evidence on factors that influence trust</article-title>
          <source>Hum Factors</source>
          <year>2015</year>
          <volume>57</volume>
          <issue>3</issue>
          <fpage>407</fpage>
          <lpage>434</lpage>
          <pub-id pub-id-type="doi">10.1177/0018720814547570</pub-id>
          <pub-id pub-id-type="medline">25875432</pub-id>
          <pub-id pub-id-type="pii">0018720814547570</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref25">
        <label>25</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Lyell</surname>
              <given-names>D</given-names>
            </name>
            <name name-style="western">
              <surname>Coiera</surname>
              <given-names>E</given-names>
            </name>
          </person-group>
          <article-title>Automation bias and verification complexity: a systematic review</article-title>
          <source>J Am Med Inform Assoc</source>
          <year>2017</year>
          <volume>24</volume>
          <issue>2</issue>
          <fpage>423</fpage>
          <lpage>431</lpage>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://europepmc.org/abstract/MED/27516495"/>
          </comment>
          <pub-id pub-id-type="doi">10.1093/jamia/ocw105</pub-id>
          <pub-id pub-id-type="medline">27516495</pub-id>
          <pub-id pub-id-type="pii">ocw105</pub-id>
          <pub-id pub-id-type="pmcid">PMC7651899</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref26">
        <label>26</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Sakamoto</surname>
              <given-names>T</given-names>
            </name>
            <name name-style="western">
              <surname>Harada</surname>
              <given-names>Y</given-names>
            </name>
            <name name-style="western">
              <surname>Shimizu</surname>
              <given-names>T</given-names>
            </name>
          </person-group>
          <article-title>Facilitating trust calibration in artificial intelligence-driven diagnostic decision support systems for determining physicians' diagnostic accuracy: quasi-experimental study</article-title>
          <source>JMIR Form Res</source>
          <year>2024</year>
          <volume>8</volume>
          <fpage>e58666</fpage>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://formative.jmir.org/2024//e58666/"/>
          </comment>
          <pub-id pub-id-type="doi">10.2196/58666</pub-id>
          <pub-id pub-id-type="medline">39602469</pub-id>
          <pub-id pub-id-type="pii">v8i1e58666</pub-id>
          <pub-id pub-id-type="pmcid">PMC11612524</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref27">
        <label>27</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Küper</surname>
              <given-names>A</given-names>
            </name>
            <name name-style="western">
              <surname>Lodde</surname>
              <given-names>GC</given-names>
            </name>
            <name name-style="western">
              <surname>Livingstone</surname>
              <given-names>E</given-names>
            </name>
            <name name-style="western">
              <surname>Schadendorf</surname>
              <given-names>D</given-names>
            </name>
            <name name-style="western">
              <surname>Krämer</surname>
              <given-names>N</given-names>
            </name>
          </person-group>
          <article-title>Psychological factors influencing appropriate reliance on AI-enabled clinical decision support systems: experimental web-based study among dermatologists</article-title>
          <source>J Med Internet Res</source>
          <year>2025</year>
          <volume>27</volume>
          <fpage>e58660</fpage>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://www.jmir.org/2025//e58660/"/>
          </comment>
          <pub-id pub-id-type="doi">10.2196/58660</pub-id>
          <pub-id pub-id-type="medline">40184614</pub-id>
          <pub-id pub-id-type="pii">v27i1e58660</pub-id>
          <pub-id pub-id-type="pmcid">PMC12008695</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref28">
        <label>28</label>
        <nlm-citation citation-type="book">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>McPhee</surname>
              <given-names>SJ</given-names>
            </name>
            <name name-style="western">
              <surname>Papadakis</surname>
              <given-names>MA</given-names>
            </name>
          </person-group>
          <source>Current Medical Diagnosis and Treatment 2010</source>
          <year>2010</year>
          <publisher-loc>New York, NY</publisher-loc>
          <publisher-name>McGraw-Hill Medical</publisher-name>
        </nlm-citation>
      </ref>
      <ref id="ref29">
        <label>29</label>
        <nlm-citation citation-type="web">
          <article-title>Audience panel and research services</article-title>
          <source>Centiment</source>
          <access-date>2026-09-10</access-date>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://www.centiment.co/">https://www.centiment.co/</ext-link>
          </comment>
        </nlm-citation>
      </ref>
      <ref id="ref30">
        <label>30</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Bandura</surname>
              <given-names>A</given-names>
            </name>
          </person-group>
          <article-title>Self-efficacy: toward a unifying theory of behavioral change</article-title>
          <source>Psychol Rev</source>
          <year>1977</year>
          <volume>84</volume>
          <issue>2</issue>
          <fpage>191</fpage>
          <lpage>215</lpage>
          <pub-id pub-id-type="doi">10.1037//0033-295x.84.2.191</pub-id>
          <pub-id pub-id-type="medline">847061</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref31">
        <label>31</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Liu</surname>
              <given-names>X</given-names>
            </name>
            <name name-style="western">
              <surname>Cruz Rivera</surname>
              <given-names>S</given-names>
            </name>
            <name name-style="western">
              <surname>Moher</surname>
              <given-names>D</given-names>
            </name>
            <name name-style="western">
              <surname>Calvert</surname>
              <given-names>MJ</given-names>
            </name>
            <name name-style="western">
              <surname>Denniston</surname>
              <given-names>AK</given-names>
            </name>
            <collab>SPIRIT-AICONSORT-AI Working Group</collab>
          </person-group>
          <article-title>Reporting guidelines for clinical trial reports for interventions involving artificial intelligence: the CONSORT-AI extension</article-title>
          <source>Nat Med</source>
          <year>2020</year>
          <volume>26</volume>
          <issue>9</issue>
          <fpage>1364</fpage>
          <lpage>1374</lpage>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://europepmc.org/abstract/MED/32908283"/>
          </comment>
          <pub-id pub-id-type="doi">10.1038/s41591-020-1034-x</pub-id>
          <pub-id pub-id-type="medline">32908283</pub-id>
          <pub-id pub-id-type="pii">10.1038/s41591-020-1034-x</pub-id>
          <pub-id pub-id-type="pmcid">PMC7598943</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref32">
        <label>32</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Kahr</surname>
              <given-names>PK</given-names>
            </name>
            <name name-style="western">
              <surname>Rooks</surname>
              <given-names>G</given-names>
            </name>
            <name name-style="western">
              <surname>Willemsen</surname>
              <given-names>MC</given-names>
            </name>
            <name name-style="western">
              <surname>Snijders</surname>
              <given-names>CCP</given-names>
            </name>
          </person-group>
          <article-title>Understanding trust and reliance development in AI advice: assessing model accuracy, model explanations, and experiences from previous interactions</article-title>
          <source>ACM Trans. Interact. Intell. Syst</source>
          <year>2024</year>
          <volume>14</volume>
          <issue>4</issue>
          <fpage>1</fpage>
          <lpage>30</lpage>
          <pub-id pub-id-type="doi">10.1145/3686164</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref33">
        <label>33</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Yang</surname>
              <given-names>XJ</given-names>
            </name>
            <name name-style="western">
              <surname>Schemanske</surname>
              <given-names>C</given-names>
            </name>
            <name name-style="western">
              <surname>Searle</surname>
              <given-names>C</given-names>
            </name>
          </person-group>
          <article-title>Toward quantifying trust dynamics: how people adjust their trust after moment-to-moment interaction with automation</article-title>
          <source>Hum Factors</source>
          <year>2023</year>
          <volume>65</volume>
          <issue>5</issue>
          <fpage>862</fpage>
          <lpage>878</lpage>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://journals.sagepub.com/doi/10.1177/00187208211034716?url_ver=Z39.88-2003&#38;rfr_id=ori:rid:crossref.org&#38;rfr_dat=cr_pub  0pubmed"/>
          </comment>
          <pub-id pub-id-type="doi">10.1177/00187208211034716</pub-id>
          <pub-id pub-id-type="medline">34459266</pub-id>
          <pub-id pub-id-type="pmcid">PMC10374998</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref34">
        <label>34</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Rittenberg</surname>
              <given-names>BSP</given-names>
            </name>
            <name name-style="western">
              <surname>Holland</surname>
              <given-names>CW</given-names>
            </name>
            <name name-style="western">
              <surname>Barnhart</surname>
              <given-names>GE</given-names>
            </name>
            <name name-style="western">
              <surname>Gaudreau</surname>
              <given-names>SM</given-names>
            </name>
            <name name-style="western">
              <surname>Neyedli</surname>
              <given-names>HF</given-names>
            </name>
          </person-group>
          <article-title>Trust with increasing and decreasing reliability</article-title>
          <source>Hum Factors</source>
          <year>2024</year>
          <volume>66</volume>
          <issue>12</issue>
          <fpage>2569</fpage>
          <lpage>2589</lpage>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://journals.sagepub.com/doi/10.1177/00187208241228636?url_ver=Z39.88-2003&#38;rfr_id=ori:rid:crossref.org&#38;rfr_dat=cr_pub  0pubmed"/>
          </comment>
          <pub-id pub-id-type="doi">10.1177/00187208241228636</pub-id>
          <pub-id pub-id-type="medline">38445652</pub-id>
          <pub-id pub-id-type="pmcid">PMC11487872</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref35">
        <label>35</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Gaube</surname>
              <given-names>S</given-names>
            </name>
            <name name-style="western">
              <surname>Suresh</surname>
              <given-names>H</given-names>
            </name>
            <name name-style="western">
              <surname>Raue</surname>
              <given-names>M</given-names>
            </name>
            <name name-style="western">
              <surname>Merritt</surname>
              <given-names>A</given-names>
            </name>
            <name name-style="western">
              <surname>Berkowitz</surname>
              <given-names>SJ</given-names>
            </name>
            <name name-style="western">
              <surname>Lermer</surname>
              <given-names>E</given-names>
            </name>
            <name name-style="western">
              <surname>Coughlin</surname>
              <given-names>JF</given-names>
            </name>
            <name name-style="western">
              <surname>Guttag</surname>
              <given-names>JV</given-names>
            </name>
            <name name-style="western">
              <surname>Colak</surname>
              <given-names>E</given-names>
            </name>
            <name name-style="western">
              <surname>Ghassemi</surname>
              <given-names>M</given-names>
            </name>
          </person-group>
          <article-title>Do as AI say: susceptibility in deployment of clinical decision-aids</article-title>
          <source>NPJ Digit Med</source>
          <year>2021</year>
          <volume>4</volume>
          <issue>1</issue>
          <fpage>31</fpage>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://doi.org/10.1038/s41746-021-00385-9"/>
          </comment>
          <pub-id pub-id-type="doi">10.1038/s41746-021-00385-9</pub-id>
          <pub-id pub-id-type="medline">33608629</pub-id>
          <pub-id pub-id-type="pii">10.1038/s41746-021-00385-9</pub-id>
          <pub-id pub-id-type="pmcid">PMC7896064</pub-id>
        </nlm-citation>
      </ref>
    </ref-list>
  </back>
</article>
