<?xml version="1.0" encoding="UTF-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Archiving and Interchange DTD v2.3 20070202//EN" "archivearticle.dtd">
<article article-type="systematic-review" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" xml:lang="EN">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Digit. Health</journal-id>
<journal-title>Frontiers in Digital Health</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Digit. Health</abbrev-journal-title>
<issn pub-type="epub">2673-253X</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="doi">10.3389/fdgth.2025.1482712</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Digital Health</subject>
<subj-group>
<subject>Systematic Review</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Comparative analysis of ChatGPT and Gemini (Bard) in medical inquiry: a scoping review</article-title>
</title-group>
<contrib-group>
<contrib contrib-type="author"><name><surname>Fattah</surname><given-names>Fattah H.</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="aff" rid="aff2"><sup>2</sup></xref><role content-type="https://credit.niso.org/contributor-roles/investigation/"/><role content-type="https://credit.niso.org/contributor-roles/methodology/"/><role content-type="https://credit.niso.org/contributor-roles/validation/"/><role content-type="https://credit.niso.org/contributor-roles/visualization/"/><role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/></contrib>
<contrib contrib-type="author"><name><surname>Salih</surname><given-names>Abdulwahid M.</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="aff" rid="aff2"><sup>2</sup></xref><role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/><role content-type="https://credit.niso.org/contributor-roles/investigation/"/><role content-type="https://credit.niso.org/contributor-roles/methodology/"/><role content-type="https://credit.niso.org/contributor-roles/resources/"/><role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/></contrib>
<contrib contrib-type="author"><name><surname>Salih</surname><given-names>Ameer M.</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="aff" rid="aff3"><sup>3</sup></xref><role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/><role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/><role content-type="https://credit.niso.org/contributor-roles/methodology/"/><role content-type="https://credit.niso.org/contributor-roles/resources/"/><role content-type="https://credit.niso.org/contributor-roles/validation/"/><role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/></contrib>
<contrib contrib-type="author"><name><surname>Asaad</surname><given-names>Saywan K.</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="aff" rid="aff2"><sup>2</sup></xref><role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/><role content-type="https://credit.niso.org/contributor-roles/resources/"/><role content-type="https://credit.niso.org/contributor-roles/software/"/><role content-type="https://credit.niso.org/contributor-roles/supervision/"/><role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/></contrib>
<contrib contrib-type="author"><name><surname>Ghafour</surname><given-names>Abdullah K.</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref><role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/><role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/><role content-type="https://credit.niso.org/contributor-roles/methodology/"/><role content-type="https://credit.niso.org/contributor-roles/resources/"/><role content-type="https://credit.niso.org/contributor-roles/validation/"/><role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/></contrib>
<contrib contrib-type="author"><name><surname>Bapir</surname><given-names>Rawa</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="aff" rid="aff4"><sup>4</sup></xref>
<xref ref-type="aff" rid="aff5"><sup>5</sup></xref><uri xlink:href="https://loop.frontiersin.org/people/2934194/overview"/><role content-type="https://credit.niso.org/contributor-roles/data-curation/"/><role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/><role content-type="https://credit.niso.org/contributor-roles/investigation/"/><role content-type="https://credit.niso.org/contributor-roles/methodology/"/><role content-type="https://credit.niso.org/contributor-roles/validation/"/><role content-type="https://credit.niso.org/contributor-roles/visualization/"/><role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/></contrib>
<contrib contrib-type="author"><name><surname>Abdalla</surname><given-names>Berun A.</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="aff" rid="aff5"><sup>5</sup></xref><role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/><role content-type="https://credit.niso.org/contributor-roles/data-curation/"/><role content-type="https://credit.niso.org/contributor-roles/visualization/"/><role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/><role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/></contrib>
<contrib contrib-type="author"><name><surname>Othman</surname><given-names>Snur</given-names></name>
<xref ref-type="aff" rid="aff5"><sup>5</sup></xref><role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/><role content-type="https://credit.niso.org/contributor-roles/methodology/"/><role content-type="https://credit.niso.org/contributor-roles/project-administration/"/><role content-type="https://credit.niso.org/contributor-roles/visualization/"/><role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/></contrib>
<contrib contrib-type="author"><name><surname>Ahmed</surname><given-names>Sasan M.</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="aff" rid="aff5"><sup>5</sup></xref><role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/><role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/><role content-type="https://credit.niso.org/contributor-roles/methodology/"/><role content-type="https://credit.niso.org/contributor-roles/visualization/"/><role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/></contrib>
<contrib contrib-type="author"><name><surname>Hasan</surname><given-names>Sabah Jalal</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref><role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/><role content-type="https://credit.niso.org/contributor-roles/methodology/"/><role content-type="https://credit.niso.org/contributor-roles/resources/"/><role content-type="https://credit.niso.org/contributor-roles/validation/"/><role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/></contrib>
<contrib contrib-type="author"><name><surname>Mahmood</surname><given-names>Yousif M.</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref><role content-type="https://credit.niso.org/contributor-roles/data-curation/"/><role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/><role content-type="https://credit.niso.org/contributor-roles/software/"/><role content-type="https://credit.niso.org/contributor-roles/supervision/"/><role content-type="https://credit.niso.org/contributor-roles/visualization/"/><role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/></contrib>
<contrib contrib-type="author" corresp="yes"><name><surname>Kakamad</surname><given-names>Fahmi H.</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="aff" rid="aff2"><sup>2</sup></xref>
<xref ref-type="aff" rid="aff5"><sup>5</sup></xref>
<xref ref-type="corresp" rid="cor1">&#x002A;</xref><uri xlink:href="https://loop.frontiersin.org/people/429700/overview" /><role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/><role content-type="https://credit.niso.org/contributor-roles/data-curation/"/><role content-type="https://credit.niso.org/contributor-roles/validation/"/><role content-type="https://credit.niso.org/contributor-roles/visualization/"/><role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/></contrib>
</contrib-group>
<aff id="aff1"><label><sup>1</sup></label><addr-line>Scientific Affairs Department</addr-line>, <institution>Smart Health Tower</institution>, <addr-line>Sulaymaniyah</addr-line>, <country>Iraq</country></aff>
<aff id="aff2"><label><sup>2</sup></label><institution>College of Medicine, University of Sulaimani</institution>, <addr-line>Sulaymaniyah</addr-line>, <country>Iraq</country></aff>
<aff id="aff3"><label><sup>3</sup></label><institution>Civil Engineering Department, College of Engineering, University of Sulaimani</institution>, <addr-line>Sulaymaniyah</addr-line>, <country>Iraq</country></aff>
<aff id="aff4"><label><sup>4</sup></label><institution>Department of Urology, Sulaimani Surgical Teaching Hospital</institution>, <addr-line>Sulaymaniyah</addr-line>, <country>Iraq</country></aff>
<aff id="aff5"><label><sup>5</sup></label><institution>Kscien Organization for Scientific Research (Middle East Office)</institution>, <addr-line>Sulaymaniyah</addr-line>, <country>Iraq</country></aff>
<author-notes>
<fn fn-type="edited-by"><p><bold>Edited by:</bold> Hosna Salmani, Iran University of Medical Sciences, Iran</p></fn>
<fn fn-type="edited-by"><p><bold>Reviewed by:</bold> Larry R. Price, Texas State University, United States</p>
<p>Thomas F. Heston, University of Washington, United States</p></fn>
<corresp id="cor1"><label>&#x002A;</label><bold>Correspondence:</bold> Fahmi H. Kakamad <email>fahmi.hussein@univsul.edu.iq</email></corresp>
</author-notes>
<pub-date pub-type="epub"><day>03</day><month>02</month><year>2025</year></pub-date>
<pub-date pub-type="collection"><year>2025</year></pub-date>
<volume>7</volume><elocation-id>1482712</elocation-id>
<history>
<date date-type="received"><day>21</day><month>08</month><year>2024</year></date>
<date date-type="accepted"><day>21</day><month>01</month><year>2025</year></date>
</history>
<permissions>
<copyright-statement>&#x00A9; 2025 Fattah, Salih, Salih, Asaad, Ghafour, Bapir, Abdalla, Othman, Ahmed, Hasan, Mahmood and Kakamad.</copyright-statement>
<copyright-year>2025</copyright-year><copyright-holder>Fattah, Salih, Salih, Asaad, Ghafour, Bapir, Abdalla, Othman, Ahmed, Hasan, Mahmood and Kakamad</copyright-holder><license license-type="open-access" xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the <ext-link ext-link-type="uri" xlink:href="http://creativecommons.org/licenses/by/4.0/">Creative Commons Attribution License (CC BY)</ext-link>. The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p></license>
</permissions>
<abstract><sec><title>Introduction</title>
<p>Artificial intelligence and machine learning are popular interconnected technologies. AI chatbots like ChatGPT and Gemini show considerable promise in medical inquiries. This scoping review aims to assess the accuracy and response length (in characters) of ChatGPT and Gemini in medical applications.</p>
</sec><sec><title>Methods</title>
<p>The eligible databases were searched to find studies published in English from January 1 to October 20, 2023. The inclusion criteria consisted of studies that focused on using AI in medicine and assessed outcomes based on the accuracy and character count (length) of ChatGPT and Gemini. Data collected from the studies included the first author&#x0027;s name, the country where the study was conducted, the type of study design, publication year, sample size, medical speciality, and the accuracy and response length.</p>
</sec><sec><title>Results</title>
<p>The initial search identified 64 papers, with 11 meeting the inclusion criteria, involving 1,177 samples. ChatGPT showed higher accuracy in radiology (87.43&#x0025; vs. Gemini&#x0027;s 71&#x0025;) and shorter responses (907 vs. 1,428 characters). Similar trends were noted in other specialties. However, Gemini outperformed ChatGPT in emergency scenarios (87&#x0025; vs. 77&#x0025;) and in renal diets with low potassium and high phosphorus (79&#x0025; vs. 60&#x0025; and 100&#x0025; vs. 77&#x0025;). Statistical analysis confirms that ChatGPT has greater accuracy and shorter responses than Gemini in medical studies, with a <italic>p</italic>-value of &#x003C;.001 for both metrics.</p>
</sec><sec><title>Conclusion</title>
<p>This Scoping review suggests that ChatGPT may demonstrate higher accuracy and provide shorter responses than Gemini in medical studies.</p>
</sec>
</abstract>
<kwd-group>
<kwd>ChatGPT</kwd>
<kwd>Google Bard</kwd>
<kwd>medical inquiries</kwd>
<kwd>comparison</kwd>
<kwd>madical AI</kwd>
</kwd-group><counts>
<fig-count count="1"/>
<table-count count="3"/><equation-count count="0"/><ref-count count="27"/><page-count count="7"/><word-count count="0"/></counts><custom-meta-wrap><custom-meta><meta-name>section-at-acceptance</meta-name><meta-value>Connected Health</meta-value></custom-meta></custom-meta-wrap>
</article-meta>
</front>
<body><sec id="s1" sec-type="intro"><title>Introduction</title>
<p>Artificial Intelligence (AI) and Machine Learning (ML) are interconnected technologies recently gaining significant popularity. AI involves creating intelligent machines capable of performing tasks that typically require human intelligence, such as visual perception, speech recognition, decision-making, and language translation. ML, a subset of AI, focuses on developing algorithms and statistical models that enable machines to learn from data and improve their performance over time without explicit programming (<xref ref-type="bibr" rid="B1">1</xref>, <xref ref-type="bibr" rid="B2">2</xref>). AI is at the forefront of transforming various aspects of our lives by altering how we analyze information and enhancing decision-making through problem-solving, reasoning, and learning (<xref ref-type="bibr" rid="B3">3</xref>).</p>
<p>In the dynamic domain of AI chatbots, the comparative analysis of ChatGPT and Gemini (formerly known as Google&#x0027;s Bard) has emerged as a focal point, particularly in medical inquiries (<xref ref-type="bibr" rid="B1">1</xref>, <xref ref-type="bibr" rid="B2">2</xref>). Recent investigations have explored the precision and effectiveness of these AI models in fielding medical questions across various specialties (<xref ref-type="bibr" rid="B1">1</xref>, <xref ref-type="bibr" rid="B4">4</xref>&#x2013;<xref ref-type="bibr" rid="B7">7</xref>). These studies demonstrate ChatGPT&#x0027;s capabilities in diagnostic imaging and clinical decision support, underscoring its potential value in healthcare settings (<xref ref-type="bibr" rid="B4">4</xref>&#x2013;<xref ref-type="bibr" rid="B6">6</xref>).</p>
<p>In recent years, AI models like ChatGPT and Gemini have significantly impacted natural language processing, particularly in healthcare. ChatGPT, developed by OpenAI, provides relevant and accurate text-based responses using a large dataset (<xref ref-type="bibr" rid="B1">1</xref>, <xref ref-type="bibr" rid="B5">5</xref>). While Gemini, from Google DeepMind, integrates multimodal capabilities, handling text, audio, and video, which is especially useful in medical imaging (<xref ref-type="bibr" rid="B5">5</xref>). However, both models face challenges. ChatGPT, for example, shows variability in psychiatric assessments and struggles with complex cases (<xref ref-type="bibr" rid="B2">2</xref>). Additionally, AI models still struggle to interpret nuanced human emotions and contexts (<xref ref-type="bibr" rid="B4">4</xref>).</p>
<p>While AI chatbots like ChatGPT and Gemini show promise in medicine, extensive research is still required to understand their capabilities properly. It is essential to address the variation in their performance across different medical scenarios and enhance their accuracy for various medical applications (<xref ref-type="bibr" rid="B8">8</xref>). The use of AI in healthcare faces several challenges, including data privacy, algorithm accuracy, adherence to ethical standards, societal acceptance, and clinical integration (<xref ref-type="bibr" rid="B9">9</xref>, <xref ref-type="bibr" rid="B10">10</xref>). These challenges make it difficult to develop precise and reliable AI systems. Privacy concerns restrict access to relevant data, and potential biases can result in inaccurate outcomes (<xref ref-type="bibr" rid="B8">8</xref>, <xref ref-type="bibr" rid="B10">10</xref>).</p>
<p>This scoping review aims to evaluate and compare the accuracy and length of ChatGPT and Gemini (Google&#x0027;s Bard) in addressing medical inquiries across diverse fields, focusing on their strengths, limitations, and practical implications for healthcare. As AI models become increasingly integrated into clinical and educational settings, understanding their performance variability is essential. Both models face challenges, including inconsistencies in complex cases, privacy concerns, and ethical issues. This review offers insights to help researchers, practitioners, and developers optimize these tools for more effective decision-making and patient care.</p>
</sec>
<sec id="s2" sec-type="methods"><title>Methods</title>
<sec id="s2a"><title>Study protocols</title>
<p>We applied a systematic approach to assess the methodological quality of our scoping review, including comprehensive literature searches, double screening, bias assessment, and evaluation of publication bias.</p>
</sec>
<sec id="s2b"><title>Data sources and search strategy</title>
<p>A systematic search was conducted in databases and search engines, including Google Scholar, PubMed/MEDLINE, Cochrane Library, Web of Science, CINAHL, and EMBASE, using keywords such as (&#x201C;ChatGPT&#x201D; OR &#x201C;GPT-3&#x201D; OR &#x201C;GPT-4&#x201D; OR &#x201C;Bard&#x201D; OR &#x201C;Gemini&#x201D;) AND (&#x201C;Medical&#x201D; OR &#x201C;Healthcare&#x201D; OR &#x201C;Clinical&#x201D; OR &#x201C;Health Inquiry&#x201D; OR &#x201C;Medical Inquiry&#x201D;) AND (&#x201C;comparison&#x201D; OR &#x201C;comparative&#x201D; OR &#x201C;analysis&#x201D; OR &#x201C;review&#x201D;) to identify studies published from January 1 to October 20, 2023. The search was restricted to studies published in English and related to human health subjects.</p>
</sec>
<sec id="s2c"><title>Eligibility criteria</title>
<p>To be included in this study, studies needed to meet the following criteria: focus on the application of ChatGPT and Gemini across different branch specialties, evaluate outcomes based on the accuracy and character count of ChatGPT and Gemini, and be verified against the most recent predatory journal list (<xref ref-type="bibr" rid="B11">11</xref>). Additionally, review articles and case reports were excluded.</p>
</sec>
<sec id="s2d"><title>Study selection process</title>
<p>The initial screening involved two researchers reviewing all titles and abstracts to check if they met the eligibility criteria. In case of disagreements, a third author was consulted to reach a final decision and resolve conflicts between the initial researchers.</p>
</sec>
<sec id="s2e"><title>Data items</title>
<p>The data collected from the studies included the first author&#x0027;s name, country of study, type of study design, publishing year, sample size, type of medical specialty, accuracy, and length (character) of ChatGPT and Gemini. Accuracy refers to the ability of ChatGPT and Gemini to provide contextually appropriate and correct responses to medical questions based on the standard guidelines specific to each medical specialty.</p>
</sec>
<sec id="s2f"><title>Data analysis and synthesis</title>
<p>Microsoft Excel (2019) was utilized to collect and organize the extracted data, while descriptive analysis was conducted using the Statistical Package for Social Sciences (SPSS) software (version 26). The data is displayed as frequencies, percentages, means, and standard deviations.</p>
</sec>
</sec>
<sec id="s3" sec-type="results"><title>Results</title>
<sec id="s3a"><title>Study selection</title>
<p>During the initial database search, a total of 64 articles were identified. Pre-screening procedures removed one duplicate, two articles in non-English languages, and eight with unretrievable data. Following a comprehensive review of titles and abstracts, 53 studies were assessed, excluding 22 for lack of relevance. The remaining 31 studies underwent full-text evaluation, excluding 19 for failing to meet the inclusion criteria. Among the 12 studies that proceeded to the eligibility assessment phase, one was excluded due to its publication in a predatory journal. Ultimately, 11 studies met the criteria for inclusion in the review (<xref ref-type="fig" rid="F1">Figure&#x00A0;1</xref>).</p>
<fig id="F1" position="float"><label>Figure 1</label>
<caption><p>Prisma flow diagram.</p></caption>
<graphic xmlns:xlink="http://www.w3.org/1999/xlink" xlink:href="fdgth-07-1482712-g001.tif"/>
</fig>
</sec>
<sec id="s3b"><title>Characteristics of the included studies</title>
<p>The summarized raw data from the included studies are all observational in <xref ref-type="table" rid="T1">Tables&#x00A0;1</xref>, <xref ref-type="table" rid="T2">2</xref>. India and the United States were the primary contributors, providing two studies. Additionally, Canada, Singapore, Turkey, Australia, Ecuador, and Iraq each contributed one study (<xref ref-type="table" rid="T1">Table&#x00A0;1</xref>).</p>
<table-wrap id="T1" position="float"><label>Table 1</label>
<caption><p>Baseline characteristics of the included studies.</p></caption>
<table frame="hsides" rules="groups">
<colgroup>
<col align="left"/>
<col align="left"/>
<col align="left"/>
<col align="center"/>
<col align="left"/>
<col align="center"/>
<col align="left"/>
</colgroup>
<thead>
<tr>
<th valign="top" align="left">No</th>
<th valign="top" align="left">Author</th>
<th valign="top" align="center">Type of study</th>
<th valign="top" align="center">Publishing years</th>
<th valign="top" align="center">Country</th>
<th valign="top" align="center">Sample size</th>
<th valign="top" align="center">Specialty</th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="left" rowspan="12">1</td>
<td valign="top" align="left" rowspan="12">Patil (<xref ref-type="bibr" rid="B12">12</xref>)</td>
<td valign="top" align="left" rowspan="12">Cross-sectional</td>
<td valign="top" align="center" rowspan="12">2023</td>
<td valign="top" align="left" rowspan="12">Canada</td>
<td valign="top" align="center">29 (9.12&#x0025;)</td>
<td valign="top" align="left">Neuroradiology</td>
</tr>
<tr>
<td valign="top" align="center">19 (5.97&#x0025;)</td>
<td valign="top" align="left">Mammography</td>
</tr>
<tr>
<td valign="top" align="center">89 (27.99&#x0025;)</td>
<td valign="top" align="left">General &#x0026; physics</td>
</tr>
<tr>
<td valign="top" align="center">30 (9.43&#x0025;)</td>
<td valign="top" align="left">Nuclear medicine</td>
</tr>
<tr>
<td valign="top" align="center">16 (5.03&#x0025;)</td>
<td valign="top" align="left">Pediatric Radiology</td>
</tr>
<tr>
<td valign="top" align="center">26 (8.18&#x0025;)</td>
<td valign="top" align="left">Interventional radiology</td>
</tr>
<tr>
<td valign="top" align="center">29 (9.12&#x0025;)</td>
<td valign="top" align="left">Gastrointestinal radiology</td>
</tr>
<tr>
<td valign="top" align="center">11 (3.46&#x0025;)</td>
<td valign="top" align="left">Genitourinary radiology</td>
</tr>
<tr>
<td valign="top" align="center">16 (5.03&#x0025;)</td>
<td valign="top" align="left">Cardiac radiology</td>
</tr>
<tr>
<td valign="top" align="center">6 (1.89&#x0025;)</td>
<td valign="top" align="left">Chest radiology</td>
</tr>
<tr>
<td valign="top" align="center">25 (7.86&#x0025;)</td>
<td valign="top" align="left">Musculoskeletal radiology</td>
</tr>
<tr>
<td valign="top" align="center">22 (6.92&#x0025;)</td>
<td valign="top" align="left">Ultrasound</td>
</tr>
<tr>
<td valign="top" align="left">2</td>
<td valign="top" align="left">Kumari (<xref ref-type="bibr" rid="B13">13</xref>)</td>
<td valign="top" align="left">Cross-sectional</td>
<td valign="top" align="center">2023</td>
<td valign="top" align="left">India</td>
<td valign="top" align="center">50</td>
<td valign="top" align="left">Hematology</td>
</tr>
<tr>
<td valign="top" align="left">3</td>
<td valign="top" align="left">Dhanvijay (<xref ref-type="bibr" rid="B14">14</xref>)</td>
<td valign="top" align="left">Cross-sectional</td>
<td valign="top" align="center">2023</td>
<td valign="top" align="left">India</td>
<td valign="top" align="center">77</td>
<td valign="top" align="left">Physiology</td>
</tr>
<tr>
<td valign="top" align="left">4</td>
<td valign="top" align="left">Muhialdeen (<xref ref-type="bibr" rid="B15">15</xref>)</td>
<td valign="top" align="left">Cross-sectional</td>
<td valign="top" align="center">2023</td>
<td valign="top" align="left">Iraq</td>
<td valign="top" align="center">20</td>
<td valign="top" align="left">Clinical Diagnosis</td>
</tr>
<tr>
<td valign="top" align="left">5</td>
<td valign="top" align="left">Koga (<xref ref-type="bibr" rid="B16">16</xref>)</td>
<td valign="top" align="left">Cross-sectional</td>
<td valign="top" align="center">2023</td>
<td valign="top" align="left">USA</td>
<td valign="top" align="center">25</td>
<td valign="top" align="left">neurodegenerative disorders</td>
</tr>
<tr>
<td valign="top" align="left">6</td>
<td valign="top" align="left">Zhi Wei Lim (<xref ref-type="bibr" rid="B17">17</xref>)</td>
<td valign="top" align="left">Cross-sectional</td>
<td valign="top" align="center">2023</td>
<td valign="top" align="left">Singapore</td>
<td valign="top" align="center">31</td>
<td valign="top" align="left">myopia care</td>
</tr>
<tr>
<td valign="top" align="left">7</td>
<td valign="top" align="left">Ilgaz (<xref ref-type="bibr" rid="B18">18</xref>)</td>
<td valign="top" align="left">Cross-sectional</td>
<td valign="top" align="center">2023</td>
<td valign="top" align="left">Turkey</td>
<td valign="top" align="center">131</td>
<td valign="top" align="left">Anatomy</td>
</tr>
<tr>
<td valign="top" align="left">8</td>
<td valign="top" align="left">Seth (<xref ref-type="bibr" rid="B19">19</xref>)</td>
<td valign="top" align="left">Cross-sectional</td>
<td valign="top" align="center">2023</td>
<td valign="top" align="left">Australia</td>
<td valign="top" align="center">6</td>
<td valign="top" align="left">Rhinoplasty</td>
</tr>
<tr>
<td valign="top" align="left" rowspan="2">9</td>
<td valign="top" align="left" rowspan="2">Salazar (<xref ref-type="bibr" rid="B20">20</xref>)</td>
<td valign="top" align="left" rowspan="2">Cross-sectional</td>
<td valign="top" align="center" rowspan="2">2023</td>
<td valign="top" align="left" rowspan="2">Ecuador</td>
<td valign="top" align="center">75</td>
<td valign="top" align="left">Emergency</td>
</tr>
<tr>
<td valign="top" align="center">101</td>
<td valign="top" align="left">Non-emergency</td>
</tr>
<tr>
<td valign="top" align="left" rowspan="3">10</td>
<td valign="top" align="left" rowspan="3">Qarajeh (<xref ref-type="bibr" rid="B21">21</xref>)</td>
<td valign="top" align="left" rowspan="3">Cross-sectional</td>
<td valign="top" align="center" rowspan="3">2023</td>
<td valign="top" align="left" rowspan="3">USA</td>
<td valign="top" align="center">81 (33.75&#x0025;)</td>
<td valign="top" align="left">Renal Diet High potassium</td>
</tr>
<tr>
<td valign="top" align="center">68 (28.33&#x0025;)</td>
<td valign="top" align="left">Renal Diet Low potassium</td>
</tr>
<tr>
<td valign="top" align="center">91 (37.91&#x0025;)</td>
<td valign="top" align="left">Renal Diet High phosphorus</td>
</tr>
<tr>
<td valign="top" align="left">11</td>
<td valign="top" align="left">Toyama (<xref ref-type="bibr" rid="B22">22</xref>)</td>
<td valign="top" align="left">Cross-sectional</td>
<td valign="top" align="center">2023</td>
<td valign="top" align="left">Japan</td>
<td valign="top" align="center">103</td>
<td valign="top" align="left">Radiology</td>
</tr>
</tbody>
</table>
</table-wrap>
<table-wrap id="T2" position="float"><label>Table 2</label>
<caption><p>Comparison between ChatGPT and Bard.</p></caption>
<table frame="hsides" rules="groups">
<colgroup>
<col align="left"/>
<col align="left"/>
<col align="left"/>
<col align="center"/>
<col align="center"/>
<col align="center"/>
<col align="center"/>
</colgroup>
<thead>
<tr>
<th valign="top" align="left">No</th>
<th valign="top" align="center">Author</th>
<th valign="top" align="center">Specialty</th>
<th valign="top" align="center">ChatGPT accurate</th>
<th valign="top" align="center">Gemini accurate</th>
<th valign="top" align="center">ChatGPT length (character)</th>
<th valign="top" align="center">Gemini length (character</th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="left" rowspan="12">1</td>
<td valign="top" align="left" rowspan="12">Patil (<xref ref-type="bibr" rid="B12">12</xref>)<xref ref-type="table-fn" rid="table-fn1"><sup>a</sup></xref></td>
<td valign="top" align="left">Neuroradiology</td>
<td valign="top" align="center">100.00&#x0025;</td>
<td valign="top" align="center">86.21&#x0025;</td>
<td valign="top" align="center">840.90 (&#x00B1;426.35)</td>
<td valign="top" align="center">1,443.52 (&#x00B1;415.88)</td>
</tr>
<tr>
<td valign="top" align="left">Mammography</td>
<td valign="top" align="center">84.21&#x0025;</td>
<td valign="top" align="center">68.42&#x0025;</td>
<td valign="top" align="center">787.63 (&#x00B1;447.38)</td>
<td valign="top" align="center">1,454.95 (&#x00B1;442.34)</td>
</tr>
<tr>
<td valign="top" align="left">General &#x0026; physics</td>
<td valign="top" align="center">85.39&#x0025;</td>
<td valign="top" align="center">68.54&#x0025;</td>
<td valign="top" align="center">1,022.38 (&#x00B1;453.50)</td>
<td valign="top" align="center">1,490.69 (&#x00B1;406.58)</td>
</tr>
<tr>
<td valign="top" align="left">Nuclear medicine</td>
<td valign="top" align="center">80.00&#x0025;</td>
<td valign="top" align="center">56.67&#x0025;</td>
<td valign="top" align="center">947.30 (&#x00B1;486.57)</td>
<td valign="top" align="center">1,321.57 (&#x00B1;374.86)</td>
</tr>
<tr>
<td valign="top" align="left">Pediatric Radiology</td>
<td valign="top" align="center">93.75&#x0025;</td>
<td valign="top" align="center">68.75&#x0025;</td>
<td valign="top" align="center">764.63 (&#x00B1;330.04)</td>
<td valign="top" align="center">1,368.88 (&#x00B1;547.91)</td>
</tr>
<tr>
<td valign="top" align="left">Interventional radiology</td>
<td valign="top" align="center">88.46&#x0025;</td>
<td valign="top" align="center">80.77&#x0025;</td>
<td valign="top" align="center">952.31 (&#x00B1;510.00)</td>
<td valign="top" align="center">1,538.31 (&#x00B1;446.90)</td>
</tr>
<tr>
<td valign="top" align="left">Gastrointestinal radiology</td>
<td valign="top" align="center">89.66&#x0025;</td>
<td valign="top" align="center">79.31&#x0025;</td>
<td valign="top" align="center">901.93 (&#x00B1;423.01)</td>
<td valign="top" align="center">1,427.66 (&#x00B1;322.21)</td>
</tr>
<tr>
<td valign="top" align="left">Genitourinary radiology</td>
<td valign="top" align="center">72.73&#x0025;</td>
<td valign="top" align="center">63.64&#x0025;</td>
<td valign="top" align="center">1,048.82 (&#x00B1;338.28)</td>
<td valign="top" align="center">1,373.09 (&#x00B1;342.2)</td>
</tr>
<tr>
<td valign="top" align="left">Cardiac radiology</td>
<td valign="top" align="center">75.00&#x0025;</td>
<td valign="top" align="center">68.75&#x0025;</td>
<td valign="top" align="center">915.50 (&#x00B1;3.48)</td>
<td valign="top" align="center">1,537.94 (&#x00B1;692.44)</td>
</tr>
<tr>
<td valign="top" align="left">Chest radiology</td>
<td valign="top" align="center">100.00&#x0025;</td>
<td valign="top" align="center">83.33&#x0025;</td>
<td valign="top" align="center">816.33 (&#x00B1;303.77)</td>
<td valign="top" align="center">1,492.33 (&#x00B1;295.28)</td>
</tr>
<tr>
<td valign="top" align="left">Musculoskeletal radiology</td>
<td valign="top" align="center">80.00&#x0025;</td>
<td valign="top" align="center">64.00&#x0025;</td>
<td valign="top" align="center">945.48 (&#x00B1;394.93)</td>
<td valign="top" align="center">1,326.40 (&#x00B1;316.33)</td>
</tr>
<tr>
<td valign="top" align="left">Ultrasound</td>
<td valign="top" align="center">100.00&#x0025;</td>
<td valign="top" align="center">63.64&#x0025;</td>
<td valign="top" align="center">944.91 (&#x00B1;518.11)</td>
<td valign="top" align="center">1,371.95 (&#x00B1;352.66)</td>
</tr>
<tr>
<td valign="top" align="left">2</td>
<td valign="top" align="left">Kumari (<xref ref-type="bibr" rid="B13">13</xref>)<xref ref-type="table-fn" rid="table-fn2"><sup>b</sup></xref></td>
<td valign="top" align="left">Hematology</td>
<td valign="top" align="center">3.15/5 (63&#x0025;)<sup>A</sup></td>
<td valign="top" align="center">2.23/5 (44&#x0025;)<sup>A</sup></td>
<td valign="top" align="center">-</td>
<td valign="top" align="center">-</td>
</tr>
<tr>
<td valign="top" align="left">3</td>
<td valign="top" align="left">Dhanvijy (<xref ref-type="bibr" rid="B14">14</xref>)<xref ref-type="table-fn" rid="table-fn2"><sup>b</sup></xref></td>
<td valign="top" align="left">Physiology</td>
<td valign="top" align="center">3.19/4 (79&#x0025;)<sup>A</sup></td>
<td valign="top" align="center">2.15/4 (53&#x0025;)<sup>A</sup></td>
<td valign="top" align="center">-</td>
<td valign="top" align="center">-</td>
</tr>
<tr>
<td valign="top" align="left">4</td>
<td valign="top" align="left">Muhialdeen (<xref ref-type="bibr" rid="B15">15</xref>)<xref ref-type="table-fn" rid="table-fn2"><sup>b</sup></xref></td>
<td valign="top" align="left">Clinical Diagnosis</td>
<td valign="top" align="center">90&#x0025;</td>
<td valign="top" align="center">80&#x0025;</td>
<td valign="top" align="center">-</td>
<td valign="top" align="center">-</td>
</tr>
<tr>
<td valign="top" align="left">5</td>
<td valign="top" align="left">Koga (<xref ref-type="bibr" rid="B16">16</xref>)<xref ref-type="table-fn" rid="table-fn1"><sup>a</sup></xref></td>
<td valign="top" align="left">neurodegenerative disorders</td>
<td valign="top" align="center">84&#x0025;</td>
<td valign="top" align="center">76&#x0025;</td>
<td valign="top" align="center">-</td>
<td valign="top" align="center">-</td>
</tr>
<tr>
<td valign="top" align="left">6</td>
<td valign="top" align="left">Zhi Wei Lim (<xref ref-type="bibr" rid="B17">17</xref>)<xref ref-type="table-fn" rid="table-fn1"><sup>a</sup></xref></td>
<td valign="top" align="left">myopia care</td>
<td valign="top" align="center">80.6&#x0025;</td>
<td valign="top" align="center">54.8&#x0025;</td>
<td valign="top" align="center">1,221.13 (&#x00B1;323.32)</td>
<td valign="top" align="center">1,275.87 (&#x00B1;393.25)</td>
</tr>
<tr>
<td valign="top" align="left">7</td>
<td valign="top" align="left">Ilgaz (<xref ref-type="bibr" rid="B18">18</xref>)<xref ref-type="table-fn" rid="table-fn2"><sup>b</sup></xref></td>
<td valign="top" align="left">Anatomy</td>
<td valign="top" align="center">44.27&#x0025;</td>
<td valign="top" align="center">41.98&#x0025;</td>
<td valign="top" align="center">-</td>
<td valign="top" align="center">-</td>
</tr>
<tr>
<td valign="top" align="left">8</td>
<td valign="top" align="left">Seth (<xref ref-type="bibr" rid="B19">19</xref>)<xref ref-type="table-fn" rid="table-fn2"><sup>b</sup></xref></td>
<td valign="top" align="left">Rhinoplasty</td>
<td valign="top" align="center">4/5 (80&#x0025;)<sup>A</sup></td>
<td valign="top" align="center">4/5 (80&#x0025;)<sup>A</sup></td>
<td valign="top" align="center">-</td>
<td valign="top" align="center">-</td>
</tr>
<tr>
<td valign="top" align="left" rowspan="2">9</td>
<td valign="top" align="left" rowspan="2">Salazar (<xref ref-type="bibr" rid="B20">20</xref>)<xref ref-type="table-fn" rid="table-fn2"><sup>b</sup></xref></td>
<td valign="top" align="left">Emergency</td>
<td valign="top" align="center">77&#x0025;</td>
<td valign="top" align="center">87&#x0025;</td>
<td valign="top" align="center">-</td>
<td valign="top" align="center">-</td>
</tr>
<tr>
<td valign="top" align="left">Non-emergency</td>
<td valign="top" align="center">36&#x0025;</td>
<td valign="top" align="center">33&#x0025;</td>
<td valign="top" align="center">-</td>
<td valign="top" align="center">-</td>
</tr>
<tr>
<td valign="top" align="left" rowspan="3">10</td>
<td valign="top" align="left" rowspan="3">Qarajeh (<xref ref-type="bibr" rid="B21">21</xref>)<xref ref-type="table-fn" rid="table-fn1"><sup>a</sup></xref></td>
<td valign="top" align="left">Renal Diet High potassium</td>
<td valign="top" align="center">99&#x0025;</td>
<td valign="top" align="center">79&#x0025;</td>
<td valign="top" align="center">-</td>
<td valign="top" align="center">-</td>
</tr>
<tr>
<td valign="top" align="left">Renal Diet Low potassium</td>
<td valign="top" align="center">60&#x0025;</td>
<td valign="top" align="center">79&#x0025;</td>
<td valign="top" align="center">-</td>
<td valign="top" align="center">-</td>
</tr>
<tr>
<td valign="top" align="left">Renal Diet high phosphorus</td>
<td valign="top" align="center">77&#x0025;</td>
<td valign="top" align="center">100&#x0025;</td>
<td valign="top" align="center">-</td>
<td valign="top" align="center">-</td>
</tr>
<tr>
<td valign="top" align="left">11</td>
<td valign="top" align="left">Toyama (<xref ref-type="bibr" rid="B22">22</xref>)<xref ref-type="table-fn" rid="table-fn1"><sup>a</sup></xref></td>
<td valign="top" align="left">Radiology</td>
<td valign="top" align="center">65&#x0025;</td>
<td valign="top" align="center">39&#x0025;</td>
<td valign="top" align="center">-</td>
<td valign="top" align="center">-</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<fn id="table-fn1"><label><sup>a</sup></label>
<p>ChatGPT-4.</p></fn>
<fn id="table-fn2"><label><sup>b</sup></label>
<p>ChatGPT-3.5.</p></fn>
</table-wrap-foot>
</table-wrap>
</sec>
<sec id="s3c"><title>Main findings</title>
<p>The research included 1,177 samples from 11 different medical specialties (<xref ref-type="bibr" rid="B12">12</xref>&#x2013;<xref ref-type="bibr" rid="B22">22</xref>). In radiology, there were 421 samples, encompassing various fields such as neuroradiology (9.12&#x0025;), mammography (5.97&#x0025;), general and physics (27.99&#x0025;), nuclear medicine (9.43&#x0025;), pediatric radiology (5.03&#x0025;), interventional radiology (8.18&#x0025;), gastrointestinal radiology (9.12&#x0025;), genitourinary radiology (3.46&#x0025;), cardiac radiology (5.03&#x0025;), chest radiology (1.89&#x0025;), musculoskeletal radiology (7.86&#x0025;), and ultrasound (6.92&#x0025;). The renal sample size was 240, divided into the renal diet with high potassium (33.75&#x0025;), the renal diet with low potassium (28.33&#x0025;), and the renal diet with high phosphorus (37.91&#x0025;). Emergency and non-emergency cases had sample sizes of 176, while the smallest samples were in clinical diagnosis and neurodegenerative disorders, with sizes of 20 and 25, respectively (<xref ref-type="table" rid="T1">Table&#x00A0;1</xref>).</p>
<p>The comparison between ChatGPT and Gemini across various specialties reveals accuracy and response length differences. ChatGPT generally may demonstrate higher accuracy than Gemini, especially in radiology specialties. The average accuracy of ChatGPT was 87.43&#x0025;, higher than Gemini 71&#x0025;. Additionally, the average response length of ChatGPT was 907 characters, shorter than Gemini&#x0027;s 1,428 characters. This indicates that ChatGPT&#x0027;s accuracy relative to response length may be more reliable and accurate than Gemini. Accuracy in the hematology specialty, ChatGPT, was 63&#x0025;, compared to Gemini&#x0027;s 44&#x0025;. Similar trends are observed in physiology, clinical diagnosis, neurodegenerative disorders, anatomy, renal diet, high potassium, and radiology. Conversely, in myopia care, the response lengths of ChatGPT and Gemini were nearly the same (1,221.13 vs. 1,275.87), with ChatGPT achieving a higher accuracy of 80.6&#x0025; compared to Gemini&#x0027;s 54.8&#x0025;. In rhinoplasty, both ChatGPT and Gemini demonstrate the same accuracy. In contrast, Gemini may be more accurate than ChatGPT in emergency scenarios, a renal diet with low potassium and a renal diet with high phosphorus (87&#x0025; vs. 77&#x0025;, 79&#x0025; vs. 60&#x0025;, and 100&#x0025; vs. 77&#x0025;, respectively) (<xref ref-type="table" rid="T2">Table&#x00A0;2</xref>).</p>
<p>The statistical analysis compared the accuracy and response length of ChatGPT and Gemini. The results indicate that ChatGPT has a higher accuracy (72.06&#x0025;) than Gemini (63.38&#x0025;), with a mean difference of 8.68, a confidence interval of 7.77&#x2013;9.58, and a statistically significant <italic>p</italic>-value of &#x003C;.001. In terms of response length, ChatGPT produces shorter responses (960.84 words) compared to Gemini (1,423.15 words), with a mean difference of 462.31 and a similarly significant <italic>p</italic>-value of &#x003C;.001. This statistical comparison emphasizes that ChatGPT may be more accurate and generates shorter responses than Gemini (<xref ref-type="table" rid="T3">Table&#x00A0;3</xref>).</p>
<table-wrap id="T3" position="float"><label>Table 3</label>
<caption><p>Statistical analysis of ChatGPT and Gemini.</p></caption>
<table frame="hsides" rules="groups">
<colgroup>
<col align="left"/>
<col align="center"/>
<col align="center"/>
<col align="center"/>
<col align="center"/>
<col align="center"/>
<col align="center"/>
</colgroup>
<thead>
<tr>
<th valign="top" align="left" rowspan="2"/>
<th valign="top" align="center" rowspan="2">Mean</th>
<th valign="top" align="center" rowspan="2">Mean difference</th>
<th valign="top" align="center" rowspan="2">Std. deviation</th>
<th valign="top" align="center" colspan="2">Confidence interval</th>
<th valign="top" align="center" rowspan="2"><italic>p</italic>-value</th>
</tr>
<tr>
<th valign="top" align="center">Lower</th>
<th valign="top" align="center">Upper</th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="left">ChatGPT Accuracy</td>
<td valign="top" align="center">72.06</td>
<td valign="top" align="center" rowspan="2">8.68</td>
<td valign="top" align="center" rowspan="2">15.82</td>
<td valign="top" align="center" rowspan="2">7.77</td>
<td valign="top" align="center" rowspan="2">9.58</td>
<td valign="top" align="center" rowspan="2">&#x003C;.001</td>
</tr>
<tr>
<td valign="top" align="left">Gemini Accuracy</td>
<td valign="top" align="center">63.38</td>
</tr>
<tr>
<td valign="top" align="left">ChatGPT Length</td>
<td valign="top" align="center">960.84</td>
<td valign="top" align="center" rowspan="2">462.31</td>
<td valign="top" align="center" rowspan="2">158.10</td>
<td valign="top" align="center" rowspan="2">445.64</td>
<td valign="top" align="center" rowspan="2">478.98</td>
<td valign="top" align="center" rowspan="2">&#x003C;.001</td>
</tr>
<tr>
<td valign="top" align="left">Gemini Length</td>
<td valign="top" align="center">1,423.15</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
</sec>
<sec id="s4" sec-type="discussion"><title>Discussion</title>
<p>Implementing large language models (LLMs) in medical education shows significant potential for transforming traditional teaching methods. Models like ChatGPT and Gemini process extensive medical literature, providing valuable, contextually relevant information for educators and students (<xref ref-type="bibr" rid="B23">23</xref>, <xref ref-type="bibr" rid="B24">24</xref>). LLMs create interactive, dynamic learning by giving students access to current medical data, clarifying complex concepts, and enhancing problem-solving. They also improve knowledge retrieval and support evidence-based decision-making. Incorporating LLMs encourages self-directed learning, critical thinking, and ongoing professional growth. However, recognizing their limitations and biases is essential for responsible, ethical use, complemented by practical training and clinical mentorship (<xref ref-type="bibr" rid="B25">25</xref>, <xref ref-type="bibr" rid="B26">26</xref>).</p>
<p>AI models have demonstrated significant potential in assisting medical professionals by enhancing efficiency in problem-solving, diagnosis, and data interpretation. For instance, ChatGPT has consistently outperformed models like Bard and Bing in accuracy when addressing medical vignettes (<xref ref-type="bibr" rid="B13">13</xref>, <xref ref-type="bibr" rid="B14">14</xref>). This underscores AI&#x0027;s pivotal role in supporting clinical decision-making, particularly in complex fields such as hematology (<xref ref-type="bibr" rid="B13">13</xref>). However, despite these encouraging results, AI models still face limitations, including inconsistencies in performance across various medical specialties, which necessitate further refinement before full integration into clinical practice (<xref ref-type="bibr" rid="B21">21</xref>).</p>
<p>The comparative analysis of ChatGPT and Gemini across various medical specialties reveals distinct patterns in their accuracy and response length performance. ChatGPT consistently demonstrates higher accuracy rates compared to Gemini in most specialties. This trend is evident in specialties such as neuroradiology (100&#x0025; vs. 86.21&#x0025;), hematology (63&#x0025; vs. 44&#x0025;), physiology (79&#x0025; vs. 53&#x0025;), clinical diagnosis (90&#x0025; vs. 80&#x0025;) and neurodegenerative disorders (84&#x0025; vs. 76&#x0025;). ChatGPT&#x0027;s superior accuracy indicates its potential as a reliable tool for medical inquiries, providing precise and dependable information across various medical fields (<xref ref-type="bibr" rid="B12">12</xref>&#x2013;<xref ref-type="bibr" rid="B16">16</xref>).</p>
<p>Despite Gemini&#x0027;s lower accuracy rates, it consistently delivers longer responses than ChatGPT. In neuroradiology, Gemini&#x0027;s responses averaged 1,443.52 characters compared to ChatGPT&#x0027;s 840.90 characters. This pattern is repeated across other specialties, such as mammography (1,454.95 vs. 787.63) and general &#x0026; physics (1,490.69 vs. 1,022.38) (<xref ref-type="bibr" rid="B12">12</xref>).</p>
<p>The longer response length of Gemini suggests that it may offer more detailed and comprehensive information, which could be beneficial in scenarios where a more exhaustive explanation is needed. While comparing ChatGPT and Gemini for accuracy and response length in chest radiology and ultrasound, ChatGPT consistently outperforms Gemini in accuracy. ChatGPT achieves 100&#x0025; accuracy for chest radiology compared to Gemini&#x0027;s 83.33&#x0025;, with a shorter average response length of 816.33 vs. Gemini&#x0027;s 1,492.33 characters. ChatGPT also has a perfect accuracy rate of 100&#x0025; in ultrasound, while Gemini&#x0027;s accuracy drops to 63.64&#x0025;. Similarly, ChatGPT&#x0027;s responses are more concise, averaging 944.91 characters compared to Gemini&#x0027;s 1,371.95 characters (<xref ref-type="bibr" rid="B12">12</xref>).</p>
<p>ChatGPT and Gemini in myopia care respond to similar lengths (1,221.13 and 1,275.87 characters, respectively). However, there is a difference in accuracy: ChatGPT achieves an accuracy of 80.6&#x0025;, whereas Gemini achieves 54.8&#x0025;. This disparity suggests that while both models may offer comparable responses in terms of content, ChatGPT tends to provide more reliable and accurate information in this specialized medical context (<xref ref-type="bibr" rid="B17">17</xref>).</p>
<p>ChatGPT and Gemini exhibit nearly identical accuracy in anatomy and rhinoplasty. ChatGPT achieves 44.27&#x0025; accuracy in anatomy, slightly higher than Gemini&#x0027;s 41.98&#x0025;. For rhinoplasty, both models perform equally well, each with an accuracy rate of 80&#x0025;. This comparison demonstrates that ChatGPT and Gemini perform similarly in these medical specialties (<xref ref-type="bibr" rid="B18">18</xref>, <xref ref-type="bibr" rid="B19">19</xref>).</p>
<p>Exceptions to this trend were observed in emergency scenarios, where Gemini achieved higher accuracy (87&#x0025;) compared to ChatGPT (77&#x0025;) (<xref ref-type="bibr" rid="B20">20</xref>). This highlights that Gemini may have strengths in specific contexts, such as emergencies where detailed information could be critical. However, both models showed lower accuracy rates in non-emergency scenarios, with ChatGPT slightly outperforming Gemini (36&#x0025; vs. 33&#x0025;) (<xref ref-type="bibr" rid="B20">20</xref>).</p>
<p>The performance of ChatGPT and Gemini in providing dietary advice for renal conditions also varied. ChatGPT excelled in high potassium contexts (99&#x0025; vs. 79&#x0025;) but was less accurate in low potassium and high phosphorus scenarios compared to Gemini (77&#x0025; vs. 100&#x0025;) (<xref ref-type="bibr" rid="B21">21</xref>). This variability suggests that each model may have specialized strengths in specific medical contexts, and their combined use could potentially enhance the quality of medical inquiry responses.</p>
<p>The comparative analysis of ChatGPT and Gemini (Bard) in medical inquiry highlights several limitations. ChatGPT may provide inaccurate medical information due to its limited understanding of complex contexts, and biases in training data can affect accuracy. Ethical concerns include the risk of outdated information and issues related to patient data privacy. Additionally, the evolving nature of large language models means that ChatGPT and Gemini are frequently updated, potentially rendering some findings obsolete as newer versions are released. The study&#x0027;s focus on specific models and predefined case vignettes may restrict its findings, as the scope of medical inquiries is limited to particular scenarios, which may not fully capture the broad range of medical topics these models could encounter. Moreover, potential biases in the responses of these language models were not fully explored, affecting the generalizability of the results. There may be limitations and potential bias in measuring accuracy, as each specialty uses different standard answers to compare with the responses of ChatGPT and Gemini across various studies. This variability makes it challenging to determine how accurately the models perform in each specialty. The findings indicate that ChatGPT generally offers more accurate and concise responses across various medical specialties, while Gemini provides more detailed but less accurate answers. The choice between these AI models should be guided by the specific needs of the medical inquiry&#x2014;whether precision or detail is prioritized. Future improvements should aim to integrate the strengths of both models, enhancing accuracy while maintaining the comprehensiveness of responses to support better clinical decision-making and patient care.</p>
</sec>
<sec id="s5" sec-type="conclusions"><title>Conclusion</title>
<p>This scoping review indicates that ChatGPT has shown promise in the included medical studies. It may demonstrate higher accuracy and a shorter response than Gemini. Therefore, further research is needed to maximize ChatGPT&#x0027;s accuracy compared to Gemini in the medical field.</p>
</sec>
</body>
<back>
<sec id="s6" sec-type="data-availability"><title>Data availability statement</title>
<p>The original contributions presented in the study are included in the article/Supplementary Material, further inquiries can be directed to the corresponding author.</p>
</sec>
<sec id="s7" sec-type="author-contributions"><title>Author contributions</title>
<p>FF: Investigation, Methodology, Validation, Visualization, Writing &#x2013; review &#x0026; editing. AbS: Conceptualization, Investigation, Methodology, Resources, Writing &#x2013; review &#x0026; editing. AmS: Conceptualization, Formal Analysis, Methodology, Resources, Validation, Writing &#x2013; review &#x0026; editing. SKA: Formal Analysis, Resources, Software, Supervision, Writing &#x2013; review &#x0026; editing. AG: Conceptualization, Formal Analysis, Methodology, Resources, Validation, Writing &#x2013; review &#x0026; editing. RB: Data curation, Formal Analysis, Investigation, Methodology, Validation, Visualization, Writing &#x2013; review &#x0026; editing. BA: Conceptualization, Data curation, Visualization, Writing &#x2013; original draft, Writing &#x2013; review &#x0026; editing. SO: Conceptualization, Methodology, Project administration, Visualization, Writing &#x2013; review &#x0026; editing. SMA: Conceptualization, Formal Analysis, Methodology, Visualization, Writing &#x2013; review &#x0026; editing. SH: Formal Analysis, Methodology, Resources, Validation, Writing &#x2013; review &#x0026; editing. YM: Data curation, Formal Analysis, Software, Supervision, Visualization, Writing &#x2013; review &#x0026; editing. FK: Conceptualization, Data curation, Validation, Visualization, Writing &#x2013; review &#x0026; editing.</p>
</sec>
<sec id="s8" sec-type="funding-information"><title>Funding</title>
<p>The author(s) declare that no financial support was received for the research, authorship, and/or publication of this article.</p>
</sec>
<sec id="s9" sec-type="COI-statement"><title>Conflict of interest</title>
<p>The authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec id="s11" sec-type="disclaimer"><title>Publisher&#x0027;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<ref-list><title>References</title>
<ref id="B1"><label>1.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Hiwa</surname><given-names>DS</given-names></name><name><surname>Abdalla</surname><given-names>SS</given-names></name><name><surname>Muhialdeen</surname><given-names>AS</given-names></name><name><surname>Hamasalih</surname><given-names>HM</given-names></name><name><surname>Karim</surname><given-names>SO</given-names></name></person-group>. <article-title>Assessment of nursing skill and knowledge of ChatGPT, Gemini, Microsoft Copilot, and Llama: a comparative study</article-title>. <source>Barw Med J</source>. (<year>2024</year>) <volume>2</volume>(<issue>2</issue>):<fpage>3</fpage>&#x2013;<lpage>6</lpage>. <pub-id pub-id-type="doi">10.58742/bmj.v2i2.87</pub-id></citation></ref>
<ref id="B2"><label>2.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Abbas</surname><given-names>YN</given-names></name><name><surname>Mahmood</surname><given-names>YM</given-names></name><name><surname>Hassan</surname><given-names>HA</given-names></name><name><surname>Hamad</surname><given-names>DQ</given-names></name><name><surname>Hasan</surname><given-names>SJ</given-names></name><name><surname>Omer</surname><given-names>DA</given-names></name><etal/></person-group> <article-title>Role of ChatGPT and google bard in the diagnosis of psychiatric disorders: a cross sectional study</article-title>. <source>Barw Med J</source>. (<year>2023</year>) <volume>1</volume>(<issue>4</issue>):<fpage>14</fpage>&#x2013;<lpage>9</lpage>. <pub-id pub-id-type="doi">10.58742/4vd6h741</pub-id></citation></ref>
<ref id="B3"><label>3.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Xu</surname><given-names>L</given-names></name><name><surname>Sanders</surname><given-names>L</given-names></name><name><surname>Li</surname><given-names>K</given-names></name><name><surname>Chow</surname><given-names>JC</given-names></name></person-group>. <article-title>Chatbot for health care and oncology applications using artificial intelligence and machine learning: systematic review</article-title>. <source>JMIR Cancer</source>. (<year>2021</year>) <volume>7</volume>(<issue>4</issue>):<fpage>e27850</fpage>. <pub-id pub-id-type="doi">10.2196/27850</pub-id><pub-id pub-id-type="pmid">34847056</pub-id></citation></ref>
<ref id="B4"><label>4.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Fuchs</surname><given-names>A</given-names></name><name><surname>Trachsel</surname><given-names>T</given-names></name><name><surname>Weiger</surname><given-names>R</given-names></name><name><surname>Eggmann</surname><given-names>F</given-names></name></person-group>. <article-title>ChatGPT&#x2019;s performance in dentistry and allergy immunology assessments: a comparative study</article-title>. <source>Swiss Dent J</source>. (<year>2024</year>) <volume>134</volume>(<issue>2</issue>):<fpage>1</fpage>&#x2013;<lpage>7</lpage>. <pub-id pub-id-type="doi">10.61872/sdj-2024-06-01</pub-id></citation></ref>
<ref id="B5"><label>5.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Masalkhi</surname><given-names>M</given-names></name><name><surname>Ong</surname><given-names>J</given-names></name><name><surname>Waisberg</surname><given-names>E</given-names></name><name><surname>Lee</surname><given-names>AG</given-names></name></person-group>. <article-title>Google DeepMind&#x2019;s Gemini AI versus ChatGPT: a comparative analysis in ophthalmology</article-title>. <source>Eye</source>. (<year>2024</year>) <volume>14</volume>:<fpage>1</fpage>&#x2013;<lpage>6</lpage>. <pub-id pub-id-type="doi">10.1038/s41433-024-02958-w</pub-id></citation></ref>
<ref id="B6"><label>6.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Salih</surname><given-names>AM</given-names></name><name><surname>Mohammed</surname><given-names>NA</given-names></name><name><surname>Mahmood</surname><given-names>YM</given-names></name><name><surname>Hasan</surname><given-names>SJ</given-names></name><name><surname>Namiq</surname><given-names>HS</given-names></name><name><surname>Ghafour</surname><given-names>AK</given-names></name><etal/></person-group> <article-title>ChatGPT insight and opinion regarding the controversies in neurogenic thoracic outlet syndrome; a case based-study</article-title>. <source>Barw Med J</source>. (<year>2023</year>) <volume>1</volume>(<issue>3</issue>):<fpage>2</fpage>&#x2013;<lpage>4</lpage>. <pub-id pub-id-type="doi">10.58742/bmj.v1i2.48</pub-id></citation></ref>
<ref id="B7"><label>7.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wei</surname><given-names>Q</given-names></name><name><surname>Yao</surname><given-names>Z</given-names></name><name><surname>Cui</surname><given-names>Y</given-names></name><name><surname>Wei</surname><given-names>B</given-names></name><name><surname>Jin</surname><given-names>Z</given-names></name><name><surname>Xu</surname><given-names>X</given-names></name></person-group>. <article-title>Evaluation of ChatGPT-generated medical responses: a systematic review and meta-analysis</article-title>. <source>J Biomed Inform</source>. (<year>2024</year>) <volume>8</volume>:<fpage>104620</fpage>. <pub-id pub-id-type="doi">10.1016/j.jbi.2024.104620</pub-id></citation></ref>
<ref id="B8"><label>8.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Farhud</surname><given-names>DD</given-names></name><name><surname>Zokaei</surname><given-names>S</given-names></name></person-group>. <article-title>Ethical issues of artificial intelligence in medicine and healthcare</article-title>. <source>Iran J Public Health</source>. (<year>2021</year>) <volume>50</volume>(<issue>11</issue>):<fpage>i</fpage>&#x2013;<lpage>v</lpage>. <pub-id pub-id-type="doi">10.18502/ijph.v50i11.7600</pub-id></citation></ref>
<ref id="B9"><label>9.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ray</surname><given-names>PP</given-names></name><name><surname>Majumder</surname><given-names>P</given-names></name></person-group>. <article-title>The potential of ChatGPT to transform healthcare and address ethical challenges in artificial intelligence-driven medicine</article-title>. <source>J Clin Neurol</source>. (<year>2023</year>) <volume>19</volume>(<issue>5</issue>):<fpage>509</fpage>. <pub-id pub-id-type="doi">10.3988/jcn.2023.0158</pub-id><pub-id pub-id-type="pmid">37635433</pub-id></citation></ref>
<ref id="B10"><label>10.</label><citation citation-type="confproc"><person-group person-group-type="author"><name><surname>Sharma</surname><given-names>D</given-names></name><name><surname>Kaushal</surname><given-names>S</given-names></name><name><surname>Kumar</surname><given-names>H</given-names></name><name><surname>Gainder</surname><given-names>S</given-names></name></person-group>. <article-title>Chatbots in healthcare: challenges, technologies and applications</article-title>. <conf-name>2022 4th International Conference on Artificial Intelligence and Speech Technology (AIST)</conf-name>. <publisher-name>IEEE</publisher-name> (<year>2022</year>). p. <fpage>1</fpage>&#x2013;<lpage>6</lpage></citation></ref>
<ref id="B11"><label>11.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Abdullah</surname><given-names>HO</given-names></name><name><surname>Abdalla</surname><given-names>BA</given-names></name><name><surname>Kakamad</surname><given-names>FH</given-names></name><name><surname>Ahmed</surname><given-names>JO</given-names></name><name><surname>Baba</surname><given-names>HO</given-names></name><name><surname>Hassan</surname><given-names>MN</given-names></name><etal/></person-group> <article-title>Predatory publishing lists: a review on the ongoing battle against fraudulent actions</article-title>. <source>Barw Med J</source>. (<year>2024</year>) <volume>2</volume>(<issue>2</issue>):<fpage>26</fpage>&#x2013;<lpage>30</lpage>. <pub-id pub-id-type="doi">10.58742/bmj.v2i2.91</pub-id></citation></ref>
<ref id="B12"><label>12.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Patil</surname><given-names>NS</given-names></name><name><surname>Huang</surname><given-names>RS</given-names></name><name><surname>van der Pol</surname><given-names>CB</given-names></name><name><surname>Larocque</surname><given-names>N</given-names></name></person-group>. <article-title>Comparative performance of ChatGPT and bard in a text-based radiology knowledge assessment</article-title>. <source>Can Assoc Radiol J</source>. (<year>2024</year>) <volume>75</volume>(<issue>2</issue>):<fpage>344</fpage>&#x2013;<lpage>50</lpage>. <pub-id pub-id-type="doi">10.1177/08465371231193716</pub-id><pub-id pub-id-type="pmid">37578849</pub-id></citation></ref>
<ref id="B13"><label>13.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Kumari</surname><given-names>A</given-names></name><name><surname>Kumari</surname><given-names>A</given-names></name><name><surname>Singh</surname><given-names>A</given-names></name><name><surname>Singh</surname><given-names>SK</given-names></name><name><surname>Juhi</surname><given-names>A</given-names></name><name><surname>Kumar</surname><given-names>A</given-names></name><etal/></person-group> <article-title>Large language models in hematology case solving: a comparative study of ChatGPT-3.5, Google Bard, and Microsoft Bing</article-title>. <source>Cureus</source>. (<year>2023</year>) <volume>21</volume>:<fpage>e43861</fpage>. <pub-id pub-id-type="doi">10.7759/cureus.43861</pub-id></citation></ref>
<ref id="B14"><label>14.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Dhanvijay</surname><given-names>AK</given-names></name><name><surname>Pinjar</surname><given-names>MJ</given-names></name><name><surname>Dhokane</surname><given-names>N</given-names></name><name><surname>Sorte</surname><given-names>SR</given-names></name><name><surname>Kumari</surname><given-names>A</given-names></name><name><surname>Mondal</surname><given-names>H</given-names></name></person-group>. <article-title>Performance of large language models (ChatGPT, Bing Search, and Google Bard) in solving case vignettes in physiology</article-title>. <source>Cureus</source>. (<year>2023</year>) <volume>15</volume>(<issue>8</issue>):<fpage>e42972</fpage>. <pub-id pub-id-type="doi">10.7759/cureus.42972</pub-id><pub-id pub-id-type="pmid">37671207</pub-id></citation></ref>
<ref id="B15"><label>15.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Muhialdeen</surname><given-names>AS</given-names></name><name><surname>Mohammed</surname><given-names>SA</given-names></name><name><surname>Ahmed</surname><given-names>NH</given-names></name><name><surname>Ahmed</surname><given-names>SF</given-names></name><name><surname>Hassan</surname><given-names>WN</given-names></name><name><surname>Asaad</surname><given-names>HR</given-names></name><etal/></person-group> <article-title>Artificial intelligence in medicine: a comparative study of ChatGPT and google bard in clinical diagnostics</article-title>. <source>Barw Med J</source>. (<year>2023</year>) <volume>6</volume>:<fpage>7</fpage>&#x2013;<lpage>13</lpage>. <pub-id pub-id-type="doi">10.58742/pry94q89</pub-id></citation></ref>
<ref id="B16"><label>16.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Koga</surname><given-names>S</given-names></name><name><surname>Martin</surname><given-names>NB</given-names></name><name><surname>Dickson</surname><given-names>DW</given-names></name></person-group>. <article-title>Evaluating the performance of large language models: ChatGPT and Google Bard in generating differential diagnoses in clinicopathological conferences of neurodegenerative disorders</article-title>. <source>Brain Pathol</source>. (<year>2024</year>) <volume>34</volume>(<issue>3</issue>):<fpage>e13207</fpage>. <pub-id pub-id-type="doi">10.1111/bpa.13207</pub-id><pub-id pub-id-type="pmid">37553205</pub-id></citation></ref>
<ref id="B17"><label>17.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Lim</surname><given-names>ZW</given-names></name><name><surname>Pushpanathan</surname><given-names>K</given-names></name><name><surname>Yew</surname><given-names>SM</given-names></name><name><surname>Lai</surname><given-names>Y</given-names></name><name><surname>Sun</surname><given-names>CH</given-names></name><name><surname>Lam</surname><given-names>JS</given-names></name><etal/></person-group> <article-title>Benchmarking large language models&#x2019; performances for myopia care: a comparative analysis of ChatGPT-3.5, ChatGPT-4.0, and google bard</article-title>. <source>EBioMedicine</source>. (<year>2023</year>) <volume>95</volume>:<fpage>104770</fpage>. <pub-id pub-id-type="doi">10.1016/j.ebiom.2023.104770</pub-id><pub-id pub-id-type="pmid">37625267</pub-id></citation></ref>
<ref id="B18"><label>18.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ilgaz</surname><given-names>HB</given-names></name><name><surname>&#x00C7;elik</surname><given-names>Z</given-names></name></person-group>. <article-title>The significance of artificial intelligence platforms in anatomy education: an experience with ChatGPT and Google Bard</article-title>. <source>Cureus</source>. (<year>2023</year>) <volume>15</volume>(<issue>9</issue>):<fpage>e45301</fpage>. <pub-id pub-id-type="doi">10.7759/cureus.45301</pub-id><pub-id pub-id-type="pmid">37846274</pub-id></citation></ref>
<ref id="B19"><label>19.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Seth</surname><given-names>I</given-names></name><name><surname>Lim</surname><given-names>B</given-names></name><name><surname>Xie</surname><given-names>Y</given-names></name><name><surname>Cevik</surname><given-names>J</given-names></name><name><surname>Rozen</surname><given-names>WM</given-names></name><name><surname>Ross</surname><given-names>RJ</given-names></name><etal/></person-group> <article-title>Comparing the efficacy of large language models ChatGPT, BARD, and Bing AI in providing information on rhinoplasty: an observational study</article-title>. <source>Aesthet Surg J Open Forum</source>. (<year>2023</year>) <volume>5</volume>:<fpage>ojad084</fpage>. <pub-id pub-id-type="doi">10.1093/asjof/ojad084</pub-id><comment>; US: Oxford University Press</comment>.<pub-id pub-id-type="pmid">37795257</pub-id></citation></ref>
<ref id="B20"><label>20.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Salazar</surname><given-names>GZ</given-names></name><name><surname>Z&#x00FA;&#x00F1;iga</surname><given-names>D</given-names></name><name><surname>Vindel</surname><given-names>CL</given-names></name><name><surname>Yoong</surname><given-names>AM</given-names></name><name><surname>Hincapie</surname><given-names>S</given-names></name><name><surname>Z&#x00FA;&#x00F1;iga</surname><given-names>AB</given-names></name><etal/></person-group> <article-title>Efficacy of AI chats to determine an emergency: a comparison between OpenAI&#x2019;s ChatGPT, Google Bard, and Microsoft Bing AI chat</article-title>. <source>Cureus</source>. (<year>2023</year>) <volume>15</volume>(<issue>9</issue>):<fpage>e45473</fpage>. <pub-id pub-id-type="doi">10.7759/cureus.45473</pub-id><pub-id pub-id-type="pmid">37727841</pub-id></citation></ref>
<ref id="B21"><label>21.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Qarajeh</surname><given-names>A</given-names></name><name><surname>Tangpanithandee</surname><given-names>S</given-names></name><name><surname>Thongprayoon</surname><given-names>C</given-names></name><name><surname>Suppadungsuk</surname><given-names>S</given-names></name><name><surname>Krisanapan</surname><given-names>P</given-names></name><name><surname>Aiumtrakul</surname><given-names>N</given-names></name><etal/></person-group> <article-title>AI-powered renal diet support: performance of ChatGPT, Bard AI, and Bing chat</article-title>. <source>Clin Pract</source>. (<year>2023</year>) <volume>13</volume>(<issue>5</issue>):<fpage>1160</fpage>&#x2013;<lpage>72</lpage>. <pub-id pub-id-type="doi">10.3390/clinpract13050104</pub-id><pub-id pub-id-type="pmid">37887080</pub-id></citation></ref>
<ref id="B22"><label>22.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Toyama</surname><given-names>Y</given-names></name><name><surname>Harigai</surname><given-names>A</given-names></name><name><surname>Abe</surname><given-names>M</given-names></name><name><surname>Nagano</surname><given-names>M</given-names></name><name><surname>Kawabata</surname><given-names>M</given-names></name><name><surname>Seki</surname><given-names>Y</given-names></name><etal/></person-group> <article-title>Performance evaluation of ChatGPT, GPT-4, and bard on the official board examination of the Japan radiology society</article-title>. <source>Jpn J Radiol</source>. (<year>2024</year>) <volume>42</volume>(<issue>2</issue>):<fpage>201</fpage>&#x2013;<lpage>7</lpage>. <pub-id pub-id-type="doi">10.1007/s11604-023-01491-2</pub-id><pub-id pub-id-type="pmid">37792149</pub-id></citation></ref>
<ref id="B23"><label>23.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Sinha</surname><given-names>RK</given-names></name><name><surname>Roy</surname><given-names>AD</given-names></name><name><surname>Kumar</surname><given-names>N</given-names></name><name><surname>Mondal</surname><given-names>H</given-names></name></person-group>. <article-title>Applicability of ChatGPT in assisting to solve higher order problems in pathology</article-title>. <source>Cureus</source>. (<year>2023</year>) <volume>15</volume>(<issue>2</issue>):<fpage>e35237</fpage>. <pub-id pub-id-type="doi">10.7759/cureus.35237</pub-id><pub-id pub-id-type="pmid">36968864</pub-id></citation></ref>
<ref id="B24"><label>24.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Das</surname><given-names>D</given-names></name><name><surname>Kumar</surname><given-names>N</given-names></name><name><surname>Longjam</surname><given-names>LA</given-names></name><name><surname>Sinha</surname><given-names>R</given-names></name><name><surname>Roy</surname><given-names>AD</given-names></name><name><surname>Mondal</surname><given-names>H</given-names></name><etal/></person-group> <article-title>Assessing the capability of ChatGPT in answering first-and second-order knowledge questions on microbiology as per competency-based medical education curriculum</article-title>. <source>Cureus</source>. (<year>2023</year>) <volume>15</volume>(<issue>3</issue>):<fpage>e36034</fpage>. <pub-id pub-id-type="doi">10.7759/cureus.36034</pub-id><pub-id pub-id-type="pmid">37056538</pub-id></citation></ref>
<ref id="B25"><label>25.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ghosh</surname><given-names>A</given-names></name><name><surname>Bir</surname><given-names>A</given-names></name></person-group>. <article-title>Evaluating ChatGPT&#x2019;s ability to solve higher-order questions on the competency-based medical education curriculum in medical biochemistry</article-title>. <source>Cureus</source>. (<year>2023</year>) <volume>15</volume>(<issue>4</issue>):<fpage>e37023</fpage>. <pub-id pub-id-type="doi">10.7759/cureus.37023</pub-id><pub-id pub-id-type="pmid">37143631</pub-id></citation></ref>
<ref id="B26"><label>26.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Gudis</surname><given-names>DA</given-names></name><name><surname>McCoul</surname><given-names>ED</given-names></name><name><surname>Marino</surname><given-names>MJ</given-names></name><name><surname>Patel</surname><given-names>ZM</given-names></name></person-group>. <article-title>Avoiding bias in artificial intelligence</article-title>. <source>Int Forum Allergy Rhinol</source>. (<year>2023</year>) <volume>13</volume>(<issue>3</issue>):<fpage>193</fpage>&#x2013;<lpage>5</lpage>. <pub-id pub-id-type="doi">10.1002/alr.23129</pub-id><pub-id pub-id-type="pmid">36573806</pub-id></citation></ref></ref-list>
</back>
</article>