<?xml version="1.0" encoding="utf-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance" article-type="research-article" dtd-version="2.3" xml:lang="EN">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Med.</journal-id>
<journal-title>Frontiers in Medicine</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Med.</abbrev-journal-title>
<issn pub-type="epub">2296-858X</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="doi">10.3389/fmed.2024.1489117</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Medicine</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Potential application of ChatGPT in <italic>Helicobacter pylori</italic> disease relevant queries</article-title>
</title-group>
<contrib-group>
<contrib contrib-type="author">
<name><surname>Gao</surname> <given-names>Zejun</given-names></name>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Ge</surname> <given-names>Jinlin</given-names></name>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Xu</surname> <given-names>Ruoshi</given-names></name>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name><surname>Chen</surname> <given-names>Xiaoyan</given-names></name>
<xref ref-type="corresp" rid="c001"><sup>&#x002A;</sup></xref>
<uri xlink:href="https://loop.frontiersin.org/people/2841592/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name><surname>Cai</surname> <given-names>Zhenzhai</given-names></name>
<xref ref-type="corresp" rid="c002"><sup>&#x002A;</sup></xref>
<uri xlink:href="https://loop.frontiersin.org/people/1512258/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
</contrib-group>
<aff><institution>Department of Gastroenterology, Second Affiliated Hospital and Yuying Children&#x2019;s Hospital of Wenzhou Medical University</institution>, <addr-line>Wenzhou</addr-line>, <country>China</country></aff>
<author-notes>
<fn fn-type="edited-by" id="fn0003">
<p>Edited by: Ponsiano Ocama, Makerere University, Uganda</p>
</fn>
<fn fn-type="edited-by" id="fn0004">
<p>Reviewed by: Raffaele Pellegrino, University of Campania Luigi Vanvitelli, Italy</p>
<p>Jonathan Soldera, University of Caxias do Sul, Brazil</p>
</fn>
<corresp id="c001">&#x002A;Correspondence: Xiaoyan Chen, <email>cxy_dr@sina.com</email></corresp>
<corresp id="c002">Zhenzhai Cai, <email>caizhenzhai@wmu.edu.cn</email></corresp>
</author-notes>
<pub-date pub-type="epub">
<day>10</day>
<month>10</month>
<year>2024</year>
</pub-date>
<pub-date pub-type="collection">
<year>2024</year>
</pub-date>
<volume>11</volume>
<elocation-id>1489117</elocation-id>
<history>
<date date-type="received">
<day>31</day>
<month>08</month>
<year>2024</year>
</date>
<date date-type="accepted">
<day>27</day>
<month>09</month>
<year>2024</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#x00A9; 2024 Gao, Ge, Xu, Chen and Cai.</copyright-statement>
<copyright-year>2024</copyright-year>
<copyright-holder>Gao, Ge, Xu, Chen and Cai</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<abstract>
<sec id="sec1">
<title>Background</title>
<p>Advances in artificial intelligence are gradually transforming various fields, but its applicability among ordinary people is unknown. This study aims to explore the ability of a large language model to address <italic>Helicobacter pylori</italic> related questions.</p>
</sec>
<sec id="sec2">
<title>Methods</title>
<p>We created several prompts on the basis of guidelines and the clinical concerns of patients. The capacity of ChatGPT on <italic>Helicobacter pylori</italic> queries was evaluated by experts. Ordinary people assessed the applicability.</p>
</sec>
<sec id="sec3">
<title>Results</title>
<p>The responses to each prompt in ChatGPT-4 were good in terms of response length and repeatability. There was good agreement in each dimension (Fleiss&#x2019; kappa ranged from 0.302 to 0.690, <italic>p</italic>&#x2009;&#x003C;&#x2009;0.05). The accuracy, completeness, usefulness, comprehension and satisfaction scores of the experts were generally high. Rated usefulness and comprehension among ordinary people were significantly lower than expert, while medical students gave a relatively positive evaluation.</p>
</sec>
<sec id="sec4">
<title>Conclusion</title>
<p>ChatGPT-4 performs well in resolving <italic>Helicobacter pylori</italic> related questions. Large language models may become an excellent tool for medical students in the future, but still requires further research and validation.</p>
</sec>
</abstract>
<kwd-group>
<kwd>
<italic>Helicobacter pylori</italic>
</kwd>
<kwd>intrafamilial transmission</kwd>
<kwd>ChatGPT</kwd>
<kwd>large language model</kwd>
<kwd>artificial intelligence</kwd>
</kwd-group>
<counts>
<fig-count count="2"/>
<table-count count="4"/>
<equation-count count="0"/>
<ref-count count="35"/>
<page-count count="7"/>
<word-count count="4703"/>
</counts>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-at-acceptance</meta-name>
<meta-value>Gastroenterology</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec sec-type="intro" id="sec5">
<label>1</label>
<title>Introduction</title>
<p><italic>Helicobacter pylori</italic> (HP) is a gram-negative bacterium transmitted through the fecal&#x2013;oral route that infects the human gastric mucosa epithelium and affects 50% of the world&#x2019;s population, especially in developing countries, due to unhealthy dietary habits1. Long-term HP infection may lead to several gastrointestinal diseases, such as chronic inflammation, peptic ulcers, gastric cancer, and mucosa-associated lymphoid tissue lymphoma (<xref ref-type="bibr" rid="ref1">1</xref>, <xref ref-type="bibr" rid="ref2">2</xref>). The World Health Organization listed it as a class I carcinogen for gastric cancer in 1994 (<xref ref-type="bibr" rid="ref3">3</xref>). In addition to traditional test-and-treat and screen-and-treat strategies, family-based control and management has been proposed as a third approach, which is not affected by HP infection rates (<xref ref-type="bibr" rid="ref4 ref5 ref6">4&#x2013;6</xref>). However, owing to differences in education among societies, the popularity of HP-related knowledge still remains a major problem.</p>
<p>Recently, with the progress of technology and the rapid development of artificial intelligence (AI), enormous changes have taken place in different areas of the world, and the medical field is no exception. The application of AI in medicine is expanding in many fields, including intelligent screening, intelligent diagnosis, risk prediction and adjuvant therapy (<xref ref-type="bibr" rid="ref7 ref8 ref9">7&#x2013;9</xref>). At the same time, there has been an interest in ChatGPT in the gastroenterology community. Gravina et al. analyzed this tool showed some attractive potential in addressing IBD related issues, while having significant limitations in updating and detailing information and providing inaccurate information in some cases (<xref ref-type="bibr" rid="ref10">10</xref>). Among them, ChatGPT, launched by OpenAI on November 30, 2022, is a new type of natural large language model (LLM) that performs well in medical education and training (<xref ref-type="bibr" rid="ref11">11</xref>). To date, ChatGPTs have successfully passed various large medical licensing exams and other medical tests (<xref ref-type="bibr" rid="ref12 ref13 ref14 ref15">12&#x2013;15</xref>).</p>
<p>Therefore, LLM could be considered an interactive information resource for patients with HP infection or their families, as well as a tool for clinicians, but its ability to guide HP management is uncertain. This study was designed to assess the potential medical capacity of LLM for HP-related questions among both experts and ordinary people.</p>
</sec>
<sec sec-type="methods" id="sec6">
<label>2</label>
<title>Methods</title>
<sec id="sec7">
<label>2.1</label>
<title>Study design</title>
<p>We formulated a series of HP-related questions on the basis of the latest relevant guidelines (<xref ref-type="bibr" rid="ref6">6</xref>, <xref ref-type="bibr" rid="ref16">16</xref>, <xref ref-type="bibr" rid="ref17">17</xref>) and clinical experience (at least 15&#x2009;years&#x2019; clinical HP work), involving life (Q1-5), test (Q6-13), and treatment guidance (Q14-22). Each question was submitted to ChatGPT-4 (<xref ref-type="bibr" rid="ref18">18</xref>) (version 3/14/2023)<xref ref-type="fn" rid="fn0001"><sup>1</sup></xref> three times during independent interactions without intervening feedback to assess repeatability. Nevertheless, the evaluation of model performance was planned to be restricted to the analysis of only the initial run (Answer 1). The whole process took place from April 28, 2024, to May 07, 2024. To prevent LLM from dodging medical questions, we inform ChatGPT-4 in advance that &#x201C;Now that you are a professional gastroenterologist, meanwhile you have mastered the latest guidelines about <italic>Helicobacter pylori</italic>. Please answer the following questions.&#x201D;</p>
</sec>
<sec id="sec8">
<label>2.2</label>
<title>Comparative analysis of answers</title>
<p>A temporary assessment team, established by three HP experts, evaluated the capacity (accuracy, completeness, usefulness, comprehension, and satisfaction) of the responses, while 14 ordinary people, as nonexperts, were divided into seven medical students groups and seven nonmedical groups and evaluated for usefulness and comprehension only. All the dimensions were scored via a Likert scale (<xref ref-type="table" rid="tab1">Table 1</xref>).</p>
<table-wrap position="float" id="tab1">
<label>Table 1</label>
<caption>
<p>Likert scales of every dimension.</p>
</caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="middle">Score</th>
<th align="left" valign="middle">Accuracy rating by a five-point Likert scale</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="middle">1</td>
<td align="left" valign="middle">Completely incorrect</td>
</tr>
<tr>
<td align="left" valign="middle">2</td>
<td align="left" valign="middle">More incorrect than correct</td>
</tr>
<tr>
<td align="left" valign="middle">3</td>
<td align="left" valign="middle">Approximately equal correct and incorrect</td>
</tr>
<tr>
<td align="left" valign="middle">4</td>
<td align="left" valign="middle">Mostly accurate, with some slight inaccuracies or irrelevant information</td>
</tr>
<tr>
<td align="left" valign="middle">5</td>
<td align="left" valign="middle">Correct</td>
</tr>
<tr>
<td align="left" valign="middle">Score</td>
<td align="left" valign="middle">Completeness rating by a three-point Likert scale</td>
</tr>
<tr>
<td align="left" valign="middle">1</td>
<td align="left" valign="middle">Incomplete, addresses some aspects of the question, but significant parts are missing or incomplete</td>
</tr>
<tr>
<td align="left" valign="middle">2</td>
<td align="left" valign="middle">Generally complete, with the minimum amount of information</td>
</tr>
<tr>
<td align="left" valign="middle">3</td>
<td align="left" valign="middle">Very complete, addresses all aspects of the question</td>
</tr>
<tr>
<td align="left" valign="middle">Score</td>
<td align="left" valign="middle">Usefullness rating by a three-point Likert scale</td>
</tr>
<tr>
<td align="left" valign="middle">1</td>
<td align="left" valign="middle">No guidance</td>
</tr>
<tr>
<td align="left" valign="middle">2</td>
<td align="left" valign="middle">Containing only generic information or guidance</td>
</tr>
<tr>
<td align="left" valign="middle">3</td>
<td align="left" valign="middle">Containing some specific guidance</td>
</tr>
<tr>
<td align="left" valign="middle">Score</td>
<td align="left" valign="middle">Comprehension rating by a three-point Likert scale</td>
</tr>
<tr>
<td align="left" valign="middle">1</td>
<td align="left" valign="middle">Difficult to understand</td>
</tr>
<tr>
<td align="left" valign="middle">2</td>
<td align="left" valign="middle">Partly difficult to understand</td>
</tr>
<tr>
<td align="left" valign="middle">3</td>
<td align="left" valign="middle">Easy to understand</td>
</tr>
<tr>
<td align="left" valign="middle">Score</td>
<td align="left" valign="middle">Satisfaction rating by a five-point Likert scale</td>
</tr>
<tr>
<td align="left" valign="middle">1</td>
<td align="left" valign="middle">Very dissatisfied</td>
</tr>
<tr>
<td align="left" valign="middle">2</td>
<td align="left" valign="middle">Dissatisfied</td>
</tr>
<tr>
<td align="left" valign="middle">3</td>
<td align="left" valign="middle">Moderate satisfied</td>
</tr>
<tr>
<td align="left" valign="middle">4</td>
<td align="left" valign="middle">Satisfied</td>
</tr>
<tr>
<td align="left" valign="middle">5</td>
<td align="left" valign="middle">Very satisfied</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec id="sec9">
<label>2.3</label>
<title>Statistical analysis</title>
<p>SPSS 26.0 software (IBM Corp.) was used for statistical analysis, and GraphPad Prism 9.5 (GraphPad Software, Inc.) was used for data visualization and graph plotting. When <italic>p</italic>&#x2009;&#x003C;&#x2009;0.05, the difference was considered statistically significant. The Kolmogorov&#x2013;Smirnov test was used to check whether the data were normally distributed, and Levene&#x2019;s test was used for homogeneity of variance. The least significant difference (LSD) test and Kruskal Walls test were used for pairwise comparisons. The consistency of scores among multiple raters was evaluated by Fless&#x2019;s kappa.</p>
</sec>
</sec>
<sec sec-type="results" id="sec10">
<label>3</label>
<title>Results</title>
<sec id="sec11">
<label>3.1</label>
<title>Response repeatability</title>
<p><xref ref-type="table" rid="tab2">Table 2</xref> and <xref rid="SM1" ref-type="supplementary-material">Supplementary Table 1</xref> list each preset question and the ChatGPT-4 answer for this project. When the same question was submitted to ChatGPT-4 independently, 86.36% (19/22) of the questions received responses that were generally consistent with the previous answers. The answers to questions 3, 12 and 19 reveal subtle inconsistency in some of the details. Regarding dietary of HP patients, ChatGPT-4 focused on the healthy diet, and the third answer of Q3 mentioned pickled foods. In terms of false negatives explanation, the last answer involved the influence of bismuth-containing compounds (Q12). For people of penicillin allergic, ChatGPT had a different advice. The last answer provided the most solutions, including the high-dose dual therapy. In addition, it focused on differences in antibiotic resistance patterns (Q19).</p>
<table-wrap position="float" id="tab2">
<label>Table 2</label>
<caption>
<p>Questions imported into ChatGPT-4.</p>
</caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="middle">Number</th>
<th align="left" valign="middle">Questions</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="middle">Q1</td>
<td align="left" valign="middle">Today I had a medical check-up at my company, and the doctor said I have HP<xref ref-type="table-fn" rid="tfn1"><sup>1</sup></xref>. What is HP?</td>
</tr>
<tr>
<td align="left" valign="middle">Q2</td>
<td align="left" valign="middle">What impact will HP have on my life, studies, or work?</td>
</tr>
<tr>
<td align="left" valign="middle">Q3</td>
<td align="left" valign="middle">I&#x2019;ve been diagnosed with HP, are there any dietary considerations I need to be aware of?</td>
</tr>
<tr>
<td align="left" valign="middle">Q4</td>
<td align="left" valign="middle">Can HP cause cancer?</td>
</tr>
<tr>
<td align="left" valign="middle">Q5</td>
<td align="left" valign="middle">I tested positive for HP, is it hereditary?</td>
</tr>
<tr>
<td align="left" valign="middle">Q6</td>
<td align="left" valign="middle">I tested positive for HP, should my family members also get tested?</td>
</tr>
<tr>
<td align="left" valign="middle">Q7</td>
<td align="left" valign="middle">I recently have acid reflux and have been taking omeprazole, should I get tested for HP?</td>
</tr>
<tr>
<td align="left" valign="middle">Q8</td>
<td align="left" valign="middle">I&#x2019;ve recently had bad breath, do I need to get tested for HP?</td>
</tr>
<tr>
<td align="left" valign="middle">Q9</td>
<td align="left" valign="middle">I had a cold last week and took cold medicine, can I still get tested for HP?</td>
</tr>
<tr>
<td align="left" valign="middle">Q10</td>
<td align="left" valign="middle">Does our whole family need to get tested for HP and treated together?</td>
</tr>
<tr>
<td align="left" valign="middle">Q11</td>
<td align="left" valign="middle">I need to get tested for HP, what tests can I do? What are the advantages and disadvantages of these tests?</td>
</tr>
<tr>
<td align="left" valign="middle">Q12</td>
<td align="left" valign="middle">My C13 breath test result was negative, does this mean I do not have HP?</td>
</tr>
<tr>
<td align="left" valign="middle">Q13</td>
<td align="left" valign="middle">My C13 breath test value is very high, does that mean it is severe?</td>
</tr>
<tr>
<td align="left" valign="middle">Q14</td>
<td align="left" valign="middle">My C13 breath test result is positive, can I avoid treatment? If not, please provide a specific treatment plan.</td>
</tr>
<tr>
<td align="left" valign="middle">Q15</td>
<td align="left" valign="middle">My blood test results show HP antibody positive, do I need treatment?</td>
</tr>
<tr>
<td align="left" valign="middle">Q16</td>
<td align="left" valign="middle">An elderly family member tested positive for HP last week, do they need treatment?</td>
</tr>
<tr>
<td align="left" valign="middle">Q17</td>
<td align="left" valign="middle">My child tested positive for HP last week, do they need treatment?</td>
</tr>
<tr>
<td align="left" valign="middle">Q18</td>
<td align="left" valign="middle">My wife is pregnant and tested positive for HP last week, does she need treatment?</td>
</tr>
<tr>
<td align="left" valign="middle">Q19</td>
<td align="left" valign="middle">I tested positive for HP, but I&#x2019;m allergic to penicillin, what should I do?</td>
</tr>
<tr>
<td align="left" valign="middle">Q20</td>
<td align="left" valign="middle">I took the prescribed antibiotics, does this mean I&#x2019;m cured?</td>
</tr>
<tr>
<td valign="middle" align="left">Q21</td>
<td align="left" valign="middle">I followed the treatment and took antibiotics, why is the follow-up test still positive? What should I do next?</td>
</tr>
<tr>
<td align="left" valign="middle">Q22</td>
<td align="left" valign="middle">I heard that potassium-competitive acid blockers is a good drug. Can this be used to treat HP infection?</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<fn id="tfn1">
<label>1</label>
<p>HP, <italic>Helicobacter pylori</italic>.</p>
</fn>
</table-wrap-foot>
</table-wrap>
</sec>
<sec id="sec12">
<label>3.2</label>
<title>Response length analysis of ChatGPT-4</title>
<p><xref ref-type="table" rid="tab3">Table 3</xref> presents the response lengths of ChatGPT-4 across each response. The average word count was 195.94&#x2009;&#x00B1;&#x2009;52.96. Among each response, the average word count and SD of answers 1 to 3 were 198.91&#x2009;&#x00B1;&#x2009;58.51, 190.82&#x2009;&#x00B1;&#x2009;52.25, and 198.09&#x2009;&#x00B1;&#x2009;69.95, respectively (<italic>p</italic>&#x2009;&#x003E;&#x2009;0.05). The average character count was 1302.30&#x2009;&#x00B1;&#x2009;382.56. Among each response, the character average count and SD of answers 1 to 3 were 1314.09&#x2009;&#x00B1;&#x2009;407.04, 1278.73&#x2009;&#x00B1;&#x2009;376.14, and 1314.09&#x2009;&#x00B1;&#x2009;493.36, respectively (<italic>p</italic>&#x2009;&#x003E;&#x2009;0.05).</p>
<table-wrap position="float" id="tab3">
<label>Table 3</label>
<caption>
<p>Length analysis of ChatGPT-4 answers.</p>
</caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="top">Response length of GPT-4</th>
<th align="center" valign="top">Word count (SD)</th>
<th align="center" valign="top">Minimum</th>
<th align="center" valign="top">Maximum</th>
<th align="center" valign="top">Character count (SD)</th>
<th align="center" valign="top">Minimum</th>
<th align="center" valign="top">Maximum</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="middle">Answer1</td>
<td align="center" valign="middle">198.91 (58.51)</td>
<td align="center" valign="middle">80</td>
<td align="center" valign="middle">312</td>
<td align="center" valign="middle">1314.09 (407.04)</td>
<td align="center" valign="middle">541</td>
<td align="center" valign="middle">2032</td>
</tr>
<tr>
<td align="left" valign="middle">Answer2</td>
<td align="center" valign="middle">190.82 (52.25)</td>
<td align="center" valign="middle">98</td>
<td align="center" valign="middle">303</td>
<td align="center" valign="middle">1278.73 (376.14)</td>
<td align="center" valign="middle">642</td>
<td align="center" valign="middle">2030</td>
</tr>
<tr>
<td align="left" valign="middle">Answer3</td>
<td align="center" valign="middle">198.09 (69.95)</td>
<td align="center" valign="middle">107</td>
<td align="center" valign="middle">363</td>
<td align="center" valign="middle">1314.09 (493.36)</td>
<td align="center" valign="middle">716</td>
<td align="center" valign="middle">2,374</td>
</tr>
<tr>
<td align="left" valign="middle">Total</td>
<td align="center" valign="middle">195.94 (52.96)</td>
<td align="center" valign="middle">95</td>
<td align="center" valign="middle">286</td>
<td align="center" valign="middle">1302.30 (382.56)</td>
<td align="center" valign="middle">633</td>
<td align="center" valign="middle">2045.67</td>
</tr>
<tr>
<td align="left" valign="middle"><italic>p</italic> value</td>
<td align="center" valign="middle">0.969</td>
<td/>
<td/>
<td align="center" valign="middle">0.991</td>
<td/>
<td/>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec id="sec13">
<label>3.3</label>
<title>Interrater reliability</title>
<p>Fleiss&#x2019; kappa was used to assess the consistency of the ratings. The Fleiss&#x2019; kappa evaluations for &#x201C;accuracy (Fleiss&#x2019; kappa: 0.690, <italic>p</italic>&#x2009;&#x003C;&#x2009;0.001)&#x201D; was rated as &#x201C;substantial agreement,&#x201D; while &#x201C;completeness (Fleiss&#x2019; kappa: 0.456, <italic>p</italic>&#x2009;&#x003C;&#x2009;0.001),&#x201D; &#x201C;usefulness (Fleiss&#x2019; kappa: 0.564, <italic>p</italic>&#x2009;&#x003C;&#x2009;0.001)&#x201D; and &#x201C;satisfaction (Fleiss&#x2019; kappa: 0.580, <italic>p</italic>&#x2009;&#x003C;&#x2009;0.001)&#x201D; were rated as &#x201C;moderate agreement,&#x201D; and &#x201C;comprehension (Fleiss&#x2019; kappa: 0.302, <italic>p</italic>&#x2009;=&#x2009;0.014)&#x201D; was rated as &#x201C;fair agreement&#x201D; (<xref rid="SM1" ref-type="supplementary-material">Supplementary Table 2</xref>).</p>
</sec>
<sec id="sec14">
<label>3.4</label>
<title>Evaluation of the ChatGPT-4 responses in each dimension by the experts</title>
<p>As shown in <xref ref-type="fig" rid="fig1">Figure 1</xref> and <xref ref-type="table" rid="tab4">Table 4</xref>, the average quality scores of accuracy, completeness, usefulness, comprehension, and satisfaction for ChatGPT-4 were 4.58&#x2009;&#x00B1;&#x2009;0.50, 2.79&#x2009;&#x00B1;&#x2009;0.41, 2.83&#x2009;&#x00B1;&#x2009;0.38, 2.95&#x2009;&#x00B1;&#x2009;0.21, and 4.55&#x2009;&#x00B1;&#x2009;0.53, respectively. The score for life guidance was higher than that for test and treat guidance, but the differences were not significant (<italic>p</italic>&#x2009;&#x003E;&#x2009;0.05).</p>
<fig position="float" id="fig1">
<label>Figure 1</label>
<caption>
<p>Different dimension analysis of ChatGPT-4 answers to HP queries by experts.</p>
</caption>
<graphic xlink:href="fmed-11-1489117-g001.tif"/>
</fig>
<table-wrap position="float" id="tab4">
<label>Table 4</label>
<caption>
<p>Each dimension scores analysis of ChatGPT-4 answers by experts (Average&#x2009;&#x00B1;&#x2009;SD).</p>
</caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="top">Dimension</th>
<th align="center" valign="top">Overall</th>
<th align="center" valign="top">Life guidance</th>
<th align="center" valign="top">Test guidance</th>
<th align="center" valign="top">Treatment guidance</th>
<th align="center" valign="top"><italic>p</italic> value</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="middle">Accuracy</td>
<td align="center" valign="middle">4.58&#x2009;&#x00B1;&#x2009;0.50</td>
<td align="center" valign="middle">4.73&#x2009;&#x00B1;&#x2009;0.46</td>
<td align="center" valign="middle">4.50&#x2009;&#x00B1;&#x2009;0.51</td>
<td align="center" valign="middle">4.56&#x2009;&#x00B1;&#x2009;0.51</td>
<td align="center" valign="middle">0.548</td>
</tr>
<tr>
<td align="left" valign="middle">Completeness</td>
<td align="center" valign="middle">2.79&#x2009;&#x00B1;&#x2009;0.41</td>
<td align="center" valign="middle">2.80&#x2009;&#x00B1;&#x2009;0.41</td>
<td align="center" valign="middle">2.79&#x2009;&#x00B1;&#x2009;0.41</td>
<td align="center" valign="middle">2.78&#x2009;&#x00B1;&#x2009;0.42</td>
<td align="center" valign="middle">0.999</td>
</tr>
<tr>
<td align="left" valign="middle">Usefulness</td>
<td align="center" valign="middle">2.83&#x2009;&#x00B1;&#x2009;0.38</td>
<td align="center" valign="middle">2.93&#x2009;&#x00B1;&#x2009;0.26</td>
<td align="center" valign="middle">2.79&#x2009;&#x00B1;&#x2009;0.41</td>
<td align="center" valign="middle">2.81&#x2009;&#x00B1;&#x2009;0.40</td>
<td align="center" valign="middle">0.697</td>
</tr>
<tr>
<td align="left" valign="middle">Comprehension</td>
<td align="center" valign="middle">2.95&#x2009;&#x00B1;&#x2009;0.21</td>
<td align="center" valign="middle">3.00&#x2009;&#x00B1;&#x2009;0.00</td>
<td align="center" valign="middle">3.00&#x2009;&#x00B1;&#x2009;0.00</td>
<td align="center" valign="middle">2.89&#x2009;&#x00B1;&#x2009;0.32</td>
<td align="center" valign="middle">0.212</td>
</tr>
<tr>
<td align="left" valign="middle">Satisfaction</td>
<td align="center" valign="middle">4.55&#x2009;&#x00B1;&#x2009;0.53</td>
<td align="center" valign="middle">4.73&#x2009;&#x00B1;&#x2009;0.46</td>
<td align="center" valign="middle">4.46&#x2009;&#x00B1;&#x2009;0.51</td>
<td align="center" valign="middle">4.52&#x2009;&#x00B1;&#x2009;0.58</td>
<td align="center" valign="middle">0.434</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec id="sec15">
<label>3.5</label>
<title>Performance of ChatGPT-4 among ordinary people</title>
<p>The average overall usefulness score for the ChatGPT-4 by experts was 2.83&#x2009;&#x00B1;&#x2009;0.38, which was significantly higher than that of nonexperts (2.42&#x2009;&#x00B1;&#x2009;0.73; <italic>p</italic>&#x2009;&#x003C;&#x2009;0.001). The scores of medical students (2.68&#x2009;&#x00B1;&#x2009;0.54) were higher than nonmedical people (2.16&#x2009;&#x00B1;&#x2009;0.79; <italic>p</italic>&#x2009;&#x003C;&#x2009;0.001; <xref ref-type="fig" rid="fig2">Figure 2A</xref>).</p>
<fig position="float" id="fig2">
<label>Figure 2</label>
<caption>
<p>Usefulness (A) and comprehension (B) score analysis of ChatGPT-4 answers between experts and nonexperts.</p>
</caption>
<graphic xlink:href="fmed-11-1489117-g002.tif"/>
</fig>
<p>In terms of comprehension scores, the average overall score for experts was 2.95&#x2009;&#x00B1;&#x2009;0.21, which was significantly higher than that for nonexperts (2.35&#x2009;&#x00B1;&#x2009;0.74; <italic>p</italic>&#x2009;&#x003C;&#x2009;0.001). Further analysis revealed that medical majors scored higher than nonmedical majors did (<italic>p</italic>&#x2009;&#x003C;&#x2009;0.01; <xref ref-type="fig" rid="fig2">Figure 2B</xref>).</p>
</sec>
</sec>
<sec sec-type="discussion" id="sec16">
<label>4</label>
<title>Discussion</title>
<p>Although the infection rate of HP in China has been slowly declining over the past three to four decades, it is still a major health threat to families and society in China (<xref ref-type="bibr" rid="ref19">19</xref>). These HP infected people develop different types and degrees of gastrointestinal and extragastrointestinal diseases, such as dyspepsia, chronic gastritis, peptic ulcers, gastric cancer, iron deficiency anemia, and idiopathic thrombocytopenic purpura (<xref ref-type="bibr" rid="ref16">16</xref>, <xref ref-type="bibr" rid="ref20">20</xref>). However, the public&#x2019;s understanding of HP is not enough. Previous studies have demonstrated the promising prospects of ChatGPT in medicine, and some have assessed the ability of ChatGPT-3.5 to address HP-related queries (<xref ref-type="bibr" rid="ref21 ref22 ref23">21&#x2013;23</xref>). However, previous studies have evaluated only the accuracy and repeatability of the ChatGPT-3.5 model, and the set of questions has neglected the family cluster characteristics associated with HP infection (<xref ref-type="bibr" rid="ref23">23</xref>). Therefore, the purpose of this project is to determine whether chaGPT-4 can solve HP-related questions and explore the potential applications of LLM among ordinary people.</p>
<p>As an important tool, the LLM of AI is gradually affecting every field among human beings. According to the HP guidelines and clinical experience, we designed several HP-related issues. We found that ChatGPT-4 performed well in terms of the repeatability of each prompt. However, there are still some subtle differences that need to be carefully identified, which is obviously better than the previous research results of ChatGPT-3.5 (<xref ref-type="bibr" rid="ref24">24</xref>). These findings suggest that LLMs, such as ChatGPT-4, have significant potential medical applications in the future.</p>
<p>We found that the ChatGPT-4 resulted in high scores on accuracy, completeness, usefulness, comprehension and satisfaction dimensions in terms of performance on HP-related questions. After the questions were classified by life, test and treatment guidance, the score for life guidance was highest, but the differences were not significant. This illustrates that the ability of ChatGPT-4 to respond to HP-related issues is good and that ChatGPT-4 has a good breadth of knowledge. Owing to its vast dataset and continuous learning ability, ChatGPT-4 performs well in processing medical information. In addition, the integration of an advanced reasoning mechanism and strict adherence to the guidelines enabled ChatGPT-4 to address complex clinical demands. Importantly, the inclusion of a substantial volume of up-to-date medical training data and the assimilation of lessons from practical application experiences collectively have improved the quality and relevance of the responses provided by ChatGPT-4 (<xref ref-type="bibr" rid="ref25">25</xref>, <xref ref-type="bibr" rid="ref26">26</xref>). ChatGPT would firstly respond according to guidelines. But there is a time limit for model training. At the same time, ChatGPT would take patient-specific factors into consideration. When mentioned the heredity (Q5), it supplemented the transmission route of HP after it denied the heredity of HP. Involved in HP treatment (Q13), it provided alternatives for penicillin allergy patients, while we did not ask about how to resolve allergy patients beforehand. Finally, ChatGPT will lead you to follow doctor&#x2019;s advice. In addition, multiple responses mentioned relevant guidelines and these contents of the repeated responses were roughly the same. ChatGPT-4 did not clearly state which literature was cited, needing a step further prompt. This means that LLMs, such as ChatGPT-4, can become excellent tools for both doctors and patients, showing considerable development prospects. However, these outputs by ChatGPT were needed suspicion, although it was rated well by the experts of our research.</p>
<p>We further collected and analyzed the masses to assess the comprehensiveness and usefulness of the data. The analysis revealed that the usefulness and comprehension results of ChatGPT-4&#x2019;s replies on HP-related queries among nonexperts were not as good as those among experts. Some individuals thought that these answers were too obscure and lacked significance. This gap widened when nonexperts were divided into medical professions and nonmedical majors. In terms of usefulness, although the scores of medical-related majors were lower than those of expert majors, the difference was not statistically significant. In contrast, perhaps owing to the inherent difficulty and threshold of medical knowledge, the average scores were moderate among nonmedical majors, with scores of only 2.08 for comprehension and 2.16 for usefulness. In fact, these results are very easy to understand. HP experts have mastered the latest advances in HP research and have pivotal positions in this field. Compared with ordinary people, those who have medical knowledge, who have a certain knowledge base, can more easily understand the answers. However, the public, especially those in developing countries, generally lack medical knowledge, and many people are still illiterate. The resolution of their problems is of utmost importance, as they constitute the main body of the world. Despite the presence of a hierarchical diagnosis model in China, it still cannot change the phenomenon whereby large hospitals in cities are full of patients and small hospitals are empty. This not only imputes the scarcity and uneven distribution of medical resources but also contributes to the imperfect knowledge system of doctors in small hospitals (<xref ref-type="bibr" rid="ref27">27</xref>). As the birth of LLM, these obsessions may be gradually resolved, which will help the knowledge acquisition of the masses and the rapid progress of medical beginners. On the other hand, LLM may be able to shorten the distance between doctors and patients and increase medical efficiency. Therefore, regardless of whether LLM replaces clinicians, it is likely to become an important consultation option for ordinary patients in the future.</p>
<p>Focusing on the familial aggregation of HP infection, we asked the corresponding questions. Considering questions 6 and 10, ChatGPT-4 suggested that family members of HP patients should only be tested and treated unless they have symptoms or a family history of gastric cancer. However, this reply was too narrow. Mounting evidence has demonstrated that the main route of transmission of HP is through the mouth and that HP infection is associated with a family cluster (<xref ref-type="bibr" rid="ref28">28</xref>, <xref ref-type="bibr" rid="ref29">29</xref>). In addition to traditional test-and-treat and screen-and-treat strategies, a third new family-based strategy has recently been proposed (<xref ref-type="bibr" rid="ref6">6</xref>). The new strategy targets HP-infected individuals within the family, and its scope of application is not affected by HP infection rates. With respect to historical traditional differences, some families in China usually share foods in the same dish or bowl, sometimes using the same utensils, which are sources of HP cross-contamination. Thus, the need for family-based test and treatment becomes the key. The &#x201C;Chinese Consensus Report on Family-Based <italic>Helicobacter pylori</italic> Infection Control and Management (2021 Edition)&#x201D; suggests that, unless there are competing considerations, family-based HP infection management and the eradication of HP infection are recommended, which is helpful for reducing the chance of transmission of infection and reinfection after its eradication (<xref ref-type="bibr" rid="ref6">6</xref>). Compared with the consensus, the answers of ChatGPT-4, which ignore cultural diversity and skip the new strategy, seem to be imperfect. Therefore, excessive care must be taken when AI models are employed in practical medical fields to ensure that inaccurate information is not generated due to model limitations.</p>
<p>Moreover, in terms of treatment regimen, we focused on potassium-competitive acid blockers (P-CABs) in HP treatment (Q22). As a new regimen, the research volume of relevant P-CAB-based HP therapy were not large. Kanu et al. found P-CAB-based therapy had a promising effect on HP eradication (<xref ref-type="bibr" rid="ref30">30</xref>). Distinguishing itself from conventional proton pump inhibitors, this class of drugs can have a more durable and stable acid control effect. P-CAB exhibits versatile clinical applications, encompassing the treatment of gastroesophageal reflux disease, peptic ulcer disease, and HP eradication therapy (<xref ref-type="bibr" rid="ref31">31</xref>, <xref ref-type="bibr" rid="ref32">32</xref>). For P-CAB aspect, ChatGPT-4&#x2019;s answers showed a good performances (Q22). However, ChatGPT-4 did not mention this treatment alternate among other treatment-related questions (Q14-21). ChatGPT will choose those widely recognized treatments, such as triple therapy, rather than those under investigation. There was no real-time access to the internet when responding to queries, so its knowledge base is fundamentally limited (<xref ref-type="bibr" rid="ref33">33</xref>). That is, ChatGPT may not derive enough accurate information until updates has been fed to the model (<xref ref-type="bibr" rid="ref34">34</xref>). This also becomes a constraint, which may not be able to keep up with the big explosion of information.</p>
<p>In addition, as a powerful peer-to-peer fast feedback interactive program, LLM can help people quickly acquire needed information, speed up life and work efficiency. However, LLMs may be addictive as drugs, causing the public to gradually lose their ability to think by themselves. Not long ago, owing to the sudden collapse of the ChatGPT website, people expressed on social media that their lives and jobs were unable to operate completely anymore (<xref ref-type="bibr" rid="ref35">35</xref>). Therefore, regardless of how AI affects our lives in the future, it is crucial to keep a clear mind to judge the progress of science.</p>
<p>This task still has several limitations. This article simply probes the potential medical applications of LLM in the future through the replies of ChatGPT-4 to HP-related questions. This study did not examine other AI models, nor did it examine responses to other clinical questions.</p>
</sec>
<sec sec-type="conclusions" id="sec17">
<label>5</label>
<title>Conclusion</title>
<p>ChatGPT-4 performs well in resolving HP related questions, which is expected to be a convenient and effective tool for people.</p>
</sec>
</body>
<back>
<sec sec-type="data-availability" id="sec18">
<title>Data availability statement</title>
<p>The original contributions presented in the study are included in the article/<xref rid="SM1" ref-type="supplementary-material">Supplementary material</xref>, further inquiries can be directed to the corresponding authors.</p>
</sec>
<sec sec-type="author-contributions" id="sec19">
<title>Author contributions</title>
<p>ZG: Data curation, Formal analysis, Visualization, Writing &#x2013; original draft, Validation. JG: Data curation, Investigation, Writing &#x2013; original draft. RX: Data curation, Validation, Writing &#x2013; original draft. XC: Supervision, Writing &#x2013; review &#x0026; editing, Project administration. ZC: Conceptualization, Methodology, Project administration, Writing &#x2013; review &#x0026; editing.</p>
</sec>
<sec sec-type="funding-information" id="sec20">
<title>Funding</title>
<p>The author(s) declare that no financial support was received for the research, authorship, and/or publication of this article.</p>
</sec>
<ack>
<p>I would like to thank my friend, H. Lin, for his proof reading to the manuscript.</p>
</ack>
<sec sec-type="COI-statement" id="sec21">
<title>Conflict of interest</title>
<p>The authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec sec-type="disclaimer" id="sec22">
<title>Publisher&#x2019;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<sec sec-type="supplementary-material" id="sec23">
<title>Supplementary material</title>
<p>The Supplementary material for this article can be found online at: <ext-link xlink:href="https://www.frontiersin.org/articles/10.3389/fmed.2024.1489117/full#supplementary-material" ext-link-type="uri">https://www.frontiersin.org/articles/10.3389/fmed.2024.1489117/full#supplementary-material</ext-link></p>
<supplementary-material xlink:href="Table_1.XLSX" id="SM1" mimetype="application/vnd.openxmlformats-officedocument.spreadsheetml.sheet" xmlns:xlink="http://www.w3.org/1999/xlink"/>
</sec>
<fn-group>
<fn id="fn0001">
<p><sup>1</sup><ext-link xlink:href="https://openai.com/gpt-4" ext-link-type="uri">https://openai.com/gpt-4</ext-link>
</p>
</fn>
</fn-group>
<ref-list>
<title>References</title>
<ref id="ref1"><label>1.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Duan</surname> <given-names>M</given-names></name> <name><surname>Li</surname> <given-names>Y</given-names></name> <name><surname>Liu</surname> <given-names>J</given-names></name> <name><surname>Zhang</surname> <given-names>W</given-names></name> <name><surname>Dong</surname> <given-names>Y</given-names></name> <name><surname>Han</surname> <given-names>Z</given-names></name> <etal/></person-group>. <article-title>Transmission routes and patterns of helicobacter pylori</article-title>. <source>Helicobacter</source>. (<year>2023</year>) <volume>28</volume>:<fpage>e12945</fpage>. doi: <pub-id pub-id-type="doi">10.1111/hel.12945</pub-id>, PMID: <pub-id pub-id-type="pmid">36645421</pub-id></citation></ref>
<ref id="ref2"><label>2.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Waldum</surname> <given-names>H</given-names></name> <name><surname>Fossmark</surname> <given-names>R</given-names></name></person-group>. <article-title>Inflammation and digestive Cancer</article-title>. <source>Int J Mol Sci</source>. (<year>2023</year>) <volume>24</volume>:<fpage>13503</fpage>. doi: <pub-id pub-id-type="doi">10.3390/ijms241713503</pub-id>, PMID: <pub-id pub-id-type="pmid">37686307</pub-id></citation></ref>
<ref id="ref3"><label>3.</label><citation citation-type="book"><person-group person-group-type="author"><collab id="coll1">IARC Working Group on the Evaluation of Carcinogenic Risks to Humans</collab></person-group>. <article-title>Infection with Helicobacter pylori</article-title> In: <source>Schistosomes, liver flukes and Helicobacter pylori</source>, vol. <volume>61</volume>. <publisher-loc>France</publisher-loc>: <publisher-name>International Agency for Research on Cancer</publisher-name> (<year>1994</year>). <fpage>1</fpage>&#x2013;<lpage>241</lpage>.</citation></ref>
<ref id="ref4"><label>4.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Shin</surname> <given-names>DW</given-names></name> <name><surname>Cho</surname> <given-names>J</given-names></name> <name><surname>Kim</surname> <given-names>SH</given-names></name> <name><surname>Kim</surname> <given-names>YJ</given-names></name> <name><surname>Choi</surname> <given-names>HC</given-names></name> <name><surname>Son</surname> <given-names>KY</given-names></name> <etal/></person-group>. <article-title>Preferences for the "screen and treat" strategy of <italic>Helicobacter pylori</italic> to prevent gastric cancer in healthy Korean populations</article-title>. <source>Helicobacter</source>. (<year>2013</year>) <volume>18</volume>:<fpage>262</fpage>&#x2013;<lpage>9</lpage>. doi: <pub-id pub-id-type="doi">10.1111/hel.12039</pub-id>, PMID: <pub-id pub-id-type="pmid">23384480</pub-id></citation></ref>
<ref id="ref5"><label>5.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Guevara</surname> <given-names>B</given-names></name> <name><surname>Cogdill</surname> <given-names>AG</given-names></name></person-group>. <article-title><italic>Helicobacter pylori</italic>: a review of current diagnostic and management strategies</article-title>. <source>Dig Dis Sci</source>. (<year>2020</year>) <volume>65</volume>:<fpage>1917</fpage>&#x2013;<lpage>31</lpage>. doi: <pub-id pub-id-type="doi">10.1007/s10620-020-06193-7</pub-id></citation></ref>
<ref id="ref6"><label>6.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ding</surname> <given-names>SZ</given-names></name> <name><surname>Du</surname> <given-names>YQ</given-names></name> <name><surname>Lu</surname> <given-names>H</given-names></name></person-group>. <article-title>Chinese consensus report on family-based <italic>Helicobacter pylori</italic> infection control and management (2021 edition)</article-title>. <source>Gut</source>. (<year>2022</year>) <volume>71</volume>:<fpage>238</fpage>&#x2013;<lpage>53</lpage>. doi: <pub-id pub-id-type="doi">10.1136/gutjnl-2021-325630</pub-id>, PMID: <pub-id pub-id-type="pmid">34836916</pub-id></citation></ref>
<ref id="ref7"><label>7.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Mitsala</surname> <given-names>A</given-names></name> <name><surname>Tsalikidis</surname> <given-names>C</given-names></name> <name><surname>Pitiakoudis</surname> <given-names>M</given-names></name> <name><surname>Simopoulos</surname> <given-names>C</given-names></name> <name><surname>Tsaroucha</surname> <given-names>AK</given-names></name></person-group>. <article-title>Artificial intelligence in colorectal Cancer screening, diagnosis and treatment</article-title>. <source>A New Era Curr Oncol</source>. (<year>2021</year>) <volume>28</volume>:<fpage>1581</fpage>&#x2013;<lpage>607</lpage>. doi: <pub-id pub-id-type="doi">10.3390/curroncol28030149</pub-id>, PMID: <pub-id pub-id-type="pmid">33922402</pub-id></citation></ref>
<ref id="ref8"><label>8.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Yang</surname> <given-names>L</given-names></name> <name><surname>Yang</surname> <given-names>J</given-names></name> <name><surname>Kleppe</surname> <given-names>A</given-names></name> <name><surname>Danielsen</surname> <given-names>HE</given-names></name> <name><surname>Kerr</surname> <given-names>DJ</given-names></name></person-group>. <article-title>Personalizing adjuvant therapy for patients with colorectal cancer</article-title>. <source>Nat Rev Clin Oncol</source>. (<year>2024</year>) <volume>21</volume>:<fpage>67</fpage>&#x2013;<lpage>79</lpage>. doi: <pub-id pub-id-type="doi">10.1038/s41571-023-00834-2</pub-id></citation></ref>
<ref id="ref9"><label>9.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Miller</surname> <given-names>RJH</given-names></name> <name><surname>Huang</surname> <given-names>C</given-names></name> <name><surname>Liang</surname> <given-names>JX</given-names></name> <name><surname>Slomka</surname> <given-names>PJ</given-names></name></person-group>. <article-title>Artificial intelligence for disease diagnosis and risk prediction in nuclear cardiology</article-title>. <source>J Nucl Cardiol</source>. (<year>2022</year>) <volume>29</volume>:<fpage>1754</fpage>&#x2013;<lpage>62</lpage>. doi: <pub-id pub-id-type="doi">10.1007/s12350-022-02977-8</pub-id></citation></ref>
<ref id="ref10"><label>10.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Gravina</surname> <given-names>AG</given-names></name> <name><surname>Pellegrino</surname> <given-names>R</given-names></name> <name><surname>Cipullo</surname> <given-names>M</given-names></name> <name><surname>Palladino</surname> <given-names>G</given-names></name> <name><surname>Imperio</surname> <given-names>G</given-names></name> <name><surname>Ventura</surname> <given-names>A</given-names></name> <etal/></person-group>. <article-title>May ChatGPT be a tool producing medical information for common inflammatory bowel disease patients' questions? An evidence-controlled analysis</article-title>. <source>World J Gastroenterol</source>. (<year>2024</year>) <volume>30</volume>:<fpage>17</fpage>&#x2013;<lpage>33</lpage>. doi: <pub-id pub-id-type="doi">10.3748/wjg.v30.i1.17</pub-id>, PMID: <pub-id pub-id-type="pmid">38293321</pub-id></citation></ref>
<ref id="ref11"><label>11.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Li</surname> <given-names>W</given-names></name> <name><surname>Zhang</surname> <given-names>Y</given-names></name> <name><surname>Chen</surname> <given-names>F</given-names></name></person-group>. <article-title>ChatGPT in colorectal surgery: a promising tool or a passing fad?</article-title> <source>Ann Biomed Eng</source>. (<year>2023</year>) <volume>51</volume>:<fpage>1892</fpage>&#x2013;<lpage>7</lpage>. doi: <pub-id pub-id-type="doi">10.1007/s10439-023-03232-y</pub-id></citation></ref>
<ref id="ref12"><label>12.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Oztermeli</surname> <given-names>AD</given-names></name> <name><surname>Oztermeli</surname> <given-names>A</given-names></name></person-group>. <article-title>ChatGPT performance in the medical specialty exam: an observational study</article-title>. <source>Medicine (Baltimore)</source>. (<year>2023</year>) <volume>102</volume>:<fpage>e34673</fpage>. doi: <pub-id pub-id-type="doi">10.1097/md.0000000000034673</pub-id>, PMID: <pub-id pub-id-type="pmid">37565917</pub-id></citation></ref>
<ref id="ref13"><label>13.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Williams</surname> <given-names>DO</given-names></name> <name><surname>Fadda</surname> <given-names>E</given-names></name></person-group>. <article-title>Can ChatGPT pass Glycobiology?</article-title> <source>Glycobiology</source>. (<year>2023</year>) <volume>33</volume>:<fpage>606</fpage>&#x2013;<lpage>14</lpage>. doi: <pub-id pub-id-type="doi">10.1093/glycob/cwad064</pub-id>, PMID: <pub-id pub-id-type="pmid">37531256</pub-id></citation></ref>
<ref id="ref14"><label>14.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Alessandri Bonetti</surname> <given-names>M</given-names></name> <name><surname>Giorgino</surname> <given-names>R</given-names></name> <name><surname>Gallo Afflitto</surname> <given-names>G</given-names></name> <name><surname>de Lorenzi</surname> <given-names>F</given-names></name> <name><surname>Egro</surname> <given-names>FM</given-names></name></person-group>. <article-title>How does ChatGPT perform on the Italian residency admission National Exam Compared to 15,869 medical graduates?</article-title> <source>Ann Biomed Eng</source>. (<year>2024</year>) <volume>52</volume>:<fpage>745</fpage>&#x2013;<lpage>9</lpage>. doi: <pub-id pub-id-type="doi">10.1007/s10439-023-03318-7</pub-id>, PMID: <pub-id pub-id-type="pmid">37490183</pub-id></citation></ref>
<ref id="ref15"><label>15.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Gravina</surname> <given-names>AG</given-names></name> <name><surname>Pellegrino</surname> <given-names>R</given-names></name> <name><surname>Palladino</surname> <given-names>G</given-names></name> <name><surname>Imperio</surname> <given-names>G</given-names></name> <name><surname>Ventura</surname> <given-names>A</given-names></name> <name><surname>Federico</surname> <given-names>A</given-names></name></person-group>. <article-title>Charting new AI education in gastroenterology: cross-sectional evaluation of ChatGPT and perplexity AI in medical residency exam</article-title>. <source>Dig Liver Dis</source>. (<year>2024</year>) <volume>56</volume>:<fpage>1304</fpage>&#x2013;<lpage>11</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.dld.2024.02.019</pub-id>, PMID: <pub-id pub-id-type="pmid">38503659</pub-id></citation></ref>
<ref id="ref16"><label>16.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Katelaris</surname> <given-names>P</given-names></name> <name><surname>Hunt</surname> <given-names>R</given-names></name> <name><surname>Bazzoli</surname> <given-names>F</given-names></name> <name><surname>Cohen</surname> <given-names>H</given-names></name> <name><surname>Fock</surname> <given-names>KM</given-names></name> <name><surname>Gemilyan</surname> <given-names>M</given-names></name> <etal/></person-group>. <article-title><italic>Helicobacter pylori</italic> world gastroenterology organization global guideline</article-title>. <source>J Clin Gastroenterol</source>. (<year>2023</year>) <volume>57</volume>:<fpage>111</fpage>&#x2013;<lpage>26</lpage>. doi: <pub-id pub-id-type="doi">10.1097/mcg.0000000000001719</pub-id>, PMID: <pub-id pub-id-type="pmid">36598803</pub-id></citation></ref>
<ref id="ref17"><label>17.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Malfertheiner</surname> <given-names>P</given-names></name> <name><surname>Megraud</surname> <given-names>F</given-names></name> <name><surname>Rokkas</surname> <given-names>T</given-names></name> <name><surname>Gisbert</surname> <given-names>JP</given-names></name> <name><surname>Liou</surname> <given-names>JM</given-names></name> <name><surname>Schulz</surname> <given-names>C</given-names></name> <etal/></person-group>. <article-title>Management of <italic>Helicobacter pylori</italic> infection: the Maastricht VI/Florence consensus report</article-title>. <source>Gut</source>. (<year>2022</year>) <volume>71</volume>:<fpage>1724</fpage>&#x2013;<lpage>62</lpage>. doi: <pub-id pub-id-type="doi">10.1136/gutjnl-2022-327745</pub-id></citation></ref>
<ref id="ref18"><label>18.</label><citation citation-type="other"><person-group person-group-type="author"><collab id="coll2">ChatGPT</collab></person-group>. OpenAI. (<year>2023</year>). Available at: <ext-link xlink:href="https://chat.openai.com/chat" ext-link-type="uri">https://chat.openai.com/chat</ext-link> (Accessed on 28 April, 2024).</citation></ref>
<ref id="ref19"><label>19.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Li</surname> <given-names>M</given-names></name> <name><surname>Sun</surname> <given-names>Y</given-names></name> <name><surname>Yang</surname> <given-names>J</given-names></name> <name><surname>de Martel</surname> <given-names>C</given-names></name> <name><surname>Charvat</surname> <given-names>H</given-names></name> <name><surname>Clifford</surname> <given-names>GM</given-names></name> <etal/></person-group>. <article-title>Time trends and other sources of variation in <italic>Helicobacter pylori</italic> infection in mainland China: a systematic review and meta-analysis</article-title>. <source>Helicobacter</source>. (<year>2020</year>) <volume>25</volume>:<fpage>e12729</fpage>. doi: <pub-id pub-id-type="doi">10.1111/hel.12729</pub-id>, PMID: <pub-id pub-id-type="pmid">32686261</pub-id></citation></ref>
<ref id="ref20"><label>20.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Santos</surname> <given-names>MLC</given-names></name> <name><surname>de Brito</surname> <given-names>BB</given-names></name> <name><surname>da Silva</surname> <given-names>FAF</given-names></name></person-group>. <article-title><italic>Helicobacter pylori</italic> infection: beyond gastric manifestations</article-title>. <source>World J Gastroenterol</source>. (<year>2020</year>) <volume>26</volume>:<fpage>4076</fpage>&#x2013;<lpage>93</lpage>. doi: <pub-id pub-id-type="doi">10.3748/wjg.v26.i28.4076</pub-id>, PMID: <pub-id pub-id-type="pmid">32821071</pub-id></citation></ref>
<ref id="ref21"><label>21.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Moazzam</surname> <given-names>Z</given-names></name> <name><surname>Cloyd</surname> <given-names>J</given-names></name> <name><surname>Lima</surname> <given-names>HA</given-names></name> <name><surname>Pawlik</surname> <given-names>TM</given-names></name></person-group>. <article-title>Quality of ChatGPT responses to questions related to pancreatic Cancer and its surgical care</article-title>. <source>Ann Surg Oncol</source>. (<year>2023</year>) <volume>30</volume>:<fpage>6284</fpage>&#x2013;<lpage>6</lpage>. doi: <pub-id pub-id-type="doi">10.1245/s10434-023-13777-w</pub-id>, PMID: <pub-id pub-id-type="pmid">37349615</pub-id></citation></ref>
<ref id="ref22"><label>22.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Henson</surname> <given-names>JB</given-names></name> <name><surname>Glissen Brown</surname> <given-names>JR</given-names></name> <name><surname>Lee</surname> <given-names>JP</given-names></name> <name><surname>Patel</surname> <given-names>A</given-names></name> <name><surname>Leiman</surname> <given-names>DA</given-names></name></person-group>. <article-title>Evaluation of the potential utility of an artificial intelligence Chatbot in gastroesophageal reflux disease management</article-title>. <source>Am J Gastroenterol</source>. (<year>2023</year>) <volume>118</volume>:<fpage>2276</fpage>&#x2013;<lpage>9</lpage>. doi: <pub-id pub-id-type="doi">10.14309/ajg.0000000000002397</pub-id>, PMID: <pub-id pub-id-type="pmid">37410934</pub-id></citation></ref>
<ref id="ref23"><label>23.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Lai</surname> <given-names>Y</given-names></name> <name><surname>Liao</surname> <given-names>F</given-names></name> <name><surname>Zhao</surname> <given-names>J</given-names></name> <name><surname>Zhu</surname> <given-names>C</given-names></name> <name><surname>Hu</surname> <given-names>Y</given-names></name> <name><surname>Li</surname> <given-names>Z</given-names></name></person-group>. <article-title>Exploring the capacities of ChatGPT: a comprehensive evaluation of its accuracy and repeatability in addressing helicobacter pylori-related queries</article-title>. <source>Helicobacter</source>. (<year>2024</year>) <volume>29</volume>:<fpage>e13078</fpage>. doi: <pub-id pub-id-type="doi">10.1111/hel.13078</pub-id>, PMID: <pub-id pub-id-type="pmid">38867649</pub-id></citation></ref>
<ref id="ref24"><label>24.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Pugliese</surname> <given-names>N</given-names></name> <name><surname>Wai-Sun Wong</surname> <given-names>V</given-names></name> <name><surname>Schattenberg</surname> <given-names>JM</given-names></name> <name><surname>Romero-Gomez</surname> <given-names>M</given-names></name> <name><surname>Sebastiani</surname> <given-names>G</given-names></name> <name><surname>Aghemo</surname> <given-names>A</given-names></name> <etal/></person-group>. <article-title>Accuracy, reliability, and comprehensibility of ChatGPT-generated medical responses for patients with nonalcoholic fatty liver disease</article-title>. <source>Clin Gastroenterol Hepatol</source>. (<year>2023</year>) <volume>22</volume>:<fpage>886</fpage>&#x2013;<lpage>889.e5</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.cgh.2023.08.033</pub-id></citation></ref>
<ref id="ref25"><label>25.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Uprety</surname> <given-names>D</given-names></name> <name><surname>Zhu</surname> <given-names>D</given-names></name> <name><surname>West</surname> <given-names>HJ</given-names></name></person-group>. <article-title>ChatGPT-A promising generative AI tool and its implications for cancer care</article-title>. <source>Cancer</source>. (<year>2023</year>) <volume>129</volume>:<fpage>2284</fpage>&#x2013;<lpage>9</lpage>. doi: <pub-id pub-id-type="doi">10.1002/cncr.34827</pub-id>, PMID: <pub-id pub-id-type="pmid">37183438</pub-id></citation></ref>
<ref id="ref26"><label>26.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Deng</surname> <given-names>L</given-names></name> <name><surname>Wang</surname> <given-names>T</given-names></name> <name><surname>Yang</surname> <given-names>Z</given-names></name></person-group>. <article-title>Evaluation of large language models in breast cancer clinical scenarios: a comparative analysis based on ChatGPT-3.5, ChatGPT-4.0, and Claude2</article-title>. <source>Int J Surg</source>. (<year>2024</year>) <volume>110</volume>:<fpage>1941</fpage>&#x2013;<lpage>50</lpage>. doi: <pub-id pub-id-type="doi">10.1097/js9.0000000000001066</pub-id>, PMID: <pub-id pub-id-type="pmid">38668655</pub-id></citation></ref>
<ref id="ref27"><label>27.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Xue</surname> <given-names>Z</given-names></name> <name><surname>Zhang</surname> <given-names>Y</given-names></name> <name><surname>Gan</surname> <given-names>W</given-names></name> <name><surname>Wang</surname> <given-names>H</given-names></name> <name><surname>She</surname> <given-names>G</given-names></name> <name><surname>Zheng</surname> <given-names>X</given-names></name></person-group>. <article-title>Quality and dependability of ChatGPT and DingXiangYuan forums for remote orthopedic consultations: comparative analysis</article-title>. <source>J Med Internet Res</source>. (<year>2024</year>) <volume>26</volume>:<fpage>e50882</fpage>. doi: <pub-id pub-id-type="doi">10.2196/50882</pub-id>, PMID: <pub-id pub-id-type="pmid">38483451</pub-id></citation></ref>
<ref id="ref28"><label>28.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Yu</surname> <given-names>XC</given-names></name> <name><surname>Shao</surname> <given-names>QQ</given-names></name> <name><surname>Ma</surname> <given-names>J</given-names></name> <name><surname>Yu</surname> <given-names>M</given-names></name> <name><surname>Zhang</surname> <given-names>C</given-names></name> <name><surname>Lei</surname> <given-names>L</given-names></name> <etal/></person-group>. <article-title>Family-based <italic>Helicobacter pylori</italic> infection status and transmission pattern in Central China, and its clinical implications for related disease prevention</article-title>. <source>World J Gastroenterol</source>. (<year>2022</year>) <volume>28</volume>:<fpage>3706</fpage>&#x2013;<lpage>19</lpage>. doi: <pub-id pub-id-type="doi">10.3748/wjg.v28.i28.3706</pub-id>, PMID: <pub-id pub-id-type="pmid">36161052</pub-id></citation></ref>
<ref id="ref29"><label>29.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Cave</surname> <given-names>DR</given-names></name></person-group>. <article-title>How is <italic>Helicobacter pylori</italic> transmitted?</article-title> <source>Gastroenterology</source>. (<year>1997</year>) <volume>113</volume>:<fpage>S9</fpage>&#x2013;<lpage>S14</lpage>. doi: <pub-id pub-id-type="doi">10.1016/s0016-5085(97)80004-2</pub-id></citation></ref>
<ref id="ref30"><label>30.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Kanu</surname> <given-names>JE</given-names></name> <name><surname>Soldera</surname> <given-names>J</given-names></name></person-group>. <article-title>Treatment of <italic>Helicobacter pylori</italic> with potassium competitive acid blockers: a systematic review and meta-analysis</article-title>. <source>World J Gastroenterol</source>. (<year>2024</year>) <volume>30</volume>:<fpage>1213</fpage>&#x2013;<lpage>23</lpage>. doi: <pub-id pub-id-type="doi">10.3748/wjg.v30.i9.1213</pub-id>, PMID: <pub-id pub-id-type="pmid">38577188</pub-id></citation></ref>
<ref id="ref31"><label>31.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Koike</surname> <given-names>T</given-names></name> <name><surname>Nakagawa</surname> <given-names>K</given-names></name> <name><surname>Kanno</surname> <given-names>T</given-names></name> <name><surname>Iijima</surname> <given-names>K</given-names></name> <name><surname>Shimosegawa</surname> <given-names>T</given-names></name></person-group>. <article-title>New trends of acid-related diseases treatment</article-title>. <source>Nihon Rinsho</source>. (<year>2015</year>) <volume>73</volume>:<fpage>1136</fpage>&#x2013;<lpage>46</lpage>. PMID: <pub-id pub-id-type="pmid">26165070</pub-id></citation></ref>
<ref id="ref32"><label>32.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Garnock-Jones</surname> <given-names>KP</given-names></name></person-group>. <article-title>Vonoprazan: first global approval</article-title>. <source>Drugs</source>. (<year>2015</year>) <volume>75</volume>:<fpage>439</fpage>&#x2013;<lpage>43</lpage>. doi: <pub-id pub-id-type="doi">10.1007/s40265-015-0368-z</pub-id>, PMID: <pub-id pub-id-type="pmid">25744862</pub-id></citation></ref>
<ref id="ref33"><label>33.</label><citation citation-type="journal"><person-group person-group-type="author"><collab id="coll3">OpenAI</collab></person-group>. <article-title>GPT-4 technical report</article-title>. <source>Preprint arXiv</source>. (<year>2024</year>). doi: <pub-id pub-id-type="doi">10.48550/arXiv.2303.08774</pub-id></citation></ref>
<ref id="ref34"><label>34.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Thirunavukarasu</surname> <given-names>AJ</given-names></name> <name><surname>Ting</surname> <given-names>DSJ</given-names></name> <name><surname>Elangovan</surname> <given-names>K</given-names></name> <name><surname>Gutierrez</surname> <given-names>L</given-names></name> <name><surname>Tan</surname> <given-names>TF</given-names></name> <name><surname>Ting</surname> <given-names>DSW</given-names></name></person-group>. <article-title>Large language models in medicine</article-title>. <source>Nat Med</source>. (<year>2023</year>) <volume>29</volume>:<fpage>1930</fpage>&#x2013;<lpage>40</lpage>. doi: <pub-id pub-id-type="doi">10.1038/s41591-023-02448-8</pub-id></citation></ref>
<ref id="ref35"><label>35.</label><citation citation-type="other"><person-group person-group-type="author"><collab id="coll4">Richard Speed</collab></person-group>. Millions forced to use brain as OpenAI's ChatGPT takes morning off. (<year>2024</year>). Available at: <ext-link xlink:href="https://www.theregister.com/2024/06/04/openai_chatgpt_outage/" ext-link-type="uri">https://www.theregister.com/2024/06/04/openai_chatgpt_outage/</ext-link> (Accessed June 10, 2024).</citation></ref>
</ref-list>
</back>
</article>