<?xml version="1.0" encoding="UTF-8"?><!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.0 20040830//EN" "journalpublishing.dtd"><article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" dtd-version="2.0" xml:lang="en" article-type="research-article"><front><journal-meta><journal-id journal-id-type="nlm-ta">J Med Internet Res</journal-id><journal-id journal-id-type="publisher-id">jmir</journal-id><journal-id journal-id-type="index">1</journal-id><journal-title>Journal of Medical Internet Research</journal-title><abbrev-journal-title>J Med Internet Res</abbrev-journal-title><issn pub-type="epub">1438-8871</issn><publisher><publisher-name>JMIR Publications</publisher-name><publisher-loc>Toronto, Canada</publisher-loc></publisher></journal-meta><article-meta><article-id pub-id-type="publisher-id">v28i1e82746</article-id><article-id pub-id-type="doi">10.2196/82746</article-id><article-categories><subj-group subj-group-type="heading"><subject>Original Paper</subject></subj-group></article-categories><title-group><article-title>Development and Usability Assessment of a Health Education Conversational Agent for Patients With Gastric Cancer: Action Research Study</article-title></title-group><contrib-group><contrib contrib-type="author"><name name-style="western"><surname>Kang</surname><given-names>YiChen</given-names></name><degrees>MS</degrees><xref ref-type="aff" rid="aff1">1</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Yan</surname><given-names>YaMin</given-names></name><degrees>MS</degrees><xref ref-type="aff" rid="aff1">1</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Wang</surname><given-names>TianXiao</given-names></name><degrees>MS</degrees><xref ref-type="aff" rid="aff2">2</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Yu</surname><given-names>ZhengHong</given-names></name><degrees>RN</degrees><xref ref-type="aff" rid="aff1">1</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Hu</surname><given-names>Yan</given-names></name><degrees>RN</degrees><xref ref-type="aff" rid="aff1">1</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Wang</surname><given-names>ZhiXun</given-names></name><degrees>MS</degrees><xref ref-type="aff" rid="aff2">2</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Zhang</surname><given-names>JiYang</given-names></name><degrees>PhD</degrees><xref ref-type="aff" rid="aff3">3</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Latour</surname><given-names>Jos M</given-names></name><degrees>PhD</degrees><xref ref-type="aff" rid="aff1">1</xref><xref ref-type="aff" rid="aff4">4</xref></contrib><contrib contrib-type="author" corresp="yes"><name name-style="western"><surname>Zhang</surname><given-names>YuXia</given-names></name><degrees>PhD</degrees><xref ref-type="aff" rid="aff1">1</xref></contrib></contrib-group><aff id="aff1"><institution>Department of Nursing, Zhongshan Hospital, Fudan University</institution><addr-line>Fenglin Road 180</addr-line><addr-line>Shanghai</addr-line><country>China</country></aff><aff id="aff2"><institution>Planning and Management Center, Zhongshan Hospital, Fudan University</institution><addr-line>Shanghai</addr-line><country>China</country></aff><aff id="aff3"><institution>Big Data and Artificial Intelligence Center, Zhongshan Hospital, Fudan University</institution><addr-line>Shanghai</addr-line><country>China</country></aff><aff id="aff4"><institution>School of Nursing and Midwifery, University of Plymouth</institution><addr-line>Plymouth</addr-line><country>United Kingdom</country></aff><contrib-group><contrib contrib-type="editor"><name name-style="western"><surname>Leung</surname><given-names>Tiffany</given-names></name></contrib><contrib contrib-type="editor"><name name-style="western"><surname>Balcarras</surname><given-names>Matthew</given-names></name></contrib></contrib-group><contrib-group><contrib contrib-type="reviewer"><name name-style="western"><surname>Carrillo</surname><given-names>Adriana</given-names></name></contrib><contrib contrib-type="reviewer"><name name-style="western"><surname>Chow</surname><given-names>James C L</given-names></name></contrib><contrib contrib-type="reviewer"><name name-style="western"><surname>Shilo</surname><given-names>Polina</given-names></name></contrib></contrib-group><author-notes><corresp>Correspondence to YuXia Zhang, PhD, Department of Nursing, Zhongshan Hospital, Fudan University, Fenglin Road 180, Shanghai, China, 86 64041990; <email>zhang.yx@aliyun.com</email></corresp></author-notes><pub-date pub-type="collection"><year>2026</year></pub-date><pub-date pub-type="epub"><day>10</day><month>9</month><year>2026</year></pub-date><volume>28</volume><elocation-id>e82746</elocation-id><history><date date-type="received"><day>21</day><month>08</month><year>2025</year></date><date date-type="rev-recd"><day>15</day><month>07</month><year>2026</year></date><date date-type="accepted"><day>31</day><month>07</month><year>2026</year></date></history><copyright-statement>&#x00A9; YiChen Kang, YaMin Yan, TianXiao Wang, ZhengHong Yu, Yan Hu, ZhiXun Wang, JiYang Zhang, Jos M Latour, YuXia Zhang. Originally published in the Journal of Medical Internet Research (<ext-link ext-link-type="uri" xlink:href="https://www.jmir.org">https://www.jmir.org</ext-link>), 10.9.2026. </copyright-statement><copyright-year>2026</copyright-year><license license-type="open-access" xlink:href="https://creativecommons.org/licenses/by/4.0/"><p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (<ext-link ext-link-type="uri" xlink:href="https://creativecommons.org/licenses/by/4.0/">https://creativecommons.org/licenses/by/4.0/</ext-link>), which permits unrestricted use, distribution, and reproduction in any medium, provided the original work, first published in the Journal of Medical Internet Research (ISSN 1438-8871), is properly cited. The complete bibliographic information, a link to the original publication on <ext-link ext-link-type="uri" xlink:href="https://www.jmir.org/">https://www.jmir.org/</ext-link>, as well as this copyright and license information must be included.</p></license><self-uri xlink:type="simple" xlink:href="https://www.jmir.org/2026/1/e82746"/><abstract><sec><title>Background</title><p>To provide patients with gastric cancer with adequate health education information for effective overall management is crucial, while traditional manners exposed certain challenges. Conversational agents have increasingly been adopted for health care use to provide innovative solutions for patient education.</p></sec><sec><title>Objective</title><p>This study aimed to develop a health education embodied conversational agent to focus on gastric cancer disease using an action research approach and test its accuracy, usability, and user experience among patients and other related stakeholders.</p></sec><sec sec-type="methods"><title>Methods</title><p>The AI-guided conversational agent was developed based on the OpenMEDLab 2.0 foundation model and the Retrieval-Augmented Generation (RAG) architecture. A 4-phase action research approach was adopted to implement this system at a gastric cancer center in China. Diagnose and plan phase used participatory observation and in-depth interviews to explore current health education models and patients&#x2019; needs for health education. Act and implement phase was used to develop and deploy the conversational agent. Evaluate phase comprised 3 rounds of alpha testing to assess accuracy and RAG knowledge hit rate, and 1 round of beta testing to assess the usability and relevance. Reflect phase conducted in-depth interviews to gain insights into users' experiences. Data collection was conducted from September 2023 through April 2025. Participants include patients, clinical nurses, nursing managers, surgeons, clinical psychologists, and dietitians. Thematic analysis and multiple-group chi-square tests were performed, respectively, for qualitative and quantitative data. A 2-sided <italic>P</italic> value of &#x003C;.05 was considered statistically significant.</p></sec><sec sec-type="results"><title>Results</title><p>A total of 44 patients, 13 nurses, 3 nursing managers, 2 surgeons, 1 clinical psychologist, and 1 dietitian were recruited during the study procedure. Favorable outcomes in terms of accuracy and usability were achieved. The accuracy of the agent in 3 rounds was 67% (37/55), 71% (44/62), and 82% (31/38), respectively; RAG knowledge hit rates reached 86% (47/55), 98% (61/62), and 100% (38/38). Significant differences (<italic>P</italic>&#x003C;.01) in RAG knowledge hit rates were observed. The mean chatbot usability questionnaire score was 91.9 (SD 3.6), while the mean content relevance score was 3.75 (SD 0.9). Three themes of user experiences were identified: perceived usefulness, ease of use, and intention to use, revealing potential in reducing staff workload and reinforcing patient education.</p></sec><sec sec-type="conclusions"><title>Conclusions</title><p>This study provided insights into how the action research approach can inform the development and usability assessment of a gastric cancer health education conversational agent, also illustrating the value of RAG technology. Additional assessments and improvements are warranted to confirm the effectiveness and safety.</p></sec></abstract><kwd-group><kwd>health education</kwd><kwd>generative artificial intelligence</kwd><kwd>digital health</kwd><kwd>user-centered design</kwd><kwd>action research</kwd></kwd-group></article-meta></front><body><sec id="s1" sec-type="intro"><title>Introduction</title><sec id="s1-1"><title>Background</title><p>Gastric cancer is the most common gastrointestinal tumor worldwide, ranking as the fifth most prevalent cancer and the fifth leading cause of cancer death globally [<xref ref-type="bibr" rid="ref1">1</xref>]. Decreases in 1 or more domains (physical, role, social, etc) of functioning and an increase in negatively associated symptoms (dysphagia, reflux, etc) were found perioperatively among the patients with gastric cancer [<xref ref-type="bibr" rid="ref2">2</xref>]. Furthermore, if patients are unfamiliar with managing these symptoms or lack adequate disease information, these problems may worsen. In addition to the impact of symptoms on the human body, related economic costs and psychological burdens also affect patients&#x2019; recovery [<xref ref-type="bibr" rid="ref3">3</xref>]. Therefore, it is crucial to provide patients with gastric cancer with adequate health education information for effective overall management.</p><p>Patients and their families frequently encounter difficulties in accessing reliable and comprehensible disease-related information [<xref ref-type="bibr" rid="ref4">4</xref>]. This gap is primarily due to the limitations of traditional health education methods, such as pamphlets, verbal instructions, and infrequent consultations, which often fail to meet the dynamic and personalized needs of patients with gastric cancer. Besides, these kinds of health education methods often include too much content with dense text, which may lead to information overload. Patients will feel confused when first facing a large amount of information and even more anxious and fearful with diagnosis and treatment. Patients with gastric cancer also face unique challenges due to the complexity of the surgery, subsequent symptoms, and perioperative care. Current health education resources and methods are insufficient to meet the demand for personalized, real-time responses, which proves to be a catalyst for innovative solutions.</p><p>With the development of technology, some patients seek online medical information when they feel uncertain or concerned. Nevertheless, such behavior sometimes may entail potential risks such as exacerbating anxiety and inappropriate self-management due to unverified information [<xref ref-type="bibr" rid="ref5">5</xref>]. Digital health interventions, particularly conversational agents&#x2014;such as chatbots or virtual assistants, offer promising solutions. Research has proved that conversational agents had the potential to significantly improve patient care by enhancing symptom management, supporting self-management, and providing patient support [<xref ref-type="bibr" rid="ref6">6</xref>-<xref ref-type="bibr" rid="ref8">8</xref>]. Gomaa et al [<xref ref-type="bibr" rid="ref9">9</xref>] developed an SMS text messaging&#x2013;integrated and chatbot-interfaced self-management program for patients with gastrointestinal cancer undergoing chemotherapy. The authors noted that the omission of natural language processing (NLP) capabilities restricted deeper interactive engagement, and future research should explore integrating NLP to enhance responsiveness to diverse patient expressions. Naseri et al [<xref ref-type="bibr" rid="ref10">10</xref>] examined the effectiveness of 2 AI chatbots, Sider Fusion AI Bot and Perplexity AI, in improving patient outcomes, alleviating anxiety, and promoting informed decision-making, highlighting the importance of tailored communication styles to enhance patient engagement and outcomes. Thus, an AI-based conversational agent is suggested to use in our study.</p><p>However, generic and unverified AI-based conversational agents may pose potential health risks such as delivering inaccurate information that seems convincing [<xref ref-type="bibr" rid="ref11">11</xref>]. Their integration into clinical practice is sometimes hindered by inherent limitations, and most refers to the phenomenon of &#x201C;hallucinations&#x201D; and unclear boundaries of responsibility and ethics [<xref ref-type="bibr" rid="ref12">12</xref>]. A growing number of AI-driven platforms have been developed for high-prevalence malignancies, such as breast, prostate, and lung cancers, but rigorously designed conversational agents specifically tailored for gastric cancer remain scarce [<xref ref-type="bibr" rid="ref13">13</xref>]. In the high-stakes context of gastric cancer care, where erroneous dietary or postoperative advice can lead to severe clinical complications, ensuring information fidelity is paramount. RAG has emerged as a robust architectural solution to mitigate these risks [<xref ref-type="bibr" rid="ref14">14</xref>]. Zhou et al [<xref ref-type="bibr" rid="ref15">15</xref>] developed a Chinese gastrointestinal disease chatbot and demonstrated the innovative potential of Retrieval-Augmented Generation (RAG) technology. The retrievable knowledge functions as a form of nonparametric memory that is easily updatable, capable of incorporating extensive long-tail knowledge, and suitable for encoding confidential information. Consequently, as the landscape of AI-based conversational agents continues to expand, the focus of development must transcend mere conversational fluency. Clinical safety must remain the cornerstone of AI deployment to ensure that these intelligent agents serve as a secure bridge between complex medical knowledge and the patient&#x2019;s daily recovery.</p><p>Based on the clinical background and gap, 4 research questions are discussed in this study as follows: (1) How to explore patients&#x2019; needs and suggestions on the current gastric health education mode? (2) How to use large language models (LLMs) and RAG to develop a conversational agent for patients with gastric cancer? (3) How to evaluate the performance of the agent in terms of accuracy and usability? (4) How to reflect the user experience among patients and other related stakeholders regarding the conversational agent as a tool for patient education?</p></sec><sec id="s1-2"><title>Objectives</title><p>The aim of this study was to develop and implement an AI-guided conversational agent for patients with gastric cancer. We adopted an action research approach to gain insights from patients and stakeholders, develop the conversational agent, evaluate its usability and accuracy, and gather user feedback by using qualitative and quantitative methods.</p></sec></sec><sec id="s2" sec-type="methods"><title>Methods</title><sec id="s2-1"><title>Study Design</title><p>An action research design was adopted, which is a collaborative, democratic approach and process where research participants are collaborators rather than participants [<xref ref-type="bibr" rid="ref16">16</xref>,<xref ref-type="bibr" rid="ref17">17</xref>]. Unlike conventional approaches, action research emphasizes collaborative knowledge building and social change to develop contextually relevant interventions [<xref ref-type="bibr" rid="ref18">18</xref>]. It has been applied by health care professionals to improve both patient experience and working conditions of those who deliver care [<xref ref-type="bibr" rid="ref19">19</xref>,<xref ref-type="bibr" rid="ref20">20</xref>]. The implementation of digital health interventions presents unique challenges in contrast to other studies [<xref ref-type="bibr" rid="ref21">21</xref>]. To gain the most from digital health products such as conversational agents in our study, we should clearly know what is really needed in practice [<xref ref-type="bibr" rid="ref20">20</xref>]. Co-design or collaboration with all stakeholders provides effective means for developing and implementing digital health interventions that suit the needs of the end users [<xref ref-type="bibr" rid="ref22">22</xref>].</p><p>Action research emphasized that each study should adapt the action research framework according to its specific research aim and content. Our study followed an iterative process adapted from Lewin&#x2019;s [<xref ref-type="bibr" rid="ref23">23</xref>] 4-step cycle of action research (<xref ref-type="fig" rid="figure1">Figure 1</xref>). The reporting guideline and checklist (<xref ref-type="supplementary-material" rid="app3">Checklist 1</xref>) have been used to report our study [<xref ref-type="bibr" rid="ref24">24</xref>].</p><fig position="float" id="figure1"><label>Figure 1.</label><caption><p>Action research cycle illustrating the iterative development of a health education conversational agent for patients with gastric cancer in this study. CUQ: chatbot usability questionnaire; RAG: Retrieval-Augmented Generation.</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="jmir_v28i1e82746_fig01.png"/></fig></sec><sec id="s2-2"><title>Action Research Team</title><p>The multidisciplinary team was composed of 3 advanced nurse specialists, 2 surgeons, 2 researchers, and an information technology engineer. Each team member played a distinct and essential role in the development and implementation of the health education conversational agent. The advanced nurse specialists and surgeons who had 7-20 years of experience in the gastrointestinal field brought rich clinical experience in gastric perioperative care and patient education. They were responsible for identifying key health education weaknesses and constructing a knowledge base for the conversational agent. The researchers led the overall research design, data collection, and analysis process. The information technology engineer was responsible for conversational agent development, program coding, and other informatics work. The action team functioned as a partnership, characterized by shared decision-making and collaborative effort throughout all phases of the research.</p></sec><sec id="s2-3"><title>Setting and Participants</title><p>This study was conducted at the Gastric Cancer Center of Zhongshan Hospital in Shanghai, China, a facility known for its advanced medical care devices and patient-centered care delivery. Patients with gastric cancer and relevant stakeholders were involved in this action research study.</p></sec><sec id="s2-4"><title>Action Research Process</title><p>Four phases were designed to realize the triangulation of research by involving multiple participants and combining qualitative and quantitative approaches [<xref ref-type="bibr" rid="ref25">25</xref>].</p><sec id="s2-4-1"><title>Phase 1: Diagnose and Plan</title><p>The first step of the action research is diagnosing and planning, which means researchers should thoroughly understand the needs of the target population and the problems for the current health education model. Therefore, participatory observation and in-depth interviews were used in the first stage of this action research.</p><p>By engaging in participatory observation, researchers were able to directly depict the structure, effectiveness, and limitations of clinical health education practices of participatory observation for patients with gastric cancer [<xref ref-type="bibr" rid="ref26">26</xref>]. Within a week in September 2023, the researcher participated as a clinical nurse in the gastric cancer center, observing the daily workflow and content of the health education for patients. Data collection included multiple methods: direct observation, audio recordings, and document analysis. The key aspects of the observation included (1) timing and triggers of health education delivery, (2) content and themes of health education, (3) delivery and interaction methods of health education, and (4) patients&#x2019; observable cognitive, emotional, and behavioral responses to health education. Face-to-face semistructured interviews with 16 patients and 12 stakeholders were conducted after the participatory observation. These interviews aimed to explore the existing problems in gastric cancer health education and to identify stakeholders&#x2019; needs and expectations for the development of a health education conversational agent. The in-depth interview outline is tailored for each participant (Table S1 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>).</p></sec><sec id="s2-4-2"><title>Phase 2: Act and Implement</title><p>Action research team collaboratively developed the conversational agent based on results from both participatory observation and in-depth interviews. An evidence-based knowledge corpus was constructed based on multiple sources of professional evidence. The AI-guided conversational agent was developed based on the OpenMEDLab2.0 foundation model and the RAG architecture. To enhance user engagement and provide a more human-like interaction experience, a nurse avatar was integrated into the chatbot interface. Electronic materials are presented in <xref ref-type="supplementary-material" rid="app2">Multimedia Appendix 2</xref>.</p></sec><sec id="s2-4-3"><title>Phase 3: Evaluate</title><p>Three rounds of alpha testing and 1 round of beta testing were conducted to evaluate the accuracy and usability of the conversational agent. Accuracy rate and RAG knowledge hit rate were measured in the alpha testing. In total, 55, 62, and 38 questions from the knowledge corpus were randomly selected for team experts to assess the accuracy of the answers in 3 rounds. Sample question used for expert evaluation is shown in Table S2 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>. A total of 28 patients were recruited to finish the chatbot usability questionnaire (CUQ) to assess the chatbot usability. They also finished a postconversation evaluation of content relevance, rating how well the chatbot&#x2019;s responses matched the questions they asked.</p></sec><sec id="s2-4-4"><title>Phase 4: Reflect</title><p>A total of 15 patients in the usability testing participated in the in-depth interview to deeply explore real experience and suggestions for the conversational agent. Besides, another 8 nurses who worked in the gastric cancer center were also invited to give their views on the new health education platform. The interview outline is shown in Table S3 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>. The results of qualitative study were used to guide the next cycle of action research.</p></sec></sec><sec id="s2-5"><title>Participant Recruitment and Data Collection</title><p>Data collection was conducted from September 2023 through April 2025 across 4 sequential phases. Participants were recruited using convenience sampling from the gastric cancer center of Zhongshan Hospital. Eligible participants were identified by action research team members including patients, clinical nurses, nursing managers, surgeons, clinical psychologists, and dietitians (recruitment criteria in each phase are shown in Table S4 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>).</p><p>Qualitative interviews were recorded using a digital voice recorder with participants&#x2019; permission. Audio recordings were transcribed verbatim into text immediately after each interview. Data collection ceased until no new themes or meaningful information emerged from subsequent interviews. Inpatients who used the conversational agent were invited to finish the validated CUQ on discharge day. CUQ is a chatbot-specific usability questionnaire that is comparable with the Systems Usability Scale [<xref ref-type="bibr" rid="ref27">27</xref>]. It consists of 16 balanced questions related to different aspects of chatbot usability. Odd-numbered questions relate to positive aspects of the chatbot, and even-numbered questions relate to negative aspects. All 16 questions are scored using a 5-point Likert-type scale. Each item was scored on a 5-point Likert scale, ranging from 1 (strongly disagree) to 5 (strongly agree). Scores are calculated out of 100. Besides the CUQ, the conversational agent was also tested by 1 more question about the content relevance. The additional question &#x201C;Do you think the responses provided by the conversational agent appropriately match your question?&#x201D; was also rated on a 5-point Likert scale as CUQ, and this score was analyzed separately. Face validity was assessed among 2 patients and 1 nurse.</p><p>The accuracy and RAG knowledge hit rate are 2 critical metrics for analyzing the performance of question-answering systems [<xref ref-type="bibr" rid="ref28">28</xref>]. The accuracy rate was evaluated by 5 team experts (3 advanced nurse specialists and 2 surgeons) independently; each response was rated as correct or incorrect, and it was considered accurate if at least 4 of the 5 experts judged it to be correct. Overall interrater agreement across all 3 rounds of evaluation was calculated to verify the consistency of the expert assessments. The RAG knowledge hit rate determines whether the system can accurately retrieve relevant medical information from the knowledge base, preventing answer errors caused by missing knowledge or retrieval failures while mitigating the risk of large-model &#x201C;hallucinations.&#x201D; It was calculated directly by the system. We combined both automated retrieval metrics and manual assessment to offer practical and adequate evaluation results.</p></sec><sec id="s2-6"><title>Statistical Analysis</title><sec id="s2-6-1"><title>Qualitative Data Analysis</title><p>NVivo (version 20; QSR International) and content analysis approach [<xref ref-type="bibr" rid="ref29">29</xref>] were used to identify themes and subthemes from the participatory observation and interview data. Two independent reviewers (YCK and YMY) launched the initial coding by immersing themselves in the data and reading the interview transcripts word by word. A tree diagram was used to organize the subthemes into themes. A discussion was held after the initial coding to resolve any discrepancies until a consensus was reached to ensure the consistency and completeness in the final results.</p></sec><sec id="s2-6-2"><title>Quantitative Statistical Analysis</title><p>EXCEL and SPSS software (version 26.0; IBM Corp) were used to complete the quantitative data analysis. A CUQ calculation tool, a Microsoft Excel spreadsheet, which was developed by the researchers from Ulster University, is available for the easy calculation of the CUQ scores [<xref ref-type="bibr" rid="ref30">30</xref>]. Outcome measures, including accuracy and usability, were evaluated through 3 calculation formulas: (1) Accuracy = (Number of correctly answered questions evaluated by health care experts/Total number of questions) &#x00D7; 100%. (2) RAG Knowledge Hit Rate = (Number of questions successfully matched with answers from the knowledge base/Total number of questions) &#x00D7; 100%. (3) <inline-formula><mml:math id="ieqn1"><mml:mstyle><mml:mrow><mml:mstyle displaystyle="false"><mml:mi>C</mml:mi><mml:mi>U</mml:mi><mml:mi>Q</mml:mi><mml:mo>=</mml:mo><mml:mo stretchy="false">(</mml:mo><mml:mo stretchy="false">(</mml:mo><mml:msubsup><mml:mi mathvariant="normal">&#x03A3;</mml:mi><mml:mrow><mml:mi>n</mml:mi><mml:mo>=</mml:mo><mml:mn>1</mml:mn></mml:mrow><mml:mrow><mml:mi>m</mml:mi></mml:mrow></mml:msubsup><mml:mn>2</mml:mn><mml:mi>n</mml:mi><mml:mo>&#x2212;</mml:mo><mml:mn>1</mml:mn><mml:mo stretchy="false">)</mml:mo><mml:mo>&#x2212;</mml:mo><mml:mn>5</mml:mn><mml:mo stretchy="false">)</mml:mo><mml:mo>+</mml:mo><mml:mo stretchy="false">(</mml:mo><mml:mn>25</mml:mn><mml:mo>&#x2212;</mml:mo><mml:mo stretchy="false">(</mml:mo><mml:msubsup><mml:mi mathvariant="normal">&#x03A3;</mml:mi><mml:mrow><mml:mi>n</mml:mi><mml:mo>=</mml:mo><mml:mn>1</mml:mn></mml:mrow><mml:mrow><mml:mi>m</mml:mi></mml:mrow></mml:msubsup><mml:mn>2</mml:mn><mml:mi>n</mml:mi><mml:mo stretchy="false">)</mml:mo><mml:mo stretchy="false">)</mml:mo><mml:mo>&#x00D7;</mml:mo><mml:mn>1.6</mml:mn></mml:mstyle></mml:mrow></mml:mstyle></mml:math></inline-formula></p><p>Descriptive statistics are presented with mean (SD), median (IQR), or frequencies (percentages), as appropriate. Comparisons among 3 rounds of accuracy rate and RAG knowledge hit rate were conducted using multiple-group chi-square tests. The interrater agreement was evaluated using the Fleiss &#x03BA;. A 2-sided <italic>P</italic> value of &#x003C;.05 was considered statistically significant.</p></sec></sec><sec id="s2-7"><title>Ethical Considerations</title><p>The use of AI-based conversational agents in health care raises important ethical issues, particularly regarding data privacy, transparency, and patient autonomy. The conversational agent was designed to support, rather than replace, professional medical advice, in order to minimize the risk of overreliance on AI-generated information. The system repeatedly emphasizes that qualified health care professionals&#x2019; advice should be regarded as the final authority. Absolute or definitive language is deliberately avoided in responses, especially in certain sensitive topics such as diagnosis and treatment. Besides, conversations were anonymized and stored in encrypted files accessible only to the research team.</p><p>This study was approved by the Medical Ethical Review Board of Zhongshan Hospital Fudan University (B2022-613R). Individual informed consent was obtained from all participants, whether in the qualitative interviews or the quantitative researches. Participants were also informed that their personal identifiers were removed from the research database and they were voluntary to withdraw from the project at any time without giving a reason. This research complies with the Declaration of Helsinki and conforms to the data protection guidelines outlined by the local governing bodies.</p></sec></sec><sec id="s3" sec-type="results"><title>Results</title><sec id="s3-1"><title>Phase 1: Diagnose and Plan</title><sec id="s3-1-1"><title>Participatory Observation Findings</title><p>The observational data revealed the traditional health education workflow, including the following: (1) patients typically receive health education at 4 key time points: upon admission, before surgery, after surgery, and at discharge, and if they have problems, they would ask health care staff anytime; (2) health education content covered topics such as preoperative preparation, postoperative recovery, dietary management, medication guidance, follow-up arrangements, complication prevention, and psychological support; (3) delivery methods primarily rely on oral instructions, printed materials, and video-based content, and educational videos were played on a continuous loop in public areas of the ward for inpatients, and some nurses use special ways such as teach-back methods to assess the understanding of patients; and (4) patients usually face information overload.</p></sec><sec id="s3-1-2"><title>In-Depth Interview Findings</title><p>Sixteen patients with gastric cancer and 12 stakeholders were recruited in semistructured interviews. Detailed information is shown in <xref ref-type="table" rid="table1">Tables 1</xref> and <xref ref-type="table" rid="table2">2</xref>.</p><table-wrap id="t1" position="float"><label>Table 1.</label><caption><p>Characteristics of patients with gastric cancer who were recruited to explore the existing problems in gastric cancer health education and needs for the development of a health education conversational agent (N=16).</p></caption><table id="table1" frame="hsides" rules="groups"><thead><tr><td align="left" valign="bottom">Patients, sex</td><td align="left" valign="bottom">Age, years</td><td align="left" valign="bottom">Education level</td><td align="left" valign="bottom">Marital status</td><td align="left" valign="bottom">Career status</td></tr></thead><tbody><tr><td align="left" valign="top">P1: female</td><td align="left" valign="top">61</td><td align="left" valign="top">Senior high school</td><td align="left" valign="top">Married</td><td align="left" valign="top">Retired</td></tr><tr><td align="left" valign="top">P2: female</td><td align="left" valign="top">77</td><td align="left" valign="top">Senior high school</td><td align="left" valign="top">Married</td><td align="left" valign="top">Retired</td></tr><tr><td align="left" valign="top">P3: female</td><td align="left" valign="top">44</td><td align="left" valign="top">Bachelor&#x2019;s degree</td><td align="left" valign="top">Married</td><td align="left" valign="top">Worker</td></tr><tr><td align="left" valign="top">P4: male</td><td align="left" valign="top">56</td><td align="left" valign="top">Master&#x2019;s degree</td><td align="left" valign="top">Married</td><td align="left" valign="top">Teacher</td></tr><tr><td align="left" valign="top">P5: female</td><td align="left" valign="top">77</td><td align="left" valign="top">Junior high school</td><td align="left" valign="top">Married</td><td align="left" valign="top">Retired</td></tr><tr><td align="left" valign="top">P6: male</td><td align="left" valign="top">51</td><td align="left" valign="top">Senior high school</td><td align="left" valign="top">Married</td><td align="left" valign="top">Farmer</td></tr><tr><td align="left" valign="top">P7: female</td><td align="left" valign="top">41</td><td align="left" valign="top">Bachelor&#x2019;s degree</td><td align="left" valign="top">Married</td><td align="left" valign="top">Worker</td></tr><tr><td align="left" valign="top">P8: male</td><td align="left" valign="top">77</td><td align="left" valign="top">Senior high school</td><td align="left" valign="top">Married</td><td align="left" valign="top">Retired</td></tr><tr><td align="left" valign="top">P9: female</td><td align="left" valign="top">81</td><td align="left" valign="top">Senior high school</td><td align="left" valign="top">Married</td><td align="left" valign="top">Retired</td></tr><tr><td align="left" valign="top">P10: female</td><td align="left" valign="top">47</td><td align="left" valign="top">Bachelor&#x2019;s degree</td><td align="left" valign="top">Married</td><td align="left" valign="top">Worker</td></tr><tr><td align="left" valign="top">P11: male</td><td align="left" valign="top">45</td><td align="left" valign="top">Bachelor&#x2019;s degree</td><td align="left" valign="top">Married</td><td align="left" valign="top">Worker</td></tr><tr><td align="left" valign="top">P12: male</td><td align="left" valign="top">78</td><td align="left" valign="top">Senior high school</td><td align="left" valign="top">Married</td><td align="left" valign="top">Retired</td></tr><tr><td align="left" valign="top">P13: male</td><td align="left" valign="top">59</td><td align="left" valign="top">Senior high school</td><td align="left" valign="top">Married</td><td align="left" valign="top">Retired</td></tr><tr><td align="left" valign="top">P14: female</td><td align="left" valign="top">38</td><td align="left" valign="top">Bachelor&#x2019;s degree</td><td align="left" valign="top">Married</td><td align="left" valign="top">Worker</td></tr><tr><td align="left" valign="top">P15: male</td><td align="left" valign="top">25</td><td align="left" valign="top">Bachelor&#x2019;s degree</td><td align="left" valign="top">Single</td><td align="left" valign="top">Worker</td></tr><tr><td align="left" valign="top">P16: male</td><td align="left" valign="top">75</td><td align="left" valign="top">Senior high school</td><td align="left" valign="top">Married</td><td align="left" valign="top">Retired</td></tr></tbody></table></table-wrap><table-wrap id="t2" position="float"><label>Table 2.</label><caption><p>Characteristics of stakeholders who were recruited to explore the existing problems in gastric cancer health education and needs for the development of a health education conversational agent (N=12).</p></caption><table id="table2" frame="hsides" rules="groups"><thead><tr><td align="left" valign="bottom">Stakeholders, sex</td><td align="left" valign="bottom">Age, years</td><td align="left" valign="bottom">Education level</td><td align="left" valign="bottom">Role</td></tr></thead><tbody><tr><td align="left" valign="top">S1: female</td><td align="left" valign="top">29</td><td align="left" valign="top">Bachelor&#x2019;s degree</td><td align="left" valign="top">Nurse</td></tr><tr><td align="left" valign="top">S2: female</td><td align="left" valign="top">34</td><td align="left" valign="top">Bachelor&#x2019;s degree</td><td align="left" valign="top">Nurse</td></tr><tr><td align="left" valign="top">S3: female</td><td align="left" valign="top">31</td><td align="left" valign="top">Bachelor&#x2019;s degree</td><td align="left" valign="top">Nurse</td></tr><tr><td align="left" valign="top">S4: male</td><td align="left" valign="top">28</td><td align="left" valign="top">Master&#x2019;s degree</td><td align="left" valign="top">Nurse</td></tr><tr><td align="left" valign="top">S5: female</td><td align="left" valign="top">35</td><td align="left" valign="top">Bachelor&#x2019;s degree</td><td align="left" valign="top">Nurse</td></tr><tr><td align="left" valign="top">S6: male</td><td align="left" valign="top">45</td><td align="left" valign="top">Master&#x2019;s degree</td><td align="left" valign="top">Nursing manager</td></tr><tr><td align="left" valign="top">S7: female</td><td align="left" valign="top">52</td><td align="left" valign="top">Master&#x2019;s degree</td><td align="left" valign="top">Nursing manager</td></tr><tr><td align="left" valign="top">S8: male</td><td align="left" valign="top">39</td><td align="left" valign="top">Master&#x2019;s degree</td><td align="left" valign="top">Nursing manager</td></tr><tr><td align="left" valign="top">S9: female</td><td align="left" valign="top">47</td><td align="left" valign="top">Doctorate</td><td align="left" valign="top">Surgeon</td></tr><tr><td align="left" valign="top">S10: female</td><td align="left" valign="top">50</td><td align="left" valign="top">Doctorate</td><td align="left" valign="top">Surgeon</td></tr><tr><td align="left" valign="top">S11: male</td><td align="left" valign="top">42</td><td align="left" valign="top">Doctorate</td><td align="left" valign="top">Clinical psychologist</td></tr><tr><td align="left" valign="top">S12: male</td><td align="left" valign="top">40</td><td align="left" valign="top">Doctorate</td><td align="left" valign="top">Dietitian</td></tr></tbody></table></table-wrap><p>Four key findings emerged from the interviews, including fragmented and inconsistent health education content, limited accessibility and personalization of education delivery, unmet needs for timely and interactive support, and expectations and concerns for new, technology-supported education manners (<xref ref-type="table" rid="table3">Table 3</xref>). Detailed analysis of each quote is provided in Text 1 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>.</p><table-wrap id="t3" position="float"><label>Table 3.</label><caption><p>Themes identified from the in-depth interview to explore the disadvantages of current health education mode and expectations for new education manners.</p></caption><table id="table3" frame="hsides" rules="groups"><thead><tr><td align="left" valign="bottom">Themes</td><td align="left" valign="bottom">Sample quotes from participants</td></tr></thead><tbody><tr><td align="left" valign="top">Fragmented and inconsistent health education content</td><td align="left" valign="top">&#x201C;I listened very carefully when the nurse was giving the health education, but afterward, I could only rely on the printed leaflet they gave me to recall the information...I feel like I didn&#x2019;t remember a lot of it.&#x201D; [P1]</td></tr><tr><td align="left" valign="top">Limited accessibility and personalization of education delivery</td><td align="left" valign="top">&#x201C;The booklet they gave me was the same for everyone. Some parts didn&#x2019;t apply to me at all, while the things I really cared about were not explained in detail.&#x201D; [P4]</td></tr><tr><td align="left" valign="top">Unmet needs for timely and interactive support</td><td align="left" valign="top">&#x201C;I was very anxious before the surgery...Although the preoperative education helped alleviate some of my anxiety and answered some of my questions, I felt it wasn&#x2019;t enough. However, I was also worried about bothering the doctors and nurses too much.&#x201D; [P3]</td></tr><tr><td align="left" valign="top">Expectations and concerns for new, technology-supported education manners</td><td align="left" valign="top">&#x201C;We&#x2019;re living in the information age now, and I hope there can be more intelligent methods and approaches for delivering health education to me.&#x201D; [P6]<break/>&#x201C;AI has indeed developed rapidly and is widely applied now, but we must also use it with caution. It should only be introduced to patients after confirming that it is reliable, practical, and provides accurate responses.&#x201D; [S8]</td></tr></tbody></table></table-wrap></sec></sec><sec id="s3-2"><title>Phase 2: Act and Implement</title><sec id="s3-2-1"><title>Evidence-Based Knowledge Corpus Development</title><p>As a prerequisite for developing the conversational agent, a corpus was established to provide a reliable and clinically validated knowledge foundation for response generation. Based on the results of phase 1, the action research team first identified 9 core knowledge domains for gastric cancer health education: disease knowledge, diagnosis, prevention, treatment, nutrition, nursing care, recovery, follow-up, and hospitalization procedures. Guided by these domains, 2 researchers (YCK and YMY) searched data from structured and unstructured resources, including clinical practice guidelines, scientific literature, hospital patient education materials, clinician-compiled frequently asked questions, and other authoritative resources. Detailed search strategies and data sources are provided in Table S5 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>. After removing duplicates, advanced nurse specialists and surgeons of our gastric cancer center screened eligible evidence and transformed the curated knowledge into structured question and answer (Q&#x0026;A) pairs. The Q&#x0026;A knowledge corpus would be integrated into the conversational agent as its knowledge base to support retrieval-augmented response generation. This corpus will be continuously maintained by chief physicians and nurses from the gastric cancer center, with per-week updates to incorporate the latest clinical evidence and practice guidelines.</p></sec><sec id="s3-2-2"><title>Conversational Agent Development and Functions</title><p>This conversational agent constructed a question-answering service system based on the OpenMEDLab2.0 foundation model and the RAG architecture (<xref ref-type="fig" rid="figure2">Figure 2</xref>). As a domain-specific LLM, the OpenMEDLab2.0 foundation model was developed based on InternLM and achieved professional capability through a 3-stage progressive medical domain&#x2013;training pipeline: (1) continual pretraining on medical literature, clinical guidelines, and diagnostic cases processed by third-generation data-cleaning technology, using masked language modeling and sentence order prediction tasks with BF16 mixed-precision training; (2) supervised fine-tuning on over 300,000 medical instruction datasets via low-rank adaptation for parameter-efficient training; and (3) 3 rounds of online reinforcement learning from human feedback (RLHF) to optimize response safety and clinical relevance. This training process enables the model to better capture medical semantic relationships and represent domain-specific knowledge, thereby improving its ability to interpret complex medical queries and generate structured responses [<xref ref-type="bibr" rid="ref31">31</xref>].</p><fig position="float" id="figure2"><label>Figure 2.</label><caption><p>A detailed operation illustration example of the conversational agent based on the OpenMEDLab2.0 foundation model and the Retrieval-Augmented Generation architecture. LLM: large language model; Q&#x0026;A: question and answer.</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="jmir_v28i1e82746_fig02.png"/></fig><p>For this study, the RAG architecture was specifically tailored to integrate with the OpenMEDLab2.0 foundation model. The RAG knowledge base was constructed by the evidence-based knowledge corpus developed in the last procedure. The retrieved information was injected into the input layer of the OpenMEDLab2.0 model as structured context. Specifically, the retrieved medical knowledge was first converted into vector data and then subjected to relevance matching via a rerank fine-ranking model to obtain the top 3 most relevant items. The most relevant knowledge was input into the OpenMEDLab2.0 model as a prompt to generate the final output. This method helps reduce the &#x201C;hallucination&#x201D; issues caused by knowledge limitations in pretrained large models and significantly improves the timeliness and accuracy of the model&#x2019;s responses, as the knowledge base mounted by RAG is synchronously updated with the latest clinical guidelines and expert consensus.</p><p>The workflow of the conversational agent begins with a patient&#x2019;s inquiry and consists of four steps:</p><list list-type="order"><list-item><p>Step 1: When a patient submits a query, the system first corrects any errors in speech recognition through hotword revision to convert misrecognized parts into accurate content. Subsequently, the user query is processed and analyzed by the safety guardrails module, which is specifically designed for sensitive content filtering. This module detects inappropriate information, including pornography, gambling, drugs, violence, and other high-risk content in patients&#x2019; questions. If any prohibited sensitive terms are identified, the system will directly refuse to generate a response. In addition, the same module will also conduct a safety review on the output answers generated by the large model; any response containing sensitive or risky content will be blocked accordingly. Only queries and corresponding outputs that fully pass the dual safety checks are allowed to enter the subsequent processing steps. If the patient&#x2019;s question-answer interaction does not comply with safety standards, the system will refuse to respond; queries that pass the safety check will proceed to the next step.</p></list-item><list-item><p>Step 2: The retrieval module converts the identified core needs into vector representations using an embedding model. The embedding model maps textual information into a low-dimensional vector space, where semantically similar texts are positioned close to each other. The vector database contains medical knowledge curated by professional health care staff and precisely broken down into rich question-answer pairs, each of which is converted into a vector and stored in the database. First, all types of medical knowledge resources, including structured medical data, unstructured clinical documents, and medical education videos, were systematically organized and converted into a standardized Q&#x0026;A pair format. For nontextual resources such as medical videos, key medical information was first extracted through manual annotation by clinical professionals and intelligent content recognition technology and then transformed into Q&#x0026;A pairs that conform to clinical consultation scenarios. Subsequently, question augmentation was performed on the constructed Q&#x0026;A pairs to improve the generalization ability of knowledge retrieval; specifically, techniques such as synonym replacement, sentence structure transformation, and clinical scenario extension were adopted to generate multiple extended question versions for each original question, ensuring that the knowledge base can effectively match various expression styles of patients&#x2019; queries. Finally, all the standardized Q&#x0026;A pairs (including original and extended ones) were uniformly subjected to vectorization processing using an encoder. During the vectorization process, the medical semantic features of the Q&#x0026;A pairs were fully extracted and converted into dense vectors, which were then stored in the vector database to support efficient dynamic retrieval of the RAG architecture. During retrieval, the system calculates the similarity between the patient&#x2019;s query vector and the vectors of the question-answer pairs in the database to efficiently identify relevant candidates.</p></list-item><list-item><p>Step 3: The candidate question-answer pairs then enter the reranking module. This module evaluates and ranks them based on their relevance to the patient&#x2019;s query using machine learning algorithms. The top 3 most relevant question-answer pairs are selected as input for the LLM.</p></list-item><list-item><p>Step 4: After retrieving the relevant information from the knowledge base, the LLM generates the final response. The reranked question-answer pairs are fed into the model according to a preconfigured prompt structure. This structured prompting, along with the top 3 matched answers, guides the model to generate responses that are accurate, professional, and aligned with ethical standards, which are then delivered to the patient.</p></list-item></list><p>The conversational agent was deployed on the intranet of Zhongshan Hospital with a Client/Server (C/S) architecture to ensure medical data security. The service is provided by a dedicated system: the client terminal includes a transparent LCD screen (for interaction) and a computer with an NVIDIA GeForce RTX 3090 graphics card (exclusively for digital human rendering), while the real-time inference service of the RAG-based large model is provided by a virtual machine on a computing power server equipped with Huawei Ascend 910B chips. Currently, it is available only in the hospital&#x2019;s gastric cancer center and not deployed online or on mobile devices. Functions are mainly divided into 2 modes: frequently asked questions and free questioning. Patients can click on frequently asked questions to view preset questions and answers, or choose the free question mode to ask questions. When the virtual person responds, its voice and gestures will mimic those of a real person as much as possible with automatic speech recognition, text-to-speech, and lip-sync technology, thereby providing patients with a more realistic Q&#x0026;A experience. Patients can adjust the volume by clicking the button on the right side of the screen.</p></sec><sec id="s3-2-3"><title>RAG Core Module Development</title><p>The RAG system in our study was developed around a vectorized knowledge base and 2 core components, namely, the retriever and the generator [<xref ref-type="bibr" rid="ref14">14</xref>]. The construction of a vector knowledge base involves transforming preprocessed text fragments and question-answer pairs into high-dimensional vectors, selecting a vector database that supports efficient semantic retrieval, importing vector data, establishing mapping relationships between vectors and question-answer pairs, configuring indexes to optimize retrieval speed, and adapting to real-time consulting needs.</p><p>The retriever module was developed using a retrieval strategy that combines sparse retrieval with a reranking model, with a primary focus on delivering the top-3 most accurate results. First, sparse retrieval based on keyword matching was used to locate core document segments and quickly generate a candidate set of 50&#x2010;100 items. Next, a reranking model adapted to the medical context was developed to perform deep ranking of the candidate results based on dimensions such as medical semantic relevance and clinical priority, precisely selecting the top-3 most relevant knowledge segments or question-answer pairs to ensure accurate alignment with core clinical knowledge. Finally, the top-3 results underwent a final validation step to filter duplicates and content exceeding specialty-specific safety boundaries, ensuring the uniqueness and safety of the output.</p><p>In the generator module, the retrieved question-answer pairs were used as contextual input to the OpenMEDLab2.0 foundation model, with explicit enforcement of medical ethics and specialty-specific safety boundaries. For complex queries, multiple knowledge segments were integrated to generate coherent responses, ensuring that the language remained both accessible and professionally accurate. The system also supported multiturn conversational memory, allowing historical consultation content to be referenced in order to optimize subsequent responses.</p></sec><sec id="s3-2-4"><title>Model Fine-Tuning and Safety Guardrail</title><p>To ensure the safety of LLMs in our conversational agent, the system uses RLHF for model fine-tuning. RLHF enables the model&#x2019;s behavior and outputs to better align with ethical norms and social values. Furthermore, through safety guardrail, incorporating interception, rewriting, and proxy-answering functionalities, the system identifies and blocks potential malicious prompts or erroneous outputs. This ensures that the system can respond reasonably when encountering latent ethical risks.</p></sec><sec id="s3-2-5"><title>Nurses as the avatar of the conversational agent</title><p>The chatbot featured a vividly embodied image&#x2014;an approachable and empathetic digital avatar designed to resemble a compassionate nurse (<xref ref-type="fig" rid="figure3">Figure 3</xref>). This design choice is grounded in both emotional appeal and clinical relevance [<xref ref-type="bibr" rid="ref32">32</xref>]. By adopting the image of a nurse, the chatbot conveys a sense of warmth, safety, and reliability that patients can easily relate to. Higher acceptance and increased adherence to treatment regimens were also found in some studies [<xref ref-type="bibr" rid="ref33">33</xref>,<xref ref-type="bibr" rid="ref34">34</xref>].</p><fig position="float" id="figure3"><label>Figure 3.</label><caption><p>Interface of the health education conversational agent including the home page and chatting screen.</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="jmir_v28i1e82746_fig03.png"/></fig><p>Both verbal and nonverbal behaviors were designed for the conversational agent. Verbal behaviors mainly include some specific language constructs such as politeness strategies, inclusive pronouns, and greeting and farewell rituals. Nonverbal behaviors mainly include pleasant facial expressions, lip movements, head nods, and hand gestures synchronized with speech. Such visual elements were proved to be beneficial for patients in establishing rapport and engagement with the chatbot [<xref ref-type="bibr" rid="ref33">33</xref>].</p></sec></sec><sec id="s3-3"><title>Phase 3: Evaluate</title><sec id="s3-3-1"><title>Alpha Testing</title><p>The accuracy of the agent in 3 rounds was 67% (37/55), 71% (44/62), and 82% (31/38), respectively; RAG knowledge hit rate reached 86% (47/55), 98% (61/62), and 100% (38/38). Significant differences (<italic>P</italic>&#x003C;.01) in RAG knowledge hit rate were observed across 3 rounds of testing, while no significant differences (<italic>P</italic>=.30) were found in accuracy. Overall interrater agreement using the Fleiss &#x03BA; (0.85; <italic>P</italic>&#x003C;.01) was revealed. RAG knowledge hit rate was automatically evaluated by the system, and interrate reliability was not applicable.</p></sec><sec id="s3-3-2"><title>Beta Testing (Usability Testing)</title><p>Twenty-eight participants were included in the chatbot usability testing. Characteristics of 28 patients are shown in <xref ref-type="table" rid="table4">Table 4</xref>.</p><table-wrap id="t4" position="float"><label>Table 4.</label><caption><p>Characteristics of patients who used the conversational agent and finished the chatbot usability questionnaire (N=28).</p></caption><table id="table4" frame="hsides" rules="groups"><thead><tr><td align="left" valign="bottom">Characteristics</td><td align="left" valign="bottom">Patients</td></tr></thead><tbody><tr><td align="left" valign="top">Sex, n (%)</td><td align="left" valign="top"/></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Male</td><td align="left" valign="top">12 (43)</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Female</td><td align="left" valign="top">16 (57)</td></tr><tr><td align="left" valign="top">Age (years), n (%)</td><td align="left" valign="top"/></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>&#x003C;40</td><td align="left" valign="top">2 (7&#xFF09;</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>40&#x2010;55</td><td align="left" valign="top">5 (18&#xFF09;</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>55&#x2010;70</td><td align="left" valign="top">8 (29&#xFF09;</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>&#x003E;70</td><td align="left" valign="top">13 (46&#xFF09;</td></tr><tr><td align="left" valign="top">Hospitalization, days, n (%)</td><td align="left" valign="top"/></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>&#x003C;3</td><td align="left" valign="top">3 (11&#xFF09;</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>3&#x2010;10</td><td align="left" valign="top">20 (71&#xFF09;</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>&#x003E;10</td><td align="left" valign="top">5 (18&#xFF09;</td></tr><tr><td align="left" valign="top">Marital status, n (%)</td><td align="left" valign="top"/></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Married</td><td align="left" valign="top">22 (79&#xFF09;</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Single</td><td align="left" valign="top">2 (7&#xFF09;</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Widowed</td><td align="left" valign="top">0 (0&#xFF09;</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Divorced</td><td align="left" valign="top">0 (0&#xFF09;</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Others</td><td align="left" valign="top">4 (14&#xFF09;</td></tr><tr><td align="left" valign="top">Medical insurance, n (%)</td><td align="left" valign="top"/></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Yes</td><td align="left" valign="top">26 (93&#xFF09;</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>No</td><td align="left" valign="top">2 (7&#xFF09;</td></tr><tr><td align="left" valign="top">Cancer disease, n (%)</td><td align="left" valign="top"/></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Yes</td><td align="left" valign="top">25 (89)</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>No</td><td align="left" valign="top">3 (11)</td></tr></tbody></table></table-wrap><p>The mean CUQ score was 91.9 (SD 3.6), with a median score of 92.2 (range 79.7&#x2010;95.3) (<xref ref-type="fig" rid="figure4">Figure 4</xref>). Among the odd questions, statements &#x201C;The chatbot was easy to navigate&#x201D; and &#x201C;The chatbot was very easy to use&#x201D; received the highest mean scores which proved that this conversational agent was easy to use. The statement &#x201C;The chatbot understood me well&#x201D; got the lowest mean scores and showed that some understanding problems still existed in this system. Among the even questions, the statements &#x201C;The chatbot seemed very unfriendly,&#x201D; &#x201C;It would be easy to get confused when using the chatbot,&#x201D; and &#x201C;The chatbot was very complex&#x201D; got the lowest mean scores, while &#x201C;The chatbot failed to recognize a lot of my inputs&#x201D; got the highest mean scores. These indicated that input and output of this conversational agent still had some problems that may confuse participants. Each score of CUQ is detailed in Table S6 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>. Besides, the mean content relevance score was 3.75 (SD 0.9), with a median score of 4 (range 3&#x2010;4).</p><fig position="float" id="figure4"><label>Figure 4.</label><caption><p>Patients&#x2019; choices and mean scores on each item of the chatbot usability questionnaire.</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="jmir_v28i1e82746_fig04.png"/></fig></sec></sec><sec id="s3-4"><title>Phase 4: Reflect</title><p>We recruited 15 patients (T1-T15) in the last phase who finished the CUQ and 8 nurses to do the in-depth interview to deeply explore the factors facilitating and hindering the conversational agent. Detailed characteristics of 8 clinical nurses are shown in <xref ref-type="table" rid="table5">Table 5</xref>.</p><p>Emerging themes from the qualitative study results were mapped on the 3 domains of the technology acceptance model (<xref ref-type="table" rid="table6">Table 6</xref>). Detailed information about the qualitative analysis is shown in Text 2 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>.</p><table-wrap id="t5" position="float"><label>Table 5.</label><caption><p>Characteristics of clinical nurses who worked in the gastrointestinal center and participated in daily health education activity (N=8).</p></caption><table id="table5" frame="hsides" rules="groups"><thead><tr><td align="left" valign="bottom">Clinical nurses, sex</td><td align="left" valign="bottom">Age, years</td><td align="left" valign="bottom">Career years</td><td align="left" valign="bottom">Education level</td></tr></thead><tbody><tr><td align="left" valign="top">N1: female</td><td align="left" valign="top">27</td><td align="left" valign="top">5</td><td align="left" valign="top">Bachelor&#x2019;s degree</td></tr><tr><td align="left" valign="top">N2: female</td><td align="left" valign="top">30</td><td align="left" valign="top">8</td><td align="left" valign="top">Bachelor&#x2019;s degree</td></tr><tr><td align="left" valign="top">N3: male</td><td align="left" valign="top">34</td><td align="left" valign="top">11</td><td align="left" valign="top">Master&#x2019;s degree</td></tr><tr><td align="left" valign="top">N4: female</td><td align="left" valign="top">26</td><td align="left" valign="top">4</td><td align="left" valign="top">Bachelor&#x2019;s degree</td></tr><tr><td align="left" valign="top">N5: male</td><td align="left" valign="top">32</td><td align="left" valign="top">10</td><td align="left" valign="top">Bachelor&#x2019;s degree</td></tr><tr><td align="left" valign="top">N6: female</td><td align="left" valign="top">29</td><td align="left" valign="top">7</td><td align="left" valign="top">Bachelor&#x2019;s degree</td></tr><tr><td align="left" valign="top">N7: female</td><td align="left" valign="top">35</td><td align="left" valign="top">13</td><td align="left" valign="top">Bachelor&#x2019;s degree</td></tr><tr><td align="left" valign="top">N8: male</td><td align="left" valign="top">31</td><td align="left" valign="top">9</td><td align="left" valign="top">Bachelor&#x2019;s degree</td></tr></tbody></table></table-wrap><table-wrap id="t6" position="float"><label>Table 6.</label><caption><p>Themes identified to explore the factors facilitating and hindering the use of conversational agent.</p></caption><table id="table6" frame="hsides" rules="groups"><thead><tr><td align="left" valign="bottom">Theme and subthemes</td><td align="left" valign="bottom">Representative quotes</td></tr></thead><tbody><tr><td align="left" valign="top">Perceived Usefulness</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>&#x201C;It was like having a knowledgeable companion by my side who could answer my questions at any time.&#x201D; [T3]</p></list-item><list-item><p>&#x201C;Sometimes I asked a very simple question, but the answer is too long and the language is quite academic, which in turn made me feel burdened.&#x201D; [T7]</p></list-item><list-item><p>&#x201C;We may not always be able to provide patients with comprehensive educational content, but the conversational agent helps us to supplement this and make up for any omissions in our work.&#x201D; [N2]</p></list-item></list></td></tr><tr><td align="left" valign="top">Perceived Ease of Use</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>&#x201C;I thought it would be difficult to use, but after the nurse showed me how, I learned it immediately. It wasn&#x2019;t difficult at all.&#x201D; [T10]</p></list-item><list-item><p>&#x201C;Do I have an accent? So when I use it, it always misunderstood what I&#x2019;m saying, which seems to affect its usability a little.&#x201D; [T8]</p></list-item><list-item><p>&#x201C;This conversational agent is indeed very convenient to use, and we don&#x2019;t really need to teach patients how to use it.&#x201D; [N5]</p></list-item></list></td></tr><tr><td align="left" valign="top">Intention to Use</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>&#x201C;It feels like I&#x2019;m really talking to a nurse, not just a machine...She looks kind and makes me feel more at ease when asking questions.&#x201D; [T9]</p></list-item><list-item><p>&#x201C;I feel that this machine is really easy to use. It would be great if you could develop a mobile app so that I can ask questions when I get home if I don&#x2019;t understand something.&#x201D; [T6]</p></list-item><list-item><p>&#x201C;This conversational agent is very useful, but we still need to improve it to make it more mature and reliable.&#x201D; [N1]</p></list-item></list></td></tr></tbody></table></table-wrap></sec></sec><sec id="s4" sec-type="discussion"><title>Discussion</title><sec id="s4-1"><title>Principal Findings</title><p>This study demonstrates how the action research framework can be a practical methodology to develop a health education conversational agent, comprising the diagnose and plan, act and implement, evaluate, and reflect phases. These 4 phases formed a continuous and iterative process in which the findings from each phase informed the subsequent stage. This user-centered approach enabled the systematic integration of patient needs, clinical expertise, technical implementation, and iterative evaluation, and offered an approach for developing digital health interventions that are both clinically relevant and acceptable to end users. By placing the diagnose and plan phase at the beginning of the process, researchers were able to identify real clinical requirements rather than technological assumptions and avoid certain barriers to adoption at the earliest stage. The act and implement phase realized outcomes that the conversational agent not only incorporated advanced technologies but also remained clinically practical and aligned with routine health care delivery by involving multidisciplinary stakeholders. The successful implementation of technologies such as RAG also relies on close collaboration between medical and informatics specialists. The evaluate and reflect phases constituted essential components of the iterative optimization process. Integration of quantitative and qualitative findings enabled developers to identify specific directions for system refinement and provided a foundation for subsequent iterative development. Overall, the action research&#x2013;based development pathway in our study may offer opportunities for improving the sustained adoption of digital health interventions in real-world clinical settings through iterative, user-centered, multidisciplinary, and mixed methods principles.</p></sec><sec id="s4-2"><title>Clinical Implications</title><p>Health education is a cornerstone of therapy and recovery for patients with physical and mental conditions [<xref ref-type="bibr" rid="ref35">35</xref>]. Conversational agents for health education are necessary because there may be a discrepancy between patients&#x2019; actual knowledge needs and what clinical practitioners consider necessary [<xref ref-type="bibr" rid="ref36">36</xref>-<xref ref-type="bibr" rid="ref38">38</xref>]. For patients, this digital human serves as a convenient and timely resource tool, providing accurate and personalized information, particularly when clinical staff are unavailable. <italic>Lancet Digital Health</italic> once mentioned the urgent concern regarding whether existing tools and digital health interventions are appropriate for the aging population [<xref ref-type="bibr" rid="ref39">39</xref>]. In our study, we emphasized this problem and found that the simplicity and convenience of this chatbot bridged this gap based on the results of phase 3 and phase 4. In usability questionnaires and in-depth interviews, almost no patients mentioned that the system was difficult to use.</p><p>The design of the nurse avatar and nonverbal behavior received many praises which proved that embodied representation did have significant impacts on patient satisfaction and health outcomes [<xref ref-type="bibr" rid="ref40">40</xref>]. This was explained in the research by Bickmore et al [<xref ref-type="bibr" rid="ref40">40</xref>] that human relationships are fundamentally formed in face-to-face conversation and nonverbal behaviors are vital for facilitating social connection and understanding. However, the input and output modalities may require further optimization. Contrary to our expectations, some end users experienced confusion and discomfort when using the voice input to initiate interactions. Patients with limited verbal expression skills or those speaking regional dialects faced difficulties in recognizing their questions accurately. In addition, the chatbot currently provides responses primarily in text format. Incorporating multimedia content, such as images and videos, may further enhance users&#x2019; comprehension and reinforce health education.</p><p>Health education content should be dynamically adapted to individual patients&#x2019; health literacy levels to maximize relevance and effectiveness. Personalization of our conversational agent remains limited. Some conversational agents integrate patients&#x2019; appointment schedules and personal health information to deliver more personalized health services. However, the incorporation of private data inevitably amplifies concerns regarding privacy and data protection. Fournier-Tombs and McHardy [<xref ref-type="bibr" rid="ref41">41</xref>] listed a list of potential ethical risks in conversational chatbots, including discrimination, stereotyping, exclusion, lack of privacy, poor data governance, stigma, error tolerance, and overconfidence. Specifically, our team set safety guardrails to remind patients that clinicians&#x2019; judgments should be regarded as the final suggestion in medical decision-making.</p><p>In our study, stakeholders were also involved among the 4 phases to provide diverse perspectives on the implementation of our conversational agent in clinical practice. Through quantitative and qualitative research, we found that for medical professionals, especially frontline staff, integrating a domain-specific conversational agent into the gastric cancer clinical workflow amidst growing clinical workloads can streamline patient-clinician communication. By alleviating such burdens, clinicians were allowed to concentrate on more complex tasks, such as observing patient conditions. Moreover, clinical experts also expressed interest in leveraging conversational agents for complex tasks such as data analysis. These findings represent important directions for future development in clinical practice.</p><p>Advancing health equity worldwide is increasingly drawing attention in recent years [<xref ref-type="bibr" rid="ref42">42</xref>]. Although medical chatbots are gaining popularity, there is still a dearth of chatbots specific to gastric cancer health education in resource-constrained regions [<xref ref-type="bibr" rid="ref43">43</xref>]. Researchers and governments have a responsibility to design and implement fair, unbiased, and easily accessible tools that prioritize the needs of patients [<xref ref-type="bibr" rid="ref44">44</xref>], and our conversational agent provided a promising solution. Once the health education conversational agent in our study can be put into formal use after effectiveness validation, shared professional knowledge bases will be available for regions&#x2019; lack of medical resources, thereby improving the health literacy of local patients, promoting health equity, and eliminating health disparities [<xref ref-type="bibr" rid="ref45">45</xref>].</p></sec><sec id="s4-3"><title>Comparison With Prior Work</title><p>Enthusiasm for conversational agents as means for health education was revealed among patients and clinical staff in the first phase of our study, which was also proved in other research [<xref ref-type="bibr" rid="ref46">46</xref>]. Multiple research studies used theory frameworks to develop such digital interventions. For example, a design science research methodology was used to develop a conversational agent for enhanced self-management after cardiothoracic surgery and verified its acceptability [<xref ref-type="bibr" rid="ref47">47</xref>]. Similarly, our study adopted an action research approach by realizing the user-centered design and usability principles.</p><p>Research gaps remain in gastric cancer&#x2013;specific chatbots although numerous mature conversational agents have been developed in the health care area. Gomaa et al [<xref ref-type="bibr" rid="ref9">9</xref>] developed a self-management program in patients with gastrointestinal cancer undergoing chemotherapy and used a mixed methods design to evaluate its effectiveness. However, this platform relied on a structured approach instead of NLP, which limited its understanding of nuanced patient contexts. Zhou et al [<xref ref-type="bibr" rid="ref15">15</xref>] developed a gastrointestinal disease chatbot using RAG and evaluated its performance with the Retrieval Augmented Generation Assessment framework, which provided an important technical foundation for our study. In contrast to the study by Zhou et al, we enriched our research perspectives by adding qualitative results. Similarly, Kim and Park [<xref ref-type="bibr" rid="ref48">48</xref>] developed a chatbot for patients with gastric cancer undergoing radical gastrectomy and evaluated both user experience and system performance. The accuracy of our conversational agent is comparable with this similar chatbot. An embodied conversational agent may potentially improve engagement by providing additional motivational and emotional support [<xref ref-type="bibr" rid="ref49">49</xref>], while also providing a foundation for multimodal interaction [<xref ref-type="bibr" rid="ref50">50</xref>]. Such design elements remain uncommon in existing health-related conversational agents, especially those targeting gastric cancer care. Besides, an evidence-based corpus, RAG technology, model fine-tuning, and safety guardrail mechanisms were implemented in our study to reduce the risk of generating misleading or unsafe information and improve the reliability of the chatbot&#x2019;s outputs.</p></sec><sec id="s4-4"><title>Limitations and Future Directions</title><p>Although the study&#x2019;s strengths lie in its novel combination of digital health intervention development and action research approach, leveraging both quantitative and qualitative data, some limitations must be acknowledged. This study reported only a single cycle of the action research framework. Additional cycles are necessary to iteratively refine the system and evaluate its effectiveness in real-world settings. Although prior studies suggested that subjective ratings and expert evaluations can provide valuable insights into revealing potential risks in real-world use [<xref ref-type="bibr" rid="ref51">51</xref>], accuracy rate, RAG knowledge hit rate, and usability score were used to test the performance of the chatbot, which was not enough. The complexity of health care contexts and variability in expert judgment may produce bias and inconsistency [<xref ref-type="bibr" rid="ref52">52</xref>]. Besides, this study was conducted with a relatively small sample from a specific center and cultural context. Stakeholder perspectives were collected during a limited period, which may not fully capture evolving needs and long-term experience over time. In this study, we focused more on the development and implementation of the conversational agent rather than its clinical effectiveness; the impact of the intervention on patient outcomes and other medical facilities remains unknown without multicenter randomized controlled trials.</p><p>Last but not least, algorithmic limitations need further attention. The conversational agent was developed using OpenMEDLab2.0 [<xref ref-type="bibr" rid="ref31">31</xref>], whose technical description is currently available primarily as a preprint rather than a peer-reviewed publication. Although the model was selected because of its medical domain-specific design, future studies should further evaluate and compare its performance with other peer-reviewed medical foundation models across diverse clinical tasks and settings. Furthermore, the current workflow requires enhancement to support functional extensions, such as follow-up questionnaire delivery. Moreover, the system primarily retrieves answers from a predefined knowledge base without incorporating expert clinical reasoning, highlighting the need for further development of chain-of-thought capabilities.</p></sec><sec id="s4-5"><title>Conclusions</title><p>This study provided insights into how the action research approach can inform the development and usability assessment of a gastric cancer health education conversational agent. It also illustrated details of how technologies such as RAG, model fine-tuning, and safety guardrail technologies can be integrated into the conversational agent. Promising results regarding accuracy and usability were demonstrated through both qualitative and quantitative research. However, additional assessments and improvements are warranted to confirm the effectiveness, safety, and long-term adoption. Enhancements such as personalization, multimodal data integration, and advanced reasoning capabilities should also be explored in further studies.</p></sec></sec></body><back><ack><p>The authors thank all patients who participated in this study for their valuable contributions and for their involvement in the design of the conversational agent. They also thank the surgeons, advanced nurse specialists, and other clinical staff of the Gastric Cancer Center for their valuable contributions to the development of the conversational agent and for sharing their professional expertise throughout the study. The authors declare the use of generative AI (GenAI) in the research and writing process. According to the GAIDeT taxonomy (2025), the following tasks were delegated to GenAI tools under full human supervision: translation. The GenAI tool used was ChatGPT-4.5. Responsibility for the final manuscript lies entirely with the authors. GenAI tools are not listed as authors and do not bear responsibility for the final outcomes. The authors used the GenAI tool ChatGPT-4.5 by OpenAI to translate part of their manuscript. The prompt they used is as follows: "help me to translate this word/sentence into fluent, native-level academic English while preserving the original meaning."</p></ack><notes><sec><title>Funding</title><p>This research was supported by Shanghai Health Science Popularization Special Program (JKKPYL-2025-A01) and Fosun Fund (FNF202513). The funders had no role in the study design; collection, analysis, and interpretation of data; writing of the paper; and/or decision to submit for publication.</p></sec><sec><title>Data Availability</title><p>The datasets generated or analyzed during this study are available from the corresponding author on reasonable request.</p></sec></notes><fn-group><fn fn-type="con"><p>YiChen KANG and YaMin YAN have contributed equally to this work and share first authorship.</p><p>Conceptualization: YXZ, YCK, YMY</p><p>Data curation: YCK, YMY, TXW</p><p>Formal analysis: YCK, YMY, JML, YXZ</p><p>Funding acquisition: YXZ</p><p>Methodology: YXZ, YCK, YMY</p><p>Project administration: YXZ, YCK, YMY</p><p>Resources: ZHY, YH</p><p>Software: TXW, ZXW, JYZ</p><p>Supervision: YXZ, JML</p><p>Validation: YCK, YMY, YXZ, JML</p><p>Visualization: YMY, TXW</p><p>Writing &#x2013; original draft: YCK, YMY</p><p>Writing &#x2013; review and editing: TXW, ZHY, YH, ZXW, JYZ, JML, YXZ.</p></fn><fn fn-type="conflict"><p>None declared.</p></fn></fn-group><glossary><title>Abbreviations</title><def-list><def-item><term id="abb1">C/S</term><def><p>Client/Server</p></def></def-item><def-item><term id="abb2">CUQ</term><def><p>chatbot usability questionnaire</p></def></def-item><def-item><term id="abb3">LLM</term><def><p>large language model</p></def></def-item><def-item><term id="abb4">NLP</term><def><p>natural language processing</p></def></def-item><def-item><term id="abb5">Q&#x0026;A</term><def><p>question and answer</p></def></def-item><def-item><term id="abb6">RAG</term><def><p>retrieval augmented generation</p></def></def-item><def-item><term id="abb7">RLHF</term><def><p>reinforcement learning from human feedback</p></def></def-item></def-list></glossary><ref-list><title>References</title><ref id="ref1"><label>1</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Bray</surname><given-names>F</given-names> </name><name name-style="western"><surname>Laversanne</surname><given-names>M</given-names> </name><name name-style="western"><surname>Sung</surname><given-names>H</given-names> </name><etal/></person-group><article-title>Global cancer statistics 2022: GLOBOCAN estimates of incidence and mortality worldwide for 36 cancers in 185 countries</article-title><source>CA Cancer J Clin</source><year>2024</year><volume>74</volume><issue>3</issue><fpage>229</fpage><lpage>263</lpage><pub-id pub-id-type="doi">10.3322/caac.21834</pub-id><pub-id pub-id-type="medline">38572751</pub-id></nlm-citation></ref><ref id="ref2"><label>2</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Vallance</surname><given-names>PC</given-names> </name><name name-style="western"><surname>Mack</surname><given-names>L</given-names> </name><name name-style="western"><surname>Bouchard-Fortier</surname><given-names>A</given-names> </name><name name-style="western"><surname>Jost</surname><given-names>E</given-names> </name></person-group><article-title>Quality of life following the surgical management of gastric cancer using patient-reported outcomes: a systematic review</article-title><source>Curr Oncol</source><year>2024</year><month>02</month><day>4</day><volume>31</volume><issue>2</issue><fpage>872</fpage><lpage>884</lpage><pub-id pub-id-type="doi">10.3390/curroncol31020065</pub-id><pub-id pub-id-type="medline">38392059</pub-id></nlm-citation></ref><ref id="ref3"><label>3</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Zhao</surname><given-names>G</given-names> </name><name name-style="western"><surname>Zhang</surname><given-names>Y</given-names> </name><name name-style="western"><surname>Liu</surname><given-names>C</given-names> </name></person-group><article-title>The effect of health education on the quality of life of postoperative patients with gastric cancer: a systematic review and meta-analysis</article-title><source>Ann Palliat Med</source><year>2021</year><month>10</month><volume>10</volume><issue>10</issue><fpage>10633</fpage><lpage>10642</lpage><pub-id pub-id-type="doi">10.21037/apm-21-2420</pub-id></nlm-citation></ref><ref id="ref4"><label>4</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Kim</surname><given-names>D</given-names> </name><name name-style="western"><surname>Lee</surname><given-names>SS</given-names> </name></person-group><article-title>Fewer feedback opportunities and health perception of gastric cancer survivors: opportunities for patient education</article-title><source>J Canc Educ</source><year>2024</year><month>08</month><volume>39</volume><issue>4</issue><fpage>455</fpage><lpage>463</lpage><pub-id pub-id-type="doi">10.1007/s13187-024-02430-z</pub-id></nlm-citation></ref><ref id="ref5"><label>5</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Jia</surname><given-names>X</given-names> </name><name name-style="western"><surname>Pang</surname><given-names>Y</given-names> </name><name name-style="western"><surname>Liu</surname><given-names>LS</given-names> </name></person-group><article-title>Online health information seeking behavior: A systematic review</article-title><source>Healthcare (Basel)</source><year>2021</year><volume>9</volume><issue>12</issue><fpage>1740</fpage><pub-id pub-id-type="doi">10.3390/healthcare9121740IF:3.4Q1B4</pub-id><pub-id pub-id-type="medline">34946466</pub-id></nlm-citation></ref><ref id="ref6"><label>6</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Laranjo</surname><given-names>L</given-names> </name><name name-style="western"><surname>Dunn</surname><given-names>AG</given-names> </name><name name-style="western"><surname>Tong</surname><given-names>HL</given-names> </name><etal/></person-group><article-title>Conversational agents in healthcare: a systematic review</article-title><source>J Am Med Inform Assoc</source><year>2018</year><month>09</month><day>1</day><volume>25</volume><issue>9</issue><fpage>1248</fpage><lpage>1258</lpage><pub-id pub-id-type="doi">10.1093/jamia/ocy072</pub-id><pub-id pub-id-type="medline">30010941</pub-id></nlm-citation></ref><ref id="ref7"><label>7</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Xing</surname><given-names>Z</given-names> </name><name name-style="western"><surname>Yu</surname><given-names>F</given-names> </name><name name-style="western"><surname>Qanir</surname><given-names>YAM</given-names> </name><name name-style="western"><surname>Guan</surname><given-names>T</given-names> </name><name name-style="western"><surname>Walker</surname><given-names>J</given-names> </name><name name-style="western"><surname>Song</surname><given-names>L</given-names> </name></person-group><article-title>Intelligent conversational agents in patient self-management: a systematic survey using multi data sources</article-title><source>Stud Health Technol Inform</source><year>2019</year><month>08</month><day>21</day><volume>264</volume><issue>1813-4</issue><fpage>1813</fpage><lpage>1814</lpage><pub-id pub-id-type="doi">10.3233/SHTI190661</pub-id><pub-id pub-id-type="medline">31438357</pub-id></nlm-citation></ref><ref id="ref8"><label>8</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Yang</surname><given-names>Q</given-names> </name><name name-style="western"><surname>Cheung</surname><given-names>K</given-names> </name><name name-style="western"><surname>Zhang</surname><given-names>Y</given-names> </name><name name-style="western"><surname>Zhang</surname><given-names>Y</given-names> </name><name name-style="western"><surname>Qin</surname><given-names>J</given-names> </name><name name-style="western"><surname>Xie</surname><given-names>YJ</given-names> </name></person-group><article-title>Conversational agents in physical and psychological symptom management: a systematic review of randomized controlled trials</article-title><source>Int J Nurs Stud</source><year>2025</year><month>03</month><volume>163</volume><fpage>104991</fpage><pub-id pub-id-type="doi">10.1016/j.ijnurstu.2024.104991</pub-id><pub-id pub-id-type="medline">39799832</pub-id></nlm-citation></ref><ref id="ref9"><label>9</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Gomaa</surname><given-names>S</given-names> </name><name name-style="western"><surname>Posey</surname><given-names>J</given-names> </name><name name-style="western"><surname>Bashir</surname><given-names>B</given-names> </name><etal/></person-group><article-title>Feasibility of a text messaging-integrated and chatbot-interfaced self-management program for symptom control in patients with gastrointestinal cancer undergoing chemotherapy: pilot mixed methods study</article-title><source>JMIR Form Res</source><year>2023</year><month>11</month><day>10</day><volume>7</volume><fpage>e46128</fpage><pub-id pub-id-type="doi">10.2196/46128</pub-id><pub-id pub-id-type="medline">37948108</pub-id></nlm-citation></ref><ref id="ref10"><label>10</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Naseri</surname><given-names>A</given-names> </name><name name-style="western"><surname>Antikchi</surname><given-names>MH</given-names> </name><name name-style="western"><surname>Barahman</surname><given-names>M</given-names> </name><etal/></person-group><article-title>AI chatbots in oncology: a comparative study of sider fusion AI and perplexity AI for gastric cancer patients</article-title><source>Indian J Surg Oncol</source><year>2025</year><month>08</month><volume>16</volume><issue>4</issue><fpage>827</fpage><lpage>836</lpage><pub-id pub-id-type="doi">10.1007/s13193-024-02145-z</pub-id><pub-id pub-id-type="medline">40949405</pub-id></nlm-citation></ref><ref id="ref11"><label>11</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Thirunavukarasu</surname><given-names>AJ</given-names> </name></person-group><article-title>Large language models will not replace healthcare professionals: curbing popular fears and hype</article-title><source>J R Soc Med</source><year>2023</year><month>05</month><volume>116</volume><issue>5</issue><fpage>181</fpage><lpage>182</lpage><pub-id pub-id-type="doi">10.1177/01410768231173123</pub-id><pub-id pub-id-type="medline">37199678</pub-id></nlm-citation></ref><ref id="ref12"><label>12</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Zhou</surname><given-names>J</given-names> </name><name name-style="western"><surname>Li</surname><given-names>H</given-names> </name><name name-style="western"><surname>Chen</surname><given-names>S</given-names> </name><name name-style="western"><surname>Chen</surname><given-names>Z</given-names> </name><name name-style="western"><surname>Han</surname><given-names>Z</given-names> </name><name name-style="western"><surname>Gao</surname><given-names>X</given-names> </name></person-group><article-title>Large language models in biomedicine and healthcare</article-title><source>NPJ Artif Intell</source><year>2025</year><volume>1</volume><issue>1</issue><fpage>44</fpage></nlm-citation></ref><ref id="ref13"><label>13</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Tudor Car</surname><given-names>L</given-names> </name><name name-style="western"><surname>Dhinagaran</surname><given-names>DA</given-names> </name><name name-style="western"><surname>Kyaw</surname><given-names>BM</given-names> </name><etal/></person-group><article-title>Conversational agents in health care: scoping review and conceptual analysis</article-title><source>J Med Internet Res</source><year>2020</year><month>08</month><day>7</day><volume>22</volume><issue>8</issue><fpage>e17158</fpage><pub-id pub-id-type="doi">10.2196/17158</pub-id><pub-id pub-id-type="medline">32763886</pub-id></nlm-citation></ref><ref id="ref14"><label>14</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Zhao</surname><given-names>P</given-names> </name><name name-style="western"><surname>Zhang</surname><given-names>H</given-names> </name><name name-style="western"><surname>Yu</surname><given-names>Q</given-names> </name><etal/></person-group><article-title>Retrieval-augmented generation for ai-generated content: a survey</article-title><source>Data Sci Eng</source><year>2026</year><month>03</month><volume>11</volume><issue>1</issue><fpage>1</fpage><lpage>29</lpage><pub-id pub-id-type="doi">10.1007/s41019-025-00335-5</pub-id></nlm-citation></ref><ref id="ref15"><label>15</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Zhou</surname><given-names>Q</given-names> </name><name name-style="western"><surname>Liu</surname><given-names>C</given-names> </name><name name-style="western"><surname>Duan</surname><given-names>Y</given-names> </name><etal/></person-group><article-title>GastroBot: a Chinese gastrointestinal disease chatbot based on the retrieval-augmented generation</article-title><source>Front Med (Lausanne)</source><year>2024</year><volume>11</volume><fpage>1392555</fpage><pub-id pub-id-type="doi">10.3389/fmed.2024.1392555</pub-id><pub-id pub-id-type="medline">38841582</pub-id></nlm-citation></ref><ref id="ref16"><label>16</label><nlm-citation citation-type="book"><person-group person-group-type="author"><name name-style="western"><surname>Williamson</surname><given-names>GR</given-names> </name><name name-style="western"><surname>Bellman</surname><given-names>L</given-names> </name><name name-style="western"><surname>Webster</surname><given-names>J</given-names> </name></person-group><source>Action Research in Nursing and Healthcare</source><year>2011</year><publisher-name>SAGE Publications</publisher-name><pub-id pub-id-type="doi">10.4135/9781446289112</pub-id><pub-id pub-id-type="other">9781446254295</pub-id></nlm-citation></ref><ref id="ref17"><label>17</label><nlm-citation citation-type="book"><person-group person-group-type="author"><name name-style="western"><surname>McNiff</surname><given-names>J</given-names> </name></person-group><source>Action Research: All You Need to Know</source><year>2017</year><publisher-name>SAGE Publications</publisher-name><pub-id pub-id-type="other">9781473967472</pub-id></nlm-citation></ref><ref id="ref18"><label>18</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Cordeiro</surname><given-names>L</given-names> </name><name name-style="western"><surname>Soares</surname><given-names>CB</given-names> </name></person-group><article-title>Action research in the healthcare field: a scoping review</article-title><source>JBI Database System Rev Implement Rep</source><year>2018</year><month>04</month><volume>16</volume><issue>4</issue><fpage>1003</fpage><lpage>1047</lpage><pub-id pub-id-type="doi">10.11124/JBISRIR-2016-003200</pub-id><pub-id pub-id-type="medline">29634517</pub-id></nlm-citation></ref><ref id="ref19"><label>19</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Bradbury</surname><given-names>H</given-names> </name><name name-style="western"><surname>Lifvergren</surname><given-names>S</given-names> </name></person-group><article-title>Action research healthcare: focus on patients, improve quality, drive down costs</article-title><source>Healthc Manage Forum</source><year>2016</year><month>11</month><volume>29</volume><issue>6</issue><fpage>269</fpage><lpage>274</lpage><pub-id pub-id-type="doi">10.1177/0840470416658905</pub-id><pub-id pub-id-type="medline">27770047</pub-id></nlm-citation></ref><ref id="ref20"><label>20</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Oberschmidt</surname><given-names>K</given-names> </name><name name-style="western"><surname>Gr&#x00FC;nloh</surname><given-names>C</given-names> </name><name name-style="western"><surname>Nijboer</surname><given-names>F</given-names> </name><name name-style="western"><surname>van Velsen</surname><given-names>L</given-names> </name></person-group><article-title>Best practices and lessons learned for action research in eHealth design and implementation: literature review</article-title><source>J Med Internet Res</source><year>2022</year><month>01</month><day>28</day><volume>24</volume><issue>1</issue><fpage>e31795</fpage><pub-id pub-id-type="doi">10.2196/31795</pub-id><pub-id pub-id-type="medline">35089158</pub-id></nlm-citation></ref><ref id="ref21"><label>21</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Verweij</surname><given-names>L</given-names> </name><name name-style="western"><surname>Metsemakers</surname><given-names>SJJPM</given-names> </name><name name-style="western"><surname>Ector</surname><given-names>GICG</given-names> </name><etal/></person-group><article-title>Improvement, implementation, and evaluation of the CMyLife digital care platform: participatory action research approach</article-title><source>J Med Internet Res</source><year>2023</year><month>09</month><day>15</day><volume>25</volume><fpage>e45259</fpage><pub-id pub-id-type="doi">10.2196/45259</pub-id><pub-id pub-id-type="medline">37713242</pub-id></nlm-citation></ref><ref id="ref22"><label>22</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Verweij</surname><given-names>L</given-names> </name><name name-style="western"><surname>Smit</surname><given-names>Y</given-names> </name><name name-style="western"><surname>Blijlevens</surname><given-names>NMA</given-names> </name><name name-style="western"><surname>Hermens</surname><given-names>RPMG</given-names> </name></person-group><article-title>A comprehensive eHealth implementation guide constructed on a qualitative case study on barriers and facilitators of the digital care platform CMyLife</article-title><source>BMC Health Serv Res</source><year>2022</year><month>06</month><day>6</day><volume>22</volume><issue>1</issue><fpage>751</fpage><pub-id pub-id-type="doi">10.1186/s12913-022-08020-3</pub-id><pub-id pub-id-type="medline">35668491</pub-id></nlm-citation></ref><ref id="ref23"><label>23</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Lewin</surname><given-names>K</given-names> </name></person-group><article-title>Action research and minority problems</article-title><source>J Soc Issues</source><year>1946</year><month>11</month><volume>2</volume><issue>4</issue><fpage>34</fpage><lpage>46</lpage><pub-id pub-id-type="doi">10.1111/j.1540-4560.1946.tb02295.x</pub-id></nlm-citation></ref><ref id="ref24"><label>24</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Casey</surname><given-names>M</given-names> </name><name name-style="western"><surname>Coghlan</surname><given-names>D</given-names> </name><name name-style="western"><surname>Carroll</surname><given-names>&#x00C1;</given-names> </name><name name-style="western"><surname>Stokes</surname><given-names>D</given-names> </name></person-group><article-title>Towards a checklist for improving action research quality in healthcare contexts</article-title><source>Syst Pract Action Res</source><year>2023</year><month>12</month><volume>36</volume><issue>6</issue><fpage>923</fpage><lpage>934</lpage><pub-id pub-id-type="doi">10.1007/s11213-023-09635-1</pub-id></nlm-citation></ref><ref id="ref25"><label>25</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Noble</surname><given-names>H</given-names> </name><name name-style="western"><surname>Heale</surname><given-names>R</given-names> </name></person-group><article-title>Triangulation in research, with examples</article-title><source>Evid Based Nurs</source><year>2019</year><month>07</month><volume>22</volume><issue>3</issue><fpage>67</fpage><lpage>68</lpage><pub-id pub-id-type="doi">10.1136/ebnurs-2019-103145</pub-id><pub-id pub-id-type="medline">31201209</pub-id></nlm-citation></ref><ref id="ref26"><label>26</label><nlm-citation citation-type="book"><person-group person-group-type="author"><name name-style="western"><surname>Blevins</surname><given-names>MD</given-names> </name></person-group><source>The SAGE Encyclopedia of Communication Research Methods</source><year>2017</year><publisher-name>SAGE Publications, Inc</publisher-name><pub-id pub-id-type="doi">10.4135/9781483381411.n293</pub-id></nlm-citation></ref><ref id="ref27"><label>27</label><nlm-citation citation-type="book"><person-group person-group-type="author"><name name-style="western"><surname>Holmes</surname><given-names>S</given-names> </name><name name-style="western"><surname>Moorhead</surname><given-names>A</given-names> </name><name name-style="western"><surname>Bond</surname><given-names>R</given-names> </name><name name-style="western"><surname>Zheng</surname><given-names>H</given-names> </name><name name-style="western"><surname>Coates</surname><given-names>V</given-names> </name><name name-style="western"><surname>Mctear</surname><given-names>M</given-names> </name></person-group><source>Usability Testing of a Healthcare Chatbot: Can We Use Conventional Methods to Assess Conversational User Interfaces?</source><year>2019</year><publisher-name>ACM</publisher-name><fpage>207</fpage><lpage>214</lpage><pub-id pub-id-type="other">978-1-4503-7166-7</pub-id></nlm-citation></ref><ref id="ref28"><label>28</label><nlm-citation citation-type="other"><person-group person-group-type="author"><name name-style="western"><surname>Gan</surname><given-names>A</given-names> </name><name name-style="western"><surname>Yu</surname><given-names>H</given-names> </name><name name-style="western"><surname>Zhang</surname><given-names>K</given-names> </name><name name-style="western"><surname>Liu</surname><given-names>Q</given-names> </name><name name-style="western"><surname>Yan</surname><given-names>W</given-names> </name><name name-style="western"><surname>Huang</surname><given-names>Z</given-names> </name><etal/></person-group><article-title>Retrieval augmented generation evaluation in the era of large language models: a comprehensive survey</article-title><source>arXiv</source><comment>Preprint posted online on  Apr 21, 2025</comment><pub-id pub-id-type="doi">10.48550/arXiv.2504.14891</pub-id></nlm-citation></ref><ref id="ref29"><label>29</label><nlm-citation citation-type="book"><person-group person-group-type="author"><name name-style="western"><surname>Bauer</surname><given-names>M</given-names> </name></person-group><article-title>Classical content analysis: a review</article-title><source>Qualitative Researching with Text, Image and Sound: A Practical Handbook</source><year>2000</year><publisher-name>Sage Publications</publisher-name><fpage>131</fpage><lpage>151</lpage><pub-id pub-id-type="doi">10.4135/9781849209731.n8</pub-id></nlm-citation></ref><ref id="ref30"><label>30</label><nlm-citation citation-type="web"><article-title>The chatbot usability questionnaire (CUQ)</article-title><source>Ulster University</source><access-date>2026-08-17</access-date><comment><ext-link ext-link-type="uri" xlink:href="https://www.ulster.ac.uk/research/topic/computer-science/artificial-intelligence/projects/cuq">https://www.ulster.ac.uk/research/topic/computer-science/artificial-intelligence/projects/cuq</ext-link></comment></nlm-citation></ref><ref id="ref31"><label>31</label><nlm-citation citation-type="other"><person-group person-group-type="author"><name name-style="western"><surname>Wang</surname><given-names>X</given-names> </name><name name-style="western"><surname>Zhang</surname><given-names>X</given-names> </name><name name-style="western"><surname>Wang</surname><given-names>G</given-names> </name><name name-style="western"><surname>He</surname><given-names>J</given-names> </name><name name-style="western"><surname>Li</surname><given-names>Z</given-names> </name><name name-style="western"><surname>Zhu</surname><given-names>W</given-names> </name><etal/></person-group><article-title>OpenMEDLab: an open-source platform for multi-modality foundation models in medicine</article-title><source>arXiv</source><comment>Preprint posted online on  Feb 28, 2024</comment><pub-id pub-id-type="doi">10.48550/arXiv.2402.18028</pub-id></nlm-citation></ref><ref id="ref32"><label>32</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>O&#x2019;Connor</surname><given-names>S</given-names> </name></person-group><article-title>Virtual reality and avatars in health care</article-title><source>Clin Nurs Res</source><year>2019</year><month>06</month><volume>28</volume><issue>5</issue><fpage>523</fpage><lpage>528</lpage><pub-id pub-id-type="doi">10.1177/1054773819845824</pub-id><pub-id pub-id-type="medline">31064283</pub-id></nlm-citation></ref><ref id="ref33"><label>33</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Griffin</surname><given-names>AC</given-names> </name><name name-style="western"><surname>Khairat</surname><given-names>S</given-names> </name><name name-style="western"><surname>Bailey</surname><given-names>SC</given-names> </name><name name-style="western"><surname>Chung</surname><given-names>AE</given-names> </name></person-group><article-title>A chatbot for hypertension self-management support: user-centered design, development, and usability testing</article-title><source>JAMIA Open</source><year>2023</year><month>10</month><volume>6</volume><issue>3</issue><fpage>ooad073</fpage><pub-id pub-id-type="doi">10.1093/jamiaopen/ooad073</pub-id><pub-id pub-id-type="medline">37693367</pub-id></nlm-citation></ref><ref id="ref34"><label>34</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Bickmore</surname><given-names>T</given-names> </name><name name-style="western"><surname>Pfeifer</surname><given-names>L</given-names> </name><name name-style="western"><surname>Yin</surname><given-names>L</given-names> </name></person-group><article-title>The role of gesture in document explanation by embodied conversational agents</article-title><source>Int J Semantic Computing</source><year>2008</year><month>03</month><volume>02</volume><issue>01</issue><fpage>47</fpage><lpage>70</lpage><pub-id pub-id-type="doi">10.1142/S1793351X08000348</pub-id></nlm-citation></ref><ref id="ref35"><label>35</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Rizvi</surname><given-names>DS</given-names> </name></person-group><article-title>Health education and global health: Practices, applications, and future research</article-title><source>J Educ Health Promot</source><year>2022</year><volume>11</volume><fpage>262</fpage><pub-id pub-id-type="doi">10.4103/jehp.jehp_218_22</pub-id><pub-id pub-id-type="medline">36325224</pub-id></nlm-citation></ref><ref id="ref36"><label>36</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Kuwabara</surname><given-names>A</given-names> </name><name name-style="western"><surname>Su</surname><given-names>S</given-names> </name><name name-style="western"><surname>Krauss</surname><given-names>J</given-names> </name></person-group><article-title>Utilizing digital health technologies for patient education in lifestyle medicine</article-title><source>Am J Lifestyle Med</source><year>2020</year><volume>14</volume><issue>2</issue><fpage>137</fpage><lpage>142</lpage><pub-id pub-id-type="doi">10.1177/1559827619892547</pub-id><pub-id pub-id-type="medline">32231478</pub-id></nlm-citation></ref><ref id="ref37"><label>37</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Schooley</surname><given-names>B</given-names> </name><name name-style="western"><surname>Singh</surname><given-names>A</given-names> </name><name name-style="western"><surname>Hikmet</surname><given-names>N</given-names> </name><name name-style="western"><surname>Brookshire</surname><given-names>R</given-names> </name><name name-style="western"><surname>Patel</surname><given-names>N</given-names> </name></person-group><article-title>Integrated digital patient education at the bedside for patients with chronic conditions: observational study</article-title><source>JMIR Mhealth Uhealth</source><year>2020</year><month>12</month><day>22</day><volume>8</volume><issue>12</issue><fpage>e22947</fpage><pub-id pub-id-type="doi">10.2196/22947</pub-id><pub-id pub-id-type="medline">33350961</pub-id></nlm-citation></ref><ref id="ref38"><label>38</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Tseng</surname><given-names>LM</given-names> </name><name name-style="western"><surname>Lien</surname><given-names>PJ</given-names> </name><name name-style="western"><surname>Huang</surname><given-names>CY</given-names> </name><name name-style="western"><surname>Tsai</surname><given-names>YF</given-names> </name><name name-style="western"><surname>Chao</surname><given-names>TC</given-names> </name><name name-style="western"><surname>Huang</surname><given-names>SM</given-names> </name></person-group><article-title>Developing a web-based shared decision-making tool for fertility preservation among reproductive-age women with breast cancer: an action research approach</article-title><source>J Med Internet Res</source><year>2021</year><month>03</month><day>17</day><volume>23</volume><issue>3</issue><fpage>e24926</fpage><pub-id pub-id-type="doi">10.2196/24926</pub-id><pub-id pub-id-type="medline">33729164</pub-id></nlm-citation></ref><ref id="ref39"><label>39</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><collab>The Lancet Digital Health</collab></person-group><article-title>Digital health equity for older populations</article-title><source>Lancet Digit Health</source><year>2023</year><month>07</month><volume>5</volume><issue>7</issue><fpage>e395</fpage><pub-id pub-id-type="doi">10.1016/S2589-7500(23)00114-0</pub-id><pub-id pub-id-type="medline">37391262</pub-id></nlm-citation></ref><ref id="ref40"><label>40</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Bickmore</surname><given-names>T</given-names> </name><name name-style="western"><surname>Gruber</surname><given-names>A</given-names> </name><name name-style="western"><surname>Picard</surname><given-names>R</given-names> </name></person-group><article-title>Establishing the computer-patient working alliance in automated health behavior change interventions</article-title><source>Patient Educ Couns</source><year>2005</year><month>10</month><volume>59</volume><issue>1</issue><fpage>21</fpage><lpage>30</lpage><pub-id pub-id-type="doi">10.1016/j.pec.2004.09.008</pub-id><pub-id pub-id-type="medline">16198215</pub-id></nlm-citation></ref><ref id="ref41"><label>41</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Fournier-Tombs</surname><given-names>E</given-names> </name><name name-style="western"><surname>McHardy</surname><given-names>J</given-names> </name></person-group><article-title>A medical ethics framework for conversational artificial intelligence</article-title><source>J Med Internet Res</source><year>2023</year><month>07</month><day>26</day><volume>25</volume><fpage>e43068</fpage><pub-id pub-id-type="doi">10.2196/43068</pub-id><pub-id pub-id-type="medline">37224277</pub-id></nlm-citation></ref><ref id="ref42"><label>42</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Kumanyika</surname><given-names>SK</given-names> </name></person-group><article-title>Health equity is the issue we have been waiting for</article-title><source>J Public Health Manag Pract</source><year>2016</year><volume>22 Suppl 1</volume><fpage>S8</fpage><lpage>S10</lpage><pub-id pub-id-type="doi">10.1097/PHH.0000000000000363</pub-id><pub-id pub-id-type="medline">26599034</pub-id></nlm-citation></ref><ref id="ref43"><label>43</label><nlm-citation citation-type="confproc"><person-group person-group-type="author"><name name-style="western"><surname>Batani</surname><given-names>J</given-names> </name><name name-style="western"><surname>Mbunge</surname><given-names>E</given-names> </name><name name-style="western"><surname>Leokana</surname><given-names>L</given-names> </name></person-group><person-group person-group-type="editor"><name name-style="western"><surname>Batani</surname><given-names>J</given-names> </name><name name-style="western"><surname>Mbunge</surname><given-names>E</given-names> </name><name name-style="western"><surname>Leokana</surname><given-names>L</given-names> </name></person-group><article-title>A deep learning-based chatbot to enhance maternal health education</article-title><conf-name>2024 Conference on Information Communications Technology and Society (ICTAS)</conf-name><conf-date>Mar 7-8, 2024</conf-date><pub-id pub-id-type="doi">10.1109/ICTAS59620.2024.10507149</pub-id></nlm-citation></ref><ref id="ref44"><label>44</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Castillo</surname><given-names>EG</given-names> </name><name name-style="western"><surname>Harris</surname><given-names>C</given-names> </name></person-group><article-title>Directing research toward health equity: a health equity research impact assessment</article-title><source>J Gen Intern Med</source><year>2021</year><month>09</month><volume>36</volume><issue>9</issue><fpage>2803</fpage><lpage>2808</lpage><pub-id pub-id-type="doi">10.1007/s11606-021-06789-3</pub-id><pub-id pub-id-type="medline">33948804</pub-id></nlm-citation></ref><ref id="ref45"><label>45</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Liburd</surname><given-names>LC</given-names> </name><name name-style="western"><surname>Hall</surname><given-names>JE</given-names> </name><name name-style="western"><surname>Mpofu</surname><given-names>JJ</given-names> </name><name name-style="western"><surname>Williams</surname><given-names>SM</given-names> </name><name name-style="western"><surname>Bouye</surname><given-names>K</given-names> </name><name name-style="western"><surname>Penman-Aguilar</surname><given-names>A</given-names> </name></person-group><article-title>Addressing health equity in public health practice: frameworks, promising strategies, and measurement considerations</article-title><source>Annu Rev Public Health</source><year>2020</year><month>04</month><day>2</day><volume>41</volume><fpage>417</fpage><lpage>432</lpage><pub-id pub-id-type="doi">10.1146/annurev-publhealth-040119-094119</pub-id><pub-id pub-id-type="medline">31900101</pub-id></nlm-citation></ref><ref id="ref46"><label>46</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Choi</surname><given-names>&#x041A;EA</given-names> </name><name name-style="western"><surname>Fitzek</surname><given-names>S</given-names> </name></person-group><article-title>User and provider experiences with health education chatbots: qualitative systematic review</article-title><source>JMIR Hum Factors</source><year>2025</year><month>06</month><day>13</day><volume>12</volume><fpage>e60205</fpage><pub-id pub-id-type="doi">10.2196/60205</pub-id><pub-id pub-id-type="medline">40513000</pub-id></nlm-citation></ref><ref id="ref47"><label>47</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Martins</surname><given-names>A</given-names> </name><name name-style="western"><surname>Velez Lap&#x00E3;o</surname><given-names>L</given-names> </name><name name-style="western"><surname>Nunes</surname><given-names>IL</given-names> </name><etal/></person-group><article-title>A conversational agent for enhanced self-management after cardiothoracic surgery</article-title><source>Int J Med Inform</source><year>2024</year><month>12</month><volume>192</volume><fpage>105640</fpage><pub-id pub-id-type="doi">10.1016/j.ijmedinf.2024.105640</pub-id><pub-id pub-id-type="medline">39321492</pub-id></nlm-citation></ref><ref id="ref48"><label>48</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Kim</surname><given-names>AR</given-names> </name><name name-style="western"><surname>Park</surname><given-names>HA</given-names> </name></person-group><article-title>A question answering chatbot for gastric cancer patients after curative gastrectomy: Development and evaluation of user experience and performance</article-title><source>Comput Inform Nurs</source><year>2024</year><month>11</month><day>1</day><volume>42</volume><issue>11</issue><fpage>829</fpage><lpage>839</lpage><pub-id pub-id-type="doi">10.1097/CIN.0000000000001153</pub-id><pub-id pub-id-type="medline">38861611</pub-id></nlm-citation></ref><ref id="ref49"><label>49</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Jiang</surname><given-names>Z</given-names> </name><name name-style="western"><surname>Huang</surname><given-names>X</given-names> </name><name name-style="western"><surname>Wang</surname><given-names>Z</given-names> </name><name name-style="western"><surname>Liu</surname><given-names>Y</given-names> </name><name name-style="western"><surname>Huang</surname><given-names>L</given-names> </name><name name-style="western"><surname>Luo</surname><given-names>X</given-names> </name></person-group><article-title>Embodied conversational agents for chronic diseases: scoping review</article-title><source>J Med Internet Res</source><year>2024</year><month>01</month><day>9</day><volume>26</volume><fpage>e47134</fpage><pub-id pub-id-type="doi">10.2196/47134</pub-id><pub-id pub-id-type="medline">38194260</pub-id></nlm-citation></ref><ref id="ref50"><label>50</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Yang</surname><given-names>FC</given-names> </name><name name-style="western"><surname>Acevedo</surname><given-names>P</given-names> </name><name name-style="western"><surname>Guo</surname><given-names>S</given-names> </name><name name-style="western"><surname>Choi</surname><given-names>M</given-names> </name><name name-style="western"><surname>Mousas</surname><given-names>C</given-names> </name></person-group><article-title>Embodied conversational agents in extended reality: a systematic review</article-title><source>IEEE Access</source><year>2025</year><volume>13</volume><fpage>79805</fpage><lpage>79824</lpage><pub-id pub-id-type="doi">10.1109/ACCESS.2025.3566698</pub-id><pub-id pub-id-type="medline">40256415</pub-id></nlm-citation></ref><ref id="ref51"><label>51</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Myllyaho</surname><given-names>L</given-names> </name><name name-style="western"><surname>Raatikainen</surname><given-names>M</given-names> </name><name name-style="western"><surname>M&#x00E4;nnist&#x00F6;</surname><given-names>T</given-names> </name><name name-style="western"><surname>Mikkonen</surname><given-names>T</given-names> </name><name name-style="western"><surname>Nurminen</surname><given-names>JK</given-names> </name></person-group><article-title>Systematic literature review of validation methods for AI systems</article-title><source>J Syst Softw</source><year>2021</year><month>11</month><volume>181</volume><fpage>111050</fpage><pub-id pub-id-type="doi">10.1016/j.jss.2021.111050</pub-id></nlm-citation></ref><ref id="ref52"><label>52</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Li</surname><given-names>L</given-names> </name></person-group><article-title>Role of chatbots on gastroenterology: Let&#x2019;s chat about the future</article-title><source>Gastrointest Endosc</source><year>2023</year><month>07</month><volume>1</volume><issue>3</issue><fpage>144</fpage><lpage>149</lpage><pub-id pub-id-type="doi">10.1016/j.gande.2023.06.002</pub-id></nlm-citation></ref></ref-list><app-group><supplementary-material id="app1"><label>Multimedia Appendix 1</label><p>In-depth interview outlines, sample questions, recruitment criteria, evidence resource of the knowledge corpus, detailed results of usability testing, and qualitative results.</p><media xlink:href="jmir_v28i1e82746_app1.docx" xlink:title="DOCX File, 49 KB"/></supplementary-material><supplementary-material id="app2"><label>Multimedia Appendix 2</label><p>Electronic materials for conversational agent.</p><media xlink:href="jmir_v28i1e82746_app2.mp4" xlink:title="MP4 File, 7029 KB"/></supplementary-material><supplementary-material id="app3"><label>Checklist 1</label><p>The reporting guideline and checklist &#x201C;Towards a Checklist for Improving Action Research Quality in Healthcare Contexts.&#x201D;</p><media xlink:href="jmir_v28i1e82746_app3.docx" xlink:title="DOCX File, 19 KB"/></supplementary-material></app-group></back></article>