<?xml version="1.0" encoding="UTF-8"?><!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.0 20040830//EN" "journalpublishing.dtd"><article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" dtd-version="2.0" xml:lang="en" article-type="research-article"><front><journal-meta><journal-id journal-id-type="nlm-ta">J Med Internet Res</journal-id><journal-id journal-id-type="publisher-id">jmir</journal-id><journal-id journal-id-type="index">1</journal-id><journal-title>Journal of Medical Internet Research</journal-title><abbrev-journal-title>J Med Internet Res</abbrev-journal-title><issn pub-type="epub">1438-8871</issn><publisher><publisher-name>JMIR Publications</publisher-name><publisher-loc>Toronto, Canada</publisher-loc></publisher></journal-meta><article-meta><article-id pub-id-type="publisher-id">v28i1e87283</article-id><article-id pub-id-type="doi">10.2196/87283</article-id><article-categories><subj-group subj-group-type="heading"><subject>Original Paper</subject></subj-group></article-categories><title-group><article-title>Readability of AI-Generated Patient Visit Summaries in Orthopedic Surgery: Retrospective Analysis</article-title></title-group><contrib-group><contrib contrib-type="author"><name name-style="western"><surname>Major III</surname><given-names>Edward Lee</given-names></name><degrees>BS</degrees><xref ref-type="aff" rid="aff1">1</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Shah</surname><given-names>Vivek P</given-names></name><degrees>MD</degrees><xref ref-type="aff" rid="aff2">2</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Carroll</surname><given-names>Amber N</given-names></name><degrees>MD</degrees><xref ref-type="aff" rid="aff3">3</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Goetz</surname><given-names>James R</given-names></name><degrees>MD</degrees><xref ref-type="aff" rid="aff1">1</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Malempati</surname><given-names>Chaitu</given-names></name><degrees>DO</degrees><xref ref-type="aff" rid="aff1">1</xref></contrib><contrib contrib-type="author" corresp="yes"><name name-style="western"><surname>Badarudeen</surname><given-names>Sameer</given-names></name><degrees>MD, MPH</degrees><xref ref-type="aff" rid="aff1">1</xref></contrib></contrib-group><aff id="aff1"><institution>Department of Orthopaedic Surgery and Sports Medicine, College of Medicine, University of Kentucky</institution><addr-line>740 S Limestone St</addr-line><addr-line>Lexington</addr-line><addr-line>KY</addr-line><country>United States</country></aff><aff id="aff2"><institution>Department of Orthopaedic Surgery, McLaren Regional Medical Center</institution><addr-line>Flint</addr-line><addr-line>MI</addr-line><country>United States</country></aff><aff id="aff3"><institution>Department of Orthopaedic Surgery, Indiana University</institution><addr-line>Indianapolis</addr-line><addr-line>IN</addr-line><country>United States</country></aff><contrib-group><contrib contrib-type="editor"><name name-style="western"><surname>Steenstra</surname><given-names>Ivan</given-names></name></contrib></contrib-group><contrib-group><contrib contrib-type="reviewer"><name name-style="western"><surname>Mateos-Garcia</surname><given-names>Daniel</given-names></name></contrib><contrib contrib-type="reviewer"><name name-style="western"><surname>Morris</surname><given-names>Kevin</given-names></name></contrib></contrib-group><author-notes><corresp>Correspondence to Sameer Badarudeen, MD, MPH, Department of Orthopaedic Surgery and Sports Medicine, College of Medicine, University of Kentucky, 740 S Limestone St, Lexington, KY, United States, 1 323 873 4940; <email>sameer.badarudeen@wellstar.org</email></corresp></author-notes><pub-date pub-type="collection"><year>2026</year></pub-date><pub-date pub-type="epub"><day>20</day><month>8</month><year>2026</year></pub-date><volume>28</volume><elocation-id>e87283</elocation-id><history><date date-type="received"><day>16</day><month>11</month><year>2025</year></date><date date-type="rev-recd"><day>22</day><month>05</month><year>2026</year></date><date date-type="accepted"><day>27</day><month>05</month><year>2026</year></date></history><copyright-statement>&#x00A9; Edward Lee Major III, Vivek P Shah, Amber N Carroll, James R Goetz, Chaitu Malempati, Sameer Badarudeen. Originally published in the Journal of Medical Internet Research (<ext-link ext-link-type="uri" xlink:href="https://www.jmir.org">https://www.jmir.org</ext-link>), 20.8.2026. </copyright-statement><copyright-year>2026</copyright-year><license license-type="open-access" xlink:href="https://creativecommons.org/licenses/by/4.0/"><p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (<ext-link ext-link-type="uri" xlink:href="https://creativecommons.org/licenses/by/4.0/">https://creativecommons.org/licenses/by/4.0/</ext-link>), which permits unrestricted use, distribution, and reproduction in any medium, provided the original work, first published in the Journal of Medical Internet Research (ISSN 1438-8871), is properly cited. The complete bibliographic information, a link to the original publication on <ext-link ext-link-type="uri" xlink:href="https://www.jmir.org/">https://www.jmir.org/</ext-link>, as well as this copyright and license information must be included.</p></license><self-uri xlink:type="simple" xlink:href="https://www.jmir.org/2026/1/e87283"/><abstract><sec><title>Background</title><p>Patient visit summaries (PVS) are patient-facing documents intended to reinforce communication and promote patient education after clinical encounters. Despite national recommendations that patient education materials be written at or below a sixth-grade reading level, most orthopedic materials substantially exceed this threshold. Current visit summaries are also time-consuming to generate, lack personalization, and often fail to meet patient literacy needs. AI-based scribes can generate personalized PVS in real time directly from patient-provider conversations, offering a potential solution.</p></sec><sec><title>Objective</title><p>This study aimed to evaluate the readability of AI-generated PVS produced by a commercial AI scribe platform in an orthopedic surgery setting and determine their alignment with established literacy standards for patient-facing materials.</p></sec><sec sec-type="methods"><title>Methods</title><p>A total of 1007 consecutive AI-generated PVS from an academic orthopedic surgery outpatient clinic between December 2023 and May 2024 were reviewed. Following standardized preprocessing, including restoration of original section headings and removal of diagnosis label headers, summaries identified as incomplete were excluded (n=25), yielding a final study cohort of 982 summaries. Readability was assessed using 5 validated indices: the Flesch-Kincaid Grade Level (FKGL), Flesch Reading Ease Score (FRES), Gunning Fog Index (GFI), Coleman-Liau Index (CLI), and Simple Measure of Gobbledygook (SMOG) Index. The proportions of PVS meeting the sixth- and eighth-grade benchmarks were calculated. Spearman&#x2019;s rank correlation (&#x03C1;) assessed associations between word count and readability metrics. The Kendall Coefficient of Concordance (<italic>W</italic>) was used to evaluate agreement among indices after reverse-coding FRES for directional alignment. Statistical significance was set at <italic>P</italic>&#x003C;.05.</p></sec><sec sec-type="results"><title>Results</title><p>Among 982 analyzed PVS, the mean FKGL was 9.3 (SD 1.2, 95% CI 9.2&#x2010;9.4) and the mean FRES was 57.6 (SD 7.9, 95% CI 57.1&#x2010;58.1), corresponding to &#x201C;fairly difficult.&#x201D; Mean scores for secondary indices were 12.4 (SD 1.6) for GFI, 11.0 (SD 1.5) for CLI, and 12.4 (SD 1.2) for SMOG. Only 0.4% (4/982) of PVS met the sixth-grade benchmark and 14.2% (139/982) met the eighth-grade threshold. Word count was not significantly correlated with FKGL (&#x03C1;=0.012, <italic>P</italic>=.70), suggesting that sentence structure and vocabulary rather than document length drive reading complexity. Readability indices demonstrated strong agreement across all 5 metrics (<italic>W</italic>=0.882, <italic>P</italic>&#x003C;.001).</p></sec><sec sec-type="conclusions"><title>Conclusions</title><p>In this orthopedic outpatient setting, AI-generated PVS consistently exceeded recommended patient literacy thresholds, with a mean ninth-grade reading level and only 14.2% of summaries meeting the eighth-grade standard. In the absence of a concurrent comparison group, the present study was not designed to evaluate improvement over traditional methods. These findings establish a quantitative readability baseline for AI scribe output in orthopedic surgery and highlight the need for algorithmic refinement, plain language optimization, and prospective patient comprehension testing to improve health communication.</p></sec></abstract><kwd-group><kwd>artificial intelligence</kwd><kwd>patient education materials</kwd><kwd>patient visit summary</kwd><kwd>readability</kwd><kwd>health literacy</kwd><kwd>AI</kwd></kwd-group></article-meta></front><body><sec id="s1" sec-type="intro"><title>Introduction</title><p>Effective communication between health care providers and patients is essential for achieving optimal health outcomes [<xref ref-type="bibr" rid="ref1">1</xref>,<xref ref-type="bibr" rid="ref2">2</xref>]. Although the average US adult reads at the eighth-grade level, a substantial proportion of patients have lower literacy skills and may struggle to comprehend medical information discussed during clinical encounters [<xref ref-type="bibr" rid="ref3">3</xref>]. Studies have shown that patients frequently forget or misinterpret key details, which can negatively affect treatment adherence, satisfaction, and safety [<xref ref-type="bibr" rid="ref4">4</xref>,<xref ref-type="bibr" rid="ref5">5</xref>]. Recognizing these challenges, national organizations such as the American Academy of Orthopedic Surgeons have recommended that patient education materials be written at or below a sixth-grade reading level [<xref ref-type="bibr" rid="ref5">5</xref>,<xref ref-type="bibr" rid="ref6">6</xref>]. However, most orthopedic patient education materials are written above the recommended reading levels, limiting their utility for many patients and providers [<xref ref-type="bibr" rid="ref7">7</xref>-<xref ref-type="bibr" rid="ref9">9</xref>].</p><p>Providing educational materials and personalized visit summaries after clinical encounters is a standard practice intended to support patient understanding and engagement [<xref ref-type="bibr" rid="ref10">10</xref>,<xref ref-type="bibr" rid="ref11">11</xref>]. These summaries are typically generated through manual entry, standardized templates, or automated extraction from electronic health records. However, each method has limitations: manual entry is time-consuming and may be inconsistently applied [<xref ref-type="bibr" rid="ref12">12</xref>]; templates often lack personalization and may include irrelevant or generic content [<xref ref-type="bibr" rid="ref13">13</xref>]; and automated extraction can result in inaccuracies, formatting inconsistencies, and limited customization for individual patient needs [<xref ref-type="bibr" rid="ref14">14</xref>,<xref ref-type="bibr" rid="ref15">15</xref>]. These challenges often result in summaries that are difficult for patients to comprehend, highlighting the need for innovative strategies that enhance readability, personalization, and accessibility for diverse patient populations [<xref ref-type="bibr" rid="ref11">11</xref>,<xref ref-type="bibr" rid="ref16">16</xref>].</p><p>AI platforms such as ambient scribes and generative language models can rapidly and accurately convert complex educational materials to recommended readability levels while maintaining content quality [<xref ref-type="bibr" rid="ref2">2</xref>,<xref ref-type="bibr" rid="ref9">9</xref>]. AI-based scribe platforms, such as the commercial Abridge system evaluated in this study, generate both structured clinical documentation and personalized visit summaries. Unlike traditional methods, AI tools use natural language processing and real-time audio recordings of patient-provider conversations to generate dynamic summaries that simplify medical jargon and integrate patient-specific concerns, potentially addressing the limitations of current materials [<xref ref-type="bibr" rid="ref9">9</xref>,<xref ref-type="bibr" rid="ref17">17</xref>,<xref ref-type="bibr" rid="ref18">18</xref>]. However, the readability of these summaries remains unexplored in the field of orthopedics, where health literacy has been shown to be lower than that in the general medical population [<xref ref-type="bibr" rid="ref8">8</xref>].</p><p>This study assessed the readability of patient visit summaries (PVS) produced by an AI-based scribe in an outpatient orthopedic clinic. The primary objectives were to determine the baseline readability of these AI-generated summaries and their alignment with the national literacy recommendations. We hypothesized that AI-generated summaries would demonstrate improved readability and closer alignment with recommended literacy standards for patient-facing materials, potentially helping address long-standing communication barriers. By situating these tools within the broader evolution of clinical documentation, this study explores the potential of AI-generated PVS to serve as a patient-centered communication tool in modern clinical care.</p></sec><sec id="s2" sec-type="methods"><title>Methods</title><sec id="s2-1"><title>Study Design and Setting</title><p>This retrospective study was conducted to assess the readability of PVS generated by an AI-based scribe platform deployed in an orthopedic surgery joint replacement clinic at a regional academic institution. All consecutive patient visits over a 6-month study period from December 2023 to May 2024 were reviewed. A commercial AI-based scribe platform, Abridge Inc, was used to generate all PVS analyzed in this study [<xref ref-type="bibr" rid="ref19">19</xref>].</p></sec><sec id="s2-2"><title>Ethical Considerations</title><p>The study protocol was approved by the Medical Center Institutional Review Board prior to initiation (number: 24-05-31-Bada-Read-Vst; approval date: May 31, 2024), and all protocols were conducted in accordance with the Declaration of Helsinki and its amendments [<xref ref-type="bibr" rid="ref20">20</xref>]. To protect patient privacy, all patient- and provider-identifying information was removed from the dataset before analysis. The AI-generated PVS were anonymized to prevent potential reidentification. The requirement for informed consent was waived by the institutional review board due to the retrospective study design and use of fully deidentified data.</p></sec><sec id="s2-3"><title>AI Scribe Patient Visit Process</title><p>During each visit, audio recordings of patient-provider conversations were processed using an AI scribe. The system used a large language model&#x2013;based real-time audio-to-text conversion and AI-driven automatic speech recognition to generate comprehensive visit transcripts. Advanced natural language processing algorithms within the AI scribe software extracted clinically relevant information from transcripts, including patient concerns, symptoms, physical examination findings, diagnoses, treatment plans, and follow-up instructions. The information was then organized into a detailed clinical note with specific sections, such as the history of present illness, physical examination, and assessment and plan.</p><p>Using the AI-generated clinical notes, the AI scribe then translated medical jargon into simple language, organized information into clear sections, and integrated the patient&#x2019;s concerns to generate a personalized, patient-facing PVS for each visit. Providers reviewed all PVS for clinical accuracy prior to release through the patient portal, email, or printout. However, all summaries analyzed in this study reflected the raw, prereview, AI-generated output prior to any provider editing. A sample of an AI-generated PVS is shown in <xref ref-type="fig" rid="figure1">Figure 1</xref>.</p><fig position="float" id="figure1"><label>Figure 1.</label><caption><p>Sample patient visit summary before and after standardized preprocessing. Panel (A) shows the raw AI-generated output, with diagnosis label headers (eg, &#x201C;-RIGHT SHOULDER PAIN:,&#x201D; &#x201C;-NECK PAIN&#x201D;:) highlighted in red within the &#x201C;YOUR PLAN&#x201D; section. Panel (B) shows the same summary following preprocessing, in which diagnosis label headers were removed as they represent structural metadata rather than patient-facing narrative prose. Section headings (&#x201C;VISIT SUMMARY,&#x201D; &#x201C;YOUR PLAN,&#x201D; and &#x201C;INSTRUCTIONS&#x201D;) and all clinical narrative content were preserved without rewriting or simplification. All readability analyses were performed on preprocessed summaries as shown in (B) prior to any provider review or editing.</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="jmir_v28i1e87283_fig01.png"/></fig></sec><sec id="s2-4"><title>Data Collection</title><p>A total of 1007 PVS were extracted from the AI scribe platform and converted into plain text for readability analyses. A standardized preprocessing protocol was applied to ensure consistent assessment while preserving patient-facing readability features. The preprocessing workflow included restoration of the original standardized section headings (&#x201C;VISIT SUMMARY,&#x201D; &#x201C;YOUR PLAN,&#x201D; and &#x201C;INSTRUCTIONS&#x201D;) to preserve the intended document organization and formatting generated by the AI scribe. Special characters were removed for software compatibility, including bullet points, degree symbols, and slash characters. Diagnosis label headers (eg, &#x201C;-RIGHT KNEE PAIN&#x201D;:) were systematically removed as these labels represented structural metadata rather than narrative patient-facing prose. The details of each diagnosis subheading were preserved as separate paragraph units to maintain the intended structure of the original summaries. Summaries missing one or more major patient-facing sections (&#x201C;VISIT SUMMARY,&#x201D; &#x201C;YOUR PLAN,&#x201D; and &#x201C;INSTRUCTIONS&#x201D;) due to incomplete export or truncation were excluded (25/1007, 2.5%), yielding a final study cohort of 982 summaries.</p><p>Following preprocessing, all PVS were uploaded into Readable.com (Readable Ltd), a commercially available readability analysis tool, as plain text and were analyzed using standardized batch text scoring. The original clinical terminology and patient-facing narrative structure were preserved throughout preprocessing to minimize artificial distortion of readability metrics while maintaining the intended patient-facing format [<xref ref-type="bibr" rid="ref5">5</xref>,<xref ref-type="bibr" rid="ref6">6</xref>,<xref ref-type="bibr" rid="ref9">9</xref>,<xref ref-type="bibr" rid="ref21">21</xref>,<xref ref-type="bibr" rid="ref22">22</xref>].</p></sec><sec id="s2-5"><title>Readability Assessment Methods</title><p>The primary readability outcomes were the Flesch-Kincaid Grade Level (FKGL) and the Flesch Reading Ease Score (FRES). FKGL estimates the US grade level required to comprehend a text, with lower scores indicating easier readability. FRES ranges from 0 to 100, with higher scores indicating easier readability and lower scores indicating more difficult readability. FRES categories were interpreted using established classifications: very easy (90-100); easy (80-89); fairly easy (70-79); standard (60-69); fairly difficult (50-59); difficult (30-49); very difficult (0&#x2010;29) [<xref ref-type="bibr" rid="ref21">21</xref>]. Both FKGL and FRES calculate readability based on sentence length (average words per sentence) and word complexity (average syllables per word), with longer sentences and multisyllabic words indicating increased reading difficulty, which is reflected by a higher FKGL and lower FRES.</p><p>To compare across multiple metrics, PVS readability was also assessed using the Gunning Fog Index (GFI), Coleman-Liau Index (CLI), and the Simple Measure of Gobbledygook (SMOG) Index [<xref ref-type="bibr" rid="ref23">23</xref>,<xref ref-type="bibr" rid="ref24">24</xref>]. The GFI estimates US grade-level readability based on sentence length (average words per sentence) and word complexity (percentage of words with 3 or more syllables), with lower scores indicating better readability. The CLI does not use sentence length but rather calculates readability based on character count per word and word count per sentence, aligning with US grade levels. The SMOG Index, designed for health-related materials, estimates readability by counting polysyllabic words (3 or more syllables) in 30 sentences, with lower values indicating an easier text. Descriptions and formulas of the readability metrics used in this study are summarized in <xref ref-type="table" rid="table1">Table 1</xref>.</p><table-wrap id="t1" position="float"><label>Table 1.</label><caption><p>Commonly used readability scales in health care.</p></caption><table id="table1" frame="hsides" rules="groups"><thead><tr><td align="left" valign="bottom">Readability scale</td><td align="left" valign="bottom">Formula</td><td align="left" valign="bottom">Range</td><td align="left" valign="bottom">Interpretation</td></tr></thead><tbody><tr><td align="left" valign="top">Flesch-Kincaid Grade Level (FKGL)</td><td align="left" valign="top">Grade level = (0.39 &#x00D7; average # words per sentence) + (11.8 &#x00D7; average # syllables per word) &#x2013; 15.59</td><td align="left" valign="top">0&#x2010;18+</td><td align="left" valign="top">&#x2193; FKGL=easier readability</td></tr><tr><td align="left" valign="top">Flesch Reading Ease Score (FRES)</td><td align="left" valign="top">Reading Ease Score = 206.835 &#x2013; (1.015 &#x00D7; average # words per sentence) &#x2013; (84.6 &#x00D7; average # syllables per word)</td><td align="left" valign="top">0&#x2010;100</td><td align="left" valign="top">&#x2191; FRES=easier readability</td></tr><tr><td align="left" valign="top">Gunning Fog Index (GFI)</td><td align="left" valign="top">Grade level = 0.4 &#x00D7; [(average # words per sentence) + (# words with &#x003E;3 syllables) &#x00D7; (100 / # of words)]</td><td align="left" valign="top">0&#x2010;20+</td><td align="left" valign="top">&#x2193; GFI=easier readability</td></tr><tr><td align="left" valign="top">Coleman-Liau Index (CLI)</td><td align="left" valign="top">Grade level = (0.0588 &#x00D7; average # of letters per 100 words) &#x2013; (0.296 &#x00D7; average # of sentences per 100 words) &#x2013; 15.8</td><td align="left" valign="top">0&#x2010;18+</td><td align="left" valign="top">&#x2193; CLI=easier readability</td></tr><tr><td align="left" valign="top">Simple Measure of Gobbledygook (SMOG) Index</td><td align="left" valign="top">Grade level = 1.043 &#x00D7; &#x221A; (# of polysyllabic words) &#x00D7; (30 / # of sentences)</td><td align="left" valign="top">0&#x2010;18+</td><td align="left" valign="top">&#x2193; SMOG=easier readability</td></tr></tbody></table></table-wrap><p>The readability metrics implemented by Readable.com (Readable Ltd), including FKGL, FRES, GFI, CLI, and SMOG, are validated measures commonly used in health communication research [<xref ref-type="bibr" rid="ref25">25</xref>-<xref ref-type="bibr" rid="ref28">28</xref>]. Both the AI scribe platform and Readable analysis tool were configured to American English conventions to maintain consistency across all analyzed summaries.</p><p>The proportions of PVS with an FKGL at or below the sixth-grade reading level and at or below the eighth-grade reading level were calculated to assess alignment with commonly recommended patient health literacy standards [<xref ref-type="bibr" rid="ref6">6</xref>] and the average adult reading grade level in the United States [<xref ref-type="bibr" rid="ref3">3</xref>]. Word count was recorded to evaluate associations between PVS length and readability metrics.</p></sec><sec id="s2-6"><title>Data Analysis</title><p>The mean, median, SD, range, and 95% CIs were computed for each readability metric to summarize the distribution of readability scores across the analyzed PVS. The Kendall Coefficient of Concordance (<italic>W</italic>) was used to evaluate agreement among the 5 readability metrics in ranking summaries by readability difficulty. Prior to concordance analysis, the FRES was reverse-coded so that higher values consistently represented greater reading difficulty across all readability metrics. Spearman rank correlation (&#x03C1;) was used to assess associations between word count and readability metrics due to the nonparametric nature of the correlation analysis. Correlation strength was classified as strong (&#x03C1;&#x2265;0.7), moderate (0.3&#x2264;&#x03C1;&#x003C;0.7), or weak (&#x03C1;&#x003C;0.3), according to established thresholds [<xref ref-type="bibr" rid="ref29">29</xref>]. All statistical tests were 2-tailed with statistical significance defined as a <italic>P</italic> value of &#x003C;.05. Statistical analyses were performed using IBM SPSS Statistics (version 31.0; IBM Corp).</p></sec></sec><sec id="s3" sec-type="results"><title>Results</title><sec id="s3-1"><title>Study Population</title><p>A total of 1007 consecutive AI-generated PVS were identified during the study period, of which 982 met the inclusion criteria for the final analysis.</p></sec><sec id="s3-2"><title>Readability of AI-Generated Summaries</title><p>Across all PVS, the mean FKGL was 9.3 (SD 1.2, 95% CI 9.2&#x2010;9.4), with a median of 9.3 and a range from 5.3 to 13.7. The mean FRES was 57.6 (SD 7.9, 95% CI 57.1&#x2010;58.1), corresponding to the &#x201C;fairly difficult&#x201D; category, with a median of 57.7 and a range of 30.3 to 84.7. The distribution of PVS across FRES categories is shown in <xref ref-type="fig" rid="figure2">Figure 2</xref>.</p><fig position="float" id="figure2"><label>Figure 2.</label><caption><p>Distribution of patient visit summaries by Flesch Reading Ease Score (FRES) category. This figure displays the distribution of 982 AI-generated patient visit summaries (PVS) across established FRES interpretation categories (range 0-100; higher scores indicate easier readability). Most PVS fell within the &#x201C;fairly difficult&#x201D; (441/982, 44.9%) and &#x201C;standard&#x201D; (321/982, 32.7%) categories. Smaller proportions were categorized as &#x201C;difficult&#x201D; (159/982, 16.2%), &#x201C;fairly easy&#x201D; (58/982, 5.9%), and &#x201C;easy&#x201D; (3/982, 0.3%). No summaries were classified as &#x201C;very easy&#x201D; or &#x201C;very difficult.&#x201D;</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="jmir_v28i1e87283_fig02.png"/></fig><p>The GFI demonstrated a mean score of 12.4 (SD 1.6, 95% CI 12.3&#x2010;12.5), ranging from 6.8 to 18.8. The CLI demonstrated a mean grade level of 11.0 (SD 1.5, 95% CI 11.0&#x2010;11.1), with a range of 6.1 to 16.0. The SMOG Index demonstrated a mean grade level of 12.4 (SD 1.2, 95% CI 12.3&#x2010;12.4), with scores ranging from 7.7 to 16.5. Comprehensive descriptive statistics for all readability indices are presented in <xref ref-type="table" rid="table2">Table 2</xref>.</p><table-wrap id="t2" position="float"><label>Table 2.</label><caption><p>Readability summary statistics for AI-generated patient visit summaries.</p></caption><table id="table2" frame="hsides" rules="groups"><thead><tr><td align="left" valign="bottom">Readability index</td><td align="left" valign="bottom">Mean (SD)</td><td align="left" valign="bottom">Median</td><td align="left" valign="bottom">Range</td><td align="left" valign="bottom">95% CI</td></tr></thead><tbody><tr><td align="left" valign="top">Flesch-Kincaid Grade Level (FKGL); range: 0&#x2010;18+ (&#x2193;=better)</td><td align="left" valign="top">9.3 (1.2)</td><td align="left" valign="top">9.3</td><td align="left" valign="top">5.3&#x2010;13.7</td><td align="left" valign="top">9.2&#x2010;9.4</td></tr><tr><td align="left" valign="top">Flesch Reading Ease Score (FRES); range: 0&#x2010;100 (&#x2191;=better)</td><td align="left" valign="top">57.6 (7.9)</td><td align="left" valign="top">57.7</td><td align="left" valign="top">30.3&#x2010;84.7</td><td align="left" valign="top">57.1&#x2010;58.1</td></tr><tr><td align="left" valign="top">Gunning Fog Index (GFI); range: 0&#x2010;20+ (&#x2193;=better)</td><td align="left" valign="top">12.4 (1.6)</td><td align="left" valign="top">12.3</td><td align="left" valign="top">6.8&#x2010;18.8</td><td align="left" valign="top">12.3&#x2010;12.5</td></tr><tr><td align="left" valign="top">Coleman-Liau Index (CLI); range: 0&#x2010;18+ (&#x2193;=better)</td><td align="left" valign="top">11.0 (1.5)</td><td align="left" valign="top">11.1</td><td align="left" valign="top">6.1&#x2010;16.0</td><td align="left" valign="top">11.0&#x2010;11.1</td></tr><tr><td align="left" valign="top">SMOG Index; range: 0&#x2010;18+ (&#x2193;=better)</td><td align="left" valign="top">12.4 (1.2)</td><td align="left" valign="top">12.4</td><td align="left" valign="top">7.7&#x2010;16.5</td><td align="left" valign="top">12.3&#x2010;12.4</td></tr></tbody></table></table-wrap><p>Kendall Coefficient of Concordance (<italic>W</italic>) analysis demonstrated strong agreement in readability rankings across the 5 metrics (<italic>W</italic>=0.882, <italic>P</italic>&#x003C;.001). The pairwise associations shown in <xref ref-type="fig" rid="figure3">Figure 3</xref> were consistent with this high degree of intermetric concordance.</p><fig position="float" id="figure3"><label>Figure 3.</label><caption><p>Scatter plot matrix illustrating pairwise relationships among 5 readability metrics for 982 AI-generated patient visit summaries. Strong linear relationships were observed across all metric pairs, consistent with high intermetric concordance.</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="jmir_v28i1e87283_fig03.png"/></fig></sec><sec id="s3-3"><title>Alignment With National Literacy Benchmarks</title><p>Only 0.4% (4/982) of PVS met the recommended sixth-grade reading level threshold for health communication materials. A higher proportion of 14.2% (n=139) of PVS met the eighth-grade benchmark. The remaining summaries were distributed across higher FKGL ranges: 588 (59.9%) between grades 8.01&#x2010;10.0, 237 (24.1%) between grades 10.01&#x2010;12.0, and 18 (1.8%) greater than grade 12. <xref ref-type="fig" rid="figure4">Figure 4</xref> illustrates the cumulative percentage of summaries that met the various grade-level thresholds.</p><fig position="float" id="figure4"><label>Figure 4.</label><caption><p>Distribution of patient visit summary readability by Flesch-Kincaid Grade Level (FKGL). This figure displays the percentage distribution of AI-generated patient visit summaries (PVS) across FKGL categories. Most summaries (588/982, 59.9%) fell within the grade 8.01&#x2010;10.0 range, whereas 237 (24.1%) were within grades 10.01&#x2010;12.0, and 135 (13.7%) were within grades 6.01&#x2010;8.0. Only 4 summaries (0.4%) met the sixth-grade benchmark, whereas 18 (1.8%) exceeded grade 12.</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="jmir_v28i1e87283_fig04.png"/></fig></sec><sec id="s3-4"><title>Relationship Between Word Count and Readability</title><p>The mean word count per PVS was 202.0 (SD 47.0, median 200.0, range 84.0&#x2010;389.0, 95% CI 199.1&#x2010;205.0). Spearman rank correlations between word count and readability metrics were uniformly weak (|&#x03C1;|&#x003C;0.15). Statistically significant but negligible correlations were identified for GFI (&#x03C1;=0.082, <italic>P</italic>=.01) and CLI (&#x03C1;=&#x2212;0.111, <italic>P</italic>&#x003C;.001). Word count was not significantly correlated with FKGL (&#x03C1;=0.012, <italic>P</italic>=.70), FRES (&#x03C1;=0.032, <italic>P</italic>=.320), or the SMOG Index (&#x03C1;=0.027, <italic>P</italic>=.40). These findings suggest that summary length alone did not meaningfully determine readability level within this cohort. The complete correlation results are summarized in <xref ref-type="table" rid="table3">Table 3</xref>.</p><table-wrap id="t3" position="float"><label>Table 3.</label><caption><p>Spearman rho correlation of word count with readability scores.</p></caption><table id="table3" frame="hsides" rules="groups"><thead><tr><td align="left" valign="bottom">Readability scale</td><td align="left" valign="bottom">Spearman rho coefficient (&#x03C1;)</td><td align="left" valign="bottom"><italic>P</italic> value</td><td align="left" valign="bottom">95% CI</td></tr></thead><tbody><tr><td align="left" valign="top">Flesch-Kincaid Grade Level</td><td align="left" valign="top">0.012</td><td align="left" valign="top">.70</td><td align="left" valign="top">&#x2212;0.054 to 0.071</td></tr><tr><td align="left" valign="top">Flesch Reading Ease Score</td><td align="left" valign="top">0.032</td><td align="left" valign="top">.32</td><td align="left" valign="top">&#x2212;0.037 to 0.108</td></tr><tr><td align="left" valign="top">Gunning Fog Index</td><td align="left" valign="top">0.082</td><td align="left" valign="top">.01</td><td align="left" valign="top">0.020 to 0.137</td></tr><tr><td align="left" valign="top">Coleman Liau Index</td><td align="left" valign="top">&#x2212;0.111</td><td align="left" valign="top">&#x003C;.001</td><td align="left" valign="top">&#x2212;0.174 to &#x2212;0.049</td></tr><tr><td align="left" valign="top">SMOG Index</td><td align="left" valign="top">0.027</td><td align="left" valign="top">.40</td><td align="left" valign="top">&#x2212;0.037 to 0.085</td></tr></tbody></table></table-wrap></sec></sec><sec id="s4" sec-type="discussion"><title>Discussion</title><sec id="s4-1"><title>Principal Findings</title><p>Effective patient-provider communication is a cornerstone of quality health care, particularly in specialized fields such as orthopedic surgery, where the complexity of information can often be overwhelming. This study evaluated 982 AI-generated PVS and found an average readability corresponding to a ninth-grade reading level. Only 0.4% (4/982) of summaries met the sixth-grade threshold and 14.2% (139/982) met the eighth-grade benchmark. In the absence of a comparison group, conclusions regarding relative improvement over traditional methods could not be established. These findings establish a quantitative readability baseline for AI scribe&#x2013;generated PVS in an orthopedic setting and highlight the need for algorithmic refinement to better align patient-facing outputs with recommended health literacy standards.</p><p>The mean FKGL of 9.3 aligns with prior studies, showing that most orthopedic educational materials exceed national literacy recommendations. For example, Karimi et al [<xref ref-type="bibr" rid="ref8">8</xref>] reported FKGL scores ranging from 10 to 13 in total joint arthroplasty education materials, while Sahhar et al [<xref ref-type="bibr" rid="ref3">3</xref>] found that 100% and 76.3% of preoperative orthopedic education articles exceeded sixth- and eighth-grade reading levels, respectively. Similarly, Gerhold et al [<xref ref-type="bibr" rid="ref30">30</xref>] reported that online orthopedic trauma materials had an average FKGL of 8.7, with only 1.25% at or below a sixth-grade level. These findings collectively illustrate that even modest deviations from national readability standards can limit comprehension, satisfaction, and engagement [<xref ref-type="bibr" rid="ref3">3</xref>,<xref ref-type="bibr" rid="ref5">5</xref>].</p><p>Beyond orthopedics, previous analyses of electronic health record-generated or online patient education materials across medical specialties (eg, internal medicine, cardiology, and radiology) have identified similar mismatches between readability and average US literacy levels, suggesting that this is a system-wide issue rather than specialty-specific [<xref ref-type="bibr" rid="ref31">31</xref>-<xref ref-type="bibr" rid="ref33">33</xref>]. By contextualizing AI-generated PVS within this broader landscape, this study contributes baseline data on the readability performance of AI-generated documentation relative to existing patient-facing materials.</p><p>Spearman correlations between word count and readability metrics were uniformly weak (|&#x03C1;|&#x003C;0.15) across all 5 indices. Although GFI (&#x03C1;=0.082, <italic>P</italic>=.01) and CLI (&#x03C1;=&#x2212;0.111, <italic>P</italic>&#x003C;.001) reached statistical significance, FKGL, FRES, and SMOG did not. The small effect sizes observed across all indices suggest that statistical significance in GFI and CLI likely reflected the large sample size rather than a clinically meaningful relationship. These findings suggest that sentence structure and vocabulary complexity, rather than overall summary length, were the primary drivers of readability in AI-generated PVS.</p><p>Beyond readability scores, limited health literacy has been consistently linked to worse surgical outcomes, higher health care utilization, and reduced adherence [<xref ref-type="bibr" rid="ref3">3</xref>,<xref ref-type="bibr" rid="ref5">5</xref>]. Inadequately readable summaries risk compounding these issues, as comprehension is essential for informed consent, medication adherence, and effective postoperative care.</p><p>AI scribe platforms represent a promising digital health intervention that could bridge these gaps. Unlike static templates, generative AI can be tailored to match a patient&#x2019;s literacy level, potentially generating multiple versions of a summary based on individual needs. This adaptability addresses the long-standing limitations of traditional visit summaries and autogenerated materials [<xref ref-type="bibr" rid="ref17">17</xref>]. Emerging AI platforms, including the one evaluated in this study, continue to evolve and have the potential to incorporate real-time feedback loops that dynamically adapt readability. Optimizing these systems with plain language frameworks, algorithmic adjustments, and targeted prompting strategies may enhance comprehension without sacrificing accuracy. The finding that 14.2% of summaries already met the eighth-grade threshold compared with only 0.4% meeting the sixth-grade standard suggests that a tiered optimization approach targeting the eighth-grade benchmark first may represent a more achievable near-term goal. This is particularly relevant given that specialty-specific medical terminology may inherently constrain readability. However, any efforts to improve readability must not compromise clinical accuracy. AI-generated content carries an inherent risk of factual errors and hallucinations, and readability optimization strategies must be validated to ensure language simplification does not introduce clinically unsafe inaccuracies. The balance between language simplification and preservation of medical precision remains a critical design challenge that must be validated through further patient-centered testing in the future.</p></sec><sec id="s4-2"><title>Strengths and Limitations</title><p>This study had several strengths, including its large sample size, focus on a specific clinical context, and the use of multiple established readability formulas to benchmark AI-generated PVS. However, this study has several limitations that should be considered.</p><p>First, traditional readability formulas are deterministic, surface-level metrics designed for static human-authored text. AI-generated content carries additional interpretive limitations. A summary may achieve a favorable FKGL through short sentences while remaining clinically incoherent, and a high grade level does not in itself indicate an inaccurate summary. The strong concordance across 5 indices (<italic>W</italic>=0.882, <italic>P</italic>&#x003C;.001) confirms consistent rank ordering of surface linguistic complexity rather than semantic coherence, clinical fidelity, or actual patient comprehension. Future research should incorporate validated qualitative instruments, such as the Patient Education Materials Assessment Tool and DISCERN tool, to evaluate comprehensibility and accuracy beyond surface readability metrics.</p><p>Second, the absence of real patients in this study limited the ability to draw conclusions about the practical impact of AI-generated summaries on patient understanding. The retrospective design precluded the collection of comprehension data or usability feedback, and the findings represent a single point in time that may evolve as AI scribe technology improves.</p><p>As all analyzed summaries reflected the raw, prereview AI output, the cognitive burden of simplifying a mean FKGL of 9.3 to the recommended sixth-grade level standard falls on the reviewing provider, a task that may partially offset the time-saving premise of AI scribes.</p><p>Third, the absence of demographic data such as education level, language proficiency, and health literacy hindered the evaluation of how patient characteristics affect summary complexity. This analysis was conducted in an academic joint replacement clinic, potentially introducing a selection bias. In this context, patients may possess higher health literacy or educational levels compared to the general population, which may limit the generalizability of our findings to other populations. However, prior studies have reported that orthopedic patients generally demonstrate lower health literacy than broader medical populations, suggesting that readability barriers may persist despite the academic setting [<xref ref-type="bibr" rid="ref8">8</xref>]. Additionally, academic centers typically manage higher complexity cases compared to community practices, potentially introducing more specialized terminology and inflation of measured readability relative to a community orthopedic setting. Fourth, the evaluation of a single commercial AI scribe system without readability-optimized prompting limits generalizability to other models or institutions [<xref ref-type="bibr" rid="ref9">9</xref>]. Although this allowed for a real-world evaluation of the platform&#x2019;s baseline output, we recognize that AI systems can be tailored through prompting or algorithmic adjustments to produce content that better aligns with national readability standards, as demonstrated in previous studies [<xref ref-type="bibr" rid="ref9">9</xref>]. Finally, the lack of a preimplementation baseline using physician-generated visit summaries restricted direct comparisons. There are no previous studies evaluating the readability of visit summaries in orthopedic surgery, as the only published benchmark for physician-generated visit summaries was conducted in the fields of internal medicine and family medicine [<xref ref-type="bibr" rid="ref31">31</xref>], which are not comparable to orthopedic surgery settings.</p><p>Future research should incorporate patient-reported outcomes and comprehension testing to assess whether enhanced readability leads to tangible improvements in understanding and adherence. Comparative evaluations across multiple generative AI systems, with and without readability optimization, are required to assess platform-level variability. Future investigations should focus on developing advanced readability scales that integrate semantic and contextual understanding, rather than depending solely on sentence structure. Future trials should examine how AI-generated summaries affect patient engagement, provider workload, and the quality of documentation. Finally, ongoing reevaluation of the sixth- and eighth-grade benchmarks is necessary to determine whether these standards remain appropriate for contemporary digital health communication.</p></sec><sec id="s4-3"><title>Conclusions</title><p>In this orthopedic outpatient setting, AI-generated PVS consistently exceeded recommended patient literacy standards, with a mean FKGL of 9.3 and only 14.2% of summaries meeting the eighth-grade reading standard. In the absence of a contemporaneous comparison group, conclusions regarding relative improvement over traditional methods could not be established. Future studies should evaluate whether tailored strategies, patient comprehension testing, and multiplatform comparisons can further enhance readability and improve health communication.</p></sec></sec></body><back><ack><p>No generative AI writing tools were used in the preparation of this manuscript.</p></ack><notes><sec><title>Funding</title><p>This study received no funding.</p></sec><sec><title>Data Availability</title><p>The datasets generated and/or analyzed during the current study are not publicly available due to restrictions related to patient privacy, institutional data governance, and the proprietary nature of the AI scribe platform but are available from the corresponding author upon reasonable request and with appropriate institutional approvals.</p></sec></notes><fn-group><fn fn-type="con"><p>Conceptualization: SB</p><p>Data curation: ELM III, VPS, JRG</p><p>Formal analysis: VPS (lead), ELM III (supporting)</p><p>Investigation: ELM III, ANC</p><p>Methodology: SB</p><p>Supervision: SB, CM</p><p>Visualization: VPS (lead), ELM III (supporting)</p><p>Writing &#x2013; original draft: ELM III (lead), VPS (equal), ANC (supporting)</p><p>Writing &#x2013; review &#x0026; editing: ELM III (lead), VPS (supporting), JRG (supporting), CM (supporting), SB (supporting)</p></fn><fn fn-type="conflict"><p>None declared.</p></fn></fn-group><glossary><title>Abbreviations</title><def-list><def-item><term id="abb1">CLI</term><def><p>Coleman-Liau Index</p></def></def-item><def-item><term id="abb2">FKGL</term><def><p>Flesch-Kincaid Grade Level</p></def></def-item><def-item><term id="abb3">FRES</term><def><p>Flesch Reading Ease Score</p></def></def-item><def-item><term id="abb4">GFI</term><def><p>Gunning Fog Index</p></def></def-item><def-item><term id="abb5">PVS</term><def><p>patient visit summary</p></def></def-item><def-item><term id="abb6">SMOG</term><def><p>Simple Measure of Gobbledygook</p></def></def-item></def-list></glossary><ref-list><title>References</title><ref id="ref1"><label>1</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Lum</surname><given-names>ZC</given-names> </name><name name-style="western"><surname>Lyles</surname><given-names>CR</given-names> </name></person-group><article-title>What&#x2019;s important: health literacy in orthopaedics</article-title><source>J Bone Joint Surg Am</source><year>2024</year><month>11</month><day>6</day><volume>106</volume><issue>21</issue><fpage>2042</fpage><lpage>2044</lpage><pub-id pub-id-type="doi">10.2106/JBJS.24.00367</pub-id><pub-id pub-id-type="medline">38896658</pub-id></nlm-citation></ref><ref id="ref2"><label>2</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Lans</surname><given-names>A</given-names> </name><name name-style="western"><surname>Schwab</surname><given-names>JH</given-names> </name></person-group><article-title>Health literacy in orthopaedics</article-title><source>J Am Acad Orthop Surg</source><year>2023</year><month>04</month><day>15</day><volume>31</volume><issue>8</issue><fpage>382</fpage><lpage>388</lpage><pub-id pub-id-type="doi">10.5435/JAAOS-D-22-01026</pub-id><pub-id pub-id-type="medline">36884220</pub-id></nlm-citation></ref><ref id="ref3"><label>3</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Sahhar</surname><given-names>M</given-names> </name><name name-style="western"><surname>Singh</surname><given-names>M</given-names> </name><name name-style="western"><surname>Mehta</surname><given-names>T</given-names> </name><etal/></person-group><article-title>Lost in translation: preoperative orthopaedic education materials significantly exceed recommended reading levels</article-title><source>JB JS Open Access</source><year>2025</year><volume>10</volume><issue>3</issue><fpage>e25.00143</fpage><pub-id pub-id-type="doi">10.2106/JBJS.OA.25.00143</pub-id><pub-id pub-id-type="medline">40777214</pub-id></nlm-citation></ref><ref id="ref4"><label>4</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Miskiewicz</surname><given-names>M</given-names> </name><name name-style="western"><surname>Capotosto</surname><given-names>S</given-names> </name><name name-style="western"><surname>Ling</surname><given-names>K</given-names> </name><name name-style="western"><surname>Hance</surname><given-names>F</given-names> </name><name name-style="western"><surname>Wang</surname><given-names>E</given-names> </name></person-group><article-title>Readability analysis of patient education material on rotator cuff injuries from the top 25 ranking orthopaedic institutions</article-title><source>J Am Acad Orthop Surg Glob Res Rev</source><year>2024</year><month>05</month><day>1</day><volume>8</volume><issue>5</issue><fpage>e24.00085</fpage><pub-id pub-id-type="doi">10.5435/JAAOSGlobal-D-24-00085</pub-id><pub-id pub-id-type="medline">38722904</pub-id></nlm-citation></ref><ref id="ref5"><label>5</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Badarudeen</surname><given-names>S</given-names> </name><name name-style="western"><surname>Sabharwal</surname><given-names>S</given-names> </name></person-group><article-title>Assessing readability of patient education materials: current role in orthopaedics</article-title><source>Clin Orthop Relat Res</source><year>2010</year><month>10</month><volume>468</volume><issue>10</issue><fpage>2572</fpage><lpage>2580</lpage><pub-id pub-id-type="doi">10.1007/s11999-010-1380-y</pub-id><pub-id pub-id-type="medline">20496023</pub-id></nlm-citation></ref><ref id="ref6"><label>6</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Eltorai</surname><given-names>AEM</given-names> </name><name name-style="western"><surname>Sharma</surname><given-names>P</given-names> </name><name name-style="western"><surname>Wang</surname><given-names>J</given-names> </name><name name-style="western"><surname>Daniels</surname><given-names>AH</given-names> </name></person-group><article-title>Most American Academy of Orthopaedic Surgeons&#x2019; online patient education material exceeds average patient reading level</article-title><source>Clinical Orthopaedics &#x0026; Related Research</source><year>2015</year><volume>473</volume><issue>4</issue><fpage>1181</fpage><lpage>1186</lpage><pub-id pub-id-type="doi">10.1007/s11999-014-4071-2</pub-id></nlm-citation></ref><ref id="ref7"><label>7</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>&#x00D3; Doinn</surname><given-names>T</given-names> </name><name name-style="western"><surname>Broderick</surname><given-names>JM</given-names> </name><name name-style="western"><surname>Clarke</surname><given-names>R</given-names> </name><name name-style="western"><surname>Hogan</surname><given-names>N</given-names> </name></person-group><article-title>Readability of patient educational materials in sports medicine</article-title><source>Orthop J Sports Med</source><year>2022</year><month>05</month><volume>10</volume><issue>5</issue><fpage>23259671221092356</fpage><pub-id pub-id-type="doi">10.1177/23259671221092356</pub-id><pub-id pub-id-type="medline">35547607</pub-id></nlm-citation></ref><ref id="ref8"><label>8</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Karimi</surname><given-names>AH</given-names> </name><name name-style="western"><surname>Shah</surname><given-names>AK</given-names> </name><name name-style="western"><surname>Hecht</surname><given-names>CJ</given-names>  <suffix>II</suffix></name><name name-style="western"><surname>Burkhart</surname><given-names>RJ</given-names> </name><name name-style="western"><surname>Acu&#x00F1;a</surname><given-names>AJ</given-names> </name><name name-style="western"><surname>Kamath</surname><given-names>AF</given-names> </name></person-group><article-title>Readability of online patient education materials for total joint arthroplasty: a systematic review</article-title><source>J Arthroplasty</source><year>2023</year><month>07</month><volume>38</volume><issue>7</issue><fpage>1392</fpage><lpage>1399</lpage><pub-id pub-id-type="doi">10.1016/j.arth.2023.01.032</pub-id><pub-id pub-id-type="medline">36716898</pub-id></nlm-citation></ref><ref id="ref9"><label>9</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Kirchner</surname><given-names>GJ</given-names> </name><name name-style="western"><surname>Kim</surname><given-names>RY</given-names> </name><name name-style="western"><surname>Weddle</surname><given-names>JB</given-names> </name><name name-style="western"><surname>Bible</surname><given-names>JE</given-names> </name></person-group><article-title>Can artificial intelligence improve the readability of patient education materials?</article-title><source>Clin Orthop Relat Res</source><year>2023</year><month>11</month><day>1</day><volume>481</volume><issue>11</issue><fpage>2260</fpage><lpage>2267</lpage><pub-id pub-id-type="doi">10.1097/CORR.0000000000002668</pub-id><pub-id pub-id-type="medline">37116006</pub-id></nlm-citation></ref><ref id="ref10"><label>10</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Salmon</surname><given-names>C</given-names> </name><name name-style="western"><surname>O&#x2019;Conor</surname><given-names>R</given-names> </name><name name-style="western"><surname>Singh</surname><given-names>S</given-names> </name><etal/></person-group><article-title>Characteristics of outpatient clinical summaries in the United States</article-title><source>Int J Med Inform</source><year>2016</year><month>10</month><volume>94</volume><fpage>75</fpage><lpage>80</lpage><pub-id pub-id-type="doi">10.1016/j.ijmedinf.2016.06.005</pub-id><pub-id pub-id-type="medline">27573314</pub-id></nlm-citation></ref><ref id="ref11"><label>11</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Federman</surname><given-names>AD</given-names> </name><name name-style="western"><surname>Sanchez-Munoz</surname><given-names>A</given-names> </name><name name-style="western"><surname>Jandorf</surname><given-names>L</given-names> </name><name name-style="western"><surname>Salmon</surname><given-names>C</given-names> </name><name name-style="western"><surname>Wolf</surname><given-names>MS</given-names> </name><name name-style="western"><surname>Kannry</surname><given-names>J</given-names> </name></person-group><article-title>Patient and clinician perspectives on the outpatient after-visit summary: a qualitative study to inform improvements in visit summary design</article-title><source>J Am Med Inform Assoc</source><year>2017</year><month>04</month><day>1</day><volume>24</volume><issue>e1</issue><fpage>e61</fpage><lpage>e68</lpage><pub-id pub-id-type="doi">10.1093/jamia/ocw106</pub-id><pub-id pub-id-type="medline">27497793</pub-id></nlm-citation></ref><ref id="ref12"><label>12</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Cai</surname><given-names>P</given-names> </name><name name-style="western"><surname>Liu</surname><given-names>F</given-names> </name><name name-style="western"><surname>Bajracharya</surname><given-names>A</given-names> </name><etal/></person-group><article-title>Generation of patient after-visit summaries to support physicians</article-title><source>Proceedings of the 29th International Conference on Computational Linguistics</source><year>2022</year><access-date>2026-08-11</access-date><fpage>6234</fpage><lpage>6247</lpage><comment><ext-link ext-link-type="uri" xlink:href="https://aclanthology.org/2022.coling-1.544/">https://aclanthology.org/2022.coling-1.544/</ext-link></comment></nlm-citation></ref><ref id="ref13"><label>13</label><nlm-citation citation-type="confproc"><person-group person-group-type="author"><name name-style="western"><surname>Acharya</surname><given-names>S</given-names> </name><name name-style="western"><surname>Boyd</surname><given-names>AD</given-names> </name><name name-style="western"><surname>Cameron</surname><given-names>R</given-names> </name><etal/></person-group><article-title>Incorporating personalization features in a hospital-stay summary generation system</article-title><year>2019</year><month>01</month><day>8</day><conf-name>Hawaii International Conference on System Sciences</conf-name><conf-date>Jan 8-11, 2019</conf-date><pub-id pub-id-type="doi">10.24251/HICSS.2019.505</pub-id></nlm-citation></ref><ref id="ref14"><label>14</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Krauss</surname><given-names>O</given-names> </name><name name-style="western"><surname>Franz</surname><given-names>B</given-names> </name><name name-style="western"><surname>Schuler</surname><given-names>A</given-names> </name></person-group><article-title>Automated on-demand generation of patient summary documents</article-title><source>International Journal of Electronics and Telecommunications</source><year>2015</year><month>06</month><volume>61</volume><issue>2</issue><fpage>151</fpage><lpage>157</lpage><pub-id pub-id-type="doi">10.1515/eletel-2015-0019</pub-id></nlm-citation></ref><ref id="ref15"><label>15</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Tremoulet</surname><given-names>P</given-names> </name><name name-style="western"><surname>Krishnan</surname><given-names>R</given-names> </name><name name-style="western"><surname>Karavite</surname><given-names>D</given-names> </name><etal/></person-group><article-title>A heuristic evaluation to assess use of after visit summaries for supporting continuity of care</article-title><source>Appl Clin Inform</source><year>2018</year><month>07</month><volume>9</volume><issue>3</issue><fpage>714</fpage><lpage>724</lpage><pub-id pub-id-type="doi">10.1055/s-0038-1668093</pub-id><pub-id pub-id-type="medline">30208496</pub-id></nlm-citation></ref><ref id="ref16"><label>16</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Murugesu</surname><given-names>L</given-names> </name><name name-style="western"><surname>Heijmans</surname><given-names>M</given-names> </name><name name-style="western"><surname>Rademakers</surname><given-names>J</given-names> </name><name name-style="western"><surname>Fransen</surname><given-names>MP</given-names> </name></person-group><article-title>Challenges and solutions in communication with patients with low health literacy: perspectives of healthcare providers</article-title><source>PLoS One</source><year>2022</year><volume>17</volume><issue>5</issue><fpage>e0267782</fpage><pub-id pub-id-type="doi">10.1371/journal.pone.0267782</pub-id><pub-id pub-id-type="medline">35507632</pub-id></nlm-citation></ref><ref id="ref17"><label>17</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Barak-Corren</surname><given-names>Y</given-names> </name><name name-style="western"><surname>Wolf</surname><given-names>R</given-names> </name><name name-style="western"><surname>Rozenblum</surname><given-names>R</given-names> </name><etal/></person-group><article-title>Harnessing the power of generative AI for clinical summaries: perspectives from emergency physicians</article-title><source>Ann Emerg Med</source><year>2024</year><month>08</month><volume>84</volume><issue>2</issue><fpage>128</fpage><lpage>138</lpage><pub-id pub-id-type="doi">10.1016/j.annemergmed.2024.01.039</pub-id><pub-id pub-id-type="medline">38483426</pub-id></nlm-citation></ref><ref id="ref18"><label>18</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Yim</surname><given-names>WW</given-names> </name><name name-style="western"><surname>Fu</surname><given-names>Y</given-names> </name><name name-style="western"><surname>Ben Abacha</surname><given-names>A</given-names> </name><name name-style="western"><surname>Snider</surname><given-names>N</given-names> </name><name name-style="western"><surname>Lin</surname><given-names>T</given-names> </name><name name-style="western"><surname>Yetisgen</surname><given-names>M</given-names> </name></person-group><article-title>Aci-bench: a novel ambient clinical intelligence dataset for benchmarking automatic visit note generation</article-title><source>Sci Data</source><year>2023</year><month>09</month><day>6</day><volume>10</volume><issue>1</issue><fpage>586</fpage><pub-id pub-id-type="doi">10.1038/s41597-023-02487-3</pub-id><pub-id pub-id-type="medline">37673893</pub-id></nlm-citation></ref><ref id="ref19"><label>19</label><nlm-citation citation-type="web"><source>Abridge for Clinicians</source><year>2024</year><access-date>2025-09-01</access-date><comment><ext-link ext-link-type="uri" xlink:href="https://www.abridge.com/">https://www.abridge.com/</ext-link></comment></nlm-citation></ref><ref id="ref20"><label>20</label><nlm-citation citation-type="journal"><article-title>World Medical Association Declaration of Helsinki</article-title><source>JAMA</source><year>2000</year><month>12</month><day>20</day><volume>284</volume><issue>23</issue><fpage>3043</fpage><pub-id pub-id-type="doi">10.1001/jama.284.23.3043</pub-id></nlm-citation></ref><ref id="ref21"><label>21</label><nlm-citation citation-type="book"><person-group person-group-type="author"><name name-style="western"><surname>Flesch</surname><given-names>R</given-names> </name></person-group><source>How to Write Plain English: A Book for Lawyers and Consumers</source><year>1979</year><publisher-name>Harper &#x0026; Row</publisher-name><pub-id pub-id-type="other">9780060112783</pub-id></nlm-citation></ref><ref id="ref22"><label>22</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Friedman</surname><given-names>DB</given-names> </name><name name-style="western"><surname>Hoffman-Goetz</surname><given-names>L</given-names> </name></person-group><article-title>A systematic review of readability and comprehension instruments used for print and web-based cancer information</article-title><source>Health Educ Behav</source><year>2006</year><month>06</month><volume>33</volume><issue>3</issue><fpage>352</fpage><lpage>373</lpage><pub-id pub-id-type="doi">10.1177/1090198105277329</pub-id><pub-id pub-id-type="medline">16699125</pub-id></nlm-citation></ref><ref id="ref23"><label>23</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Costantini</surname><given-names>H</given-names> </name><name name-style="western"><surname>Fuse</surname><given-names>R</given-names> </name></person-group><article-title>Health information on COVID-19 vaccination: readability of online sources and newspapers in Singapore, Hong Kong, and the Philippines</article-title><source>Journalism and Media</source><year>2022</year><volume>3</volume><issue>1</issue><fpage>228</fpage><lpage>237</lpage><pub-id pub-id-type="doi">10.3390/journalmedia3010017</pub-id></nlm-citation></ref><ref id="ref24"><label>24</label><nlm-citation citation-type="other"><person-group person-group-type="author"><name name-style="western"><surname>Lee</surname><given-names>BW</given-names> </name><name name-style="western"><surname>Lee</surname><given-names>JJ</given-names> </name></person-group><article-title>Traditional readability formulas compared for English</article-title><source>arXiv</source><comment>Preprint posted online on  Jan 8, 2023</comment><pub-id pub-id-type="doi">10.48550/arXiv.2301.02975</pub-id></nlm-citation></ref><ref id="ref25"><label>25</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Perni</surname><given-names>S</given-names> </name><name name-style="western"><surname>Rooney</surname><given-names>MK</given-names> </name><name name-style="western"><surname>Horowitz</surname><given-names>DP</given-names> </name><etal/></person-group><article-title>Assessment of use, specificity, and readability of written clinical informed consent forms for patients with cancer undergoing radiotherapy</article-title><source>JAMA Oncol</source><year>2019</year><month>08</month><day>1</day><volume>5</volume><issue>8</issue><fpage>e190260</fpage><pub-id pub-id-type="doi">10.1001/jamaoncol.2019.0260</pub-id><pub-id pub-id-type="medline">31046122</pub-id></nlm-citation></ref><ref id="ref26"><label>26</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Onder</surname><given-names>CE</given-names> </name><name name-style="western"><surname>Koc</surname><given-names>G</given-names> </name><name name-style="western"><surname>Gokbulut</surname><given-names>P</given-names> </name><name name-style="western"><surname>Taskaldiran</surname><given-names>I</given-names> </name><name name-style="western"><surname>Kuskonmaz</surname><given-names>SM</given-names> </name></person-group><article-title>Evaluation of the reliability and readability of ChatGPT-4 responses regarding hypothyroidism during pregnancy</article-title><source>Sci Rep</source><year>2024</year><month>01</month><day>2</day><volume>14</volume><issue>1</issue><fpage>243</fpage><pub-id pub-id-type="doi">10.1038/s41598-023-50884-w</pub-id><pub-id pub-id-type="medline">38167988</pub-id></nlm-citation></ref><ref id="ref27"><label>27</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>van Ballegooie</surname><given-names>C</given-names> </name><name name-style="western"><surname>Hoang</surname><given-names>P</given-names> </name></person-group><article-title>Assessment of the readability of online patient education material from major geriatric associations</article-title><source>J Am Geriatr Soc</source><year>2021</year><month>04</month><volume>69</volume><issue>4</issue><fpage>1051</fpage><lpage>1056</lpage><pub-id pub-id-type="doi">10.1111/jgs.16960</pub-id><pub-id pub-id-type="medline">33236778</pub-id></nlm-citation></ref><ref id="ref28"><label>28</label><nlm-citation citation-type="web"><source>Readable</source><access-date>2026-08-12</access-date><comment><ext-link ext-link-type="uri" xlink:href="https://app.readable.com/text">https://app.readable.com/text</ext-link></comment></nlm-citation></ref><ref id="ref29"><label>29</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Schober</surname><given-names>P</given-names> </name><name name-style="western"><surname>Boer</surname><given-names>C</given-names> </name><name name-style="western"><surname>Schwarte</surname><given-names>LA</given-names> </name></person-group><article-title>Correlation coefficients: appropriate use and interpretation</article-title><source>Anesth Analg</source><year>2018</year><month>05</month><volume>126</volume><issue>5</issue><fpage>1763</fpage><lpage>1768</lpage><pub-id pub-id-type="doi">10.1213/ANE.0000000000002864</pub-id><pub-id pub-id-type="medline">29481436</pub-id></nlm-citation></ref><ref id="ref30"><label>30</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Gerhold</surname><given-names>C</given-names> </name><name name-style="western"><surname>Bassin</surname><given-names>J</given-names> </name><name name-style="western"><surname>Alam</surname><given-names>H</given-names> </name><name name-style="western"><surname>Vemulapalli</surname><given-names>KC</given-names> </name></person-group><article-title>Elucidating orthopaedic trauma procedures: are the current patient educational materials useful?</article-title><source>Journal of Orthopaedic Experience &#x0026; Innovation</source><year>2025</year><pub-id pub-id-type="doi">10.60118/001c.132262</pub-id></nlm-citation></ref><ref id="ref31"><label>31</label><nlm-citation citation-type="thesis"><person-group person-group-type="author"><name name-style="western"><surname>Amundsen</surname><given-names>T</given-names> </name></person-group><article-title>Readability of after visit summaries: comparing the level of information in after visit summaries from internal medicine and family medicine residencies [dissertation]</article-title><year>2019</year><access-date>2026-08-12</access-date><publisher-name>The University of Arizona</publisher-name><comment><ext-link ext-link-type="uri" xlink:href="http://hdl.handle.net/10150/633413">http://hdl.handle.net/10150/633413</ext-link></comment></nlm-citation></ref><ref id="ref32"><label>32</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Sharma</surname><given-names>S</given-names> </name><name name-style="western"><surname>Latif</surname><given-names>Z</given-names> </name><name name-style="western"><surname>Makuvire</surname><given-names>TT</given-names> </name><etal/></person-group><article-title>Readability and accessibility of patient-education materials for heart failure in the United States</article-title><source>J Card Fail</source><year>2025</year><month>01</month><volume>31</volume><issue>1</issue><fpage>154</fpage><lpage>157</lpage><pub-id pub-id-type="doi">10.1016/j.cardfail.2024.06.015</pub-id><pub-id pub-id-type="medline">39094729</pub-id></nlm-citation></ref><ref id="ref33"><label>33</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Bange</surname><given-names>M</given-names> </name><name name-style="western"><surname>Huh</surname><given-names>E</given-names> </name><name name-style="western"><surname>Novin</surname><given-names>SA</given-names> </name><name name-style="western"><surname>Hui</surname><given-names>FK</given-names> </name><name name-style="western"><surname>Yi</surname><given-names>PH</given-names> </name></person-group><article-title>Readability of patient education materials from RadiologyInfo.org: has there been progress over the past 5 years?</article-title><source>AJR Am J Roentgenol</source><year>2019</year><month>10</month><volume>213</volume><issue>4</issue><fpage>875</fpage><lpage>879</lpage><pub-id pub-id-type="doi">10.2214/AJR.18.21047</pub-id><pub-id pub-id-type="medline">31386570</pub-id></nlm-citation></ref></ref-list></back></article>