<?xml version="1.0" encoding="UTF-8" standalone="no"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Archiving and Interchange DTD v2.3 20070202//EN" "archivearticle.dtd">
<article xml:lang="EN" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" article-type="methods-article">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Artif. Intell.</journal-id>
<journal-title>Frontiers in Artificial Intelligence</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Artif. Intell.</abbrev-journal-title>
<issn pub-type="epub">2624-8212</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="doi">10.3389/frai.2024.1375419</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Artificial Intelligence</subject>
<subj-group>
<subject>Methods</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>SATS: simplification aware text summarization of scientific documents</article-title>
</title-group>
<contrib-group>
<contrib contrib-type="author">
<name><surname>Zaman</surname> <given-names>Farooq</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/2636683/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Kamiran</surname> <given-names>Faisal</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name><surname>Shardlow</surname> <given-names>Matthew</given-names></name>
<xref ref-type="aff" rid="aff2"><sup>2</sup></xref>
<xref ref-type="corresp" rid="c001"><sup>&#x0002A;</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/1699330/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Hassan</surname> <given-names>Saeed-Ul</given-names></name>
<xref ref-type="aff" rid="aff2"><sup>2</sup></xref>
<xref ref-type="author-notes" rid="fn002"><sup>&#x02020;</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Karim</surname> <given-names>Asim</given-names></name>
<xref ref-type="aff" rid="aff3"><sup>3</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Aljohani</surname> <given-names>Naif Radi</given-names></name>
<xref ref-type="aff" rid="aff4"><sup>4</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
</contrib-group>
<aff id="aff1"><sup>1</sup><institution>Scientometrics Lab, Information Technology University</institution>, <addr-line>Lahore</addr-line>, <country>Pakistan</country></aff>
<aff id="aff2"><sup>2</sup><institution>Department of Computing and Mathematics, Manchester Metropolitan University</institution>, <addr-line>Manchester</addr-line>, <country>United Kingdom</country></aff>
<aff id="aff3"><sup>3</sup><institution>Department of Computer Science, Syed Babar Ali School of Science and Engineering (SBASSE), Lahore University of Management Sciences</institution>, <addr-line>Lahore</addr-line>, <country>Pakistan</country></aff>
<aff id="aff4"><sup>4</sup><institution>Information Systems Department, Faculty of Computing and Information Technology, King Abdulaziz University</institution>, <addr-line>Jeddah</addr-line>, <country>Saudi Arabia</country></aff>
<author-notes>
<fn fn-type="edited-by"><p>Edited by: Arkaitz Zubiaga, Queen Mary University of London, United Kingdom</p></fn>
<fn fn-type="edited-by"><p>Reviewed by: Christina Niklaus, University of St. Gallen, Switzerland</p>
<p>Diego Molla, Macquarie University, Australia</p></fn>
<corresp id="c001">&#x0002A;Correspondence: Matthew Shardlow <email>m.shardlow&#x00040;mmu.ac.uk</email></corresp>
<fn fn-type="equal" id="fn002"><p>&#x02020;Deceased</p></fn></author-notes>
<pub-date pub-type="epub">
<day>10</day>
<month>07</month>
<year>2024</year>
</pub-date>
<pub-date pub-type="collection">
<year>2024</year>
</pub-date>
<volume>7</volume>
<elocation-id>1375419</elocation-id>
<history>
<date date-type="received">
<day>23</day>
<month>01</month>
<year>2024</year>
</date>
<date date-type="accepted">
<day>06</day>
<month>06</month>
<year>2024</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#x000A9; 2024 Zaman, Kamiran, Shardlow, Hassan, Karim and Aljohani.</copyright-statement>
<copyright-year>2024</copyright-year>
<copyright-holder>Zaman, Kamiran, Shardlow, Hassan, Karim and Aljohani</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/"><p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p></license>
</permissions>
<abstract>
<p>Simplifying summaries of scholarly publications has been a popular method for conveying scientific discoveries to a broader audience. While text summarization aims to shorten long documents, simplification seeks to reduce the complexity of a document. To accomplish these tasks collectively, there is a need to develop machine learning methods to shorten and simplify longer texts. This study presents a new Simplification Aware Text Summarization model (SATS) based on future n-gram prediction. The proposed SATS model extends ProphetNet, a text summarization model, by enhancing the objective function using a word frequency lexicon for simplification tasks. We have evaluated the performance of SATS on a recently published text summarization and simplification corpus consisting of 5,400 scientific article pairs. Our results in terms of automatic evaluation demonstrate that SATS outperforms state-of-the-art models for simplification, summarization, and joint simplification-summarization across two datasets on ROUGE, SARI, and <bold>CSS<sub>1</sub></bold>. We also provide human evaluation of summaries generated by the SATS model. We evaluated 100 summaries from eight annotators for grammar, coherence, consistency, fluency, and simplicity. The average human judgment for all evaluated dimensions lies between 4.0 and 4.5 on a scale from 1 to 5 where 1 means low and 5 means high.</p></abstract>
<kwd-group>
<kwd>scientific documents</kwd>
<kwd>summarization</kwd>
<kwd>simplification</kwd>
<kwd>transformer model</kwd>
<kwd>deep learning</kwd>
</kwd-group>
<counts>
<fig-count count="4"/>
<table-count count="7"/>
<equation-count count="6"/>
<ref-count count="99"/>
<page-count count="15"/>
<word-count count="12352"/>
</counts>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-at-acceptance</meta-name>
<meta-value>Natural Language Processing</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec sec-type="intro" id="s1">
<title>1 Introduction</title>
<p>The automated generation of simplified summaries of scholarly articles is a popular mechanism for the public dissemination of scientific discoveries (Hou et al., <xref ref-type="bibr" rid="B34">2022</xref>). This task can be performed by leveraging text summarization (Tomer and Kumar, <xref ref-type="bibr" rid="B83">2022</xref>) and text simplification (Al-Thanyyan and Azmi, <xref ref-type="bibr" rid="B4">2021</xref>), which are well-defined tasks in Natural Language Processing (NLP) (Iqbal et al., <xref ref-type="bibr" rid="B35">2021</xref>). Text summarization is the task of reducing text size while maintaining the information presented in the text. Using text summarization, a long text is provided as input and by automated means a shorter and more concise version is generated (Cai et al., <xref ref-type="bibr" rid="B15">2021</xref>). Text simplification reduces the complexity of a text while retaining the original meaning by producing an easy-to-read version of the input source text (Shardlow, <xref ref-type="bibr" rid="B76">2014b</xref>; Alva-Manchego et al., <xref ref-type="bibr" rid="B7">2020b</xref>). Text simplification typically operates at the sentence level, with each complex sentence being transformed into a simpler alternative.</p>
<p>At the intersection of text simplification and summarization, lies the combination of these two tasks. The desired output is summarized (shorter in length) and simplified (simple to read and understand). We argue that, if a text is only summarized, technical words may remain, impeding a reader. If it is only simplified, the resulting text might be too long and repetitive. Therefore, having both simplification and summarization is vital for true understanding in heterogeneous areas such as public health literature, legal texts, or scientific communications.</p>
<p>As an example, consider two pairs of texts in <xref ref-type="table" rid="T1">Table 1</xref>. In source 1 and source 2, there are technical terms such as &#x0201C;locus,&#x0201D; &#x0201C;perturb,&#x0201D; &#x0201C;microbial,&#x0201D; &#x0201C;apex predator,&#x0201D; &#x0201C;juxtacrine axo-glial,&#x0201D; and &#x0201C;myelination.&#x0201D; These technical terms make it difficult for a lay reader to digest; on the other hand, both the summaries are clear, concise, and easy to consume and digest.</p>
<table-wrap position="float" id="T1">
<label>Table 1</label>
<caption><p>Examples of simple summaries.</p></caption>
<table frame="box" rules="all">
<thead>
<tr>
<th valign="top" align="left"><bold>Source 1</bold></th>
<th valign="top" align="left"><bold>Single gene locus changes perturb complex microbial communities</bold></th>
</tr>
<tr>
<th/>
<th valign="top" align="left"><bold>as much as apex predator loss</bold></th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="left">Summary 1</td>
<td valign="top" align="left">Genetic mutants alter entire biological communities</td>
</tr> <tr>
<td valign="top" align="left">Source 2</td>
<td valign="top" align="left">Spatial mapping of juxtacrine axoglial interactions identifies novel molecules</td>
</tr>
<tr>
<td/>
<td valign="top" align="left">in peripheral myelination</td>
</tr>
<tr>
<td valign="top" align="left">Summary 2</td>
<td valign="top" align="left">New technique lets scientists see and study the interface where two cells touch</td>
</tr></tbody>
</table>
</table-wrap>


<p>The key focus of this study is to model the task of summarization and simplification to be executed simultaneously using deep learning methods (Zerva et al., <xref ref-type="bibr" rid="B93">2020</xref>); for this objective, we have extended the ProphetNet (Qi et al., <xref ref-type="bibr" rid="B68">2020</xref>) architecture to perform simplification and summarization.</p>
<p>The main contributions of this study are as follows: (a) we propose a simplification aware loss function for the joint task of simplification and summarization as explained in Section 3.1; (b) we evaluate the proposed model against two datasets; Eureka (Zaman et al., <xref ref-type="bibr" rid="B92">2020</xref>) and CNN-Daily Mail (See et al., <xref ref-type="bibr" rid="B74">2017</xref>); and (c) we demonstrate that our model outperforms the previous state of the art on combined summarization and simplification.</p>
<p>The rest of the study is organized as follows: Section 2 presents related work on recent summarization, simplification, and evaluation indices. Section 3 presents data and methods, including state-of-the-art models. While Section 4 discusses the results of this study, Section 5 presents concluding remarks.</p></sec>
<sec id="s2">
<title>2 Related work</title>
<p>Text summarization is the task of compressing a text by automated means while preserving meaning. Various NLP approaches exist for text summarization (Mao et al., <xref ref-type="bibr" rid="B53">2021</xref>; Suleiman and Awajan, <xref ref-type="bibr" rid="B80">2022</xref>). There are two ways of achieving text summarization: <italic>extractive</italic> and <italic>abstractive</italic>.</p>
<sec>
<title>2.1 Extractive summarization</title>
<p>In extractive text summarization, the essential sections of the original document are first marked, and then those important sections are arranged in a document to produce the summary. The first attempt at text summarization introduced automated ways of producing summaries (or abstracts) of the original documents (Baxendale, <xref ref-type="bibr" rid="B12">1958</xref>; Luhn, <xref ref-type="bibr" rid="B50">1958</xref>). It was found that the count of occurrence determines the importance of a word; hence, the importance of a sentence was obtained by counting the important words in that sentence (Luhn, <xref ref-type="bibr" rid="B50">1958</xref>). It was also found that the location of a sentence in a document determines its importance. It was further explored that in 85% of paragraphs, the important sentence is located at the start of the paragraph and 7% of the paragraphs it is located at the end (Baxendale, <xref ref-type="bibr" rid="B12">1958</xref>). Later, Edmundson (<xref ref-type="bibr" rid="B23">1969</xref>) manually wrote the summaries of 400 documents and used the word frequencies from Luhn (<xref ref-type="bibr" rid="B50">1958</xref>) and location-based sentence importance from Baxendale (<xref ref-type="bibr" rid="B12">1958</xref>) to produce summaries.</p>
<p>A supervised approach has been taken for extractive summarization (Collins et al., <xref ref-type="bibr" rid="B20">2017</xref>), and to extract essential sentences, regression was used as sentence scoring (Zopf et al., <xref ref-type="bibr" rid="B99">2018</xref>).</p>
<p>Extractive text summarization has been extended to multi-document level (Sanchez-Gomez et al., <xref ref-type="bibr" rid="B70">2020</xref>). In this study, the authors explored the Artificial Bee Colony algorithm for extractive summarization, and for feature extraction, TF-IDF approach was tailored. The authors formulate the summarization task as an optimization problem and set two objectives: (1) coverage and (2) redundancy. The coverage covers the main contents; redundancy is concerned with avoiding and controlling the redundant sentence selection in the output summary. Extractive text summarization has also been explored through the lens of query-based sentiment analysis (Sanchez-Gomez et al., <xref ref-type="bibr" rid="B71">2021</xref>). In this study, the authors explored multi-document summarization with query-based and sentiment score-based approaches; that is, the output summary will have an identical sentiment score to the input user query over a set of source documents.</p>
<p>Recently, deep learning has also been studied for extractive (Kinugawa and Tsuruoka, <xref ref-type="bibr" rid="B40">2017</xref>) and abstractive (Liu et al., <xref ref-type="bibr" rid="B46">2015</xref>; Chen and Bansal, <xref ref-type="bibr" rid="B17">2018</xref>) summarization. A study Mackie et al. (<xref ref-type="bibr" rid="B52">2014</xref>) revealed that sumBasic algorithm (Nenkova and Vanderwende, <xref ref-type="bibr" rid="B62">2005</xref>) performs well for the task of microblogs summarization. This approach is good for producing long and fluent text passages but produces non-factual text in the output summaries (See et al., <xref ref-type="bibr" rid="B74">2017</xref>). To address this problem, Generative Adversarial Neural Networks (Goodfellow et al., <xref ref-type="bibr" rid="B31">2014</xref>) have been applied (Liu et al., <xref ref-type="bibr" rid="B47">2018</xref>).</p>
<p>The performance of extractive summarization may be greatly improved by identifying structured elements in text (Filatova and Hatzivassiloglou, <xref ref-type="bibr" rid="B26">2004</xref>), across genres such as news articles (Thompson et al., <xref ref-type="bibr" rid="B82">2017</xref>) and medical text (Shardlow et al., <xref ref-type="bibr" rid="B77">2018</xref>).</p></sec>
<sec>
<title>2.2 Abstractive summarization</title>
<p>Abstractive summarization exploits deep learning methods such as sequence-to-sequence models to generate the output summary based on the input text (Liu et al., <xref ref-type="bibr" rid="B46">2015</xref>; Chen and Bansal, <xref ref-type="bibr" rid="B17">2018</xref>). The authors in Azmi and Altmami (<xref ref-type="bibr" rid="B9">2018</xref>) extended abstractive summarization with the granularity level of the generated summary to be used and controlled by the end-user. The authors in Mehta and Majumder (<xref ref-type="bibr" rid="B57">2018</xref>) studied the aggregation of multiple models for text summarization tasks. In a recent study (Barros et al., <xref ref-type="bibr" rid="B11">2019</xref>), the idea of using multi-document as input for text summarization has been studied, the authors applied text summarization for news articles, the idea used in this study is the narrative approach, events from different news articles were first extracted and then sorted over a timeline, after sorting, and for each event, one-line output has been generated based on the contextual information available in news articles. The authors reported that their proposed narrative approach generates better summaries than other state-of-the-art methods. Abstractive text summarization is also been studied in the biomedical domain (Van Veen et al., <xref ref-type="bibr" rid="B84">2023</xref>; Yang et al., <xref ref-type="bibr" rid="B91">2023</xref>). Large language models (LLM) have been studied for text summarization (Liu et al., <xref ref-type="bibr" rid="B48">2023</xref>), and it was found that the LLM generates very fluent summaries but commits mistakes such as generation of non-factual text in the output summary. Abstractive text summarization has been explored for the Greek Language (Giarelis et al., <xref ref-type="bibr" rid="B29">2023</xref>). More recently, a study Rehman et al. (<xref ref-type="bibr" rid="B69">2023</xref>) adp summarization.</p>
<p>The combination of abstractive methods with extractive methods has been explored recently (Cho et al., <xref ref-type="bibr" rid="B18">2019</xref>). Such work leads to selecting some parts from the source text and generating new content for diversification. The authors designed their setup with three decoders to generate three output summaries that are different from each other; for selecting contents from the source text, each chunk of the input text was scored; they call it &#x0201C;focus.&#x0201D; Instead of training both encoder and decoder, utilization of pre-trained encoder such as BERT (Devlin et al., <xref ref-type="bibr" rid="B21">2019</xref>) for the task of summarization both abstractive and extractive has been explored (Liu and Lapata, <xref ref-type="bibr" rid="B49">2019</xref>). In this study, the authors reported that setting a pipeline of the pre-trained encoder and a fine-tuned decoder improves the performance of text summarization methods. As pre-training improves the performance of text summarization methods, in Dong et al. (<xref ref-type="bibr" rid="B22">2019</xref>) self-supervised pre-training was explored. The authors used a transformer for summarization with the supervised objective in this study. Pre-training and self-supervised objectives were further explored for the task of summarization (Zhang et al., <xref ref-type="bibr" rid="B94">2019a</xref>), and the authors investigated the use of masking for extractive methods, that is, masking some tokens in the input text and treating them as targets to make the source-target pairs for supervised training. The self-supervised pre-training has been recently explored in Yan et al. (<xref ref-type="bibr" rid="B90">2020</xref>) for summarization. The authors used masking; some input tokens are marked as masks and presented as missing. The objective is set to predict those missing tokens, the masked input, and the masked token as target presented to the training model. The authors further explored the prediction/generation of more than one token at a time on the decoder side.</p>
<p>In See et al. (<xref ref-type="bibr" rid="B74">2017</xref>), the authors used pointer generator network for text summarization (Vinyals et al., <xref ref-type="bibr" rid="B86">2015</xref>). The output summary of this model is both abstractive and extractive. A mechanism called coverage was proposed to avoid the generation of repetitive tokens in the generated summary.</p></sec>
<sec>
<title>2.3 Text simplification</title>
<p>Automated text simplification is a new problem that lies under the domain of NLP, that is, a text modification process aiming at making the text easy to understand. In Hoard et al. (<xref ref-type="bibr" rid="B33">1992</xref>), the authors worked on writing technical manuals and assisting stroke survivors to read (Carroll et al., <xref ref-type="bibr" rid="B16">1998</xref>). In the academic literature, text simplification can be found in three categories: neural, syntactic, and lexical simplification.</p>
<p>In syntactic text simplification, the structure of the grammar is rewritten such that it transforms its constituents such as voices and narration or long sentences into small understandable pieces (Siddharthan, <xref ref-type="bibr" rid="B79">2014</xref>). This approach is different from the lexical approach (Shardlow, <xref ref-type="bibr" rid="B76">2014b</xref>). In the lexical approach, complex words are identified at first; then, their substitutes are generated. After that, the step of word sense disambiguation is performed. Finally, the substitutes are re-ranked, and a selection is performed for the final synonym of the original word. This approach produces the output with significant errors (Shardlow, <xref ref-type="bibr" rid="B75">2014a</xref>).</p>
<p>Neural machine translation models can be modified to generate a hybrid of syntactic and lexical text simplification. In Wubben et al. (<xref ref-type="bibr" rid="B88">2012</xref>) and Li et al. (<xref ref-type="bibr" rid="B42">2018</xref>) statistical machine translation model was used. Recently, neural machine translation has been applied for text simplification (Nisioi et al., <xref ref-type="bibr" rid="B64">2017</xref>). Further work has shown that the level of complexity of the output can be controlled (Agrawal and Carpuat, <xref ref-type="bibr" rid="B2">2019</xref>; Marchisio et al., <xref ref-type="bibr" rid="B54">2019</xref>; Nishihara et al., <xref ref-type="bibr" rid="B63">2019</xref>). Recent work on simplification has focused on the improvement of datasets (Alva-Manchego et al., <xref ref-type="bibr" rid="B5">2020a</xref>) and evaluation measures (Alva-Manchego et al., <xref ref-type="bibr" rid="B6">2019</xref>). In Macdonald and Siddharthan (<xref ref-type="bibr" rid="B51">2016</xref>) simplification was used as a preprocessing step prior to summarization of children&#x00027;s stories. Simplification was also used similarly more recently for clinical data summaries generation (Acharya et al., <xref ref-type="bibr" rid="B1">2019</xref>).</p>
<sec>
<title>2.3.1 Sentence compression</title>
<p>Sentence compression eliminates repeated information from a sentence while preserving key contents present in the original sentence. The first attempt without using linguistic features for sentence compression was made in 2015 by Katja et al. at Google (Filippova et al., <xref ref-type="bibr" rid="B27">2015</xref>). In this study, the authors proposed an approach based on LSTM for token deletion, which led to omitting the redundant information present in the sentence. The authors tuned and evaluated their proposed LSTM-based method against sentence compression dataset (Filippova and Altun, <xref ref-type="bibr" rid="B28">2013</xref>). The authors found that the simple LSTM-based method produces readable and more informative compression without using any syntactic information. The authors furthermore reported that syntactic information does not contribute to performance. Later, Wang et al. (<xref ref-type="bibr" rid="B87">2017</xref>) proposed an extended LSTM-based model which incorporates syntactic information of Part of speech and dependency parsing in the form of embedding, and the authors further imposed minimum and maximum sentence length constrained over the compressed output. Sentence compression methods perform well on short-length sentences (Kamigaito et al., <xref ref-type="bibr" rid="B37">2018</xref>). In recent work, sentence compression has been explored to improve the capability of sentence compression methods for long sentence lengths. The authors used attention-based weights and dependency trees to capture syntactic information. The proposed method is evaluated using token level F1 and Rouge scores (Lin and Och, <xref ref-type="bibr" rid="B45">2004</xref>) and reported that their proposed method outperforms state-of-the-art methods for sentence compression task (Kamigaito et al., <xref ref-type="bibr" rid="B37">2018</xref>). Sentence compression has also been studied in an unsupervised manner (Zhao et al., <xref ref-type="bibr" rid="B97">2018</xref>). In this study, a reinforcement-based approach is used for sentence compression. The authors used a language model as an evaluator; a series of deletion and evaluation passes were executed consecutively, such as first a deletion operation is performed, then an evaluation is performed to assess the correctness of the resultant sentence. The sentence compression task is studied to improve the grammar of the compressed sentences (Kamigaito and Okumura, <xref ref-type="bibr" rid="B38">2020</xref>). In this study, the author enhanced the decoder part of the LSTM-based model with extra information of the parent word and child word from the dependency tree.</p></sec>
<sec>
<title>2.3.2 Sentence splitting</title>
<p>Sentence splitting is the task of breaking longer sentences into shorter pieces so that each individual piece is a complete sentence; the aim is to reduce the complexity of longer sentences; this comes under text simplification. Sentence splitting can further be used in other NLP tasks such as machine translation. Splitting can also help second language learners to read a longer sentence in small pieces (Narayan et al., <xref ref-type="bibr" rid="B61">2017</xref>). Recently, Aharoni and Goldberg explored that the performance of sentence splitting can be improved by employing a copy mechanism; a copy mechanism is used in other tasks such as abstractive text summarization (See et al., <xref ref-type="bibr" rid="B74">2017</xref>). It was further observed that a unique train-validation-test split is required to validate the performance of sentence splitting (Aharoni and Goldberg, <xref ref-type="bibr" rid="B3">2018</xref>). Aharoni and Goldberg explored the dataset for sentence splitting should have train-validation-test sets to be unique and crafted carefully, they observed that in the previous studies, the dataset splitting was not carefully carried out, they further analyzed that validation and test sets contained more than 89% of unique simpler sentences from the train set, and this was making the sequence-to-sequence models to memorize the simpler sentences and thus leads to high-performance scores. Later, a new corpus for the task of sentence splitting is built (Botha et al., <xref ref-type="bibr" rid="B13">2018</xref>), with this new corpus, machine learning models are able to capture more information present in the corpus than before with the previous benchmark by Narayan et al. (<xref ref-type="bibr" rid="B61">2017</xref>), and using this, the models were trained with unique and disjointed train-validation-test sets. The authors established a new state of the art with almost double the previous results.</p></sec></sec>
<sec>
<title>2.4 Evaluation metrics for text simplification and summarization</title>
<p>To quantify the performance of any machine learning, deep learning or NLP task, one or more metrics are used, such as accuracy, precision, recall, and F1 measure. A good metric captures all the critical aspects of a task it is designed to evaluate. For text simplification, a widely used metric known as SARI (Xu et al., <xref ref-type="bibr" rid="B89">2016</xref>) is used, although some metrics such as BLEU are adopted from other NLP tasks such as machine translation (Papineni et al., <xref ref-type="bibr" rid="B67">2002</xref>). For the evaluation of summarization systems, the ROUGE metric (Lin and Hovy, <xref ref-type="bibr" rid="B44">2003</xref>) is used widely; however, due to the limitations of the existing ROUGE metric, it must be addressed carefully (Schluter, <xref ref-type="bibr" rid="B73">2017</xref>). Another evaluation method has been proposed in Jia et al. (<xref ref-type="bibr" rid="B36">2023</xref>) which evaluates consistency and faithfulness aspects. In the following subsections, we will discuss some appropriate evaluation metrics.</p>
<sec>
<title>2.4.1 BLEU</title>
<p>The best way to measure the performance of machine translation tasks is human judgment, but human judgment requires time and effort, which is a time-consuming process. To automate this process, the BLEU metric was introduced (Papineni et al., <xref ref-type="bibr" rid="B67">2002</xref>). The advantages of using BLEU are that it is fast and easy to compute, language-independent, and correlates with human judgements. As BLEU work by matching and counting the n-grams where <italic>n</italic>&#x02208;{1, 2, 3, 4} between reference text and candidate text generated by the machine translation system. BLEU does not take word order into account and uses the modified precision with brevity penalty. For full mathematical detail, the reader is referred to the work presented in Papineni et al. (<xref ref-type="bibr" rid="B67">2002</xref>). BLEU is also used to measure the performance of many text generations tasks such as abstractive text summarization, automatic image caption generation, and question answering tasks.</p></sec>
<sec>
<title>2.4.2 ROUGE</title>
<p>ROUGE is a measure used for the evaluation of text summarization systems, and ROUGE stands for Recall-Oriented Understudy for Gisting Evaluation (Lin, <xref ref-type="bibr" rid="B43">2004</xref>). Similar to BLEU (Papineni et al., <xref ref-type="bibr" rid="B67">2002</xref>), the ROUGE metric also measures the overlap between system-generated summaries and the gold standard reference summaries. One difference between ROUGE and BLEU is that BLEU is precision-based where ROUGE is recall-based. Similar to BLEU, ROUGE also used multiple references to match with the generated summary; in case of multiple references, ROUGE measure computes the pairwise score with each reference and then take the maximum of all the pairs.</p></sec>
<sec>
<title>2.4.3 METEOR</title>
<p>METEOR is a metric used to evaluate the performance of machine translation systems (Banerjee and Lavie, <xref ref-type="bibr" rid="B10">2005</xref>), Meteor is also a unigram matching-based metric, which finds the overlap between the system-generated output summary and the reference ground truth summary, similar to BLEU (Papineni et al., <xref ref-type="bibr" rid="B67">2002</xref>) and ROUGE (Lin, <xref ref-type="bibr" rid="B43">2004</xref>). METEOR is both unigram precision-based and unigram recall-based, similar to BLEU and Rouge. The authors&#x00027; enhanced capability is that METEOR accounts for the order of the matched unigrams.</p></sec>
<sec>
<title>2.4.4 SARI</title>
<p>SARI is a text simplification metric (Xu et al., <xref ref-type="bibr" rid="B89">2016</xref>), used to measure performance of text simplification systems. Text simplification systems are based on the three basic approaches: (1) Simplification by splitting longer sentences into their constituent counterparts, (2) Simplification by deletion of complex words and reordering, (3) simplification by paraphrasing and replacing complex words with their simple counterparts (Feng, <xref ref-type="bibr" rid="B25">2008</xref>). To capture the three types of operations discussed earlier, SARI considers the addition of tokens, deletion of tokens, and retention of tokens and aggregates the count of these three operations into one final score. SARI compares the system-generated output with both single or multiple references and the input source text.</p></sec>
<sec>
<title>2.4.5 SAMSA</title>
<p>SAMSA is a text simplification metric (Sulem et al., <xref ref-type="bibr" rid="B81">2018</xref>), used to evaluate text simplification systems. Unlike other metrics, SAMSA is a reference-less evaluation metric that does not compare the system-generated output with the ground truth references. Unlike SARI (Xu et al., <xref ref-type="bibr" rid="B89">2016</xref>) SAMSA considers the structural aspects of the generated text. It finds scenes in a sentence and then look for a single scene to be present in a single sentence and thus have the mapping between the scenes present in the input and the splits of the scenes in the output.</p></sec>
<sec>
<title>2.4.6 FKGL</title>
<p>Flesch Kincaid Grade Level (FKGL) is an evaluation metric widely used to measure the readability scores and grade level of the input text (Kincaid et al., <xref ref-type="bibr" rid="B39">1975</xref>). The FKGL metric translates the readability score into a U.S. school grade level, making it easier to understand the reading difficulty of a given text. The formula to compute FKGL score is given in <xref ref-type="disp-formula" rid="E1">Equation (1)</xref>.</p>
<disp-formula id="E1"><label>(1)</label><mml:math id="M1"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:mi>F</mml:mi><mml:mi>K</mml:mi><mml:mi>G</mml:mi><mml:mi>L</mml:mi><mml:mo>=</mml:mo><mml:mn>0</mml:mn><mml:mo>.</mml:mo><mml:mn>39</mml:mn><mml:mo>&#x000D7;</mml:mo><mml:mfrac><mml:mrow><mml:msub><mml:mrow><mml:mi>N</mml:mi></mml:mrow><mml:mrow><mml:mi>w</mml:mi><mml:mi>o</mml:mi><mml:mi>r</mml:mi><mml:mi>d</mml:mi><mml:mi>s</mml:mi></mml:mrow></mml:msub></mml:mrow><mml:mrow><mml:msub><mml:mrow><mml:mi>N</mml:mi></mml:mrow><mml:mrow><mml:mi>s</mml:mi><mml:mi>e</mml:mi><mml:mi>n</mml:mi><mml:mi>t</mml:mi><mml:mi>e</mml:mi><mml:mi>n</mml:mi><mml:mi>c</mml:mi><mml:mi>e</mml:mi><mml:mi>s</mml:mi></mml:mrow></mml:msub></mml:mrow></mml:mfrac><mml:mo>&#x0002B;</mml:mo><mml:mn>11</mml:mn><mml:mo>.</mml:mo><mml:mn>8</mml:mn><mml:mo>&#x000D7;</mml:mo><mml:mfrac><mml:mrow><mml:msub><mml:mrow><mml:mi>N</mml:mi></mml:mrow><mml:mrow><mml:mi>s</mml:mi><mml:mi>y</mml:mi><mml:mi>l</mml:mi><mml:mi>l</mml:mi><mml:mi>a</mml:mi><mml:mi>b</mml:mi><mml:mi>l</mml:mi><mml:mi>e</mml:mi><mml:mi>s</mml:mi></mml:mrow></mml:msub></mml:mrow><mml:mrow><mml:msub><mml:mrow><mml:mi>N</mml:mi></mml:mrow><mml:mrow><mml:mi>w</mml:mi><mml:mi>o</mml:mi><mml:mi>r</mml:mi><mml:mi>d</mml:mi><mml:mi>s</mml:mi></mml:mrow></mml:msub></mml:mrow></mml:mfrac><mml:mo>-</mml:mo><mml:mn>15</mml:mn><mml:mo>.</mml:mo><mml:mn>59</mml:mn></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula></sec>
<sec>
<title>2.4.7 BERT SCORE</title>
<p>BERTScore is used to evaluate text summarization, machine translation, and other text generation models. Unlike BLEU (Papineni et al., <xref ref-type="bibr" rid="B67">2002</xref>) and ROUGE (Lin, <xref ref-type="bibr" rid="B43">2004</xref>), it computes scores and evaluates model outputs even if there are no lexical matching words but there are semantically matching words in the gold reference text and the generated candidate text (Zhang et al., <xref ref-type="bibr" rid="B95">2019b</xref>). BERTScore uses contextual embedding to compute semantic similarity between reference and generated text. BERTScore is also word order invariant, unlike BLEU it computes semantic similarity between generated text and the gold standard reference without word order. BERTScore has shown high correlation with human judgments.</p></sec>
<sec>
<title>2.4.8 CSS1</title>
<p>To measure the effectiveness of the hybrid model for text summarization and simplification, a metric has been proposed in Zaman et al. (<xref ref-type="bibr" rid="B92">2020</xref>). This metric measures both the summarization and simplification and then takes the F1 measure of both to compute the once single score, which can provide insights for both the aspects (simplification and summarization) of the hybrid models. CSS1 can be formulated using <italic>ROUGE</italic><sub>1</sub> and <italic>SARI</italic> score as in <xref ref-type="disp-formula" rid="E2">Equation (2)</xref></p>
<disp-formula id="E2"><label>(2)</label><mml:math id="M2"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:mi>C</mml:mi><mml:mi>S</mml:mi><mml:msub><mml:mrow><mml:mi>S</mml:mi></mml:mrow><mml:mrow><mml:mn>1</mml:mn></mml:mrow></mml:msub><mml:mo>=</mml:mo><mml:mn>2</mml:mn><mml:mo>&#x000D7;</mml:mo><mml:mfrac><mml:mrow><mml:msub><mml:mrow><mml:mi>R</mml:mi></mml:mrow><mml:mrow><mml:mn>1</mml:mn></mml:mrow></mml:msub><mml:mo>&#x000D7;</mml:mo><mml:mi>S</mml:mi><mml:mi>A</mml:mi><mml:mi>R</mml:mi><mml:mi>I</mml:mi></mml:mrow><mml:mrow><mml:msub><mml:mrow><mml:mi>R</mml:mi></mml:mrow><mml:mrow><mml:mn>1</mml:mn></mml:mrow></mml:msub><mml:mo>&#x0002B;</mml:mo><mml:mi>S</mml:mi><mml:mi>A</mml:mi><mml:mi>R</mml:mi><mml:mi>I</mml:mi></mml:mrow></mml:mfrac></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula></sec></sec>
<sec>
<title>2.5 Resources for text simplification</title>
<p>Automated text simplification is transforming a complex text into a simpler one that should be easy to understand. This task can be handled in many ways, as discussed in Section 2.3. Different kinds of resources are required to perform text simplification, depending upon the type of simplification, such as lexical (Shardlow, <xref ref-type="bibr" rid="B75">2014a</xref>,<xref ref-type="bibr" rid="B76">b</xref>) simplification or structural simplification and the method used such as sequence-to-sequence learning (Liu et al., <xref ref-type="bibr" rid="B46">2015</xref>; Chen and Bansal, <xref ref-type="bibr" rid="B17">2018</xref>), machine translation. Considering the machine-translation method for text simplification (Zhu et al., <xref ref-type="bibr" rid="B98">2010</xref>; Xu et al., <xref ref-type="bibr" rid="B89">2016</xref>) requires a vast amount of parallel monolingual corpus and a considerable number of parameters to be optimized during training. So heavy computational resources are required to accommodate such large corpora in memory alongside the model parameters. For instance, working in the domain of text simplification, one can build and train state-of-the-art models only if rich computational resources are available. The availability of such resources is a big problem that the text simplification community is facing. Developing optimized resource algorithms for text simplification are required to advance the domain&#x00027;s further state of the art. Below are some benchmark datasets used to train text simplification models; further detail is given in subsections.</p>
<sec>
<title>2.5.1 PWKP corpus/WikiSmall</title>
<p>In 2010, Zhu et al. collected a monolingual parallel corpus containing 108k sentences and its simplified pairs, and this dataset was harvested from around 65k Wikipedia articles and its corresponding simple Wikipedia articles (Zhu et al., <xref ref-type="bibr" rid="B98">2010</xref>). The authors aligned the sentences in 1-to-1 and 1-to-many fashion. In the later alignment, the authors used sentence splitting&#x02014;the corpus can be found at the link.<xref ref-type="fn" rid="fn0001"><sup>1</sup></xref></p></sec>
<sec>
<title>2.5.2 Turk corpus/WikiLarge</title>
<p>Xu et al. (<xref ref-type="bibr" rid="B89">2016</xref>) explored a statistical machine translation model for text simplification purposes; along with this, they introduced a new metric for text simplification evaluation and a new dataset named Turk corpus; they used some filtered data from PWKP corpus (see section PWKP), along with this modified PWKP, they used multiple 8 reference sentences of each original sentence, and this was done through Amazon Mechanical Turk. There are 2,000 sentences for tuning and 350 sentences for testing; for training the simplification models, most researchers use WikiLarge (Zhang and Lapata, <xref ref-type="bibr" rid="B96">2017</xref>), or WikiSmall dataset (see section 2.5.1).</p></sec>
<sec>
<title>2.5.3 ASSET corpus</title>
<p>Alva-Manchego et al. (<xref ref-type="bibr" rid="B5">2020a</xref>) prepared a new dataset for text simplification task, which is based on Turk corpus (see Section 2.5.2). The authors selected the same source sentences with modified reference sentences. The idea behind maintaining multiple references is that multiple paraphrase operations can simplify. ASSET corpus contains both 1-to-1 and 1-to-many alignments. Unlike Turk corpus, ASSET contains 10 references for each source sentence. The tuning and test set size is the same as Turk corpus that is 2,000 tuning examples and 350 test examples. The ASSET dataset can be downloaded from the link.<xref ref-type="fn" rid="fn0002"><sup>2</sup></xref></p></sec></sec>
<sec>
<title>2.6 Resources for text summarization</title>
<p>Different kinds of resources are required to perform text summarization, depending upon the approach and method employed. Consider the two broad methods of text summarization, Extractive and Abstractive; some Extractive summarization methods only rank parts of the input text and then output the filtered portion of the input text such as text rank (Mihalcea and Tarau, <xref ref-type="bibr" rid="B58">2004</xref>), and such approaches do not require many resources as compared to the Deep Neural Network-based methods such as Kinugawa and Tsuruoka (<xref ref-type="bibr" rid="B40">2017</xref>). Similarly, Abstractive summarization methods such as See et al. (<xref ref-type="bibr" rid="B74">2017</xref>) and Yan et al. (<xref ref-type="bibr" rid="B90">2020</xref>) are based on sequence-to-sequence learning which requires millions of parameters to be tuned; in addition to a vast number of parameters and computational resources, these methods required a large corpus of parallel pairs of input long documents and reference summaries. The requirement of such resources becomes a bottleneck for the research communities. To further advance state of the art in the domain, more advanced algorithms are required that are resource optimized. In the next subsections, we discuss the datasets available for summarization research.</p>
<sec>
<title>2.6.1 Gigaword</title>
<p>In 2003, Graff et al. compiled a dataset named Gigaword (Graff et al., <xref ref-type="bibr" rid="B32">2003</xref>) that was mainly used for sentence summarization tasks or generation of headlines tasks. The size of the source and target summaries in terms of tokens is short; the statistics of the dataset are as follows: there are 3.8M training pairs, 189k validation pairs, and 1951 test pairs. The model trained and developed with this dataset are primarily evaluated with ROUGE (1, 2, L) scores.</p></sec>
<sec>
<title>2.6.2 CNN-daily mail dataset</title>
<p>For text summarization research CNN-daily mail prepared by Nallapati et al. (<xref ref-type="bibr" rid="B59">2016</xref>) is a large dataset containing 287,226 training article pairs, 13,368 validation article pairs, and 11,490 test article pairs. Researchers tune and validate their developed models with CNN-daily mail dataset and report ROUGE-1, ROUGE-2, and ROUGE-L scores. There are multiple versions available such as the entity-anonymized version (Nallapati et al., <xref ref-type="bibr" rid="B59">2016</xref>) and the non-anonymized version (See et al., <xref ref-type="bibr" rid="B74">2017</xref>). A pre-processed version (See et al., <xref ref-type="bibr" rid="B74">2017</xref>) of the dataset can be downloaded.<xref ref-type="fn" rid="fn0003"><sup>3</sup></xref></p></sec>
<sec>
<title>2.6.3 X-Sum</title>
<p>Narayan et al. (<xref ref-type="bibr" rid="B60">2018</xref>) prepared a new dataset for the task of pure abstractive summarization. This dataset is not suitable for extractive methods. The authors collected news articles and their corresponding one-line new summary from BBC website. The statistics of the dataset are as follows: 204,045 training article-pairs, 11,332 tuning article pairs, and 11,334 test article pairs. The source document size is on average 430 tokens, and target summary size is on average 23 tokens, which leads to extreme summarization; hence, the name X-sum was given to this dataset. Models trained and developed with this dataset are evaluated using ROUGE (1, 2, L) scores. The dataset can be downloaded.<xref ref-type="fn" rid="fn0004"><sup>4</sup></xref></p></sec>
<sec>
<title>2.6.4 Sentence compression Google dataset</title>
<p>In 2013, Google developed a dataset for the task of sentence compression (Filippova and Altun, <xref ref-type="bibr" rid="B28">2013</xref>), which may be considered as summarization at the sentence level. The first version of this dataset was released with 10<italic>k</italic> pairs; recently, Google released an updated version of this dataset which contains 200k pairs. The dataset can be downloaded from the link.<xref ref-type="fn" rid="fn0005"><sup>5</sup></xref></p></sec></sec>
<sec>
<title>2.7 Resources for text summarization and simplification</title>
<p>Beside from the available resources for summarization alone and simplification alone, in this section we discuss the availability of corpus for the combined task of summarization and simplification.</p>
<sec>
<title>2.7.1 PLOS and eLife datasets</title>
<p>Two datasets PLOS and eLife were introduced for the task of summarization and simplification, and both datasets focus on scientific documents from the biomedical domain and their corresponding summaries in plain English. The datasets were created by parsing XML articles in Python, the datasets are further structured into sections such as abstracts and article text, and article text is further formatted into subsections as per the headings presented in the original articles (Goldsack et al., <xref ref-type="bibr" rid="B30">2022</xref>)</p></sec></sec></sec>
<sec id="s3">
<title>3 Data and method</title>
<p>Hybrid text simplification and summarization models require data in the form of a parallel corpus containing complex-simple pairs. To create such a corpus manually, we need domain knowledge and extensive time. We used the publicly available dataset from Zaman et al. (<xref ref-type="bibr" rid="B92">2020</xref>), and this dataset consists of 5,204 article and summary pairs. The corpus links full-text scientific articles and their abstracts to simplified summaries taken from EurekAlert.</p>
<sec>
<title>3.1 Simplification aware text summarization (SATS)</title>
<p>We have so far described our baseline models, taken from the literature. Each of these models is designed for either simplification or summarization, except HTTS no one model is designed to simultaneously complete both tasks. HTSS is the only model that is a hybrid of summarization and simplification. To address the limitations of HTSS such as general language errors and grammar issues, we propose an adaptation to the ProphetNet architecture, which allows a summarization model to prioritize simple terms during generation. First, we need to understand what makes a word complex. While significant work has been done on automated prediction of lexical complexity (Shardlow et al., <xref ref-type="bibr" rid="B78">2022</xref>), it is a well-established fact in the literature that lexical frequency is a strong indicator of how difficult a reader will find a word to understand (Paetzold and Specia, <xref ref-type="bibr" rid="B66">2016</xref>; Martin et al., <xref ref-type="bibr" rid="B55">2020a</xref>; North et al., <xref ref-type="bibr" rid="B65">2023</xref>). As such, we take a corpus of 13,588,391 words, each associated with a frequency value which was derived as the result of counting word frequencies from over 1 trillion words of English web-text (Brants, <xref ref-type="bibr" rid="B14">2006</xref>).</p>
<p>In our proposed model which is depicted in <xref ref-type="fig" rid="F1">Figure 1</xref>, lookup-difficulty module is the contribution that we add to the summarization model to enhance its capability for the task of simplification. This module is responsible for guiding the generation of the underlying transformer model and ensuring that the generated output summary is simple to read by contributing values of difficulty scores to the loss function. The loss function is further used to update the parameters of the entire network.</p>
<fig id="F1" position="float">
<label>Figure 1</label>
<caption><p>Proposed architecture of SATS model.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="frai-07-1375419-g0001.tif"/>
</fig>


<p>We use this data and reformulate this idea into difficulty scores as in <xref ref-type="disp-formula" rid="E3">Equation (3)</xref>:</p>
<disp-formula id="E3"><label>(3)</label><mml:math id="M3"><mml:mrow><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>s</mml:mi><mml:mi>c</mml:mi><mml:mi>o</mml:mi><mml:mi>r</mml:mi></mml:mstyle><mml:msub><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>e</mml:mi></mml:mstyle><mml:mrow><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>l</mml:mi><mml:mi>o</mml:mi><mml:mi>g</mml:mi></mml:mstyle></mml:mrow></mml:msub><mml:mo>=</mml:mo><mml:mfrac><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mn>1</mml:mn></mml:mstyle><mml:mrow><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>l</mml:mi><mml:mi>o</mml:mi><mml:mi>g</mml:mi></mml:mstyle><mml:mo stretchy='false'>(</mml:mo><mml:msub><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>w</mml:mi></mml:mstyle><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>i</mml:mi></mml:mstyle></mml:msub><mml:mo stretchy='false'>)</mml:mo></mml:mrow></mml:mfrac></mml:mrow></mml:math></disp-formula>
<p>Here, <italic>score</italic><sub><italic>log</italic></sub> is the intermediate score which is further normalized between 0 and 1 as shown in <xref ref-type="disp-formula" rid="E4">Equation (4)</xref>:</p>
<disp-formula id="E4"><label>(4)</label><mml:math id="M4"><mml:mrow><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>s</mml:mi><mml:mi>c</mml:mi><mml:mi>o</mml:mi><mml:mi>r</mml:mi></mml:mstyle><mml:msub><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>e</mml:mi></mml:mstyle><mml:mrow><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>d</mml:mi><mml:mi>i</mml:mi><mml:mi>f</mml:mi><mml:mi>f</mml:mi><mml:mi>i</mml:mi><mml:mi>c</mml:mi><mml:mi>u</mml:mi><mml:mi>l</mml:mi><mml:mi>t</mml:mi><mml:mi>y</mml:mi></mml:mstyle></mml:mrow></mml:msub><mml:mo>=</mml:mo><mml:mfrac><mml:mrow><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>s</mml:mi><mml:mi>c</mml:mi><mml:mi>o</mml:mi><mml:mi>r</mml:mi></mml:mstyle><mml:msub><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>e</mml:mi></mml:mstyle><mml:mrow><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>l</mml:mi><mml:mi>o</mml:mi><mml:mi>g</mml:mi></mml:mstyle></mml:mrow></mml:msub><mml:mo>&#x02212;</mml:mo><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>m</mml:mi><mml:mi>i</mml:mi><mml:mi>n</mml:mi></mml:mstyle><mml:mo stretchy='false'>(</mml:mo><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>s</mml:mi><mml:mi>c</mml:mi><mml:mi>o</mml:mi><mml:mi>r</mml:mi></mml:mstyle><mml:msub><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>e</mml:mi></mml:mstyle><mml:mrow><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>l</mml:mi><mml:mi>o</mml:mi><mml:mi>g</mml:mi></mml:mstyle></mml:mrow></mml:msub><mml:mo stretchy='false'>)</mml:mo></mml:mrow><mml:mrow><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>m</mml:mi><mml:mi>a</mml:mi><mml:mi>x</mml:mi></mml:mstyle><mml:mo stretchy='false'>(</mml:mo><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>s</mml:mi><mml:mi>c</mml:mi><mml:mi>o</mml:mi><mml:mi>r</mml:mi></mml:mstyle><mml:msub><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>e</mml:mi></mml:mstyle><mml:mrow><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>l</mml:mi><mml:mi>o</mml:mi><mml:mi>g</mml:mi></mml:mstyle></mml:mrow></mml:msub><mml:mo stretchy='false'>)</mml:mo><mml:mo>&#x02212;</mml:mo><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>m</mml:mi><mml:mi>i</mml:mi><mml:mi>n</mml:mi></mml:mstyle><mml:mo stretchy='false'>(</mml:mo><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>s</mml:mi><mml:mi>c</mml:mi><mml:mi>o</mml:mi><mml:mi>r</mml:mi></mml:mstyle><mml:msub><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>e</mml:mi></mml:mstyle><mml:mrow><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>l</mml:mi><mml:mi>o</mml:mi><mml:mi>g</mml:mi></mml:mstyle></mml:mrow></mml:msub><mml:mo stretchy='false'>)</mml:mo></mml:mrow></mml:mfrac></mml:mrow></mml:math></disp-formula>
<p>Here <italic>score</italic><sub><italic>difficulty</italic></sub> is the final difficulty score that we use for computing simplification. This score is in the range 0&#x02013;1, and scales with the log of the frequency to avoid very common words overly influencing the simplification. We pre-compute <italic>score</italic><sub><italic>log</italic></sub>, and <italic>score</italic><sub><italic>difficulty</italic></sub> for every word in our word frequency dataset and create a lookup table, which can be accessed during the generation phase.</p>
<p>We reformulate the difficulty score as a loss term in the following manner, relying on the lookup table to identify the difficulty score of the given token as demonstrated in <xref ref-type="disp-formula" rid="E5">Equation (5)</xref>.</p>
<disp-formula id="E5"><label>(5)</label><mml:math id="M5"><mml:mrow><mml:msub><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>L</mml:mi></mml:mstyle><mml:mrow><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>s</mml:mi><mml:mi>i</mml:mi><mml:mi>m</mml:mi><mml:mi>p</mml:mi></mml:mstyle></mml:mrow></mml:msub><mml:mo>=</mml:mo><mml:mfrac><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mn>1</mml:mn></mml:mstyle><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>T</mml:mi></mml:mstyle></mml:mfrac><mml:mstyle displaystyle='true'><mml:munderover><mml:mo>&#x02211;</mml:mo><mml:mrow><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>i</mml:mi></mml:mstyle><mml:mo>=</mml:mo><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mn>1</mml:mn></mml:mstyle></mml:mrow><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>T</mml:mi></mml:mstyle></mml:munderover><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>l</mml:mi></mml:mstyle></mml:mstyle><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>o</mml:mi><mml:mi>o</mml:mi><mml:mi>k</mml:mi><mml:mi>u</mml:mi></mml:mstyle><mml:msub><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>p</mml:mi></mml:mstyle><mml:mrow><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>s</mml:mi><mml:mi>c</mml:mi><mml:mi>o</mml:mi><mml:mi>r</mml:mi></mml:mstyle><mml:msub><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>e</mml:mi></mml:mstyle><mml:mrow><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>d</mml:mi><mml:mi>i</mml:mi><mml:mi>f</mml:mi><mml:mi>f</mml:mi><mml:mi>i</mml:mi><mml:mi>c</mml:mi><mml:mi>u</mml:mi><mml:mi>l</mml:mi><mml:mi>t</mml:mi><mml:mi>y</mml:mi></mml:mstyle></mml:mrow></mml:msub></mml:mrow></mml:msub><mml:mo stretchy='false'>(</mml:mo><mml:msub><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>w</mml:mi></mml:mstyle><mml:mstyle mathvariant='bold' mathsize='normal'><mml:mi>i</mml:mi></mml:mstyle></mml:msub><mml:mo stretchy='false'>)</mml:mo></mml:mrow></mml:math></disp-formula>
<p>In <xref ref-type="disp-formula" rid="E5">Equation (5)</xref>, <italic>T</italic> is the size of the predictions output of ProphetNet. We updated the objective function by adding simplification loss from <xref ref-type="disp-formula" rid="E5">Equation (5)</xref>, thus, the updated objective can be written as in shown in <xref ref-type="disp-formula" rid="E6">Equation (6)</xref>.</p>
<disp-formula id="E6"><label>(6)</label><mml:math id="M6"><mml:mtable columnalign='left'><mml:mtr><mml:mtd><mml:mtext>&#x02009;&#x02009;&#x02009;&#x02009;</mml:mtext><mml:mi>L</mml:mi><mml:mo>=</mml:mo><mml:munder><mml:munder><mml:mrow><mml:mo>&#x02212;</mml:mo><mml:msub><mml:mi>&#x003B1;</mml:mi><mml:mn>0</mml:mn></mml:msub><mml:mo>&#x000B7;</mml:mo><mml:mo stretchy='false'>(</mml:mo><mml:mstyle displaystyle='true'><mml:munderover><mml:mo>&#x02211;</mml:mo><mml:mrow><mml:mi>t</mml:mi><mml:mo>=</mml:mo><mml:mn>1</mml:mn></mml:mrow><mml:mi>T</mml:mi></mml:munderover><mml:mrow><mml:mi>log</mml:mi></mml:mrow></mml:mstyle><mml:mi>P</mml:mi><mml:mi>&#x003B8;</mml:mi><mml:mo stretchy='false'>(</mml:mo><mml:msub><mml:mi>y</mml:mi><mml:mi>t</mml:mi></mml:msub><mml:mo>&#x0007C;</mml:mo><mml:msub><mml:mi>y</mml:mi><mml:mrow><mml:mo>&#x0003C;</mml:mo><mml:mi>t</mml:mi></mml:mrow></mml:msub><mml:mo>,</mml:mo><mml:mi>x</mml:mi><mml:mo stretchy='false'>)</mml:mo><mml:mo stretchy='false'>)</mml:mo></mml:mrow><mml:mo stretchy='true'>&#x0FE38;</mml:mo></mml:munder><mml:mrow><mml:mi>L</mml:mi><mml:msub><mml:mi>M</mml:mi><mml:mrow><mml:mi>l</mml:mi><mml:mi>o</mml:mi><mml:mi>s</mml:mi><mml:mi>s</mml:mi></mml:mrow></mml:msub></mml:mrow></mml:munder></mml:mtd></mml:mtr><mml:mtr><mml:mtd><mml:munder><mml:munder><mml:mrow><mml:mo>&#x02212;</mml:mo><mml:mstyle displaystyle='true'><mml:munderover><mml:mo>&#x02211;</mml:mo><mml:mrow><mml:mi>j</mml:mi><mml:mo>=</mml:mo><mml:mn>1</mml:mn></mml:mrow><mml:mrow><mml:mi>m</mml:mi><mml:mo>&#x02212;</mml:mo><mml:mn>1</mml:mn></mml:mrow></mml:munderover><mml:mrow><mml:msub><mml:mi>&#x003B1;</mml:mi><mml:mi>t</mml:mi></mml:msub></mml:mrow></mml:mstyle><mml:mo>&#x000B7;</mml:mo><mml:mo stretchy='false'>(</mml:mo><mml:mstyle displaystyle='true'><mml:munderover><mml:mo>&#x02211;</mml:mo><mml:mrow><mml:mi>t</mml:mi><mml:mo>=</mml:mo><mml:mn>1</mml:mn></mml:mrow><mml:mrow><mml:mi>T</mml:mi><mml:mo>&#x02212;</mml:mo><mml:mi>j</mml:mi></mml:mrow></mml:munderover><mml:mrow><mml:mi>log</mml:mi></mml:mrow></mml:mstyle><mml:mi>P</mml:mi><mml:mi>&#x003B8;</mml:mi><mml:mo stretchy='false'>(</mml:mo><mml:msub><mml:mi>y</mml:mi><mml:mrow><mml:mi>t</mml:mi><mml:mo>+</mml:mo><mml:mi>j</mml:mi></mml:mrow></mml:msub><mml:mo>&#x0007C;</mml:mo><mml:msub><mml:mi>y</mml:mi><mml:mrow><mml:mo>&#x0003C;</mml:mo><mml:mi>t</mml:mi></mml:mrow></mml:msub><mml:mo>,</mml:mo><mml:mi>x</mml:mi><mml:mo stretchy='false'>)</mml:mo><mml:mo stretchy='false'>)</mml:mo></mml:mrow><mml:mo stretchy='true'>&#x0FE38;</mml:mo></mml:munder><mml:mrow><mml:mi>F</mml:mi><mml:mi>u</mml:mi><mml:mi>t</mml:mi><mml:mi>u</mml:mi><mml:mi>r</mml:mi><mml:mi>e</mml:mi><mml:mtext>&#x02009;</mml:mtext><mml:mi>n</mml:mi><mml:mtext>-</mml:mtext><mml:mi>g</mml:mi><mml:mi>r</mml:mi><mml:mi>a</mml:mi><mml:msub><mml:mi>m</mml:mi><mml:mrow><mml:mi>l</mml:mi><mml:mi>o</mml:mi><mml:mi>s</mml:mi><mml:mi>s</mml:mi></mml:mrow></mml:msub></mml:mrow></mml:munder></mml:mtd></mml:mtr><mml:mtr><mml:mtd><mml:mtext>&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;&#x02009;</mml:mtext><mml:munder><mml:munder><mml:mrow><mml:mo>&#x02212;</mml:mo><mml:mi>&#x003BB;</mml:mi><mml:mo>&#x000B7;</mml:mo><mml:msub><mml:mi>L</mml:mi><mml:mrow><mml:mi>s</mml:mi><mml:mi>i</mml:mi><mml:mi>m</mml:mi><mml:mi>p</mml:mi></mml:mrow></mml:msub></mml:mrow><mml:mo stretchy='true'>&#x0FE38;</mml:mo></mml:munder><mml:mrow><mml:mtext>simplification&#x000A0;loss</mml:mtext></mml:mrow></mml:munder></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<p>In <xref ref-type="disp-formula" rid="E6">Equation (6)</xref>, <italic>L</italic><sub><italic>simp</italic></sub> is the simplification loss and &#x003BB; is a hyperparameter used to control the degree of the effect of simplification loss as compared to the overall loss. Note that we tuned lambda in our experiments using the validation set and set it to 0.8. An example of the calculation of our simplification loss is shown in <xref ref-type="table" rid="T2">Table 2</xref>.</p>
<table-wrap position="float" id="T2">
<label>Table 2</label>
<caption><p>Sample of words and their frequencies along with log score and simplification loss.</p></caption>
<table frame="box" rules="all">
<thead>
<tr style="background-color:#919498;color:#ffffff">
<th valign="top" align="left"><bold>Vocab Index</bold></th>
<th valign="top" align="left"><bold>Word</bold></th>
<th valign="top" align="center"><bold>Frequency</bold></th>
<th valign="top" align="center"><bold><italic>score</italic><sub><italic>log</italic></sub></bold></th>
<th valign="top" align="center"><bold><italic>L</italic><sub><italic>simp</italic></sub></bold></th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="left">1</td>
<td valign="top" align="left">Where</td>
<td valign="top" align="center">282,489,721</td>
<td valign="top" align="center">0.80212354</td>
<td valign="top" align="center">0.05503001</td>
</tr> <tr>
<td valign="top" align="left">2</td>
<td valign="top" align="left">Business</td>
<td valign="top" align="center">280,687,568</td>
<td valign="top" align="center">0.801758097</td>
<td valign="top" align="center">0.055149779</td>
</tr> <tr>
<td valign="top" align="left">3</td>
<td valign="top" align="left">Must</td>
<td valign="top" align="center">269,659,445</td>
<td valign="top" align="center">0.799469364</td>
<td valign="top" align="center">0.055901677</td>
</tr> <tr>
<td valign="top" align="left">&#x022EE;</td>
<td valign="top" align="left">&#x022EE;</td>
<td valign="top" align="center">&#x022EE;</td>
<td valign="top" align="center">&#x022EE;</td>
<td valign="top" align="center">&#x022EE;</td>
</tr> <tr>
<td valign="top" align="left">5,725</td>
<td valign="top" align="left">Barometers</td>
<td valign="top" align="center">59,714</td>
<td valign="top" align="center">0.318946422</td>
<td valign="top" align="center">0.335138044</td>
</tr> <tr>
<td valign="top" align="left">5,726</td>
<td valign="top" align="left">Scribing</td>
<td valign="top" align="center">59,712</td>
<td valign="top" align="center">0.318944509</td>
<td valign="top" align="center">0.335140006</td>
</tr> <tr>
<td valign="top" align="left">5,727</td>
<td valign="top" align="left">Splattering</td>
<td valign="top" align="center">57,154</td>
<td valign="top" align="center">0.31644443</td>
<td valign="top" align="center">0.337714815</td>
</tr> <tr>
<td valign="top" align="left">&#x022EE;</td>
<td valign="top" align="left">&#x022EE;</td>
<td valign="top" align="center">&#x022EE;</td>
<td valign="top" align="center">&#x022EE;</td>
<td valign="top" align="center">&#x022EE;</td>
</tr> <tr>
<td valign="top" align="left">77,530</td>
<td valign="top" align="left">Antrum</td>
<td valign="top" align="center">37,668</td>
<td valign="top" align="center">0.292636921</td>
<td valign="top" align="center">0.363306085</td>
</tr> <tr>
<td valign="top" align="left">77,531</td>
<td valign="top" align="left">Kudu</td>
<td valign="top" align="center">37,640</td>
<td valign="top" align="center">0.29259446</td>
<td valign="top" align="center">0.363353536</td>
</tr>
<tr>
<td valign="top" align="left">77,532</td>
<td valign="top" align="left">Victimizing</td>
<td valign="top" align="center">6,857</td>
<td valign="top" align="center">0.195363413</td>
<td valign="top" align="center">0.492969086</td>
</tr></tbody>
</table>
</table-wrap>

</sec></sec>
<sec sec-type="results" id="s4">
<title>4 Results</title>
<p>For our experiments, we designed a computational environment using Python version 3.7 as a scripting language and PyTorch as a deep learning framework; in our experiments, we leveraged two datasets: (1) CNN-Daily mail specialized for text summarization and (2) Eureka Alert dataset specialized for hybrid text summarization and text simplification. For the evaluation of our proposed SATS model and comparison with baseline models, we used ROUGE as an automatic metric for text summarization, SARI and FKGL as an automatic metric for text simplification, BERTScore for semantic similarity and overall goodness of the generated output, and CSS as an automatic metric for combined text summarization and text simplification.</p>
<p>To generate the automatic evaluation scores for the baseline systems against the datasets that we used in our experiments, we retrained the implementation of three baseline systems and applied these to the Eureka dataset specialized for summarization and simplification. The first, HTSS, is a state-of-the-art joint simplification-summarization model. It is a hybrid model that implements a binary simplification loss to the pointer-generator architecture. The second is ACCESS, which provides controllable text simplification through a sequence-to-sequence architecture. The third is MUSS. Finally, we have implemented ProphetNet as described above and included the results with our adjusted loss function. The results are in <xref ref-type="table" rid="T3">Tables 3</xref>, <xref ref-type="table" rid="T4">4</xref>.</p>




<table-wrap position="float" id="T3">
<label>Table 3</label>
<caption><p>Compare SATS (proposed) with HTSS, ACCESS, MUSS, and ProphetNet using the Eureka Dataset.</p></caption>
<table frame="box" rules="all">
<thead>
<tr style="background-color:#919498;color:#ffffff">
<th valign="top" align="left" colspan="0"><bold>Model</bold></th>
<th valign="top" align="left"><bold>Type</bold></th>
<th valign="top" align="center"><bold>BERTscore</bold></th>
<th valign="top" align="center"><bold>FKGL</bold></th>
<th valign="top" align="center"><bold>ROUGE1</bold></th>
<th valign="top" align="center"><bold>ROUGE2</bold></th>
<th valign="top" align="center"><bold>ROUGEL</bold></th>
<th valign="top" align="center"><bold>SARI</bold></th>
<th valign="top" align="center"><bold>CSS<sub>1</sub></bold></th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="left">ACCESS (Martin et al., <xref ref-type="bibr" rid="B55">2020a</xref>)</td>
<td valign="top" align="left">Simplification</td>
<td valign="top" align="center">82.34</td>
<td valign="top" align="center">7.52</td>
<td valign="top" align="center">9.641</td>
<td valign="top" align="center">1.177</td>
<td valign="top" align="center">7.893</td>
<td valign="top" align="center">38.3834</td>
<td valign="top" align="center">15.411</td>
</tr> <tr>
<td valign="top" align="left">MUSS (Martin et al., <xref ref-type="bibr" rid="B56">2020b</xref>)</td>
<td valign="top" align="left">Simplification</td>
<td valign="top" align="center"><bold>84.15</bold></td>
<td valign="top" align="center">7.34</td>
<td valign="top" align="center">22.714</td>
<td valign="top" align="center">7.425</td>
<td valign="top" align="center">19.646</td>
<td valign="top" align="center">32.093204</td>
<td valign="top" align="center">26.60106</td>
</tr> <tr>
<td valign="top" align="left">ProphetNet (Qi et al., <xref ref-type="bibr" rid="B68">2020</xref>)</td>
<td valign="top" align="left">Summarization</td>
<td valign="top" align="center">83.42</td>
<td valign="top" align="center">7.25</td>
<td valign="top" align="center">33.414</td>
<td valign="top" align="center">9.443</td>
<td valign="top" align="center">18.962</td>
<td valign="top" align="center">40.44</td>
<td valign="top" align="center">36.5987</td>
</tr> <tr>
<td valign="top" align="left">HTSS (Zaman et al., <xref ref-type="bibr" rid="B92">2020</xref>)</td>
<td valign="top" align="left">Hybrid</td>
<td valign="top" align="center">81.73</td>
<td valign="top" align="center">8.42</td>
<td valign="top" align="center">21.938</td>
<td valign="top" align="center">03.21</td>
<td valign="top" align="center">17.171</td>
<td valign="top" align="center">37.650</td>
<td valign="top" align="center">27.7225</td>
</tr>
<tr>
<td valign="top" align="left">SATS</td>
<td valign="top" align="left">Hybrid</td>
<td valign="top" align="center">83.38</td>
<td valign="top" align="center">7.35</td>
<td valign="top" align="center"><bold>34.24</bold></td>
<td valign="top" align="center"><bold>10.15</bold></td>
<td valign="top" align="center"><bold>20.21</bold></td>
<td valign="top" align="center"><bold>40.83</bold></td>
<td valign="top" align="center"><bold>37.24</bold></td>
</tr></tbody>
</table>
<table-wrap-foot>
<p>The bold values indicate the system with the highest performing score.</p>
</table-wrap-foot>
</table-wrap>

<table-wrap position="float" id="T4">
<label>Table 4</label>
<caption><p>Compare SATS with HTSS, ACCESS, MUSS, and ProphetNet using CNN-Daily Mail Dataset.</p></caption>
<table frame="box" rules="all">
<thead>
<tr style="background-color:#919498;color:#ffffff">
<th valign="top" align="left" colspan="0"><bold>Model</bold></th>
<th valign="top" align="left"><bold>Type</bold></th>
<th valign="top" align="left"><bold>ROUGE1</bold></th>
<th valign="top" align="left"><bold>ROUGE2</bold></th>
<th valign="top" align="left"><bold>ROUGEL</bold></th>
<th valign="top" align="left"><bold>SARI</bold></th>
<th valign="top" align="center"><bold>CSS<sub>1</sub></bold></th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="left">ACCESS (Martin et al., <xref ref-type="bibr" rid="B55">2020a</xref>)</td>
<td valign="top" align="left">Simplification</td>
<td valign="top" align="left">15.557</td>
<td valign="top" align="left">2.787</td>
<td valign="top" align="left">11.668</td>
<td valign="top" align="left">34.3696</td>
<td valign="top" align="center">21.481</td>
</tr> <tr>
<td valign="top" align="left">MUSS (Martin et al., <xref ref-type="bibr" rid="B56">2020b</xref>)</td>
<td valign="top" align="left">Simplification</td>
<td valign="top" align="left">22.933</td>
<td valign="top" align="left">9.575</td>
<td valign="top" align="left">15.334</td>
<td valign="top" align="left">36.501</td>
<td valign="top" align="center">28.165</td>
</tr> <tr>
<td valign="top" align="left">ProphetNet (Qi et al., <xref ref-type="bibr" rid="B68">2020</xref>)</td>
<td valign="top" align="left">Summarization</td>
<td valign="top" align="left"><bold>44.36</bold></td>
<td valign="top" align="left">21.262</td>
<td valign="top" align="left">30.723</td>
<td valign="top" align="left">42.033</td>
<td valign="top" align="center">43.167</td>
</tr> <tr>
<td valign="top" align="left">HTSS (Zaman et al., <xref ref-type="bibr" rid="B92">2020</xref>)</td>
<td valign="top" align="left">Hybrid</td>
<td valign="top" align="left">37.1011</td>
<td valign="top" align="left">15.91226</td>
<td valign="top" align="left">32.522</td>
<td valign="top" align="left">35.420429</td>
<td valign="top" align="center">36.24128</td>
</tr> <tr>
<td valign="top" align="left">SATS</td>
<td valign="top" align="left">Hybrid</td>
<td valign="top" align="left">44.26</td>
<td valign="top" align="left"><bold>21.48</bold></td>
<td valign="top" align="left"><bold>30.93</bold></td>
<td valign="top" align="left"><bold>43.40</bold></td>
<td valign="top" align="center"><bold>43.82</bold></td>
</tr></tbody>
</table>
<table-wrap-foot>
<p>The bold values indicate the system with the highest performing score.</p>
</table-wrap-foot>
</table-wrap>



<sec>
<title>4.1 Baseline models</title>
<p>To compare and validate our method, some baseline methods are required, since our method is a hybrid of simplification and summarization, so we used baseline methods for simplification, summarization, and hybrid, for simplification we used two models ACCESS and MUSS, and we ran their code to reproduce the results, and for text summarization we ran ProphetNet model and reproduce the results, whereas, for hybrid model comparison we used HTSS. Details of each baseline method are given below.</p>
<sec>
<title>4.1.1 ACCESS</title>
<p>For text simplification, we choose ACCESS (Martin et al., <xref ref-type="bibr" rid="B55">2020a</xref>) as a baseline. ACCESS is a text simplification model that uses extra control tokens to control the generated simplified text. ACCESS uses a sequence-to-sequence architecture which is based on the transformer model (Vaswani et al., <xref ref-type="bibr" rid="B85">2017</xref>); furthermore, ACCESS is based on BART (Lewis et al., <xref ref-type="bibr" rid="B41">2020</xref>).</p></sec>
<sec>
<title>4.1.2 MUSS</title>
<p>MUSS (Martin et al., <xref ref-type="bibr" rid="B56">2020b</xref>) is an unsupervised multilingual text simplification model that is based on ACCESS (Martin et al., <xref ref-type="bibr" rid="B55">2020a</xref>) and BART (Lewis et al., <xref ref-type="bibr" rid="B41">2020</xref>), from ACCESS it adapts the capability of controllable generation, and from BART, it adapts the sequence-to-sequence multilingual capability. MUSS uses mined sequences and paraphrasing to build the training dataset. There are two variations of MUSS, one is trained with mined sequences only. There is another version of MUSS which is supervised, this version is available for English only, and the supervised version uses the WikiLarge parallel corpus for training. We used the parallel version as a baseline for our study.</p></sec>
<sec>
<title>4.1.3 HTSS</title>
<p>HTSS (Zaman et al., <xref ref-type="bibr" rid="B92">2020</xref>) is a hybrid model for text simplification and text summarization. HTSS is based on the Pointer Generator model (See et al., <xref ref-type="bibr" rid="B74">2017</xref>). We used HTSS implementation as it is without modification. To the best of our knowledge, HTSS is the only hybrid method for the combined task of summarization and simplification; thus, we consider it as a baseline model.</p></sec>
<sec>
<title>4.1.4 ProphetNet</title>
<p>ProphetNet introduces an innovative pre-training model for sequence-to-sequence tasks, incorporating a unique self-supervised objective termed future n-gram prediction along with a newly proposed n-stream self-attention mechanism. In contrast to traditional sequence-to-sequence models that focus on one-step ahead prediction, ProphetNet optimizes for n-step ahead prediction. This entails predicting the next n tokens simultaneously based on the context tokens at each time step. This distinctive approach explicitly encourages the model to anticipate future tokens, thereby addressing concerns related to overfitting on strong local correlations. The model underwent pre-training using both a base-scale dataset (16GB) and an extensive large-scale dataset (160GB). Then, the model was fine-tuned for downstream tasks such as text summarization on CNN/Daily-Mail dataset and question answering on SQuAD 1.1 dataset.</p></sec></sec>
<sec>
<title>4.2 Automatic evaluation</title>
<p>For the automatic evaluation of our proposed model and the baseline models, we used the ROUGE score as a metric for summarization SARI as a metric for simplification and CSS as a metric for combined summarization and simplification.</p>
<p>The results in terms of ROUGE scores SARI and CSS demonstrate that our newly proposed model: Simplification Aware Text summarization (SATS) significantly (<italic>p</italic> &#x0003C; 0.001 according to a paired <italic>t</italic>-test), outperforms ProphetNet on all metrics, although the gains in performance are reasonably small, with the majority of the improvement over other systems coming from the ProphetNet architecture itself. Our results demonstrate that ACCESS performs poorly on the summarization task showing that a simplification system alone is insufficient. ProphetNet gains a better SARI score than ACCESS (Martin et al., <xref ref-type="bibr" rid="B55">2020a</xref>) and MUSS (Martin et al., <xref ref-type="bibr" rid="B56">2020b</xref>), despite not being tuned for simplification. In our tuning, &#x003BB; was set to 0.8, indicating its usefulness. Our model is better than all other models in terms of the Combined Simplification and Summarization metric (<italic>CSS</italic><sub>1</sub>) score, which gives the harmonic mean of ROUGE1 and SARI, giving a new state of the art for the joint simplification and summarization task with our model.</p>
<p>We demonstrate that SATS outperforms ProphetNet and ACCESS on both a simplification dataset (Eureka) and a summarization dataset (CNN-Daily Mail). This demonstrates that the system we have developed is effective for the joint task of Simplification Aware Text Summarization that we set out to achieve. We also compared to the HTSS system, which is the former state of the art in terms of the CSS-1 score on the Eureka dataset, and we demonstrate that our new model outperforms HTSS not only on the Eureka Dataset but also on the CNN-Daily Mail dataset.</p>
<p>SATS is an adaptable architecture that can be fine-tuned for new datasets, and we envision that the results we have shown on the datasets we selected will be transferable to other datasets with the appropriate training being first undertaken. We have documented the model development and training process throughout our study to enable future researchers to replicate our results where needed. We also release our code and model checkpoints for future researchers to work with.</p>
<p>Further to the analysis that we have presented in <xref ref-type="table" rid="T3">Tables 3</xref>, <xref ref-type="table" rid="T4">4</xref>, we also present examples of the best and worst summaries provided by our system (SATS) in <xref ref-type="table" rid="T5">Tables 5</xref>, <xref ref-type="table" rid="T6">6</xref>. It is clear that when the system is performing well, the summaries are of a high quality and reflect the semantic content of the references, while not directly corresponding to the lexical and syntactic choices made by the original authors. The poor summaries are typically single lines that contain technical content. It is possible that this is emblematic of the model not being able to provide summaries for such content, which has led to a poor final output.</p>
<table-wrap position="float" id="T5">
<label>Table 5</label>
<caption><p>Examples of good summaries generated by SATS.</p></caption>
<table frame="box" rules="all">
<thead>
<tr style="background-color:#919498;color:#ffffff">
<th valign="top" align="left"><bold>S. No</bold></th>
<th valign="top" align="left"><bold>Reference</bold></th>
<th valign="top" align="left"><bold>Generated</bold></th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="left">1</td>
<td valign="top" align="left">&#x0201C;Women in Sub-Saharan Africa who use modern contraceptives are more likely to be tested for HIV than those who do not according to a study published April 25 2016 in the open-access journal Plos One by Katherine Center from the University of Arizona and colleagues. HIV/AIDS remains a major public health challenge for women in Sub-saharan Africa and medical appointments for use of modern contraceptives could provide an opportunity to offer confidential HIV testing&#x0201D;</td>
<td valign="top" align="left">&#x0201C;A new study published in Plos medicine finds that women who go to the doctor for contraception are more likely to be tested for HIV than those who use traditional methods of contraception according to a study published April 25 2016 in the open-access journal Plos One by Marie-Louise Derez from the University of Oxford UK and colleagues. HIV/AIDS is a major public health problem in sub-saharan Africa and the world health organization&#x0201D;</td>
</tr> <tr>
<td valign="top" align="left">2</td>
<td valign="top" align="left">&#x0201C;A study published today in the journal Plos Medicine has identified the five genetic variants associated with higher levels of the branched-chain amino acids isoleucine leucine and valine. the researchers also found that these genetic variants were associated with an increased risk of type 2 diabetes. The researchers led by the medical research council (MRC) epidemiology unit at the University of Cambridge used large-scale genetic data together with detailed measurements of the branched-chain amino acids and&#x0201D;</td>
<td valign="top" align="left">&#x0201C;A study led by the max planck institute for the science of human genetics (LMU) in collaboration with colleagues in Germany and the united states has found that higher levels of the branched-chain amino acids isoleucine leucine and valine are associated with a higher risk of type 2 diabetes. The study was published in Plos Genetics. High levels of these amino acids have been linked to increased risk of diabetes in people&#x0201D;</td>
</tr>
<tr>
<td valign="top" align="left">3</td>
<td valign="top" align="left">&#x0201C;Polar bears, use of land during substantial summer sea ice loss in the Chukchi Sea increased by 30 days according to a study published November 18 2015 in the open-access journal Plos One by Karyn rode from the U.S. Geological Survey Alaska and colleagues. recent observations suggest that polar bears increasingly use land habitats where they have minimal access to their preferred prey likely in response to loss of their sea ice habitat associated with climatic warming. The authors of this study used location data&#x0201D;</td>
<td valign="top" align="left">&#x0201C;Polar bears are increasingly using land habitats in some parts of their range where they have minimal access to their preferred prey according to new research published in the open-access journal Plos one. Polar bears in the Chukchi Sea in Russia have increasingly used land habitats during the summer and in the winter according to a study published November 18 2015 in the peer-reviewed journal Plos One by Alexander Kravchenko from the Russian academy of sciences and colleagues in the summer sea-ice loss&#x0201D;</td>
</tr></tbody>
</table>
</table-wrap>

<table-wrap position="float" id="T6">
<label>Table 6</label>
<caption><p>Examples of bad summaries generated by SATS.</p></caption>
<table frame="box" rules="all">
<thead>
<tr style="background-color:#919498;color:#ffffff">
<th valign="top" align="left"><bold>S. No</bold></th>
<th valign="top" align="left"><bold>Reference</bold></th>
<th valign="top" align="left"><bold>Generated</bold></th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="left">1</td>
<td valign="top" align="left">&#x0201C;Massachusetts general hospital (MGH) investigators may have uncovered a novel mechanism behind the ability of the diabetes drug Metformin to inhibit the progression of pancreatic cancer. In their report that has been published in the open access journal Plos One the research team describes finding that Metformin decreases the inflammation and fibrosis characteristic of the most common form of pancreatic cancer. Their findings in cellular and animal models and in patient tumor samples also indicate that this beneficial effect may be&#x0201D;</td>
<td valign="top" align="left">&#x0201C;Pancreatic ductal adenocarcinoma (PDAC) is a highly desmoplastic tumor with a dismal prognosis&#x0201D;</td>
</tr> <tr>
<td valign="top" align="left">2</td>
<td valign="top" align="left">&#x0201C;Studies find that female meerkats can produce up to twice as much testosterone as their male counterparts. Their unusual hormone levels fuel their bullying behavior but at a cost photo courtesy of Charli Davies Duke University. Testosterone. It&#x00027;s often lauded as the hormone that makes males bigger bolder stronger. Now a pair of Duke University studies has identified one group of animals the meerkats of the southern tip of Africa in which females can produce&#x0003E;</td>
<td valign="top" align="left">&#x0003C; this is a female meerkat (Suricata Suricatta)&#x0201D;</td>
</tr>
<tr>
<td valign="top" align="left">3</td>
<td valign="top" align="left">&#x0201C;Taking human immune diversity into account increases the clinical relevance of immunological research. La Jolla institute for allergy and immunology La Jolla CA &#x02014;mouse models have advanced our understanding of immune function and disease in many ways but they have failed to account for the natural diversity in human immune responses. As a result insights gained in the lab may be lost in translation. In their latest study researchers at La Jolla institute for allergy&#x0201D;</td>
<td valign="top" align="left">"This is a hybrid mouse diversity panel (HMDP) of 83 inbred strains of lipopolysaccharide (LPS)&#x0201D;</td>
</tr></tbody>
</table>
</table-wrap>
</sec>


<sec>
<title>4.3 Human evaluation</title>
<p>The quality of machine-generated summaries can be judged by automatic means such as ROUGE scores. These automatic metrics have their own limitations, such as the ROUGE metric favors short summaries and sometimes scores higher than expected in cases of extractive summarization. In the case of abstractive summarization, ROUGE score may be low for semantically identical summaries with high lexical differences. Therefore, we evaluated the summaries generated by SATS, and ProphetNet using a group of seven PhD students and one Masters student, all studying in computer science. We randomly selected 100 summaries from our test set. We then pose five questions about the quality of each summary following the study of Fabbri et al. (<xref ref-type="bibr" rid="B24">2021</xref>). Detail about these 5 questions is presented in <xref ref-type="table" rid="T7">Table 7</xref>, and we asked each question on a LIKERT scale from 1 to 5 where 5 means very good and 1 means very low. We carefully analyze responses from eight respondents and box plot responses from eight respondents and present it in <xref ref-type="fig" rid="F2">Figure 2</xref>, and the mean in the box plot for each of our evaluation questions shows a positive ranking overall across 100 selected summaries. To investigate statistical significance, We further perform a Friedman test (Sawilowsky and Fahoome, <xref ref-type="bibr" rid="B72">2005</xref>). We observed that statistic = 15.43 and <italic>p</italic> &#x0003D; 0.031, to interpret this <italic>P</italic>-value less than 0.05 means that the null hypothesis which states that the means of our observation are the same, can be rejected.</p>
<table-wrap position="float" id="T7">
<label>Table 7</label>
<caption><p>Evaluation questions.</p></caption>
<table frame="box" rules="all">
<thead>
<tr style="background-color:#919498;color:#ffffff">
<th valign="top" align="left"><bold>S. No</bold></th>
<th valign="top" align="left"><bold>Question</bold></th>
<th valign="top" align="left"><bold>Explanation</bold></th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="left">Q1: Coherence</td>
<td valign="top" align="left">How would you rate the Coherency of the Generated summary? On a scale from 1 to 5</td>
<td valign="top" align="left">The summary should be well-structured and well-organized. The summary should not just be a heap of related information, but should build from sentence to sentence to a coherent body of information about a topic; all the sentences should be well connected and have an overall theme or topic</td>
</tr> <tr>
<td valign="top" align="left">Q2: Consistency</td>
<td valign="top" align="left">How would you rate the Consistency of the Generated summary? On a scale from 1 to 5</td>
<td valign="top" align="left">The factual alignment between the summary and the summarized source. A factually consistent summary contains only statements that are entailed in the source document. penalize summaries that contain hallucinated (non-existent) facts. A consistent summary should contain all the facts and the correct information.</td>
</tr> <tr>
<td valign="top" align="left">Q3: Fluency and Grammatical</td>
<td valign="top" align="left">How would you rate the fluency and grammar of the Generated summary? On a scale from 1 to 5</td>
<td valign="top" align="left">The quality of individual sentences. Sentences in the summary should have no formatting problems, capitalization errors, or ungrammatical sentences (e.g., fragments, missing components) that make the text difficult to read. A fluent and grammatically correct summary should be easy to follow and have a natural flow in its sentences.</td>
</tr> <tr>
<td valign="top" align="left">Q4: Relevance with ground truth</td>
<td valign="top" align="left">Is the generated summary Relevant to the ground truth summary? On a scale from 1 to 5</td>
<td valign="top" align="left">The generated Summary by the model is easy to understand for non-native speakers. And is relevant to the gold standard summary provided. the information and key idea presented in the summary should match the key idea of the gold standard summary.</td>
</tr>
<tr>
<td valign="top" align="left">Q5: Simplicity</td>
<td valign="top" align="left">Is the generated summary simple and easy to understand? On a scale from 1 to 5</td>
<td valign="top" align="left">Selection of important content from the source. The summary should include only important information from the source document. Penalize summaries that contain redundancies and excess information. the generated summary should be as easy to understand for a level of high school students.</td>
</tr></tbody>
</table>
</table-wrap>

<fig id="F2" position="float">
<label>Figure 2</label>
<caption><p>Box plot of five questions responses for 100 summaries. <bold>(A)</bold> SATS. <bold>(B)</bold> ProphetNet.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="frai-07-1375419-g0002.tif"/>
</fig>



<p>To further analyze the responses, we planned to compute inter-rater agreement using Cohen kappa score (Cohen, <xref ref-type="bibr" rid="B19">1960</xref>; Artstein and Poesio, <xref ref-type="bibr" rid="B8">2008</xref>); using this, we calculate the Cohen Kappa score for each question of one responded to each question of the rest of the <italic>n</italic>&#x02212;1 responded and then average across five questions for each respondent; thus, we obtain an average score of five questions for each responded. Finally, we compute and present this analysis in <xref ref-type="fig" rid="F3">Figure 3</xref> where all possible agreements between the rater can be observed.</p>
<fig id="F3" position="float">
<label>Figure 3</label>
<caption><p>Overall summarization and simplification Interrater agreement: Cohen Kappa score between respondents. <bold>(A)</bold> SATS. (<bold>B)</bold> ProphetNet.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="frai-07-1375419-g0003.tif"/>
</fig>


<p>We investigated the simplification and summarization aspect simultaneously. In our human evaluation responses, we have a separate question for evaluating the easiness of our generated summaries. We analyzed the responses to question 5 and computed the Cohen Kappa score between the raters and presented our analysis in <xref ref-type="fig" rid="F4">Figure 4</xref>. In <xref ref-type="fig" rid="F4">Figure 4</xref>, we observe the majority of the raters show significant agreement that the generated summaries are simple and easy to understand.</p>


<fig id="F4" position="float">
<label>Figure 4</label>
<caption><p>Interratter agreement: Cohen Kappa score for simplification between respondents. <bold>(A)</bold> SATS. <bold>(B)</bold> ProphetNet.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="frai-07-1375419-g0004.tif"/>
</fig>

</sec></sec>
<sec sec-type="discussion" id="s5">
<title>5 Discussion</title>
<p>During our experiments and evaluation, we explored and found that summarization models perform better than the simplification models on the CNN-daily mail dataset. One of the possible reasons could be the structure and characteristics of the summarization dataset which is designed for the summarization task only. We also observed that summarization models perform better than the simplification model on the Eureka dataset. This is due to the easy word replacement of the simplification model which leads to generating different phrases as compared to the phrases present in the gold standard and thus leads to low scores as compared to the summarization task only. Our proposed SATS models improves the state-of-the-art HTSS model on the Eureka dataset and CNN-Daily Mail dataset. The improvement of the SATS model over ProphetNet is marginally low, this implies that adding a simplification module hurts the summarization capability of the model to some extent, the fact behind this is that summarization and simplification tasks often contradict each other, and such summarization aims at reducing the size of text and compressing while simplification aims at explaining and expanding for easiness in understanding. The CSS scores show overall improvements. In addition to this, the Rouge-L scores of the proposed SATS model are low on CNN-daily mail dataset as compared to prophetNet model due to the breaking of long phrases into simpler structures, and thus, long pieces could not match the gold set. Finally, human annotators found that SATS model produces a more simplified summary as compared to the ProphetNet model.</p></sec>
<sec id="s6">
<title>6 Concluding remarks</title>
<sec>
<title>6.1 Limitations</title>
<p>During our experiments and evaluation, we observed that our proposed SATS model has some limitations such as it is unable to understand and interpret mathematical formulas presented in the given source document. Besides from this, the proposed model is also unable to interpret and incorporate information presented in the form of figures. Processing long input text is also an issue and needs to be addressed.</p></sec>
<sec>
<title>6.2 Future work</title>
<p>Efforts were made to propose a model that can combine the task of summarization with simplification and improve the previously established state of the art for the combined task, but the field of summarization and especially simplification requires more attention to address the limitation discussed in Section 6.1.</p></sec>
<sec>
<title>6.3 Conclusion</title>
<p>Our study has explored a new method of adapting ProphetNet or other loss-based sequence-to-sequence generation methods to produce simplified summaries of scholarly documents. We evaluated our summaries using automatic metrics and human judgments, and we found that our generated summaries are up to the mark. We have demonstrated that this leads to a significant improvement in the advancement of research in scientific communications. Our study shows that hybrid simplification and summarization are possible and that models can produce high-fidelity simplified summaries compared to reference texts. Future work incorporating larger model architectures and advances in simplification and summarization will doubtlessly lead to improvements on this helpful task in future - under the umbrella of science communication, that is, making science understandable for everyone.</p></sec></sec>
<sec sec-type="data-availability" id="s7">
<title>Data availability statement</title>
<p>Publicly available datasets were analyzed in this study. This data can be found at: <ext-link ext-link-type="uri" xlink:href="https://github.com/slab-itu/HTSS/">https://github.com/slab-itu/HTSS/</ext-link>.</p></sec>
<sec sec-type="author-contributions" id="s8">
<title>Author contributions</title>
<p>FZ: Methodology, Writing &#x02013; original draft, Writing &#x02013; review &#x00026; editing, Conceptualization, Data curation, Software. FK: Project administration, Supervision, Validation, Writing &#x02013; original draft, Writing &#x02013; review &#x00026; editing. MS: Conceptualization, Data curation, Formal analysis, Investigation, Writing &#x02013; original draft, Writing &#x02013; review &#x00026; editing. S-UH: Conceptualization, Investigation, Methodology, Supervision, Writing &#x02013; original draft, Writing &#x02013; review &#x00026; editing. AK: Writing &#x02013; original draft, Writing &#x02013; review &#x00026; editing. NA: Writing &#x02013; original draft, Writing &#x02013; review &#x00026; editing.</p></sec>
</body>
<back>
<sec sec-type="funding-information" id="s9">
<title>Funding</title>
<p>The author(s) declare that no financial support was received for the research, authorship, and/or publication of this article.</p>
</sec>
<sec sec-type="COI-statement" id="conf1">
<title>Conflict of interest</title>
<p>The authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec sec-type="disclaimer" id="s10">
<title>Publisher&#x00027;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<fn-group>
<fn id="fn0001"><p><sup>1</sup><ext-link ext-link-type="uri" xlink:href="https://github.com/louismartin/dress-data/">https://github.com/louismartin/dress-data/</ext-link></p></fn>
<fn id="fn0002"><p><sup>2</sup><ext-link ext-link-type="uri" xlink:href="https://github.com/facebookresearch/asset">https://github.com/facebookresearch/asset</ext-link></p></fn>
<fn id="fn0003"><p><sup>3</sup><ext-link ext-link-type="uri" xlink:href="https://github.com/abisee/cnn-dailymail">https://github.com/abisee/cnn-dailymail</ext-link></p></fn>
<fn id="fn0004"><p><sup>4</sup><ext-link ext-link-type="uri" xlink:href="https://github.com/EdinburghNLP/XSum">https://github.com/EdinburghNLP/XSum</ext-link></p></fn>
<fn id="fn0005"><p><sup>5</sup><ext-link ext-link-type="uri" xlink:href="https://github.com/google-research-datasets/sentence-compression">https://github.com/google-research-datasets/sentence-compression</ext-link></p></fn>
</fn-group>
<ref-list>
<title>References</title>
<ref id="B1">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Acharya</surname> <given-names>S.</given-names></name> <name><surname>Boyd</surname> <given-names>A. D.</given-names></name> <name><surname>Cameron</surname> <given-names>R.</given-names></name> <name><surname>Lopez</surname> <given-names>K. D.</given-names></name> <name><surname>Martyn-Nemeth</surname> <given-names>P.</given-names></name> <name><surname>Dickens</surname> <given-names>C.</given-names></name> <etal/></person-group>. (<year>2019</year>). <article-title>&#x0201C;Incorporating personalization features in a hospital-stay summary generation system,&#x0201D;</article-title> in <source>Proceedings of the 52nd Hawaii International Conference on System Sciences</source> (<publisher-loc>Hawaii</publisher-loc>), <fpage>4175</fpage>&#x02013;<lpage>4184</lpage>. <pub-id pub-id-type="doi">10.24251/HICSS.2019.505</pub-id></citation>
</ref>
<ref id="B2">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Agrawal</surname> <given-names>S.</given-names></name> <name><surname>Carpuat</surname> <given-names>M.</given-names></name></person-group> (<year>2019</year>). <article-title>&#x0201C;Controlling text complexity in neural machine translation,&#x0201D;</article-title> in <source>Proceedings of the 2019 Conference on Empirical Methods in Natural Language Processing and the 9th International Joint Conference on Natural Language Processing (EMNLP-IJCNLP)</source> (<publisher-loc>Hong Kong</publisher-loc>: <publisher-name>Association for Computational Linguistics</publisher-name>), <fpage>1549</fpage>&#x02013;<lpage>1564</lpage>. <pub-id pub-id-type="doi">10.18653/v1/D19-1166</pub-id><pub-id pub-id-type="pmid">36568019</pub-id></citation></ref>
<ref id="B3">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Aharoni</surname> <given-names>R.</given-names></name> <name><surname>Goldberg</surname> <given-names>Y.</given-names></name></person-group> (<year>2018</year>). <article-title>&#x0201C;Split and rephrase: better evaluation and stronger baselines,&#x0201D;</article-title> in <source>Proceedings of the 56th Annual Meeting of the Association for Computational Linguistics (Volume 2: Short Papers)</source> (<publisher-loc>Melbourne, VIC</publisher-loc>: <publisher-name>Association for Computational Linguistics</publisher-name>), <fpage>719</fpage>&#x02013;<lpage>724</lpage>. <pub-id pub-id-type="doi">10.18653/v1/P18-2114</pub-id><pub-id pub-id-type="pmid">36568019</pub-id></citation></ref>
<ref id="B4">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Al-Thanyyan</surname> <given-names>S. S.</given-names></name> <name><surname>Azmi</surname> <given-names>A. M.</given-names></name></person-group> (<year>2021</year>). <article-title>Automated text simplification: a survey</article-title>. <source>ACM Comput. Surv</source>. <volume>54</volume>, <fpage>1</fpage>&#x02013;<lpage>36</lpage>. <pub-id pub-id-type="doi">10.1145/3442695</pub-id><pub-id pub-id-type="pmid">36083212</pub-id></citation></ref>
<ref id="B5">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Alva-Manchego</surname> <given-names>F.</given-names></name> <name><surname>Martin</surname> <given-names>L.</given-names></name> <name><surname>Bordes</surname> <given-names>A.</given-names></name> <name><surname>Scarton</surname> <given-names>C.</given-names></name> <name><surname>Sagot</surname> <given-names>B.</given-names></name> <name><surname>Specia</surname> <given-names>L.</given-names></name> <etal/></person-group>. (<year>2020a</year>). <article-title>&#x0201C;ASSET: a dataset for tuning and evaluation of sentence simplification models with multiple rewriting transformations,&#x0201D;</article-title> in <source>Proceedings of the 58th Annual Meeting of the Association for Computational Linguistics</source> (<publisher-loc>Association for Computational Linguistics</publisher-loc>), <fpage>4668</fpage>&#x02013;<lpage>4679</lpage>. <pub-id pub-id-type="doi">10.18653/v1/2020.acl-main.424</pub-id><pub-id pub-id-type="pmid">36568019</pub-id></citation></ref>
<ref id="B6">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Alva-Manchego</surname> <given-names>F.</given-names></name> <name><surname>Martin</surname> <given-names>L.</given-names></name> <name><surname>Scarton</surname> <given-names>C.</given-names></name> <name><surname>Specia</surname> <given-names>L.</given-names></name></person-group> (<year>2019</year>). <article-title>&#x0201C;EASSE: easier automatic sentence simplification evaluation,&#x0201D;</article-title> in <source>Proceedings of the 2019 Conference on Empirical Methods in Natural Language Processing and the 9th International Joint Conference on Natural Language Processing (EMNLP-IJCNLP): System Demonstrations</source> (<publisher-loc>Hong Kong</publisher-loc>: <publisher-name>Association for Computational Linguistics</publisher-name>), <fpage>49</fpage>&#x02013;<lpage>54</lpage>. <pub-id pub-id-type="doi">10.18653/v1/D19-3009</pub-id><pub-id pub-id-type="pmid">36568019</pub-id></citation></ref>
<ref id="B7">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Alva-Manchego</surname> <given-names>F.</given-names></name> <name><surname>Scarton</surname> <given-names>C.</given-names></name> <name><surname>Specia</surname> <given-names>L.</given-names></name></person-group> (<year>2020b</year>). <article-title>Data-driven sentence simplification: survey and benchmark</article-title>. <source>Comput. Linguist</source>. <volume>46</volume>, <fpage>135</fpage>&#x02013;<lpage>187</lpage>. <pub-id pub-id-type="doi">10.1162/coli_a_00370</pub-id></citation>
</ref>
<ref id="B8">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Artstein</surname> <given-names>R.</given-names></name> <name><surname>Poesio</surname> <given-names>M.</given-names></name></person-group> (<year>2008</year>). <article-title>Inter-coder agreement for computational linguistics</article-title>. <source>Comput. Linguist</source>. <volume>34</volume>, <fpage>555</fpage>&#x02013;<lpage>596</lpage>. <pub-id pub-id-type="doi">10.1162/coli.07-034-R2</pub-id></citation>
</ref>
<ref id="B9">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Azmi</surname> <given-names>A. M.</given-names></name> <name><surname>Altmami</surname> <given-names>N. I.</given-names></name></person-group> (<year>2018</year>). <article-title>An abstractive arabic text summarizer with user controlled granularity</article-title>. <source>Inform. Process. Manag</source>. <volume>54</volume>, <fpage>903</fpage>&#x02013;<lpage>921</lpage>. <pub-id pub-id-type="doi">10.1016/j.ipm.2018.06.002</pub-id></citation>
</ref>
<ref id="B10">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Banerjee</surname> <given-names>S.</given-names></name> <name><surname>Lavie</surname> <given-names>A.</given-names></name></person-group> (<year>2005</year>). <article-title>&#x0201C;METEOR: an automatic metric for MT evaluation with improved correlation with human judgments,&#x0201D;</article-title> in <source>Proceedings of the ACL Workshop on Intrinsic and Extrinsic Evaluation Measures for Machine Translation and/or Summarization</source> (<publisher-loc>Ann Arbor, MI</publisher-loc>: <publisher-name>Association for Computational Linguistics</publisher-name>), <fpage>65</fpage>&#x02013;<lpage>72</lpage>.</citation>
</ref>
<ref id="B11">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Barros</surname> <given-names>C.</given-names></name> <name><surname>Lloret</surname> <given-names>E.</given-names></name> <name><surname>Saquete</surname> <given-names>E.</given-names></name> <name><surname>Navarro-Colorado</surname> <given-names>B.</given-names></name></person-group> (<year>2019</year>). <article-title>Natsum: narrative abstractive summarization through cross-document timeline generation</article-title>. <source>Inform. Process. Manag</source>. <volume>56</volume>, <fpage>1775</fpage>&#x02013;<lpage>1793</lpage>. <pub-id pub-id-type="doi">10.1016/j.ipm.2019.02.010</pub-id></citation>
</ref>
<ref id="B12">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Baxendale</surname> <given-names>P. B.</given-names></name></person-group> (<year>1958</year>). <article-title>Machine-made index for technical literature-an experiment</article-title>. <source>IBM J. Res. Dev</source>. <volume>2</volume>, <fpage>354</fpage>&#x02013;<lpage>361</lpage>. <pub-id pub-id-type="doi">10.1147/rd.24.0354</pub-id><pub-id pub-id-type="pmid">33813791</pub-id></citation></ref>
<ref id="B13">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Botha</surname> <given-names>J. A.</given-names></name> <name><surname>Faruqui</surname> <given-names>M.</given-names></name> <name><surname>Alex</surname> <given-names>J.</given-names></name> <name><surname>Baldridge</surname> <given-names>J.</given-names></name> <name><surname>Das</surname> <given-names>D.</given-names></name></person-group> (<year>2018</year>). <article-title>&#x0201C;Learning to split and rephrase from Wikipedia edit history,&#x0201D;</article-title> in <source>Proceedings of the 2018 Conference on Empirical Methods in Natural Language Processing</source> (<publisher-loc>Brussels</publisher-loc>: <publisher-name>Association for Computational Linguistics</publisher-name>), <fpage>732</fpage>&#x02013;<lpage>737</lpage>. <pub-id pub-id-type="doi">10.18653/v1/D18-1080</pub-id><pub-id pub-id-type="pmid">36568019</pub-id></citation></ref>
<ref id="B14">
<citation citation-type="web"><person-group person-group-type="author"><name><surname>Brants</surname> <given-names>T.</given-names></name></person-group> (<year>2006</year>). <source>Web 1t 5-gram version 1</source>. Available online at: <ext-link ext-link-type="uri" xlink:href="https://catalog.ldc.upenn.edu/LDC2006T13">https://catalog.ldc.upenn.edu/LDC2006T13</ext-link> (accessed June 24, 2024).</citation>
</ref>
<ref id="B15">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Cai</surname> <given-names>X.</given-names></name> <name><surname>Shi</surname> <given-names>K.</given-names></name> <name><surname>Jiang</surname> <given-names>Y.</given-names></name> <name><surname>Yang</surname> <given-names>L.</given-names></name> <name><surname>Liu</surname> <given-names>S.</given-names></name></person-group> (<year>2021</year>). <article-title>Hits-based attentional neural model for abstractive summarization</article-title>. <source>Knowl.-Based Syst</source>. <volume>222</volume>:<fpage>106996</fpage>. <pub-id pub-id-type="doi">10.1016/j.knosys.2021.106996</pub-id></citation>
</ref>
<ref id="B16">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Carroll</surname> <given-names>J.</given-names></name> <name><surname>Minnen</surname> <given-names>G.</given-names></name> <name><surname>Canning</surname> <given-names>Y.</given-names></name> <name><surname>Devlin</surname> <given-names>S.</given-names></name> <name><surname>Tait</surname> <given-names>J.</given-names></name></person-group> (<year>1998</year>). <article-title>&#x0201C;Practical simplification of english newspaper text to assist aphasic readers,&#x0201D;</article-title> in <source>Proc. of AAAI-98 Workshop on Integrating Artificial Intelligence and Assistive Technology</source> (<publisher-loc>Madison, WI</publisher-loc>), <fpage>7</fpage>&#x02013;<lpage>10</lpage>.</citation>
</ref>
<ref id="B17">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Chen</surname> <given-names>Y.-C.</given-names></name> <name><surname>Bansal</surname> <given-names>M.</given-names></name></person-group> (<year>2018</year>). <article-title>&#x0201C;Fast abstractive summarization with reinforce-selected sentence rewriting,&#x0201D;</article-title> in <source>Proceedings of the 56th Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers)</source> (<publisher-loc>Melbourne, VIC</publisher-loc>: <publisher-name>Association for Computational Linguistics</publisher-name>), <fpage>675</fpage>&#x02013;<lpage>686</lpage>. <pub-id pub-id-type="doi">10.18653/v1/P18-1063</pub-id><pub-id pub-id-type="pmid">36568019</pub-id></citation></ref>
<ref id="B18">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Cho</surname> <given-names>J.</given-names></name> <name><surname>Seo</surname> <given-names>M.</given-names></name> <name><surname>Hajishirzi</surname> <given-names>H.</given-names></name></person-group> (<year>2019</year>). <article-title>&#x0201C;Mixture content selection for diverse sequence generation,&#x0201D;</article-title> in <source>Proceedings of the 2019 Conference on Empirical Methods in Natural Language Processing and the 9th International Joint Conference on Natural Language Processing (EMNLP-IJCNLP)</source> (<publisher-loc>Hong Kong</publisher-loc>), <fpage>3112</fpage>&#x02013;<lpage>3122</lpage>. <pub-id pub-id-type="doi">10.18653/v1/D19-1308</pub-id><pub-id pub-id-type="pmid">36568019</pub-id></citation></ref>
<ref id="B19">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Cohen</surname> <given-names>J.</given-names></name></person-group> (<year>1960</year>). <article-title>A coefficient of agreement for nominal scales</article-title>. <source>Educ. Psychol. Meas</source>. <volume>20</volume>, <fpage>37</fpage>&#x02013;<lpage>46</lpage>. <pub-id pub-id-type="doi">10.1177/001316446002000104</pub-id></citation>
</ref>
<ref id="B20">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Collins</surname> <given-names>E.</given-names></name> <name><surname>Augenstein</surname> <given-names>I.</given-names></name> <name><surname>Riedel</surname> <given-names>S.</given-names></name></person-group> (<year>2017</year>). <article-title>&#x0201C;A supervised approach to extractive summarisation of scientific papers,&#x0201D;</article-title> in <source>Proceedings of the 21st Conference on Computational Natural Language Learning (CoNLL 2017)</source> (<publisher-loc>Vancouver, BC</publisher-loc>: <publisher-name>Association for Computational Linguistics</publisher-name>), <fpage>195</fpage>&#x02013;<lpage>205</lpage>. <pub-id pub-id-type="doi">10.18653/v1/K17-1021</pub-id><pub-id pub-id-type="pmid">36568019</pub-id></citation></ref>
<ref id="B21">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Devlin</surname> <given-names>J.</given-names></name> <name><surname>Chang</surname> <given-names>M.-W.</given-names></name> <name><surname>Lee</surname> <given-names>K.</given-names></name> <name><surname>Toutanova</surname> <given-names>K.</given-names></name></person-group> (<year>2019</year>). <article-title>&#x0201C;Bert: pre-training of deep bidirectional transformers for language understanding,&#x0201D;</article-title> in <source>Proceedings of the 2019 Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Technologies, Volume 1 (Long and Short Papers)</source> (<publisher-loc>Minneapolis, MN</publisher-loc>: <publisher-name>Association for Computational Linguistics</publisher-name>), <fpage>4171</fpage>&#x02013;<lpage>4186</lpage>.</citation>
</ref>
<ref id="B22">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Dong</surname> <given-names>L.</given-names></name> <name><surname>Yang</surname> <given-names>N.</given-names></name> <name><surname>Wang</surname> <given-names>W.</given-names></name> <name><surname>Wei</surname> <given-names>F.</given-names></name> <name><surname>Liu</surname> <given-names>X.</given-names></name> <name><surname>Wang</surname> <given-names>Y.</given-names></name> <etal/></person-group>. (<year>2019</year>). <article-title>&#x0201C;Unified language model pre-training for natural language understanding and generation,&#x0201D;</article-title> in <source>Advances in Neural Information Processing Systems</source> (<publisher-loc>Vancouver, BC</publisher-loc>), <fpage>13042</fpage>&#x02013;<lpage>13054</lpage>.</citation>
</ref>
<ref id="B23">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Edmundson</surname> <given-names>H. P.</given-names></name></person-group> (<year>1969</year>). <article-title>New methods in automatic extracting</article-title>. <source>J. ACM</source> <volume>16</volume>, <fpage>264</fpage>&#x02013;<lpage>285</lpage>. <pub-id pub-id-type="doi">10.1145/321510.321519</pub-id></citation>
</ref>
<ref id="B24">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Fabbri</surname> <given-names>A. R.</given-names></name> <name><surname>Kry&#x0015B;ci&#x00144;ski</surname> <given-names>W.</given-names></name> <name><surname>McCann</surname> <given-names>B.</given-names></name> <name><surname>Xiong</surname> <given-names>C.</given-names></name> <name><surname>Socher</surname> <given-names>R.</given-names></name> <name><surname>Radev</surname> <given-names>D.</given-names></name></person-group> (<year>2021</year>). <article-title>Summeval: re-evaluating summarization evaluation</article-title>. <source>Trans. Assoc. Comput. Linguist</source>. <volume>9</volume>, <fpage>391</fpage>&#x02013;<lpage>409</lpage>. <pub-id pub-id-type="doi">10.1162/tacl_a_00373</pub-id></citation>
</ref>
<ref id="B25">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Feng</surname> <given-names>L.</given-names></name></person-group> (<year>2008</year>). <source>Text simplification: A survey</source>. <publisher-loc>Technical Report. New York, NY</publisher-loc>: <publisher-name>The City University of New York</publisher-name>.</citation>
</ref>
<ref id="B26">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Filatova</surname> <given-names>E.</given-names></name> <name><surname>Hatzivassiloglou</surname> <given-names>V.</given-names></name></person-group> (<year>2004</year>). <article-title>&#x0201C;Event-based extractive summarization,&#x0201D;</article-title> in <source>Text Summarization Branches Out</source> (<publisher-loc>Barcelona</publisher-loc>: <publisher-name>Association for Computational Linguistics</publisher-name>), <fpage>104</fpage>&#x02013;<lpage>111</lpage>.</citation>
</ref>
<ref id="B27">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Filippova</surname> <given-names>K.</given-names></name> <name><surname>Alfonseca</surname> <given-names>E.</given-names></name> <name><surname>Colmenares</surname> <given-names>C. A.</given-names></name> <name><surname>Kaiser</surname> <given-names>&#x00141;.</given-names></name> <name><surname>Vinyals</surname> <given-names>O.</given-names></name></person-group> (<year>2015</year>). <article-title>&#x0201C;Sentence compression by deletion with lstms,&#x0201D;</article-title> in <source>Proceedings of the 2015 Conference on Empirical Methods in Natural Language Processing</source> (<publisher-loc>Lisbon</publisher-loc>), <fpage>360</fpage>&#x02013;<lpage>368</lpage>. <pub-id pub-id-type="doi">10.18653/v1/D15-1042</pub-id><pub-id pub-id-type="pmid">36568019</pub-id></citation></ref>
<ref id="B28">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Filippova</surname> <given-names>K.</given-names></name> <name><surname>Altun</surname> <given-names>Y.</given-names></name></person-group> (<year>2013</year>). <article-title>&#x0201C;Overcoming the lack of parallel data in sentence compression,&#x0201D;</article-title> in <source>Proceedings of the 2013 Conference on Empirical Methods in Natural Language Processing</source> (<publisher-loc>Seattle, WA</publisher-loc>), <fpage>1481</fpage>&#x02013;<lpage>1491</lpage>.</citation>
</ref>
<ref id="B29">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Giarelis</surname> <given-names>N.</given-names></name> <name><surname>Mastrokostas</surname> <given-names>C.</given-names></name> <name><surname>Karacapilidis</surname> <given-names>N.</given-names></name></person-group> (<year>2023</year>). <article-title>Greekt5: a series of greek sequence-to-sequence models for news summarization</article-title>. <source>arXiv</source> [Preprint]. arXiv:2311.07767. <pub-id pub-id-type="doi">10.48550/arXiv.2311.07767</pub-id></citation>
</ref>
<ref id="B30">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Goldsack</surname> <given-names>T.</given-names></name> <name><surname>Zhang</surname> <given-names>Z.</given-names></name> <name><surname>Lin</surname> <given-names>C.</given-names></name> <name><surname>Scarton</surname> <given-names>C.</given-names></name></person-group> (<year>2022</year>). <article-title>&#x0201C;Making science simple: corpora for the lay summarisation of scientific literature,&#x0201D;</article-title> in <source>Proceedings of the 2022 Conference on Empirical Methods in Natural Language Processing</source> (<publisher-loc>Abu Dhabi</publisher-loc>), <fpage>10589</fpage>&#x02013;<lpage>10604</lpage>. <pub-id pub-id-type="doi">10.18653/v1/2022.emnlp-main.724</pub-id><pub-id pub-id-type="pmid">36568019</pub-id></citation></ref>
<ref id="B31">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Goodfellow</surname> <given-names>I.</given-names></name> <name><surname>Pouget-Abadie</surname> <given-names>J.</given-names></name> <name><surname>Mirza</surname> <given-names>M.</given-names></name> <name><surname>Xu</surname> <given-names>B.</given-names></name> <name><surname>Warde-Farley</surname> <given-names>D.</given-names></name> <name><surname>Ozair</surname> <given-names>S.</given-names></name> <etal/></person-group>. (<year>2014</year>). <article-title>&#x0201C;Generative adversarial nets,&#x0201D;</article-title> in <source>Advances in neural information processing systems</source> (<publisher-loc>Montreal, QC</publisher-loc>), <fpage>2672</fpage>&#x02013;<lpage>2680</lpage>.</citation>
</ref>
<ref id="B32">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Graff</surname> <given-names>D.</given-names></name> <name><surname>Kong</surname> <given-names>J.</given-names></name> <name><surname>Chen</surname> <given-names>K.</given-names></name> <name><surname>Maeda</surname> <given-names>K.</given-names></name></person-group> (<year>2003</year>). <source>English gigaword</source>, 4th Edn. Philadelphia, PA: Linguistic Data Consortium, 34.</citation>
</ref>
<ref id="B33">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Hoard</surname> <given-names>J. E.</given-names></name> <name><surname>Wojcik</surname> <given-names>R.</given-names></name> <name><surname>Holzhauser</surname> <given-names>K.</given-names></name></person-group> (<year>1992</year>). <article-title>&#x0201C;An automated grammar and style checker for writers of simplified english,&#x0201D;</article-title> in Computers <italic>and Writing: State of the Art</italic> (Dordrecht: Springer Netherlands), <fpage>278</fpage>&#x02013;<lpage>296</lpage>. <pub-id pub-id-type="doi">10.1007/978-94-011-2854-4_19</pub-id></citation>
</ref>
<ref id="B34">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Hou</surname> <given-names>J.</given-names></name> <name><surname>Wang</surname> <given-names>Y.</given-names></name> <name><surname>Zhang</surname> <given-names>Y.</given-names></name> <name><surname>Wang</surname> <given-names>D.</given-names></name></person-group> (<year>2022</year>). <article-title>How do scholars and non-scholars participate in dataset dissemination on twitter</article-title>. <source>J. Informetr</source>. <volume>16</volume>:<fpage>101223</fpage>. <pub-id pub-id-type="doi">10.1016/j.joi.2021.101223</pub-id></citation>
</ref>
<ref id="B35">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Iqbal</surname> <given-names>S.</given-names></name> <name><surname>Hassan</surname> <given-names>S.-U.</given-names></name> <name><surname>Aljohani</surname> <given-names>N. R.</given-names></name> <name><surname>Alelyani</surname> <given-names>S.</given-names></name> <name><surname>Nawaz</surname> <given-names>R.</given-names></name> <name><surname>Bornmann</surname> <given-names>L.</given-names></name> <etal/></person-group>. (<year>2021</year>). <article-title>A decade of in-text citation analysis based on natural language processing and machine learning techniques: an overview of empirical studies</article-title>. <source>Scientometrics</source> <volume>126</volume>, <fpage>6551</fpage>&#x02013;<lpage>6599</lpage>. <pub-id pub-id-type="doi">10.1007/s11192-021-04055-1</pub-id></citation>
</ref>
<ref id="B36">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Jia</surname> <given-names>Q.</given-names></name> <name><surname>Ren</surname> <given-names>S.</given-names></name> <name><surname>Liu</surname> <given-names>Y.</given-names></name> <name><surname>Zhu</surname> <given-names>K. Q.</given-names></name></person-group> (<year>2023</year>). <article-title>Zero-shot faithfulness evaluation for text summarization with foundation language model</article-title>. <source>arXiv</source> [Preprint]. arXiv:2310.11648. <pub-id pub-id-type="doi">10.48550/arXiv.2310.11648</pub-id></citation>
</ref>
<ref id="B37">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Kamigaito</surname> <given-names>H.</given-names></name> <name><surname>Hayashi</surname> <given-names>K.</given-names></name> <name><surname>Hirao</surname> <given-names>T.</given-names></name> <name><surname>Nagata</surname> <given-names>M.</given-names></name></person-group> (<year>2018</year>). <article-title>&#x0201C;Higher-order syntactic attention network for longer sentence compression,&#x0201D;</article-title> in <source>Proceedings of the 2018 Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Technologies, Volume 1 (Long Papers)</source> (<publisher-loc>New Orleans, LA</publisher-loc>: <publisher-name>Association for Computational Linguistics</publisher-name>), <fpage>1716</fpage>&#x02013;<lpage>1726</lpage>. <pub-id pub-id-type="doi">10.18653/v1/N18-1155</pub-id><pub-id pub-id-type="pmid">36568019</pub-id></citation></ref>
<ref id="B38">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Kamigaito</surname> <given-names>H.</given-names></name> <name><surname>Okumura</surname> <given-names>M.</given-names></name></person-group> (<year>2020</year>). <article-title>&#x0201C;Syntactically look-ahead attention network for sentence compression,&#x0201D;</article-title> in <source>Proceedings of the AAAI Conference on Artificial Intelligence, volume</source> 34 (New York, NY), <fpage>8050</fpage>&#x02013;<lpage>8057</lpage>. <pub-id pub-id-type="doi">10.1609/aaai.v34i05.6315</pub-id></citation>
</ref>
<ref id="B39">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Kincaid</surname> <given-names>J. P.</given-names></name> <name><surname>Fishburne Jr</surname> <given-names>R. P.</given-names></name> <name><surname>Rogers</surname> <given-names>R. L.</given-names></name> <name><surname>Chissom</surname> <given-names>B. S.</given-names></name></person-group> (<year>1975</year>). <source>Derivation of new readability formulas (automated readability index, fog count and flesch reading ease formula) for navy enlisted personnel</source>. <publisher-loc>Technical report. Millington, TN</publisher-loc>: <publisher-name>Naval Technical Training Command</publisher-name>. <pub-id pub-id-type="doi">10.21236/ADA006655</pub-id></citation>
</ref>
<ref id="B40">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Kinugawa</surname> <given-names>K.</given-names></name> <name><surname>Tsuruoka</surname> <given-names>Y.</given-names></name></person-group> (<year>2017</year>). <article-title>&#x0201C;A hierarchical neural extractive summarizer for academic papers,&#x0201D;</article-title> in <source>JSAI International Symposium on Artificial Intelligence</source> (<publisher-loc>New York, NY</publisher-loc>: <publisher-name>Springer</publisher-name>), <fpage>339</fpage>&#x02013;<lpage>354</lpage>. <pub-id pub-id-type="doi">10.1007/978-3-319-93794-6_25</pub-id></citation>
</ref>
<ref id="B41">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Lewis</surname> <given-names>M.</given-names></name> <name><surname>Liu</surname> <given-names>Y.</given-names></name> <name><surname>Goyal</surname> <given-names>N.</given-names></name> <name><surname>Ghazvininejad</surname> <given-names>M.</given-names></name> <name><surname>Mohamed</surname> <given-names>A.</given-names></name> <name><surname>Levy</surname> <given-names>O.</given-names></name> <etal/></person-group>. (<year>2020</year>). <article-title>&#x0201C;BART: denoising sequence-to-sequence pre-training for natural language generation, translation, and comprehension,&#x0201D;</article-title> in <source>Proceedings of the 58th Annual Meeting of the Association for Computational Linguistics</source> (<publisher-loc>Association for Computational Linguistics</publisher-loc>), <fpage>7871</fpage>&#x02013;<lpage>7880</lpage>. <pub-id pub-id-type="doi">10.18653/v1/2020.acl-main.703</pub-id><pub-id pub-id-type="pmid">36568019</pub-id></citation></ref>
<ref id="B42">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Li</surname> <given-names>T.</given-names></name> <name><surname>Li</surname> <given-names>Y.</given-names></name> <name><surname>Qiang</surname> <given-names>J.</given-names></name> <name><surname>Yuan</surname> <given-names>Y.-H.</given-names></name></person-group> (<year>2018</year>). <article-title>&#x0201C;Text simplification with self-attention-based pointer-generator networks,&#x0201D;</article-title> in <source>Neural Information Processing</source>, eds. L. Cheng, A. C. S. Leung, and S. Ozawa (Cham: Springer International Publishing), <fpage>537</fpage>&#x02013;<lpage>545</lpage>. <pub-id pub-id-type="doi">10.1007/978-3-030-04221-9_48</pub-id></citation>
</ref>
<ref id="B43">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Lin</surname> <given-names>C.-Y.</given-names></name></person-group> (<year>2004</year>). &#x0201C;ROUGE: a package for automatic evaluation of summaries,&#x0201D; <italic>in Text Summarization Branches Out</italic> (Barcelona: Association for Computational Linguistics), <fpage>74</fpage>&#x02013;<lpage>81</lpage>.</citation>
</ref>
<ref id="B44">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Lin</surname> <given-names>C.-Y.</given-names></name> <name><surname>Hovy</surname> <given-names>E.</given-names></name></person-group> (<year>2003</year>). <article-title>&#x0201C;Automatic evaluation of summaries using n-gram co-occurrence statistics,&#x0201D;</article-title> in <source>Proceedings of the 2003 Human Language Technology Conference of the North American Chapter of the Association for Computational Linguistics</source> (<publisher-loc>Stroudsburg, PA</publisher-loc>: <publisher-name>ACM</publisher-name>), <fpage>150</fpage>&#x02013;<lpage>157</lpage>. <pub-id pub-id-type="doi">10.3115/1073445.1073465</pub-id></citation>
</ref>
<ref id="B45">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Lin</surname> <given-names>C.-Y.</given-names></name> <name><surname>Och</surname> <given-names>F. J.</given-names></name></person-group> (<year>2004</year>). <article-title>&#x0201C;Automatic evaluation of machine translation quality using longest common subsequence and skip-bigram statistics,&#x0201D;</article-title> in <source>Proceedings of the 42nd Annual Meeting of the Association for Computational Linguistics (ACL-04)</source> (<publisher-loc>Barcelona</publisher-loc>: <publisher-name>ACL</publisher-name>), <fpage>605</fpage>&#x02013;<lpage>612</lpage>. <pub-id pub-id-type="doi">10.3115/1218955.1219032</pub-id></citation>
</ref>
<ref id="B46">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Liu</surname> <given-names>F.</given-names></name> <name><surname>Flanigan</surname> <given-names>J.</given-names></name> <name><surname>Thomson</surname> <given-names>S.</given-names></name> <name><surname>Sadeh</surname> <given-names>N.</given-names></name> <name><surname>Smith</surname> <given-names>N. A.</given-names></name></person-group> (<year>2015</year>). <article-title>&#x0201C;Toward abstractive summarization using semantic representations,&#x0201D;</article-title> in <source>Proceedings of the 2015 Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Technologies</source> (<publisher-loc>Denver, CO</publisher-loc>: <publisher-name>Association for Computational Linguistics</publisher-name>), <fpage>1077</fpage>&#x02013;<lpage>1086</lpage>. <pub-id pub-id-type="doi">10.3115/v1/N15-1114</pub-id></citation>
</ref>
<ref id="B47">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Liu</surname> <given-names>L.</given-names></name> <name><surname>Lu</surname> <given-names>Y.</given-names></name> <name><surname>Yang</surname> <given-names>M.</given-names></name> <name><surname>Qu</surname> <given-names>Q.</given-names></name> <name><surname>Zhu</surname> <given-names>J.</given-names></name> <name><surname>Li</surname> <given-names>H.</given-names></name> <etal/></person-group>. (<year>2018</year>). <article-title>Generative adversarial network for abstractive text summarization</article-title>. <source>Proc. AAAI Conf. Artif. Intell</source>. <volume>32</volume>, <fpage>8109</fpage>&#x02013;<lpage>8110</lpage>. <pub-id pub-id-type="doi">10.1609/aaai.v32i1.12141</pub-id><pub-id pub-id-type="pmid">32701451</pub-id></citation></ref>
<ref id="B48">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Liu</surname> <given-names>Y.</given-names></name> <name><surname>Fabbri</surname> <given-names>A. R.</given-names></name> <name><surname>Chen</surname> <given-names>J.</given-names></name> <name><surname>Zhao</surname> <given-names>Y.</given-names></name> <name><surname>Han</surname> <given-names>S.</given-names></name> <name><surname>Joty</surname> <given-names>S.</given-names></name> <etal/></person-group>. (<year>2023</year>). <article-title>Benchmarking generation and evaluation capabilities of large language models for instruction controllable summarization</article-title>. <source>arXiv</source> [Preprint]. arXiv:2311.09184. <pub-id pub-id-type="doi">10.48550/arXiv.2311.09184</pub-id></citation>
</ref>
<ref id="B49">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Liu</surname> <given-names>Y.</given-names></name> <name><surname>Lapata</surname> <given-names>M.</given-names></name></person-group> (<year>2019</year>). <article-title>&#x0201C;Text summarization with pretrained encoders,&#x0201D;</article-title> in <source>Proceedings of the 2019 Conference on Empirical Methods in Natural Language Processing and the 9th International Joint Conference on Natural Language Processing (EMNLP-IJCNLP)</source> (<publisher-loc>Hong Kong</publisher-loc>), <fpage>3721</fpage>&#x02013;<lpage>3731</lpage>. <pub-id pub-id-type="doi">10.18653/v1/D19-1387</pub-id><pub-id pub-id-type="pmid">36568019</pub-id></citation></ref>
<ref id="B50">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Luhn</surname> <given-names>H. P.</given-names></name></person-group> (<year>1958</year>). <article-title>The automatic creation of literature abstracts</article-title>. <source>IBM J. Res. Dev</source>. <volume>2</volume>, <fpage>159</fpage>&#x02013;<lpage>165</lpage>. <pub-id pub-id-type="doi">10.1147/rd.22.0159</pub-id><pub-id pub-id-type="pmid">33813791</pub-id></citation></ref>
<ref id="B51">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Macdonald</surname> <given-names>I.</given-names></name> <name><surname>Siddharthan</surname> <given-names>A.</given-names></name></person-group> (<year>2016</year>). <article-title>&#x0201C;Summarising news stories for children,&#x0201D;</article-title> in <source>Proceedings of the 9th International Natural Language Generation conference</source> (<publisher-loc>Edinburgh</publisher-loc>: <publisher-name>Association for Computational Linguistics</publisher-name>), <fpage>1</fpage>&#x02013;<lpage>10</lpage>. <pub-id pub-id-type="doi">10.18653/v1/W16-6601</pub-id><pub-id pub-id-type="pmid">36568019</pub-id></citation></ref>
<ref id="B52">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Mackie</surname> <given-names>S.</given-names></name> <name><surname>McCreadie</surname> <given-names>R.</given-names></name> <name><surname>Macdonald</surname> <given-names>C.</given-names></name> <name><surname>Ounis</surname> <given-names>I.</given-names></name></person-group> (<year>2014</year>). <article-title>&#x0201C;Comparing algorithms for microblog summarisation,&#x0201D;</article-title> in <source>Information Access Evaluation. Multilinguality, Multimodality, and Interaction</source>, eds. E. Kanoulas, M. Lupu, P. Clough, M. Sanderson, M. Hall, A. Hanbury, et al.(Cham: Springer International Publishing), <fpage>153</fpage>&#x02013;<lpage>159</lpage>. <pub-id pub-id-type="doi">10.1007/978-3-319-11382-1_15</pub-id><pub-id pub-id-type="pmid">29107976</pub-id></citation></ref>
<ref id="B53">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Mao</surname> <given-names>X.</given-names></name> <name><surname>Huang</surname> <given-names>S.</given-names></name> <name><surname>Shen</surname> <given-names>L.</given-names></name> <name><surname>Li</surname> <given-names>R.</given-names></name> <name><surname>Yang</surname> <given-names>H.</given-names></name></person-group> (<year>2021</year>). <article-title>Single document summarization using the information from documents with the same topic</article-title>. <source>Knowl.-Based Syst</source>. <volume>228</volume>:<fpage>107265</fpage>. <pub-id pub-id-type="doi">10.1016/j.knosys.2021.107265</pub-id></citation>
</ref>
<ref id="B54">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Marchisio</surname> <given-names>K.</given-names></name> <name><surname>Guo</surname> <given-names>J.</given-names></name> <name><surname>Lai</surname> <given-names>C.-I.</given-names></name> <name><surname>Koehn</surname> <given-names>P.</given-names></name></person-group> (<year>2019</year>). <article-title>&#x0201C;Controlling the reading level of machine translation output,&#x0201D;</article-title> in <source>Proceedings of Machine Translation Summit XVII Volume 1: Research Track</source> (<publisher-loc>Dublin</publisher-loc>: <publisher-name>European Association for Machine Translation</publisher-name>), <fpage>193</fpage>&#x02013;<lpage>203</lpage>.</citation>
</ref>
<ref id="B55">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Martin</surname> <given-names>L.</given-names></name> <name><surname>de la Clergerie</surname> <given-names>&#x000C9;.</given-names></name> <name><surname>Sagot</surname> <given-names>B.</given-names></name> <name><surname>Bordes</surname> <given-names>A.</given-names></name></person-group> (<year>2020a</year>). <article-title>&#x0201C;Controllable sentence simplification,&#x0201D;</article-title> in <source>Proceedings of the 12th Language Resources and Evaluation Conference</source> (<publisher-loc>Marseille</publisher-loc>: <publisher-name>European Language Resources Association</publisher-name>), <fpage>4689</fpage>&#x02013;<lpage>4698</lpage>.</citation>
</ref>
<ref id="B56">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Martin</surname> <given-names>L.</given-names></name> <name><surname>Fan</surname> <given-names>A.</given-names></name> <name><surname>de la Clergerie</surname> <given-names>&#x000C9;.</given-names></name> <name><surname>Bordes</surname> <given-names>A.</given-names></name> <name><surname>Sagot</surname> <given-names>B.</given-names></name></person-group> (<year>2020b</year>). <article-title>Multilingual unsupervised sentence simplification</article-title>. <source>arXiv</source> [Preprint]. arXiv:2005.00352. <pub-id pub-id-type="doi">10.48550/arXiv.2005.00352</pub-id></citation>
</ref>
<ref id="B57">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Mehta</surname> <given-names>P.</given-names></name> <name><surname>Majumder</surname> <given-names>P.</given-names></name></person-group> (<year>2018</year>). <article-title>Effective aggregation of various summarization techniques</article-title>. <source>Inform. Process. Manag</source>. <volume>54</volume>, <fpage>145</fpage>&#x02013;<lpage>158</lpage>. <pub-id pub-id-type="doi">10.1016/j.ipm.2017.11.002</pub-id></citation>
</ref>
<ref id="B58">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Mihalcea</surname> <given-names>R.</given-names></name> <name><surname>Tarau</surname> <given-names>P.</given-names></name></person-group> (<year>2004</year>). <article-title>&#x0201C;TextRank: bringing order into text,&#x0201D;</article-title> in <source>Proceedings of the 2004 Conference on Empirical Methods in Natural Language Processing</source> (<publisher-loc>Barcelona</publisher-loc>: <publisher-name>Association for Computational Linguistics</publisher-name>), <fpage>404</fpage>&#x02013;<lpage>411</lpage>.</citation>
</ref>
<ref id="B59">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Nallapati</surname> <given-names>R.</given-names></name> <name><surname>Zhou</surname> <given-names>B.</given-names></name> <name><surname>dos Santos</surname> <given-names>C.</given-names></name> <name><surname>Gulcehre</surname> <given-names>C.</given-names></name> <name><surname>Xiang</surname> <given-names>B.</given-names></name></person-group> (<year>2016</year>). <article-title>&#x0201C;Abstractive text summarization using sequence-to-sequence RNNs and beyond,&#x0201D;</article-title> in <source>Proceedings of The 20th SIGNLL Conference on Computational Natural Language Learning</source> (<publisher-loc>Berlin</publisher-loc>: <publisher-name>Association for Computational Linguistics</publisher-name>), <fpage>280</fpage>&#x02013;<lpage>290</lpage>. <pub-id pub-id-type="doi">10.18653/v1/K16-1028</pub-id><pub-id pub-id-type="pmid">36568019</pub-id></citation></ref>
<ref id="B60">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Narayan</surname> <given-names>S.</given-names></name> <name><surname>Cohen</surname> <given-names>S. B.</given-names></name> <name><surname>Lapata</surname> <given-names>M.</given-names></name></person-group> (<year>2018</year>). <article-title>&#x0201C;Don&#x00027;t give me the details, just the summary! topic-aware convolutional neural networks for extreme summarization,&#x0201D;</article-title> in <source>Proceedings of the 2018 Conference on Empirical Methods in Natural Language Processing</source> (<publisher-loc>Brussels</publisher-loc>: <publisher-name>Association for Computational Linguistics</publisher-name>), <fpage>1797</fpage>&#x02013;<lpage>1807</lpage>. <pub-id pub-id-type="doi">10.18653/v1/D18-1206</pub-id><pub-id pub-id-type="pmid">36568019</pub-id></citation></ref>
<ref id="B61">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Narayan</surname> <given-names>S.</given-names></name> <name><surname>Gardent</surname> <given-names>C.</given-names></name> <name><surname>Cohen</surname> <given-names>S. B.</given-names></name> <name><surname>Shimorina</surname> <given-names>A.</given-names></name></person-group> (<year>2017</year>). <article-title>&#x0201C;Split and rephrase,&#x0201D;</article-title> in <source>Proceedings of the 2017 Conference on Empirical Methods in Natural Language Processing</source> (<publisher-loc>Copenhagen</publisher-loc>: <publisher-name>Association for Computational Linguistics</publisher-name>), <fpage>606</fpage>&#x02013;<lpage>616</lpage>. <pub-id pub-id-type="doi">10.18653/v1/D17-1064</pub-id><pub-id pub-id-type="pmid">36568019</pub-id></citation></ref>
<ref id="B62">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Nenkova</surname> <given-names>A.</given-names></name> <name><surname>Vanderwende</surname> <given-names>L.</given-names></name></person-group> (<year>2005</year>). <source>The impact of frequency on summarization</source>. <publisher-loc>Tech. Rep. MSR-TR-2005. Redmond, WA</publisher-loc>: <publisher-name>Microsoft Research, 101</publisher-name>.</citation>
</ref>
<ref id="B63">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Nishihara</surname> <given-names>D.</given-names></name> <name><surname>Kajiwara</surname> <given-names>T.</given-names></name> <name><surname>Arase</surname> <given-names>Y.</given-names></name></person-group> (<year>2019</year>). <article-title>&#x0201C;Controllable text simplification with lexical constraint loss,&#x0201D;</article-title> in <source>Proceedings of the 57th Annual Meeting of the Association for Computational Linguistics: Student Research Workshop</source> (<publisher-loc>Florence</publisher-loc>: <publisher-name>Association for Computational Linguistics</publisher-name>), <fpage>260</fpage>&#x02013;<lpage>266</lpage>. <pub-id pub-id-type="doi">10.18653/v1/P19-2036</pub-id><pub-id pub-id-type="pmid">36568019</pub-id></citation></ref>
<ref id="B64">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Nisioi</surname> <given-names>S.</given-names></name> <name><surname>&#x00160;tajner</surname> <given-names>S.</given-names></name> <name><surname>Ponzetto</surname> <given-names>S. P.</given-names></name> <name><surname>Dinu</surname> <given-names>L. P.</given-names></name></person-group> (<year>2017</year>). <article-title>&#x0201C;Exploring neural text simplification models,&#x0201D;</article-title> in <source>Proceedings of the 55th Annual Meeting of the Association for Computational Linguistics (Volume 2: Short Papers)</source> (<publisher-loc>Vancouver, BC</publisher-loc>: <publisher-name>Association for Computational Linguistics</publisher-name>), <fpage>85</fpage>&#x02013;<lpage>91</lpage>. <pub-id pub-id-type="doi">10.18653/v1/P17-2014</pub-id><pub-id pub-id-type="pmid">36083212</pub-id></citation></ref>
<ref id="B65">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>North</surname> <given-names>K.</given-names></name> <name><surname>Zampieri</surname> <given-names>M.</given-names></name> <name><surname>Shardlow</surname> <given-names>M.</given-names></name></person-group> (<year>2023</year>). <article-title>Lexical complexity prediction: an overview</article-title>. <source>ACM Comput. Surv</source>. <volume>55</volume>, <fpage>1</fpage>&#x02013;<lpage>42</lpage>. <pub-id pub-id-type="doi">10.1145/3557885</pub-id><pub-id pub-id-type="pmid">21489393</pub-id></citation></ref>
<ref id="B66">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Paetzold</surname> <given-names>G.</given-names></name> <name><surname>Specia</surname> <given-names>L.</given-names></name></person-group> (<year>2016</year>). <article-title>&#x0201C;SemEval 2016 task 11: complex word identification,&#x0201D;</article-title> in <source>Proceedings of the 10th International Workshop on Semantic Evaluation (SemEval-2016)</source> (<publisher-loc>San Diego, CA</publisher-loc>: <publisher-name>Association for Computational Linguistics</publisher-name>), <fpage>560</fpage>&#x02013;<lpage>569</lpage>. <pub-id pub-id-type="doi">10.18653/v1/S16-1085</pub-id><pub-id pub-id-type="pmid">36568019</pub-id></citation></ref>
<ref id="B67">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Papineni</surname> <given-names>K.</given-names></name> <name><surname>Roukos</surname> <given-names>S.</given-names></name> <name><surname>Ward</surname> <given-names>T.</given-names></name> <name><surname>Zhu</surname> <given-names>W.-J.</given-names></name></person-group> (<year>2002</year>). <article-title>&#x0201C;Bleu: a method for automatic evaluation of machine translation,&#x0201D;</article-title> in <source>Proceedings of the 40th Annual Meeting of the Association for Computational Linguistics</source> (<publisher-loc>Philadelphia, PA</publisher-loc>: <publisher-name>Association for Computational Linguistics</publisher-name>), <fpage>311</fpage>&#x02013;<lpage>318</lpage>. <pub-id pub-id-type="doi">10.3115/1073083.1073135</pub-id></citation>
</ref>
<ref id="B68">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Qi</surname> <given-names>W.</given-names></name> <name><surname>Yan</surname> <given-names>Y.</given-names></name> <name><surname>Gong</surname> <given-names>Y.</given-names></name> <name><surname>Liu</surname> <given-names>D.</given-names></name> <name><surname>Duan</surname> <given-names>N.</given-names></name> <name><surname>Chen</surname> <given-names>J.</given-names></name> <etal/></person-group>. (<year>2020</year>). <article-title>&#x0201C;ProphetNet: predicting future n-gram for sequence-to-SequencePre-training,&#x0201D;</article-title> in <source>Findings of the Association for Computational Linguistics: EMNLP 2020</source> (<publisher-loc>Association for Computational Linguistics</publisher-loc>), <fpage>2401</fpage>&#x02013;<lpage>2410</lpage>. <pub-id pub-id-type="doi">10.18653/v1/2020.findings-emnlp.217</pub-id><pub-id pub-id-type="pmid">36568019</pub-id></citation></ref>
<ref id="B69">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Rehman</surname> <given-names>T.</given-names></name> <name><surname>Mandal</surname> <given-names>R.</given-names></name> <name><surname>Agarwal</surname> <given-names>A.</given-names></name> <name><surname>Sanyal</surname> <given-names>D. K.</given-names></name></person-group> (<year>2023</year>). <article-title>Hallucination reduction in long input text summarization</article-title>. <source>arXiv</source> [Preprint]. arXiv:2309.16781. <pub-id pub-id-type="doi">10.48550/arXiv.2309.16781</pub-id></citation>
</ref>
<ref id="B70">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Sanchez-Gomez</surname> <given-names>J. M.</given-names></name> <name><surname>Vega-Rodr&#x000ED;guez</surname> <given-names>M. A.</given-names></name> <name><surname>Perez</surname> <given-names>C. J.</given-names></name></person-group> (<year>2020</year>). <article-title>A decomposition-based multi-objective optimization approach for extractive multi-document text summarization</article-title>. <source>Appl. Soft Comput</source>. <volume>91</volume>:<fpage>106231</fpage>. <pub-id pub-id-type="doi">10.1016/j.asoc.2020.106231</pub-id></citation>
</ref>
<ref id="B71">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Sanchez-Gomez</surname> <given-names>J. M.</given-names></name> <name><surname>Vega-Rodr&#x000ED;guez</surname> <given-names>M. A.</given-names></name> <name><surname>P&#x000E9;rez</surname> <given-names>C. J.</given-names></name></person-group> (<year>2021</year>). <article-title>Sentiment-oriented query-focused text summarization addressed with a multi-objective optimization approach</article-title>. <source>Appl. Soft Comput</source>. <volume>113</volume>:<fpage>107915</fpage>. <pub-id pub-id-type="doi">10.1016/j.asoc.2021.107915</pub-id></citation>
</ref>
<ref id="B72">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Sawilowsky</surname> <given-names>S.</given-names></name> <name><surname>Fahoome</surname> <given-names>G.</given-names></name></person-group> (<year>2005</year>). <article-title>Friedman&#x00027;s test</article-title>. <source>Encycl. Stat. Behav. Sci</source>. <pub-id pub-id-type="doi">10.1002/0470013192.bsa385</pub-id></citation>
</ref>
<ref id="B73">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Schluter</surname> <given-names>N.</given-names></name></person-group> (<year>2017</year>). <article-title>&#x0201C;The limits of automatic summarisation according to ROUGE,&#x0201D;</article-title> in <source>Proceedings of the 15th Conference of the European Chapter of the Association for Computational Linguistics: Volume 2, Short Papers</source> (<publisher-loc>Valencia</publisher-loc>: <publisher-name>Association for Computational Linguistics</publisher-name>), <fpage>41</fpage>&#x02013;<lpage>45</lpage>. <pub-id pub-id-type="doi">10.18653/v1/E17-2007</pub-id><pub-id pub-id-type="pmid">36568019</pub-id></citation></ref>
<ref id="B74">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>See</surname> <given-names>A.</given-names></name> <name><surname>Liu</surname> <given-names>P. J.</given-names></name> <name><surname>Manning</surname> <given-names>C. D.</given-names></name></person-group> (<year>2017</year>). <article-title>&#x0201C;Get to the point: summarization with pointer-generator networks,&#x0201D;</article-title> in <source>Proceedings of the 55th Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers)</source> (<publisher-loc>Vancouver</publisher-loc>: <publisher-name>Association for Computational Linguistics</publisher-name>), <fpage>1073</fpage>&#x02013;<lpage>1083</lpage>. <pub-id pub-id-type="doi">10.18653/v1/P17-1099</pub-id><pub-id pub-id-type="pmid">36568019</pub-id></citation></ref>
<ref id="B75">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Shardlow</surname> <given-names>M.</given-names></name></person-group> (<year>2014a</year>). <article-title>&#x0201C;Out in the open: finding and categorising errors in the lexical simplification pipeline,&#x0201D;</article-title> in <source>Proceedings of the Ninth International Conference on Language Resources and Evaluation (LREC-2014)</source> (<publisher-loc>Reykjavik</publisher-loc>: <publisher-name>European Languages Resources Association</publisher-name>), <fpage>1583</fpage>&#x02013;<lpage>1590</lpage>.</citation>
</ref>
<ref id="B76">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Shardlow</surname> <given-names>M.</given-names></name></person-group> (<year>2014b</year>). <article-title>A survey of automated text simplification</article-title>. <source>Int. J. Adv. Comput. Sci. Appl</source>. <volume>4</volume>, <fpage>58</fpage>&#x02013;<lpage>70</lpage>. <pub-id pub-id-type="doi">10.14569/SpecialIssue.2014.040109</pub-id><pub-id pub-id-type="pmid">36083212</pub-id></citation></ref>
<ref id="B77">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Shardlow</surname> <given-names>M.</given-names></name> <name><surname>Batista-Navarro</surname> <given-names>R.</given-names></name> <name><surname>Thompson</surname> <given-names>P.</given-names></name> <name><surname>Nawaz</surname> <given-names>R.</given-names></name> <name><surname>McNaught</surname> <given-names>J.</given-names></name> <name><surname>Ananiadou</surname> <given-names>S.</given-names></name> <etal/></person-group>. (<year>2018</year>). <article-title>Identification of research hypotheses and new knowledge from scientific literature</article-title>. <source>BMC Med. Inform. Decis. Mak</source>. <volume>18</volume>:<fpage>46</fpage>. <pub-id pub-id-type="doi">10.1186/s12911-018-0639-1</pub-id><pub-id pub-id-type="pmid">29940927</pub-id></citation></ref>
<ref id="B78">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Shardlow</surname> <given-names>M.</given-names></name> <name><surname>Evans</surname> <given-names>R.</given-names></name> <name><surname>Zampieri</surname> <given-names>M.</given-names></name></person-group> (<year>2022</year>). <article-title>Predicting lexical complexity in english texts: the complex 2.0 dataset</article-title>. <source>Lang. Resour. Eval</source>. <volume>56</volume>, <fpage>1153</fpage>&#x02013;<lpage>1194</lpage>. <pub-id pub-id-type="doi">10.1007/s10579-022-09588-2</pub-id></citation>
</ref>
<ref id="B79">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Siddharthan</surname> <given-names>A.</given-names></name></person-group> (<year>2014</year>). <article-title>A survey of research on text simplification</article-title>. <source>ITL-Int. J. Appl. Linguist</source>. <volume>165</volume>, <fpage>259</fpage>&#x02013;<lpage>298</lpage>. <pub-id pub-id-type="doi">10.1075/itl.165.2.06sid</pub-id><pub-id pub-id-type="pmid">33486653</pub-id></citation></ref>
<ref id="B80">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Suleiman</surname> <given-names>D.</given-names></name> <name><surname>Awajan</surname> <given-names>A.</given-names></name></person-group> (<year>2022</year>). <article-title>Multilayer encoder and single-layer decoder for abstractive arabic text summarization</article-title>. <source>Knowl.-Based Syst</source>. <volume>237</volume>:<fpage>107791</fpage>. <pub-id pub-id-type="doi">10.1016/j.knosys.2021.107791</pub-id></citation>
</ref>
<ref id="B81">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Sulem</surname> <given-names>E.</given-names></name> <name><surname>Abend</surname> <given-names>O.</given-names></name> <name><surname>Rappoport</surname> <given-names>A.</given-names></name></person-group> (<year>2018</year>). <article-title>&#x0201C;Semantic structural evaluation for text simplification,&#x0201D;</article-title> in <source>Proceedings of the 2018 Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Technologies, Volume 1 (Long Papers)</source> (<publisher-loc>New Orleans, LA</publisher-loc>: <publisher-name>Association for Computational Linguistics</publisher-name>), <fpage>685</fpage>&#x02013;<lpage>696</lpage>. <pub-id pub-id-type="doi">10.18653/v1/N18-1063</pub-id><pub-id pub-id-type="pmid">36568019</pub-id></citation></ref>
<ref id="B82">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Thompson</surname> <given-names>P.</given-names></name> <name><surname>Nawaz</surname> <given-names>R.</given-names></name> <name><surname>McNaught</surname> <given-names>J.</given-names></name> <name><surname>Ananiadou</surname> <given-names>S.</given-names></name></person-group> (<year>2017</year>). <article-title>Enriching news events with meta-knowledge information</article-title>. <source>Lang. Resour. Eval</source>. <volume>51</volume>, <fpage>409</fpage>&#x02013;<lpage>438</lpage>. <pub-id pub-id-type="doi">10.1007/s10579-016-9344-9</pub-id></citation>
</ref>
<ref id="B83">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Tomer</surname> <given-names>M.</given-names></name> <name><surname>Kumar</surname> <given-names>M.</given-names></name></person-group> (<year>2022</year>). <article-title>&#x0201C;STV-BEATS: skip thought vector and bi-encoder based automatic text summarizer</article-title>. <source>Knowl.-Based Syst</source>. <volume>240</volume>:<fpage>108108</fpage>. <pub-id pub-id-type="doi">10.1016/j.knosys.2021.108108</pub-id></citation>
</ref>
<ref id="B84">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Van Veen</surname> <given-names>D.</given-names></name> <name><surname>Van Uden</surname> <given-names>C.</given-names></name> <name><surname>Blankemeier</surname> <given-names>L.</given-names></name> <name><surname>Delbrouck</surname> <given-names>J.-B.</given-names></name> <name><surname>Aali</surname> <given-names>A.</given-names></name> <name><surname>Bluethgen</surname> <given-names>C.</given-names></name> <etal/></person-group>. (<year>2023</year>). <article-title>Clinical text summarization: adapting large language models can outperform human experts</article-title>. <source>arXiv</source> [Preprint]. arXiv:2309,07430. <pub-id pub-id-type="doi">10.48550/arXiv.2309.07430</pub-id></citation>
</ref>
<ref id="B85">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Vaswani</surname> <given-names>A.</given-names></name> <name><surname>Shazeer</surname> <given-names>N.</given-names></name> <name><surname>Parmar</surname> <given-names>N.</given-names></name> <name><surname>Uszkoreit</surname> <given-names>J.</given-names></name> <name><surname>Jones</surname> <given-names>L.</given-names></name> <name><surname>Gomez</surname> <given-names>A. N.</given-names></name> <etal/></person-group>. (<year>2017</year>). <article-title>&#x0201C;Attention is all you need,&#x0201D;</article-title> in <source>Advances in neural information processing systems</source> (<publisher-loc>Long Beach, CA</publisher-loc>), <fpage>5998</fpage>&#x02013;<lpage>6008</lpage>.</citation>
</ref>
<ref id="B86">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Vinyals</surname> <given-names>O.</given-names></name> <name><surname>Fortunato</surname> <given-names>M.</given-names></name> <name><surname>Jaitly</surname> <given-names>N.</given-names></name></person-group> (<year>2015</year>). <article-title>&#x0201C;Pointer networks,&#x0201D;</article-title> in <source>Advances in Neural Information Processing Systems, pages</source> (<publisher-loc>Montreal, QC</publisher-loc>), <fpage>2692</fpage>&#x02013;<lpage>2700</lpage>.</citation>
</ref>
<ref id="B87">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Wang</surname> <given-names>L.</given-names></name> <name><surname>Jiang</surname> <given-names>J.</given-names></name> <name><surname>Chieu</surname> <given-names>H. L.</given-names></name> <name><surname>Ong</surname> <given-names>C. H.</given-names></name> <name><surname>Song</surname> <given-names>D.</given-names></name> <name><surname>Liao</surname> <given-names>L.</given-names></name> <etal/></person-group>. (<year>2017</year>). <article-title>&#x0201C;Can syntax help? improving an LSTM-based sentence compression model for new domains,&#x0201D;</article-title> in <source>Proceedings of the 55th Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers)</source> (<publisher-loc>Vancouver, BC</publisher-loc>: <publisher-name>Association for Computational Linguistics</publisher-name>), <fpage>1385</fpage>&#x02013;<lpage>1393</lpage>. <pub-id pub-id-type="doi">10.18653/v1/P17-1127</pub-id><pub-id pub-id-type="pmid">36568019</pub-id></citation></ref>
<ref id="B88">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Wubben</surname> <given-names>S.</given-names></name> <name><surname>van den Bosch</surname> <given-names>A.</given-names></name> <name><surname>Krahmer</surname> <given-names>E.</given-names></name></person-group> (<year>2012</year>). <article-title>&#x0201C;Sentence simplification by monolingual machine translation,&#x0201D;</article-title> in <source>Proceedings of the 50th Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers)</source> (<publisher-loc>Jeju Island</publisher-loc>: <publisher-name>Association for Computational Linguistics</publisher-name>), <fpage>1015</fpage>&#x02013;<lpage>1024</lpage>.</citation>
</ref>
<ref id="B89">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Xu</surname> <given-names>W.</given-names></name> <name><surname>Napoles</surname> <given-names>C.</given-names></name> <name><surname>Pavlick</surname> <given-names>E.</given-names></name> <name><surname>Chen</surname> <given-names>Q.</given-names></name> <name><surname>Callison-Burch</surname> <given-names>C.</given-names></name></person-group> (<year>2016</year>). <article-title>Optimizing statistical machine translation for text simplification</article-title>. <source>Trans. Assoc. Comput. Linguist</source>. <volume>4</volume>, <fpage>401</fpage>&#x02013;<lpage>415</lpage>. <pub-id pub-id-type="doi">10.1162/tacl_a_00107</pub-id></citation>
</ref>
<ref id="B90">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Yan</surname> <given-names>Y.</given-names></name> <name><surname>Qi</surname> <given-names>W.</given-names></name> <name><surname>Gong</surname> <given-names>Y.</given-names></name> <name><surname>Liu</surname> <given-names>D.</given-names></name> <name><surname>Duan</surname> <given-names>N.</given-names></name> <name><surname>Chen</surname> <given-names>J.</given-names></name> <etal/></person-group>. (<year>2020</year>). <article-title>Prophetnet: predicting future n-gram for sequence-to-sequence pre-training</article-title>. <source>arXiv</source> [Preprint] arXiv:2001.04063. <pub-id pub-id-type="doi">10.48550/arXiv.2001.04063</pub-id></citation>
</ref>
<ref id="B91">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Yang</surname> <given-names>R.</given-names></name> <name><surname>Zeng</surname> <given-names>Q.</given-names></name> <name><surname>You</surname> <given-names>K.</given-names></name> <name><surname>Qiao</surname> <given-names>Y.</given-names></name> <name><surname>Huang</surname> <given-names>L.</given-names></name> <name><surname>Hsieh</surname> <given-names>C.-C.</given-names></name> <etal/></person-group>. (<year>2023</year>). <article-title>Medgen: a python natural language processing toolkit for medical text processing</article-title>. <source>arXiv</source> [Preprint]. arXiv:2311.16588. <pub-id pub-id-type="doi">10.48550/arXiv.2311.16588</pub-id></citation>
</ref>
<ref id="B92">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zaman</surname> <given-names>F.</given-names></name> <name><surname>Shardlow</surname> <given-names>M.</given-names></name> <name><surname>Hassan</surname> <given-names>S.-U.</given-names></name> <name><surname>Aljohani</surname> <given-names>N. R.</given-names></name> <name><surname>Nawaz</surname> <given-names>R.</given-names></name></person-group> (<year>2020</year>). <article-title>HTSS: a novel hybrid text summarisation and simplification architecture</article-title>. <source>Inform. Process. Manag</source>. <volume>57</volume>:<fpage>102351</fpage>. <pub-id pub-id-type="doi">10.1016/j.ipm.2020.102351</pub-id></citation>
</ref>
<ref id="B93">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zerva</surname> <given-names>C.</given-names></name> <name><surname>Nghiem</surname> <given-names>M.-Q.</given-names></name> <name><surname>Nguyen</surname> <given-names>N. T.</given-names></name> <name><surname>Ananiadou</surname> <given-names>S.</given-names></name></person-group> (<year>2020</year>). <article-title>Cited text span identification for scientific summarisation using pre-trained encoders</article-title>. <source>Scientometrics</source> <volume>125</volume>, <fpage>3109</fpage>&#x02013;<lpage>3137</lpage>. <pub-id pub-id-type="doi">10.1007/s11192-020-03455-z</pub-id></citation>
</ref>
<ref id="B94">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zhang</surname> <given-names>J.</given-names></name> <name><surname>Zhao</surname> <given-names>Y.</given-names></name> <name><surname>Saleh</surname> <given-names>M.</given-names></name> <name><surname>Liu</surname> <given-names>P. J.</given-names></name></person-group> (<year>2019a</year>). <article-title>Pegasus: pre-training with extracted gap-sentences for abstractive summarization</article-title>. <source>arXiv</source> [Preprint]. arXiv:1912.08777. <pub-id pub-id-type="doi">10.48550/arXiv.1912.08777</pub-id></citation>
</ref>
<ref id="B95">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Zhang</surname> <given-names>T.</given-names></name> <name><surname>Kishore</surname> <given-names>V.</given-names></name> <name><surname>Wu</surname> <given-names>F.</given-names></name> <name><surname>Weinberger</surname> <given-names>K. Q.</given-names></name> <name><surname>Artzi</surname> <given-names>Y.</given-names></name></person-group> (<year>2019b</year>). <article-title>&#x0201C;Bertscore: evaluating text generation with Bert,&#x0201D;</article-title> in <source>International Conference on Learning Representations</source> (<publisher-loc>New Orleans, LA</publisher-loc>).</citation>
</ref>
<ref id="B96">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Zhang</surname> <given-names>X.</given-names></name> <name><surname>Lapata</surname> <given-names>M.</given-names></name></person-group> (<year>2017</year>). <article-title>&#x0201C;Sentence simplification with deep reinforcement learning,&#x0201D;</article-title> in <source>Proceedings of the 2017 Conference on Empirical Methods in Natural Language Processing</source> (<publisher-loc>Copenhagen</publisher-loc>), <fpage>584</fpage>&#x02013;<lpage>594</lpage>. <pub-id pub-id-type="doi">10.18653/v1/D17-1062</pub-id><pub-id pub-id-type="pmid">36568019</pub-id></citation></ref>
<ref id="B97">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Zhao</surname> <given-names>Y.</given-names></name> <name><surname>Luo</surname> <given-names>Z.</given-names></name> <name><surname>Aizawa</surname> <given-names>A.</given-names></name></person-group> (<year>2018</year>). <article-title>&#x0201C;A language model based evaluator for sentence compression,&#x0201D;</article-title> in <source>Proceedings of the 56th Annual Meeting of the Association for Computational Linguistics (Volume 2: Short Papers)</source> (<publisher-loc>Melbourne, VIC</publisher-loc>: <publisher-name>Association for Computational Linguistics</publisher-name>), <fpage>170</fpage>&#x02013;<lpage>175</lpage>. <pub-id pub-id-type="doi">10.18653/v1/P18-2028</pub-id><pub-id pub-id-type="pmid">36568019</pub-id></citation></ref>
<ref id="B98">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Zhu</surname> <given-names>Z.</given-names></name> <name><surname>Bernhard</surname> <given-names>D.</given-names></name> <name><surname>Gurevych</surname> <given-names>I.</given-names></name></person-group> (<year>2010</year>). <article-title>&#x0201C;A monolingual tree-based translation model for sentence simplification,&#x0201D;</article-title> in <source>Proceedings of the 23rd International Conference on Computational Linguistics (Coling 2010)</source> (<publisher-loc>Beijing</publisher-loc>: <publisher-name>Coling 2010 Organizing Committee</publisher-name>), <fpage>1353</fpage>&#x02013;<lpage>1361</lpage>.</citation>
</ref>
<ref id="B99">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Zopf</surname> <given-names>M.</given-names></name> <name><surname>Loza Menc&#x000ED;a</surname> <given-names>E.</given-names></name> <name><surname>F&#x000FC;rnkranz</surname> <given-names>J.</given-names></name></person-group> (<year>2018</year>). <article-title>&#x0201C;Which scores to predict in sentence regression for text summarization?&#x0201D;</article-title> in <source>Proceedings of the 2018 Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Technologies, Volume 1 (Long Papers)</source> (<publisher-loc>New Orleans, LA</publisher-loc>: <publisher-name>Association for Computational Linguistics</publisher-name>), <fpage>1782</fpage>&#x02013;<lpage>1791</lpage>. <pub-id pub-id-type="doi">10.18653/v1/N18-1161</pub-id><pub-id pub-id-type="pmid">17910536</pub-id></citation></ref>
</ref-list>
</back>
</article>