<?xml version="1.0" encoding="UTF-8" standalone="no"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance" article-type="research-article" dtd-version="2.3" xml:lang="EN">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Mar. Sci.</journal-id>
<journal-title>Frontiers in Marine Science</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Mar. Sci.</abbrev-journal-title>
<issn pub-type="epub">2296-7745</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="doi">10.3389/fmars.2024.1368356</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Marine Science</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>ChatBBNJ: a question&#x2013;answering system for acquiring knowledge on biodiversity beyond national jurisdiction</article-title>
</title-group>
<contrib-group>
<contrib contrib-type="author">
<name>
<surname>Wang</surname>
<given-names>Xiaowei</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1851303"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Zhang</surname>
<given-names>Mingdan</given-names>
</name>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Liu</surname>
<given-names>Hao</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1439934"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/resources/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Ma</surname>
<given-names>Xiaodong</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1688074"/>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Liu</surname>
<given-names>Yingchao</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2627489"/>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Chen</surname>
<given-names>Yitong</given-names>
</name>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="author-notes" rid="fn001">
<sup>*</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1736141"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/funding-acquisition/"/>
<role content-type="https://credit.niso.org/contributor-roles/resources/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
</contrib-group>
<aff id="aff1">
<sup>1</sup>
<institution>College of Computer Science and Technology, Ocean University of China</institution>, <addr-line>Qingdao</addr-line>, <country>China</country>
</aff>
<aff id="aff2">
<sup>2</sup>
<institution>Law School, Ocean University of China</institution>, <addr-line>Qingdao</addr-line>, <country>China</country>
</aff>
<author-notes>
<fn fn-type="edited-by">
<p>Edited by: Yen-Chiang Chang, Dalian Maritime University, China</p>
</fn>
<fn fn-type="edited-by">
<p>Reviewed by: Nadia Khadam, Fatima Jinnah Women University, Pakistan</p>
<p>Wen Duan, Hainan University, China</p>
</fn>
<fn fn-type="corresp" id="fn001">
<p>*Correspondence: Yitong Chen, <email xlink:href="mailto:chenyitong@ouc.edu.cn">chenyitong@ouc.edu.cn</email>
</p>
</fn>
</author-notes>
<pub-date pub-type="epub">
<day>08</day>
<month>04</month>
<year>2024</year>
</pub-date>
<pub-date pub-type="collection">
<year>2024</year>
</pub-date>
<volume>11</volume>
<elocation-id>1368356</elocation-id>
<history>
<date date-type="received">
<day>10</day>
<month>01</month>
<year>2024</year>
</date>
<date date-type="accepted">
<day>26</day>
<month>03</month>
<year>2024</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#xa9; 2024 Wang, Zhang, Liu, Ma, Liu and Chen</copyright-statement>
<copyright-year>2024</copyright-year>
<copyright-holder>Wang, Zhang, Liu, Ma, Liu and Chen</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<abstract>
<p>The marine biodiversity in Areas beyond national jurisdiction (ABNJ), encompassing approximately two-thirds of the global ocean, is persistently declining. In 2023, the agreement on the Conservation and Sustainable Use of Marine Biodiversity of Areas Beyond National Jurisdiction (BBNJ) was officially adopted. Implementing the BBNJ Agreement has the potential to effectively meet global needs for preserving marine biodiversity. Nevertheless, the implementation requires dealing with thousands of legal clauses, and the parties participating in the process lack adequate means to acquire knowledge connected to BBNJ. This paper introduces ChatBBNJ, a highly efficient question-answering system that combines a novel data engineering technique with large language models (LLMs) of Natural Language Processing (NLP). The system aims to efficiently provide stakeholders with BBNJ-related knowledge, thereby facilitating and enhancing their comprehension and involvement with the subject matter. The experimental results demonstrate that the proposed ChatBBNJ exhibits superior expertise in the BBNJ domain, outperforming baseline models in terms of precision, recall, and F1-scores. The successful deployment of the suggested system is expected to greatly assist stakeholders in acquiring BBNJ knowledge and facilitating the effective implementation of the BBNJ Agreement. Therefore, this is expected to contribute to the conservation and sustainable use of marine biodiversity in ABNJ.</p>
</abstract>
<kwd-group>
<kwd>BBNJ agreement</kwd>
<kwd>ABNJ</kwd>
<kwd>LLMS</kwd>
<kwd>nlp</kwd>
<kwd>intelligent question-answering</kwd>
</kwd-group>
<counts>
<fig-count count="9"/>
<table-count count="6"/>
<equation-count count="3"/>
<ref-count count="47"/>
<page-count count="14"/>
<word-count count="6698"/>
</counts>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-in-acceptance</meta-name>
<meta-value>Marine Affairs and Policy</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec id="s1" sec-type="intro">
<label>1</label>
<title>Introduction</title>
<p>Areas beyond national jurisdiction (ABNJ) face persistent degradation of marine biodiversity (<xref ref-type="bibr" rid="B16">Humphries and Harden-Davies, 2020</xref>). A United Nations agreement on the conservation and sustainable use of marine biodiversity in areas beyond national jurisdiction (the BBNJ Agreement) was formally adopted in August 2023 (<xref ref-type="bibr" rid="B34">United Nations, 2023</xref>). The BBNJ Agreement regulates four key elements concerning ocean governance: marine genetic resources; area-based management tools, including marine protected areas; environmental impact assessments; and capacity building and marine technology transfer (<xref ref-type="bibr" rid="B30">Tessnow-von Wysocki and Vadrot, 2020</xref>). It will act as a governance mechanism to achieve conservation and sustainable use of marine biodiversity in ABNJ (<xref ref-type="bibr" rid="B31">Tiller et&#xa0;al., 2023</xref>).</p>
<p>However, the implementation of the treaty is facing a lot of resistance. First, as a package deal, the BBNJ Agreement provides a framework for reaching consensus; it still lacks specifics, which needs to be discussed in future Conference of Parties (COP) (<xref ref-type="bibr" rid="B9">Deasy, 2023</xref>). Second, learning from the implementation of past international law (<xref ref-type="bibr" rid="B3">Bodansky, 2011</xref>), industry may seek to weaken implementation measures in order to reduce its adjustment costs. To solve these problems, countries, organizations, and other stakeholders need to submit implementation reports and discuss these issues in the COP. However, many of the stakeholders were not involved in the BBNJ negotiations. They needed to rapidly comprehend thousands of clauses and four interdisciplinary key knowledge points before the discussion. This presents a significant challenge to the stakeholders.</p>
<p>Even though stakeholders can acquire BBNJ-related knowledge through search engines, existing search engines provide information retrieval based on keywords within relevant documents. Users have to deal with the burden of browsing and filtering out results to find the candidate passages (<xref ref-type="bibr" rid="B19">Lau et&#xa0;al., 2005</xref>). Thus, providing a convenient information acquisition system is necessary, which is already called for by the BBNJ Agreement (<xref ref-type="bibr" rid="B34">United Nations, 2023</xref>). Knowledge question-answering (Q-A) systems take in natural language questions and provide accurate answers, which can reduce the burden of users reading a large number of irrelevant documents to obtain answers (<xref ref-type="bibr" rid="B47">Zhu et&#xa0;al., 2021</xref>). At present, it has been widely used in many fields (<xref ref-type="bibr" rid="B46">Zhong et&#xa0;al., 2020</xref>; <xref ref-type="bibr" rid="B8">Dai et&#xa0;al., 2022</xref>; <xref ref-type="bibr" rid="B20">Lee et&#xa0;al., 2023</xref>), but few Q-A systems popularize professional knowledge of international law of the sea. Therefore, this study aims to develop a Q-A system and to efficiently provide stakeholders with BBNJ-related knowledge, thereby facilitating and enhancing their comprehension and involvement with the subject matter, prompting the effective implementation of the BBNJ Agreement.</p>
<p>Researchers can use various methods to develop a Q-A system for BBNJ knowledge popularization. Early Q-A systems heavily relied on rule-based methods (<xref ref-type="bibr" rid="B26">Riloff and Thelen, 2000</xref>), in which linguistic experts needed to manually formulate rules based on the characteristics of BBNJ texts. This method lacks generalization and makes it difficult to cover all scenarios. With rapid advancements in artificial intelligence technologies, several Q-A models have been developed applying statistical language models (<xref ref-type="bibr" rid="B27">Rosenfeld, 2000</xref>). These models do not require manual rules, they automatically learn statistical language patterns to predict the correct answers to questions. However, this method cannot obtain semantic information. The BBNJ Agreement contains many repetitive terminologies, which makes it challenging for most statistical language models to differentiate subtle situations. Neural language models (<xref ref-type="bibr" rid="B2">Bengio et&#xa0;al., 2003</xref>) characterize the probability of word sequences generated by neural networks. These models generate a contextual representation that encodes semantic and syntactic information. However, in addition to repetition, the language of the BBNJ Agreement shows long-term dependency. Limited by the context window, these models cannot model global semantics. Pretrained language models (<xref ref-type="bibr" rid="B10">Devlin et&#xa0;al., 2019</xref>) capture context-aware word representations by fine-tuning the networks according to specific downstream tasks. Q-A systems built upon these models exhibit better performance.</p>
<p>Recently, researchers find that scaling pretrained language models often leads to improved model capacity on downstream tasks (i.e., following the scaling law (<xref ref-type="bibr" rid="B18">Kaplan et&#xa0;al., 2020</xref>)). With the significant success of large language models (LLMs) like ChatGPT (<xref ref-type="bibr" rid="B25">Ouyang et&#xa0;al., 2022</xref>) in tasks related to understanding and generating human-like responses (<xref ref-type="bibr" rid="B12">Eloundou et&#xa0;al., 2023</xref>), applying LLMs to Q-A systems has become a popular choice among researchers (<xref ref-type="bibr" rid="B7">Cui et&#xa0;al., 2023</xref>; <xref ref-type="bibr" rid="B15">Huang et&#xa0;al., 2023</xref>; <xref ref-type="bibr" rid="B21">Li et&#xa0;al., 2023</xref>; <xref ref-type="bibr" rid="B35">Vaghefi et&#xa0;al., 2023</xref>; <xref ref-type="bibr" rid="B37">Wang et&#xa0;al., 2023</xref>; <xref ref-type="bibr" rid="B42">Xiong et&#xa0;al., 2023</xref>; <xref ref-type="bibr" rid="B43">Yang et&#xa0;al., 2023</xref>). At present, LLMs have significantly improved their performance in open-domain Q-A (<xref ref-type="bibr" rid="B22">Li et&#xa0;al., 2023</xref>). However, when applying LLMs to the BBNJ domain, it is difficult to fully utilize its advantages (<xref ref-type="bibr" rid="B38">Wang et&#xa0;al., 2023</xref>). LLMs are trained on general corpora, such as Common Crawl (<xref ref-type="bibr" rid="B6">Common Crawl, 2023</xref>) and Wikipedia (<xref ref-type="bibr" rid="B41">Wikipedia, 2023</xref>). BBNJ-related knowledge is complex and multidisciplinary, involving legal, scientific, and international relations considerations (<xref ref-type="bibr" rid="B17">Humphries et&#xa0;al., 2021</xref>). These specialized areas differ significantly from the pre-trained data of LLMs. Therefore, efficiently adapting LLMs to the BBNJ domain, fully utilizing LLMs&#x2019; understanding abilities remains a challenging problem.</p>
<p>Recent studies have shown (<xref ref-type="bibr" rid="B1">Amer-Yahia et&#xa0;al., 2023</xref>) that fine-tuning LLMs by using high-quality domain-specific data can improve their domain adaptation ability. However, developing a Q-A system for BBNJ knowledge popularization by fine-tuning LLMs with domain-specific data presents challenges. First, there is a lack of available Q-A datasets for fine-tuning LLMs in the BBNJ domain. Second, LLMs suffer from outdated information after fine-tuning has concluded. Ensuring stakeholders rapidly acquire the latest implementation situations and recommendations through the Q-A system is urgently needed in practice. Hence, providing accurate and up-to-date responses is paramount. Such accurate responses can help stakeholders understand the complex and dynamically updated BBNJ knowledge and prompt the implementation of the BBNJ Agreement. Thus, this study (<xref ref-type="fig" rid="f1">
<bold>Figure&#xa0;1</bold>
</xref>) aims to solve these problems with the following contributions:</p>
<fig id="f1" position="float">
<label>Figure&#xa0;1</label>
<caption>
<p>Overview of ChatBBNJ<bold>&#x2019;</bold>s work.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fmars-11-1368356-g001.tif"/>
</fig>
<p>(1) We proposed a data engineering method, called PDGC, to generate a higher-quality Q-A dataset for the BBNJ domain. PDGC contains two-stage data-generation and iterative correction. The two-stage data-generation method enables the model to generate higher-quality data based on BBNJ Q-A examples. Moreover, our iterative correction improves correction quality by following the human learning pattern based on easy-to-difficult.</p>
<p>(2) We developed a BBNJ domain language model called ChatBBNJ, which is fine-tuned by utilizing the United Nations Convention on the Law of the Sea (UNCLOS) and its annexes. BBNJ Agreement is developed under the UNCLOS. Furthermore, we introduced a domain knowledge-based prompt engineering. The domain knowledge base is constructed by utilizing the BBNJ Agreement. The agreement offers the latest regulations for the management of BBNJ. It ensures that ChatBBNJ obtains the latest BBNJ domain information quickly.</p>
</sec>
<sec id="s2">
<label>2</label>
<title>Methods</title>
<p>An overview of the proposed method framework for BBNJ-related knowledge popularization is shown in <xref ref-type="fig" rid="f2">
<bold>Figure&#xa0;2</bold>
</xref>. The method framework is composed of three parts: dataset construction, model fine-tuning, and domain knowledge-based prompt engineering. At the dataset construction stage, we applied PDGC to construct a high-quality dataset for fine-tuning. Then, at the model fine-tuning stage, we applied the LoRA (<xref ref-type="bibr" rid="B14">Hu et&#xa0;al., 2021</xref>) method to fine-tune the base model, which improved the domain adaptation ability of the base model. Finally, at the domain knowledge-based prompt engineering stage, a domain knowledge base was used when constructing prompts, which ensured the timeliness of the model&#x2019;s answers.</p>
<fig id="f2" position="float">
<label>Figure&#xa0;2</label>
<caption>
<p>Overview of biodiversity of areas beyond national jurisdiction knowledge question&#x2013;answering.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fmars-11-1368356-g002.tif"/>
</fig>
<sec id="s2_1">
<label>2.1</label>
<title>Q-A dataset construction</title>
<p>The proposed PDGC method comprises three modules: text preprocessing, data generation, and data correction (<xref ref-type="fig" rid="f3">
<bold>Figure&#xa0;3</bold>
</xref>). During the text preprocessing, we performed reference completion as the BBNJ-related documents have many reference expressions. Then, during the data generation, since the BBNJ-related documents are multidisciplinary, involving legal, scientific, and international relations considerations (<xref ref-type="bibr" rid="B17">Humphries et&#xa0;al., 2021</xref>), single data generation will result in a significant amount of noisy data. To solve this problem, a two-stage data generation method was applied. Finally, during the data correction, existing consistency validation methods reduce data quantity and diversity. Therefore, we proposed a similarity-based data division and iterative correction method.</p>
<fig id="f3" position="float">
<label>Figure&#xa0;3</label>
<caption>
<p>Dataset construction method framework. &#x201c;UNCLOS&#x201d; represents the United Nations Convention on the Law of the Sea. &#x201c;BBNJ Draft&#x201d; represents the draft agreement under the United Nations Convention on the Law of the Sea on the conservation and sustainable use of marine biological diversity of areas beyond national jurisdiction. &#x201c;BBNJ Agreement&#x201d; represents the agreement under the United Nations Convention on the Law of the Sea on the conservation and sustainable use of marine biological diversity of areas beyond national jurisdiction. &#x201c;Vicuna-13B&#x201d; is a large language model used to generate data. &#x201c; <inline-formula>
<mml:math display="inline" id="im1">
<mml:mrow>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mstyle mathvariant="bold" mathsize="normal">
<mml:mi>t</mml:mi>
</mml:mstyle>
<mml:mo>,</mml:mo>
<mml:mstyle mathvariant="bold" mathsize="normal">
<mml:mi>q</mml:mi>
</mml:mstyle>
<mml:mo>,</mml:mo>
<mml:mstyle mathvariant="bold" mathsize="normal">
<mml:mi>a</mml:mi>
</mml:mstyle>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula> &#x201c; contains three parts: &#x201c; <inline-formula>
<mml:math display="inline" id="im2">
<mml:mstyle mathvariant="bold" mathsize="normal">
<mml:mi>t</mml:mi>
</mml:mstyle>
</mml:math>
</inline-formula>&#x201c; is the text of clause obtained in text preprocessing, &#x201c; <inline-formula>
<mml:math display="inline" id="im3">
<mml:mstyle mathvariant="bold" mathsize="normal">
<mml:mi>q</mml:mi>
</mml:mstyle>
</mml:math>
</inline-formula>&#x201c; and &#x201c; <inline-formula>
<mml:math display="inline" id="im4">
<mml:mstyle mathvariant="bold" mathsize="normal">
<mml:mi>a</mml:mi>
</mml:mstyle>
</mml:math>
</inline-formula>&#x201c; represent the question and answer generated by Vicuna-13B in data generation respectively. &#x201c; <inline-formula>
<mml:math display="inline" id="im5">
<mml:mrow>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mstyle mathvariant="bold" mathsize="normal">
<mml:mi>t</mml:mi>
</mml:mstyle>
<mml:mo>,</mml:mo>
<mml:mstyle mathvariant="bold" mathsize="normal">
<mml:mi>q</mml:mi>
</mml:mstyle>
<mml:mo>,</mml:mo>
<mml:mover accent="true">
<mml:mi>a</mml:mi>
<mml:mo>^</mml:mo>
</mml:mover>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula>&#x201c; contains three parts: &#x201c; <inline-formula>
<mml:math display="inline" id="im6">
<mml:mstyle mathvariant="bold" mathsize="normal">
<mml:mi>t</mml:mi>
</mml:mstyle>
</mml:math>
</inline-formula>&#x201c; and &#x201c; <inline-formula>
<mml:math display="inline" id="im7">
<mml:mstyle mathvariant="bold" mathsize="normal">
<mml:mi>q</mml:mi>
</mml:mstyle>
</mml:math>
</inline-formula>&#x201c; are the same as &#x201c; <inline-formula>
<mml:math display="inline" id="im8">
<mml:mstyle mathvariant="bold" mathsize="normal">
<mml:mi>t</mml:mi>
</mml:mstyle>
</mml:math>
</inline-formula>&#x201c; and &#x201c; <inline-formula>
<mml:math display="inline" id="im9">
<mml:mstyle mathvariant="bold" mathsize="normal">
<mml:mi>q</mml:mi>
</mml:mstyle>
</mml:math>
</inline-formula>&#x201c; in &#x201c; <inline-formula>
<mml:math display="inline" id="im10">
<mml:mrow>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mstyle mathvariant="bold" mathsize="normal">
<mml:mi>t</mml:mi>
</mml:mstyle>
<mml:mo>,</mml:mo>
<mml:mstyle mathvariant="bold" mathsize="normal">
<mml:mi>q</mml:mi>
</mml:mstyle>
<mml:mo>,</mml:mo>
<mml:mstyle mathvariant="bold" mathsize="normal">
<mml:mi>a</mml:mi>
</mml:mstyle>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula>&#x201c;, but &#x201c; <inline-formula>
<mml:math display="inline" id="im11">
<mml:mover accent="true">
<mml:mi>a</mml:mi>
<mml:mo>^</mml:mo>
</mml:mover>
</mml:math>
</inline-formula>&#x201c; represents the new answer generated using &#x201c;Vicuna-13B Q-A prompt&#x201d; for data correction. &#x201c;ChatGLM&#x201d; is a large language model used to revise data.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fmars-11-1368356-g003.tif"/>
</fig>
<sec id="s2_1_1">
<label>2.1.1</label>
<title>Text preprocessing</title>
<p>BBNJ-related documents were collected to build our dataset. The details are shown in <xref ref-type="table" rid="T1">
<bold>Table&#xa0;1</bold>
</xref>. In our study, text preprocessing consists of two steps: (1) division of regulations; and (2) reference completion. Division of regulations divides regulations into multiple paragraphs according to their clauses. As some clauses list multiple contents when describing &#x201c;requirements&#x201d; and &#x201c;steps&#x201d;, they often go beyond the input limit of the data generation model. Therefore, we took each part of the clause as a paragraph and supplemented Q-A pairs to ensure completeness. In addition, all the section titles were merged with any sentence within the corresponding section to retain semantic information. Reference completion ensures that each clause retains complete semantic information. Lots of clauses use referential terms such as &#x201c;above&#x201d; and &#x201c;this section&#x201d; in their descriptions, they lose complete semantic information after the division of regulations. Therefore, we completed the reference to the terms according to the context. Finally, 3,089 BBNJ-related paragraphs were obtained and used for Q-A data generation.</p>
<table-wrap id="T1" position="float">
<label>Table&#xa0;1</label>
<caption>
<p>Documents contained in the dataset.</p>
</caption>
<table frame="hsides">
<thead>
<tr>
<th valign="top" align="left">No.</th>
<th valign="top" align="center">BBNJ-related documents</th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="center">1</td>
<td valign="top" align="center">Draft agreement under the United Nations Convention on the Law of the Sea on the conservation and sustainable use of marine biological diversity of areas beyond national jurisdiction</td>
</tr>
<tr>
<td valign="top" align="center">2</td>
<td valign="top" align="center">Agreement under the United Nations Convention on the Law of the Sea on the conservation and sustainable use of marine biological diversity of areas beyond national jurisdiction</td>
</tr>
<tr>
<td valign="top" align="center">3</td>
<td valign="top" align="center">United Nations Convention on the Law of the Sea</td>
</tr>
<tr>
<td valign="top" align="center">4</td>
<td valign="top" align="center">United Nations Convention on the Law of the Sea Annex I</td>
</tr>
<tr>
<td valign="top" align="center">5</td>
<td valign="top" align="center">United Nations Convention on the Law of the Sea Annex II</td>
</tr>
<tr>
<td valign="top" align="center">6</td>
<td valign="top" align="center">United Nations Convention on the Law of the Sea Annex III</td>
</tr>
<tr>
<td valign="top" align="center">7</td>
<td valign="top" align="center">United Nations Convention on the Law of the Sea Annex IV</td>
</tr>
<tr>
<td valign="top" align="center">8</td>
<td valign="top" align="center">United Nations Convention on the Law of the Sea Annex V</td>
</tr>
<tr>
<td valign="top" align="center">9</td>
<td valign="top" align="center">United Nations Convention on the Law of the Sea Annex VI</td>
</tr>
<tr>
<td valign="top" align="center">10</td>
<td valign="top" align="center">United Nations Convention on the Law of the Sea Annex VII</td>
</tr>
<tr>
<td valign="top" align="center">11</td>
<td valign="top" align="center">United Nations Convention on the Law of the Sea Annex VIII</td>
</tr>
<tr>
<td valign="top" align="center">12</td>
<td valign="top" align="center">United Nations Convention on the Law of the Sea Annex IX</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec id="s2_1_2">
<label>2.1.2</label>
<title>Data generation</title>
<p>Preliminary evaluation using GPT-4 as a judge shows that Vicuna-13B achieves more than 90% quality of OpenAI&#x2019;s ChatGPT (<xref ref-type="bibr" rid="B5">Chiang et&#xa0;al., 2023</xref>). Since ChatGPT (<xref ref-type="bibr" rid="B25">Ouyang et&#xa0;al., 2022</xref>) is in a non-open source state, we applied open source Vicuna-13B as the base model for data generation to minimize experimental costs.</p>
<p>Prompting is a method for guiding the LLMs toward desired outputs. To achieve the best performance of LLMs in data generation, proper design of prompts is essential. We designed prompts for generating data at different stages. In the first stage, we used BBNJ-related paragraphs as input to generate Q-A pairs based on the content of the Phase 1 data generation prompt shown in <xref ref-type="fig" rid="f4">
<bold>Figure&#xa0;4</bold>
</xref>. In the second stage, following the concept of in-context learning (<xref ref-type="bibr" rid="B11">Dong et&#xa0;al., 2023</xref>), BBNJ-related paragraphs and Q-A pairs generated in Phase 1 were used as input to generate higher-quality Q-A pairs based on the content of the Phase 2 data generation prompt shown in <xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5</bold>
</xref>. Although it has been clearly stated in the prompt that LLMs need to generate new questions different from the examples, it was found that there were still occurrences of repetitive Q-A pairs during the experiment. Therefore, after the question generation process, it is necessary to carry out a deduplication operation on the generated Q-A pairs. After deduplication, we obtained 29,273 Q-A pairs, <inline-formula>
<mml:math display="inline" id="im12">
<mml:mrow>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mstyle mathsize="normal">
<mml:mi>q</mml:mi>
</mml:mstyle>
<mml:mo>,</mml:mo>
<mml:mi>a</mml:mi>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula>.</p>
<fig id="f4" position="float">
<label>Figure&#xa0;4</label>
<caption>
<p>Phase 1 data generation prompt.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fmars-11-1368356-g004.tif"/>
</fig>
<fig id="f5" position="float">
<label>Figure&#xa0;5</label>
<caption>
<p>Phase 2 data generation prompt.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fmars-11-1368356-g005.tif"/>
</fig>
</sec>
<sec id="s2_1_3">
<label>2.1.3</label>
<title>Data correction</title>
<p>To further improve the quality of the generated data, the data correction module depicted in <xref ref-type="fig" rid="f6">
<bold>Figure&#xa0;6</bold>
</xref> is applied to revise the generated Q-A pairs. Since the LLMs will generate out-of-scope answers, we obtained new model-generated answers <inline-formula>
<mml:math display="inline" id="im23">
<mml:mover accent="true">
<mml:mi>a</mml:mi>
<mml:mo>^</mml:mo>
</mml:mover>
</mml:math>
</inline-formula> based on the content of the Vicuna-13B Q-A prompt shown in <xref ref-type="fig" rid="f7">
<bold>Figure 7</bold>
</xref>. Then, we removed Q-A pairs with &#x201c;no answer&#x201d; and obtained new pairs <inline-formula>
<mml:math display="inline" id="im24">
<mml:mrow>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mtext>q</mml:mtext>
<mml:mo>,</mml:mo>
<mml:mover accent="true">
<mml:mi>a</mml:mi>
<mml:mo>^</mml:mo>
</mml:mover>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula>. The data were divided into three levels: easy, medium, and difficult by calculating the similarity between <inline-formula>
<mml:math display="inline" id="im25">
<mml:mtext>a</mml:mtext>
</mml:math>
</inline-formula> and <inline-formula>
<mml:math display="inline" id="im26">
<mml:mover accent="true">
<mml:mi>a</mml:mi>
<mml:mo>^</mml:mo>
</mml:mover>
</mml:math>
</inline-formula>.</p>
<fig id="f6" position="float">
<label>Figure&#xa0;6</label>
<caption>
<p>Data correction module. &#x201c; <inline-formula>
<mml:math display="inline" id="im13">
<mml:mrow>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mstyle mathvariant="bold" mathsize="normal">
<mml:mi>t</mml:mi>
</mml:mstyle>
<mml:mo>,</mml:mo>
<mml:mstyle mathvariant="bold" mathsize="normal">
<mml:mi>q</mml:mi>
</mml:mstyle>
<mml:mo>,</mml:mo>
<mml:mstyle mathvariant="bold" mathsize="normal">
<mml:mi>a</mml:mi>
</mml:mstyle>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula>&#x201c; is obtained from data generation. &#x201c; <inline-formula>
<mml:math display="inline" id="im14">
<mml:mstyle mathvariant="bold" mathsize="normal">
<mml:mi>t</mml:mi>
</mml:mstyle>
</mml:math>
</inline-formula>&#x201c; and &#x201c; <inline-formula>
<mml:math display="inline" id="im15">
<mml:mstyle mathvariant="bold" mathsize="normal">
<mml:mi>q</mml:mi>
</mml:mstyle>
</mml:math>
</inline-formula>&#x201c; in &#x201c; <inline-formula>
<mml:math display="inline" id="im16">
<mml:mrow>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mstyle mathvariant="bold" mathsize="normal">
<mml:mi>t</mml:mi>
</mml:mstyle>
<mml:mo>,</mml:mo>
<mml:mstyle mathvariant="bold" mathsize="normal">
<mml:mi>q</mml:mi>
</mml:mstyle>
<mml:mo>,</mml:mo>
<mml:mover accent="true">
<mml:mi>a</mml:mi>
<mml:mo>^</mml:mo>
</mml:mover>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula>&#x201c; are the same as &#x201c; <inline-formula>
<mml:math display="inline" id="im17">
<mml:mstyle mathvariant="bold" mathsize="normal">
<mml:mi>t</mml:mi>
</mml:mstyle>
</mml:math>
</inline-formula>&#x201c; and &#x201c; <inline-formula>
<mml:math display="inline" id="im18">
<mml:mstyle mathvariant="bold" mathsize="normal">
<mml:mi>q</mml:mi>
</mml:mstyle>
</mml:math>
</inline-formula>&#x201c; in &#x201c; <inline-formula>
<mml:math display="inline" id="im19">
<mml:mrow>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mstyle mathvariant="bold" mathsize="normal">
<mml:mi>t</mml:mi>
</mml:mstyle>
<mml:mo>,</mml:mo>
<mml:mstyle mathvariant="bold" mathsize="normal">
<mml:mi>q</mml:mi>
</mml:mstyle>
<mml:mo>,</mml:mo>
<mml:mstyle mathvariant="bold" mathsize="normal">
<mml:mi>a</mml:mi>
</mml:mstyle>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula>&#x201c;, but &#x201c; <inline-formula>
<mml:math display="inline" id="im20">
<mml:mover accent="true">
<mml:mi>a</mml:mi>
<mml:mo>^</mml:mo>
</mml:mover>
</mml:math>
</inline-formula>&#x201c; is the new model-generated answer using the &#x201c;Vicuna-13B Q-A prompt&#x201d;. During data division, first, we use the &#x201c;Vicuna-13B Q-A prompt&#x201d; to obtain<inline-formula>
<mml:math display="inline" id="im21">
<mml:mrow>
<mml:mstyle mathvariant="bold" mathsize="normal">
<mml:mi>a</mml:mi>
</mml:mstyle>
</mml:mrow>
</mml:math>
</inline-formula>&#xa0;and <inline-formula>
<mml:math display="inline" id="im22">
<mml:mover accent="true">
<mml:mi>a</mml:mi>
<mml:mo>^</mml:mo>
</mml:mover>
</mml:math>
</inline-formula>; finally, we obtain the divided data including easy, medium and difficult levels. During iterative correction, first, we fine-tune the ChatGLM using only easy data; then, we use the supervised fine-tuned model to infer pseudo-labels on medium data and form the pseudo-labels in memory; then, we use the model, which accesses the pseudo-labels in memory to infer pseudo-labels on difficult data; finally, we obtain an annotated biodiversity of areas beyond national jurisdiction domain dataset.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fmars-11-1368356-g006.tif"/>
</fig>
<fig id="f7" position="float">
<label>Figure&#xa0;7</label>
<caption>
<p>Vicuna-13B Q-A prompt.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fmars-11-1368356-g007.tif"/>
</fig>
<p>To minimize the computational resources used for model fine-tuning during the data correction, ChatGLM (<xref ref-type="bibr" rid="B44">Zeng et&#xa0;al., 2022</xref>) is used as the base model for data correction, which has only 6.2 billion parameters. Inspired by (<xref ref-type="bibr" rid="B36">Wang et&#xa0;al., 2021</xref>), we optimized the data correction model iteratively by increasing the difficulty of the Q-A pairs fed to the model gradually. First, we used easy Q-A pairs to conduct the initial fine-tuning of ChatGLM, and used the fine-tuned model to reannotate the medium Q-A pairs. Then, we further fine-tuned the model using the reannotated data and used the model obtained after this fine-tuning to reannotate the difficult Q-A pairs. Thus, the abilities of Vicuna-13B are transferred to ChatGLM for data correction in low-resource settings. Finally, 18,296 annotated Q-A pairs of BBNJ domain were obtained.</p>
</sec>
</sec>
<sec id="s2_2">
<label>2.2</label>
<title>Model fine-tuning</title>
<p>To improve the domain adaptation ability of LLMs applied to BBNJ domain, we fine-tuned the LLMs by using the BBNJ domain Q-A dataset constructed in Section 2.1. ChatGLM is applied as the base model. The open-source nature of the model is an important consideration.</p>
<p>To minimize the computational resources used for model fine-tuning, we applied the commonly used LoRA technique (<xref ref-type="bibr" rid="B14">Hu et&#xa0;al., 2021</xref>), which has been shown to effectively adapt LLMs to specific domain tasks and improve their performance (<xref ref-type="bibr" rid="B23">Lukichev et&#xa0;al., 2023</xref>). LoRA applies a simple linear design that allows the trainable matrix to be combined with frozen weights during model deployment. Compared with fully fine-tuned models, this approach does not create inference latencies, which is necessary for a knowledge Q-A system.</p>
<p>Data processing is necessary when fine-tuning ChatGLM. Research shows that ChatGLM and other LLMs can generalize well to unseen tasks and follow task descriptions after instruction tuning (<xref ref-type="bibr" rid="B40">Wei et&#xa0;al., 2021</xref>). However, LLMs have some weaknesses when lacking instructions, such as repetitive output and difficulty in fulfilling researchers&#x2019; expected task types. Therefore, it is necessary to construct the training set based on the organization of data in instruction tuning before fine-tuning. Each training data used for instruction tuning consists of three parts: instruction, question, and answer. <xref ref-type="table" rid="T2">
<bold>Table 2</bold>
</xref> shows an example of the BBNJ domain Q-A dataset.</p>
<table-wrap id="T2" position="float">
<label>Table&#xa0;2</label>
<caption>
<p>Example for instruction tuning.</p>
</caption>
<table frame="hsides">
<thead>
<tr>
<th valign="top" align="center">Instruction</th>
<th valign="top" align="center">The Conference of the Parties shall ordinarily meet at the seat of the secretariat or at United Nations Headquarters. Answer the following question based on the text given.</th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="center">Question</td>
<td valign="top" align="center">Where should the Conference of the Parties meet?</td>
</tr>
<tr>
<td valign="top" align="center">Answer</td>
<td valign="top" align="center">At the seat of the secretariat or at United Nations Headquarters.</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>The ChatBBNJ model was obtained by using the LoRA technique to fine-tune ChatGLM for the BBNJ knowledge Q-A task.</p>
</sec>
<sec id="s2_3">
<label>2.3</label>
<title>Domain knowledge-based prompt engineering</title>
<p>The training data for the ChatBBNJ model is limited to a specific time period, the model cannot provide time-sensitive knowledge. To solve this problem, we applied the framework depicted in <xref ref-type="fig" rid="f8">
<bold>Figure&#xa0;8</bold>
</xref> to the ChatBBNJ model.</p>
<fig id="f8" position="float">
<label>Figure&#xa0;8</label>
<caption>
<p>Domain knowledge-based prompt engineering framework. &#x201c;BBNJ Draft&#x201d; represents the draft agreement under the United Nations Convention on the Law of the Sea on the conservation and sustainable use of marine biological diversity of areas beyond national jurisdiction. &#x201c;BBNJ Agreement&#x201d; represents the Agreement under the United Nations Convention on the Law of the Sea on the conservation and sustainable use of marine biological diversity of areas beyond national jurisdiction.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fmars-11-1368356-g008.tif"/>
</fig>
<p>First, we create the BBNJ domain knowledge base using BBNJ Agreement and its draft. Second, we applied ERNIE 2.0 (<xref ref-type="bibr" rid="B29">Sun et&#xa0;al., 2020</xref>) to obtain vectorized representations of the domain knowledge, which was stored in our vector database. When a user poses a question, it is first embedded and then indexed using semantic similarity to find the top-k nearest vectors corresponding to the inquiry. The dot product of two vectors is utilized to analyze the similarity between vector embeddings, which is obtained by multiplying their respective components and summing the results. After identifying the nearest vectors to the query vector, we decode the numeric vectors to text and retrieve the corresponding text from the database. The textual information and user question are used to improve ChatBBNJ&#x2019;s prompt. This method enables users to receive reliable and up-to-date answers. The framework can be extended to other domains that require periodic knowledge updates.</p>
</sec>
</sec>
<sec id="s3">
<label>3</label>
<title>Experiment</title>
<sec id="s3_1">
<label>3.1</label>
<title>Baseline</title>
<p>To assess the performance of the proposed ChatBBNJ, we performed a comparative analysis using its base model ChatGLM and two additional language models.</p>
<p>LLaMA (<xref ref-type="bibr" rid="B32">Touvron et&#xa0;al., 2023a</xref>) is a collection of foundation language models ranging from 7 to 65 billion parameters, these models are trained using publicly available data. The experimental results demonstrate that LLaMA-7B outperforms GPT-3 in several natural language processing benchmark tests without relying on domain datasets.</p>
<p>Vicuna-13B (<xref ref-type="bibr" rid="B5">Chiang et&#xa0;al., 2023</xref>) is an open-source chatbot trained by fine-tuning LLaMA on user-shared conversations collected from ShareGPT. Preliminary evaluation using GPT-4 as a judge showed that Vicuna-13B achieved more than 90% quality of OpenAI&#x2019;s ChatGPT and Google&#x2019;s Bard.</p>
</sec>
<sec id="s3_2">
<label>3.2</label>
<title>Evaluation metrics</title>
<p>The models used in this study provide knowledge answers using a generative method. Traditional generative task evaluation metrics only consider word matching, they cannot provide a reasonable assessment for expressions with the same semantics. Thus, we applied the BERTScore (<xref ref-type="bibr" rid="B45">Zhang et&#xa0;al., 2019</xref>) to evaluate semantic equivalence. The precision, recall, and F1-scores were computed. <xref ref-type="disp-formula" rid="eq1">Equations (1</xref>-<xref ref-type="disp-formula" rid="eq3">3)</xref> provide the calculation methods for the evaluation metrics, where <inline-formula>
<mml:math display="inline" id="im27">
<mml:mrow>
<mml:mi>x</mml:mi>
<mml:mo>=</mml:mo>
<mml:mrow>
<mml:mo>&#x2329;</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mn>1</mml:mn>
</mml:msub>
<mml:mo>,</mml:mo>
<mml:mi>&#x2026;</mml:mi>
<mml:mo>,</mml:mo>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mi>k</mml:mi>
</mml:msub>
</mml:mrow>
<mml:mo>&#x232a;</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula> represents the ground-truth answers in the test set and <inline-formula>
<mml:math display="inline" id="im28">
<mml:mrow>
<mml:mover accent="true">
<mml:mi>x</mml:mi>
<mml:mo>^</mml:mo>
</mml:mover>
<mml:mo>=</mml:mo>
<mml:mrow>
<mml:mo>&#x2329;</mml:mo>
<mml:mrow>
<mml:mover accent="true">
<mml:mrow>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mn>1</mml:mn>
</mml:msub>
</mml:mrow>
<mml:mo stretchy="true">^</mml:mo>
</mml:mover>
<mml:mo>,</mml:mo>
<mml:mi>&#x2026;</mml:mi>
<mml:mo>,</mml:mo>
<mml:mover accent="true">
<mml:mrow>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mi>l</mml:mi>
</mml:msub>
</mml:mrow>
<mml:mo stretchy="true">^</mml:mo>
</mml:mover>
</mml:mrow>
<mml:mo>&#x232a;</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula> represents the answers from ChatBBNJ and baselines.</p>
<disp-formula id="eq1">
<label>(1)</label>
<mml:math display="block" id="M1">
<mml:mrow>
<mml:msub>
<mml:mtext mathvariant="bold-italic">P</mml:mtext>
<mml:mrow>
<mml:mtext mathvariant="bold-italic">BERT</mml:mtext>
</mml:mrow>
</mml:msub>
<mml:mo>=</mml:mo>
<mml:mfrac>
<mml:mn mathvariant="bold">1</mml:mn>
<mml:mrow>
<mml:mrow>
<mml:mo>|</mml:mo>
<mml:mover accent="true">
<mml:mtext mathvariant="bold-italic">x</mml:mtext>
<mml:mo>^</mml:mo>
</mml:mover>
<mml:mo>|</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:mfrac>
<mml:mstyle displaystyle="true">
<mml:munder>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mover accent="true">
<mml:mrow>
<mml:msub>
<mml:mtext mathvariant="bold-italic">x</mml:mtext>
<mml:mtext mathvariant="bold-italic">h</mml:mtext>
</mml:msub>
</mml:mrow>
<mml:mo stretchy="false">^</mml:mo>
</mml:mover>
<mml:mo>&#x2208;</mml:mo>
<mml:mover accent="true">
<mml:mtext mathvariant="bold-italic">x</mml:mtext>
<mml:mo>^</mml:mo>
</mml:mover>
</mml:mrow>
</mml:munder>
</mml:mstyle>
<mml:munder>
<mml:mrow><mml:mtext mathvariant="bold-italic">max</mml:mtext></mml:mrow>
<mml:mrow>
<mml:msub>
<mml:mtext mathvariant="bold-italic">x</mml:mtext>
<mml:mi>i</mml:mi>
</mml:msub>
<mml:mo>&#x2208;</mml:mo>
<mml:mtext mathvariant="bold-italic">x</mml:mtext>
</mml:mrow>
</mml:munder>
<mml:mo>&#xa0;</mml:mo>
<mml:msubsup>
<mml:mtext mathvariant="bold-italic">x</mml:mtext>
<mml:mrow><mml:mtext mathvariant="bold-italic">i</mml:mtext></mml:mrow>
<mml:mrow><mml:mtext mathvariant="bold-italic">T</mml:mtext></mml:mrow>
</mml:msubsup>
<mml:mover accent="true">
<mml:mrow>
<mml:msub>
<mml:mtext mathvariant="bold-italic">x</mml:mtext>
<mml:mtext mathvariant="bold-italic">h</mml:mtext>
</mml:msub>
</mml:mrow>
<mml:mo stretchy="true">^</mml:mo>
</mml:mover>
</mml:mrow>
</mml:math>
</disp-formula>
<disp-formula id="eq2">
<label>(2)</label>
<mml:math display="block" id="M2">
<mml:mrow>
<mml:msub>
<mml:mtext mathvariant="bold-italic">R</mml:mtext>
<mml:mrow>
<mml:mtext mathvariant="bold-italic">BERT</mml:mtext>
</mml:mrow>
</mml:msub>
<mml:mo>=</mml:mo>
<mml:mfrac>
<mml:mn mathvariant="bold">1</mml:mn>
<mml:mrow>
<mml:mrow>
<mml:mo>|</mml:mo>
<mml:mtext mathvariant="bold-italic">x</mml:mtext>
<mml:mo>|</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:mfrac>
<mml:mstyle displaystyle="true">
<mml:munder>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mtext mathvariant="bold-italic">x</mml:mtext>
<mml:mtext mathvariant="bold-italic">i</mml:mtext>
</mml:msub>
<mml:mo>&#x2208;</mml:mo>
<mml:mtext mathvariant="bold-italic">x</mml:mtext>
</mml:mrow>
</mml:munder>
</mml:mstyle>
<mml:munder>
<mml:mrow>
<mml:mtext mathvariant="bold-italic">max</mml:mtext>
<mml:mo>&#xa0;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mover accent="true">
<mml:mrow>
<mml:msub>
<mml:mtext mathvariant="bold-italic">x</mml:mtext>
<mml:mtext mathvariant="bold-italic">h</mml:mtext>
</mml:msub>
</mml:mrow>
<mml:mo stretchy="false">^</mml:mo>
</mml:mover>
<mml:mo>&#x2208;</mml:mo>
<mml:mover accent="true">
<mml:mtext mathvariant="bold-italic">x</mml:mtext>
<mml:mo>^</mml:mo>
</mml:mover>
</mml:mrow>
</mml:munder>
<mml:msubsup>
<mml:mtext mathvariant="bold-italic">x</mml:mtext>
<mml:mrow><mml:mtext mathvariant="bold-italic">i</mml:mtext></mml:mrow>
<mml:mrow><mml:mtext mathvariant="bold-italic">T</mml:mtext></mml:mrow>
</mml:msubsup>
<mml:mover accent="true">
<mml:mrow>
<mml:msub>
<mml:mtext mathvariant="bold-italic">x</mml:mtext>
<mml:mtext mathvariant="bold-italic">h</mml:mtext>
</mml:msub>
</mml:mrow>
<mml:mo stretchy="true">^</mml:mo>
</mml:mover>
</mml:mrow>
</mml:math>
</disp-formula>
<disp-formula id="eq3">
<label>(3)</label>
<mml:math display="block" id="M3">
<mml:mrow>
<mml:msub>
<mml:mtext mathvariant="bold-italic">F</mml:mtext>
<mml:mrow>
<mml:mtext mathvariant="bold-italic">BERT</mml:mtext>
</mml:mrow>
</mml:msub>
<mml:mo>=</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mn>2</mml:mn>
<mml:mo>*</mml:mo>
<mml:mo>&#xa0;</mml:mo>
<mml:msub>
<mml:mtext mathvariant="bold-italic">P</mml:mtext>
<mml:mrow>
<mml:mtext mathvariant="bold-italic">BERT</mml:mtext>
</mml:mrow>
</mml:msub>
<mml:mo>*</mml:mo>
<mml:mo>&#xa0;</mml:mo>
<mml:msub>
<mml:mtext mathvariant="bold-italic">R</mml:mtext>
<mml:mrow>
<mml:mtext mathvariant="bold-italic">BERT</mml:mtext>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mrow>
<mml:msub>
<mml:mtext mathvariant="bold-italic">P</mml:mtext>
<mml:mrow>
<mml:mtext mathvariant="bold-italic">BERT</mml:mtext>
</mml:mrow>
</mml:msub>
<mml:mo>+</mml:mo>
<mml:msub>
<mml:mtext mathvariant="bold-italic">R</mml:mtext>
<mml:mrow>
<mml:mtext mathvariant="bold-italic">BERT</mml:mtext>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:math>
</disp-formula>
</sec>
<sec id="s3_3" sec-type="results">
<label>3.3</label>
<title>Results</title>
<p>We demonstrated the effectiveness of our method framework, the efficacy of its individual modules, and the Q-A performance of ChatBBNJ quantitatively through experiments. We divided the BBNJ domain dataset constructed in Section 2.1 into training and testing sets with ratios of 77 and 23%, respectively.</p>
<p>First, we compared ChatBBNJ with several other LLMs in Q-A tasks. To quantitatively evaluate the performance of ChatBBNJ, we calculated the metrics for both baselines in Section 3.1 and ChatBBNJ. The results are shown in <xref ref-type="fig" rid="f9">
<bold>Figure&#xa0;9</bold>
</xref>. The results show that ChatBBNJ achieves significantly higher precision, recall, and F1-scores, compared to baselines. The results show that ChatBBNJ achieves 0.097 precision, 0.025 recall, and 0.062 F1-scores improvement over its base model ChatGLM. These results demonstrate that ChatBBNJ exhibits superior expertise in the BBNJ domain.</p>
<fig id="f9" position="float">
<label>Figure&#xa0;9</label>
<caption>
<p>Comparison of the experimental results for ChatBBNJ.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fmars-11-1368356-g009.tif"/>
</fig>
<p>Second, we conducted ablation experiments to demonstrate the efficacy of individual modules in our method framework. To quantitatively evaluate the effectiveness of the three proposed modules, namely PDGC for BBNJ domain dataset construction, model fine-tuning, and domain knowledge-based prompt engineering, we also calculated the metrics under different scenarios. The complete evaluation results are shown in <xref ref-type="table" rid="T3">
<bold>Table&#xa0;3</bold>
</xref>. The quality of training data influences how well the model learns. Therefore, to evaluate the effectiveness of PDGC, we fine-tuned the model for Q-A task before and after using PDGC, then tested the model&#x2019;s performance. The results indicate that the model&#x2019;s Q-A performance improves after fine-tuning the model with data generated by PDGC. Additionally, fine-tuning the model helps enhance its performance. We find that the model achieves improvement over ChatGLM. To evaluate the effectiveness of the domain knowledge-based prompt engineering, we also compared the Q-A performance before and after using this module. The results indicate that the model performs better after using the domain knowledge-based prompt engineering when other modules are the same. The results indicate that the model incorporating with each of our module achieves higher precision, recall and F1-scores in all cases.</p>
<table-wrap id="T3" position="float">
<label>Table&#xa0;3</label>
<caption>
<p>Ablation experiment results for the three modules of ChatBBNJ.</p>
</caption>
<table frame="hsides">
<thead>
<tr>
<th valign="top" align="center">Models</th>
<th valign="top" align="center">Precision</th>
<th valign="top" align="center">Recall</th>
<th valign="top" align="center">F1</th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="center">ChatGLM (without PDGC, model fine-tuning, and DKB prompt engineering)</td>
<td valign="top" align="center">0.822</td>
<td valign="top" align="center">0.876</td>
<td valign="top" align="center">0.848</td>
</tr>
<tr>
<td valign="top" align="center">ChatBBNJ (without PDGC, with model fine-tuning, without DKB prompt engineering)</td>
<td valign="top" align="center">0.865</td>
<td valign="top" align="center">0.891</td>
<td valign="top" align="center">0.873</td>
</tr>
<tr>
<td valign="top" align="center">ChatBBNJ (with PDGC, with model fine-tuning, without DKB prompt engineering)</td>
<td valign="top" align="center">0.872</td>
<td valign="top" align="center">0.894</td>
<td valign="top" align="center">0.888</td>
</tr>
<tr>
<td valign="top" align="center">ChatBBNJ (without PDGC, with model fine-tuning and DKB prompt engineering)</td>
<td valign="top" align="center">0.897</td>
<td valign="top" align="center">0.900</td>
<td valign="top" align="center">0.898</td>
</tr>
<tr>
<td valign="top" align="center">ChatBBNJ (with PDGC, model fine-tuning and DKB prompt engineering)</td>
<td valign="top" align="center">
<bold>0.919</bold>
</td>
<td valign="top" align="center">
<bold>0.901</bold>
</td>
<td valign="top" align="center">
<bold>0.910</bold>
</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<fn>
<p>DKB prompt engineering represents domain knowledge-based prompt engineering.</p>
</fn>
<fn>
<p>The bold values represent the best results.</p>
</fn>
</table-wrap-foot>
</table-wrap>
<p>Finally, we demonstrate the effectiveness of ChatBBNJ through specific Q-A examples. The examples generated by ChatBBNJ and baseline models are listed in <xref ref-type="table" rid="T4">
<bold>Table&#xa0;4</bold>
</xref>, <xref ref-type="table" rid="T5">
<bold>Table&#xa0;5</bold>
</xref> and <xref ref-type="supplementary-material" rid="SM1">
<bold>supplementary materials</bold>
</xref>. In the tables, &#x201c;Q&#x201d; represents the question provided to the models. &#x201c;ChatGLM&#x201d;, &#x201c;LLaMA-7B&#x201d;, &#x201c;Vicuna-13B&#x201d;, and &#x201c;ChatBBNJ&#x201d; represent the answers generated by the models. &#x201c;Answer&#x201d; represents the ground-truth answer. We underline the key points of the models&#x2019; answers.</p>
<table-wrap id="T4" position="float">
<label>Table&#xa0;4</label>
<caption>
<p>Q-A example in BBNJ domain terminology.</p>
</caption>
<table frame="hsides">
<thead>
<tr>
<th valign="top" align="left">Q: What is the purpose of area-based management tool?</th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="center">
<bold>ChatGLM</bold>: An area-based management tool is a software application that allows organizations to manage their resources and operations across different locations, regions, or areas. The purpose of an area-based management tool is to help organizations more effectively manage their resources, improve communication, and increase productivity. One of the main benefits of using an area-based management tool is that it allows organizations to manage their resources and operations across different locations, regions, or areas more efficiently. This can help to reduce costs by reducing the need for travel and accommodation, and improving communication by allowing for real-time collaboration between teams in different locations. Additionally, area-based management tools can help to improve productivity by allowing organizations to manage their resources and operations more effectively across different locations. This can help to reduce the time it takes for teams to work together, and improve the overall performance of the organization. Overall, area-based management tools are an important tool for organizations looking to improve their resource management, communication, and productivity across different locations.</td>
</tr>
<tr>
<td valign="top" align="center">
<bold>LLaMA-7B</bold>: Area based Management Tool (ABMT) was developed to help managers and employees in making decisions about how best to manage their resources, such as people or money. It helps them identify which areas are most important for improvement by identifying what needs attention first. The ABM system can be used with any type of resource but it works particularly well when you have limited time and/or budget available so that you need to prioritize your efforts effectively. The main steps involved include: 1. Identify all possible options;2. Priority ranking each option according to its potential impact on performance goals;3. Selecting one or more priority actions from those ranked highest;4. Implement selected action(s);5. Evaluate results against original objectives;6. Repeat process if necessary until desired level of effectiveness has been achieved.</td>
</tr>
<tr>
<td valign="top" align="center">
<bold>Vicuna-13B</bold>: Area-based management tools are used to help organizations and individuals manage specific areas or regions. These tools can be used for a variety of purposes, such as to track and manage resources, monitor activity, and identify trends and patterns. Some examples of area-based management tools include geographic information systems (GIS), which can be used to create maps and visualizations of an area, and resource management software, which can be used to track and manage resources such as fish stocks or forests. Area-based management tools can be useful for a wide range of applications, including natural resource management, emergency response, and urban planning.</td>
</tr>
<tr>
<td valign="top" align="center">
<bold>ChatBBNJ</bold>: The purpose of an area-based management tool is to <underline>manage a Geographically defined area through which one or several sectors or activities are managed with the aim of achieving particular conservation.</underline> The tool may <underline>include a marine protected area</underline> as one of its components. The establishment of area-based management tools, including marine protected areas, is to prevent the inclusion of any areas within national jurisdiction in order to prevent asserts or denying any claims to sovereignty.</td>
</tr>
<tr>
<td valign="top" align="center">
<bold>Answer</bold>: &#x201c;Area-based management tool&#x201d; means a tool, including a marine protected area, for a geographically defined area through which one or several sectors or activities are managed with the aim of achieving particular conservation and sustainable use objectives in accordance with this Agreement.</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<fn>
<p>ChatGLM', 'LLaMA-7B', 'Vicuna-13B', and 'ChatBBNJ' represent the answers generated by the models. 'Answer' represents the ground-truth answer.</p>
<p>The underline words represent the key points of the models' answers.</p>
</fn>
</table-wrap-foot>
</table-wrap>
<table-wrap id="T5" position="float">
<label>Table&#xa0;5</label>
<caption>
<p>Time-sensitive Q-A example in the BBNJ domain.</p>
</caption>
<table frame="hsides">
<thead>
<tr>
<th valign="top" align="left">Q: Where should the Conference of the Parties meet?</th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="center">
<bold>ChatGLM</bold>: The Conference of the Parties (CP) is an international organization that&#x81f4;&#x529b;&#x4e8e; developing and implementing policies to reduce climate change. It meets regularly, usually in locations that are convenient for its members. TheCP&#x2019;s current meeting, which took place in Madrid, Spain in 2020, was the 11th meeting of theCP since its formation in 1995. The next meeting of theCP will take place in late 2021 or early 2022, and the location of the meeting will be announced at a later date. TheCP&#x2019;s meetings are held in a variety of locations, including &#x4f1a;&#x8bae;&#x5ba4; (&#x4f1a;&#x8bae;&#x5ba4;&#x662f;&#x6307;&#x4f1a;&#x8bae;&#x6240;&#x9700;&#x7684;&#x4f1a;&#x8bae;&#x5ba4;&#x3001;&#x4f1a;&#x8bae;&#x684c;&#x3001;&#x6905;&#x5b50;&#x7b49;&#x4f1a;&#x8bae;&#x8bbe;&#x65bd;), which are designed to provide a comfortable and productive environment for conference attendees. The location of theCP&#x2019;s next meeting will be determined by the needs and preferences of its members, as well as the availability of&#x4f1a;&#x8bae;&#x5ba4; and other &#x4f1a;&#x8bae;&#x8bbe;&#x65bd;. TheCP meets regularly to discuss and develop policies to reduce climate change, including policies to reduce carbon emissions, promote renewable energy, and protect the environment. The organization is committed to working towards a goal of reducing climate change by 2C by the end of the century.</td>
</tr>
<tr>
<td valign="top" align="center">
<bold>LLaMA-7B</bold>: The conference will be held in Paris, France because it has been suggested by many countries and they have agreed to host this meeting there for 2 weeks starting on December 1st until January 3rd at Hotel de Ville (City Hall). This location was chosen due to its centrality within Europe which makes traveling easier as well as being able to accommodate all delegates from around the world with ease. The city also offers great food options such as baguettes or crepes along with other delicious cuisines making sure everyone&#x2019;s taste buds are satisfied! It would not only benefit those who attend but also help boost tourism throughout the country bringing more money into their economy.</td>
</tr>
<tr>
<td valign="top" align="center">
<bold>Vicuna-13B</bold>: The Conference of the Parties (COP) is a meeting of the parties to the United Nations Framework Convention on Climate Change (UNFCCC). The COP is the supreme decision-making body of the UNFCCC, and it is responsible for reviewing the implementation of the Convention and taking decisions to further its objectives. The COP is typically held annually, and it is attended by representatives from the governments of the parties to the Convention, as well as experts, civil society organizations, and the media. The location of the COP is typically decided by the UNFCCC secretariat and the government of the host country. The COP has been held in a variety of locations, including developed and developing countries, and it has typically been hosted by countries that are parties to the Convention and have the capacity to host a large international meeting. In the past, the COP has been held in countries such as Germany, France, and South Africa. The next COP is scheduled to be held in Glasgow, Scotland in 2021.</td>
</tr>
<tr>
<td valign="top" align="center">
<bold>ChatBBNJ</bold>: <underline>At the seat of the secretariat or at United Nations Headquarters</underline>.</td>
</tr>
<tr>
<td valign="top" align="center">
<bold>Answer</bold>: The Conference of the Parties shall ordinarily meet at the seat of the secretariat or at United Nations Headquarters.</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<fn>
<p>ChatGLM', 'LLaMA-7B', 'Vicuna-13B', and 'ChatBBNJ' represent the answers generated by the models. 'Answer' represents the ground-truth answer.</p>
<p>The underline words represent the key points of the models' answers.</p>
</fn>
</table-wrap-foot>
</table-wrap>
<p>In <xref ref-type="table" rid="T4">
<bold>Table&#xa0;4</bold>
</xref>, we provide a question related to BBNJ domain terminology. ChatGLM considers &#x201c;area-based management tool&#x201d; to be a software application, LLaMA-7B categorizes as manager and employee resource management, their answers deviate from the BBNJ domain. Vicuna-13B&#x2019;s answer is more inclusive, but it lacks domain specificity. Our ChatBBNJ provides an answer that reflects the BBNJ domain specificities and includes most of the key points from the ground-truth answer.</p>
<p>In <xref ref-type="table" rid="T5">
<bold>Table&#xa0;5</bold>
</xref>, we provide a time-sensitive question related to the BBNJ domain. This question is regulated first in the BBNJ Agreement, which was formally adopted in August 2023. The answers from ChatGLM, LLaMA-7B, and Vicuna-13B are related to the COP but not to the BBNJ domain, and ChatGLM displayed language inconsistencies. It can be seen that ChatGLM's answer is a mixture of Chinese and English format, the answer includes "&#x4f1a;&#x8bae;&#x5ba4;"(conference room)&#x3001;"&#x4f1a;&#x8bae;&#x8bbe;&#x65bd;"(conference facilities) and other Chinese format. ChatGLM&#x2019;s knowledge is cut off in 2022. However, ChatBBNJ, which is based on ChatGLM, provides correct BBNJ domain answers based on knowledge in 2023.</p>
<p>In <xref ref-type="supplementary-material" rid="SM1">
<bold>Supplementary Table&#xa0;1</bold>
</xref>, we provide two Q-A examples for complex application scenario. The answers of ChatGLM, LLaMA-7B and Vicuna-13B use terms from the BBNJ Agreement rarely, especially ChatGLM expresses ABMT as ABM tools in Q2, which is not conform to the treaty. Besides, ChatGLM's answer is a mixture of Chinese and English format, it includes "&#x5bfb;&#x6c42;&#x56fd;&#x9645;&#x534f;&#x4f5c;, &#x4ee5;&#x89e3;&#x51b3;&#x8fd9;&#x4e00;&#x95ee;&#x9898;" (seek international cooperation to solve this problem), "&#x91c7;&#x53d6;&#x53ef;&#x6301;&#x7eed;&#x7684;&#x63aa;&#x65bd;" (take sustainable measures) and other Chinese format. However, ChatBBNJ not only provides the original text from the Agreement in Q2, but also emphasizes the necessity of the relevant reports and the role of the United Nations in Q1, which make the answers more comprehensive. It can be seen that compared to other models, our ChatBBNJ can provide answers more aligned with the BBNJ Agreement in complex scenarios, especially in Q2, where ChatBBNJ provides almost all requirements related to ABMTs in the BBNJ Agreement.</p>
<p>In <xref ref-type="supplementary-material" rid="SM1">
<bold>Supplementary Table&#xa0;2</bold>
</xref>, we provide two Q-A examples for the interpretation of the BBNJ Agreement under Article 31 of the Vienna Convention on the Law of Treaties (VCLT). Article 31 of the VCLT can be summarized as: be interpreted in good faith, be interpreted in the light of treaty&#x2019;s object and purpose, be interpreted with supplementary means of interpretation, be interpreted in accordance with the ordinary meaning and be interpreted with the terms of the treaty in their context. In Q1, ChatBBNJ&#x2019;s response highlights the importance of international cooperation in marine scientific research and technology development in first sentence, which complies with the principle of good faith interpretation. And the first sentence also links to the objective of the treaty. Moreover, ChatBBNJ&#x2019;s response provides relevant contents from Article 143 of the UNCLOS, reflecting supplementing interpretation based on external materials. In Q2, ChatBBNJ provides an ordinary explanation of &#x201c;transparency&#x201d;, supplemented it with interpretations of the term in different contexts, and provide the role of transparency as well as the challenges it faces. This response reflects an ordinary meaning interpretation. Besides, ChatBBNJ provides some measures involved in the BBNJ Agreement, including making decisions and documents open to the public, open meeting practices, publishing and maintaining a public record of decisions and publishing decision information. These measures come from different articles in the agreement, which reflects interpreting with the terms of the treaty in their context.</p>
<p>The latest versions of the baseline models (e.g., ChatGLM2 and LLaMA2-7B (<xref ref-type="bibr" rid="B33">Touvron et&#xa0;al., 2023b</xref>)) had improved abilities over ChatGLM and LLaMA-7B used in this study. Since we applied the ChatGLM as a base model, we compared ChatBBNJ with the latest baseline models, ChatGLM2 and LLaMA2-7B. The results are shown in <xref ref-type="table" rid="T6">
<bold>Table&#xa0;6</bold>
</xref>. The results show that our ChatBBNJ outperforms the more advanced ChatGLM2 and LLaMA2-7B. These findings suggest that our method framework for BBNJ domain Q-A, significantly enhance the performance of the LLMs in the BBNJ domain.</p>
<table-wrap id="T6" position="float">
<label>Table&#xa0;6</label>
<caption>
<p>Supplementary comparisons of experimental results for ChatBBNJ.</p>
</caption>
<table frame="hsides">
<thead>
<tr>
<th valign="top" align="center">Models</th>
<th valign="top" align="center">Precision</th>
<th valign="top" align="center">Recall</th>
<th valign="top" align="center">F1</th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="center">ChatGLM2</td>
<td valign="top" align="center">0.846</td>
<td valign="top" align="center">0.883</td>
<td valign="top" align="center">0.864</td>
</tr>
<tr>
<td valign="top" align="center">LLaMA2-7B</td>
<td valign="top" align="center">0.852</td>
<td valign="top" align="center">0.889</td>
<td valign="top" align="center">0.870</td>
</tr>
<tr>
<td valign="top" align="center">ChatBBNJ</td>
<td valign="top" align="center">
<bold>0.919</bold>
</td>
<td valign="top" align="center">
<bold>0.901</bold>
</td>
<td valign="top" align="center">
<bold>0.910</bold>
</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<fn>
<p>The bold values represent the best results.</p>
</fn>
</table-wrap-foot>
</table-wrap>
</sec>
</sec>
<sec id="s4" sec-type="discussion">
<label>4</label>
<title>Discussion</title>
<p>Over the past 50 years, the degradation of marine biodiversity in ABNJ persists, with a lack of effective governance mechanisms to halt this decline (<xref ref-type="bibr" rid="B39">Ward et&#xa0;al., 2022</xref>). Industrial fishing affects about 50% of oceans (<xref ref-type="bibr" rid="B28">Sala et&#xa0;al., 2018</xref>), leading to overexploitation of 31% of marine fish stocks (<xref ref-type="bibr" rid="B13">FAO, 2016</xref>), and ecosystem-level alterations in high seas (<xref ref-type="bibr" rid="B24">Ortu&#xf1;o Crespo and Dunn, 2017</xref>). The international maritime order is dynamic and evolving, and UNCLOS&#x2019; authoritative, comprehensive, and extended nature does not imply its perfection. Due to the game of interests and compromise among countries, many UNCLOS rules are principled and articulated, creating legal ambiguities that often need to be further addressed in practice. The current international ocean order is in a state of rapid transition, necessitating a reaction to several earth system changes such as sea-level rise, plastic pollution of the seas, acidification, and destruction of marine biodiversity (<xref ref-type="bibr" rid="B4">Chen and Liu, 2023</xref>). BBNJ agreement is intended to serve as a governance mechanism for the protection and sustainable use of marine biodiversity in the ABNJ.</p>
<p>The present study introduces ChatBBNJ, a question-answering system, that designed to enhance the treaty implementation. First, we proposed PDGC, a data engineering method, that constructs a high-quality Q-A dataset specific to the BBNJ domain. Second, we applied this dataset to fine-tune the ChatGLM to obtain the BBNJ domain model, ChatBBNJ. Finally, we introduced a domain knowledge-based prompt engineering. We demonstrated improvements by testing ChatBBNJ on Q-A data related to the BBNJ Agreement and its draft. Experiment results demonstrate that ChatBBNJ outperforms baseline LLMs across three Q-A metrics. Additionally, hybrid ChatBBNJ, which introduces a domain knowledge-based prompt framework outperforms standalone ChatBBNJ. The main findings of our work are summarized as follows:</p>
<list list-type="simple">
<list-item>
<p>(1) The quality of LLMs text generation can be enhanced through appropriate prompt engineering and data correction. The effectiveness of model training is to some extent related to the quality of training data. Therefore, <xref ref-type="table" rid="T3">
<bold>Table 3</bold>
</xref> compares the Q-A performance of ChatBBNJ before and after using PDGC. It can be seen that the Q-A performance of ChatBBNJ improves when the model is fine-tuned with data generated by PDGC.</p>
</list-item>
<list-item>
<p>(2) The domain adaptation ability of LLMs in Q-A tasks can be improved by fine-tuning the model on domain-specific Q-A datasets. Analyzing the experimental results in <xref ref-type="table" rid="T3">
<bold>Table 3</bold>
</xref>, regardless of the method used for data generation, ChatBBNJ&#x2019;s Q-A performance improves after model fine-tuning.</p>
</list-item>
<list-item>
<p>(3) The outdated issues of LLMs can be refined by giving the model access to the knowledge beyond its fine-tuning phase time and instructing LLMs on how to utilize that knowledge. In <xref ref-type="table" rid="T5">
<bold>Table 5</bold>
</xref>, the model was asked a time-sensitive question, with relevant information not within the model&#x2019;s training data. However, our ChatBBNJ, equipped with the domain knowledge-based prompt engineering, accessed external knowledge and provided the correct answer.</p>
</list-item>
<list-item>
<p>(4) This study not only demonstrates the feasibility of BBNJ knowledge acquisition by using Q-A system but also provides a method framework to apply Q-A system in BBNJ domain. The method framework has three parts. Firstly, domain datasets are acquired using data engineering method. Secondly, LLMs are fine-tuned for Q-A task based on the generated datasets. Finally, a domain knowledge-based prompt engineering is employed to assist the model in accessing the latest information, ensuring the timeliness of its responses.</p>
</list-item>
</list>
<p>In the future, the successful application of ChatBBNJ can bring several benefits. From the perspective of the representatives who have not participated in BBNJ negotiations, they can acquire BBNJ knowledge and avoid reading a large number of irrelevant documents. Simultaneously, governments can benefit from the Q-A system. The BBNJ Agreement calls for parties to take the necessary legislative measures to ensure the implementation of this agreement (<xref ref-type="bibr" rid="B34">United Nations, 2023</xref>). ChatBBNJ can provide popularization and interpretation of the agreement, which is helpful to assist governments in enacting implementation legislation. The successful implementation of the proposed ChatBBNJ will make a substantial contribution towards prompting the effectie implementation of the BBNJ Agreement, which will hopefully result in the conservation and sustainable use of marine biodiversity in ABNJ.</p>
<p>However, the limitations of this study remain. First, the accuracy of ChatBBNJ needs to be improved. The present system has achieved reasonable accuracy in the current research field, but there is still room for improvement. Second, the sources of knowledge need to be expanded. The present study mainly focused on promoting the implementation of the BBNJ Agreement, but the methods can be applied to more implementation in international agreements in the law of the sea domain. In future research, more international law of the sea data will be added to the study to obtain richer knowledge and expand the scope of applications.</p>
</sec>
<sec id="s5" sec-type="data-availability">
<title>Data availability statement</title>
<p>The raw data supporting the conclusions of this article will be made available by the authors, without undue reservation.</p>
</sec>
<sec id="s6" sec-type="author-contributions">
<title>Author contributions</title>
<p>XW: Conceptualization, Data curation, Formal Analysis, Investigation, Methodology, Project administration, Software, Validation, Visualization, Writing &#x2013; original draft, Writing &#x2013; review &amp; editing. MZ: Data curation, Formal analysis, Writing &#x2013; review &amp; editing. HL: Conceptualization, Resources, Supervision, Writing &#x2013; review &amp; editing. XM: Project administration, Validation, Writing &#x2013; original draft. YL: Project administration, Writing &#x2013; original draft. YC: Conceptualization, Data curation, Formal Analysis, Funding acquisition, Resources, Supervision, Writing &#x2013; review &amp; editing.</p>
</sec>
</body>
<back>
<sec id="s7" sec-type="funding-information">
<title>Funding</title>
<p>The author(s) declare financial support was received for the research, authorship, and/or publication of this article. This research was funded by the National Social Science Fund of China, Grant No. 20CFX082.</p>
</sec>
<ack>
<title>Acknowledgments</title>
<p>We would like to thank the computing resources and the professional technical guidance provided by Hao Liu. Thanks to Mittens and Panda.</p>
</ack>
<sec id="s8" sec-type="COI-statement">
<title>Conflict of interest</title>
<p>The authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec id="s9" sec-type="disclaimer">
<title>Publisher&#x2019;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<sec id="s10" sec-type="supplementary-material">
<title>Supplementary material</title>
<p>The Supplementary Material for this article can be found online at: <ext-link ext-link-type="uri" xlink:href="https://www.frontiersin.org/articles/10.3389/fmars.2024.1368356/full#supplementary-material">https://www.frontiersin.org/articles/10.3389/fmars.2024.1368356/full#supplementary-material</ext-link>
</p>
<supplementary-material xlink:href="Table_1.docx" id="SM1" mimetype="application/vnd.openxmlformats-officedocument.wordprocessingml.document"/>
</sec>
<ref-list>
<title>References</title>
<ref id="B1">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Amer-Yahia</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Bonifati</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Chen</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>G. L.</given-names>
</name>
<name>
<surname>Shim</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Xu</surname> <given-names>J. L.</given-names>
</name>
<etal/>
</person-group>. (<year>2023</year>). <article-title>From large language models to databases and back: A discussion on research and education</article-title>. <source>SIGMOD Rec.</source> <volume>52</volume>, <fpage>49</fpage>&#x2013;<lpage>56</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1145/3631504.3631518</pub-id>
</citation>
</ref>
<ref id="B2">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Bengio</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Ducharme</surname> <given-names>R.</given-names>
</name>
<name>
<surname>Vincent</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Jauvin</surname> <given-names>C.</given-names>
</name>
</person-group> (<year>2003</year>). <article-title>A neural probabilistic language model</article-title>. <source>JMLR</source> <volume>3</volume>, <fpage>1137</fpage>&#x2013;<lpage>1155</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1162/153244303322533223</pub-id>
</citation>
</ref>
<ref id="B3">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Bodansky</surname> <given-names>D.</given-names>
</name>
</person-group> (<year>2011</year>). <article-title>Implementation of international environmental law</article-title>. <source>Jpn. Yearb. Int. Law.</source> <volume>54</volume>, <fpage>62</fpage>&#x2013;<lpage>96</lpage>.</citation>
</ref>
<ref id="B4">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Chen</surname> <given-names>Y. T.</given-names>
</name>
<name>
<surname>Liu</surname> <given-names>H. R.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Critical perspectives on the new situation of global ocean governance</article-title>. <source>Sustainability</source> <volume>15</volume>, <elocation-id>10921</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.3390/su151410921</pub-id>
</citation>
</ref>
<ref id="B5">
<citation citation-type="web">
<person-group person-group-type="author">
<name>
<surname>Chiang</surname> <given-names>W.-L.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Lin</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Sheng</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Wu</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Zhang</surname> <given-names>H.</given-names>
</name>
<etal/>
</person-group>. (<year>2023</year>) <source>Vicuna: an open-source chatbot impressing GPT-4 with 90%* chatGPT quality</source>. Available online at: <uri xlink:href="https://lmsys.org/blog/2023-03-30-vicuna/">https://lmsys.org/blog/2023-03-30-vicuna/</uri> (Accessed <access-date>November 4, 2023</access-date>).</citation>
</ref>
<ref id="B6">
<citation citation-type="web">
<person-group person-group-type="author">
<collab>Common Crawl</collab>
</person-group> (<year>2023</year>). Available online at: <uri xlink:href="https://commoncrawl.org/">https://commoncrawl.org/</uri> (Accessed <access-date>November 4, 2023</access-date>).</citation>
</ref>
<ref id="B7">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Cui</surname> <given-names>J. X.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>Z. J.</given-names>
</name>
<name>
<surname>Yan</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Chen</surname> <given-names>B. H.</given-names>
</name>
<name>
<surname>Yuan</surname> <given-names>L.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>ChatLaw: open-source legal large language model with integrated external knowledge bases</article-title>. <source>arXiv</source>. doi:&#xa0;<pub-id pub-id-type="doi">10.48550/arXiv.2306.16092</pub-id>
</citation>
</ref>
<ref id="B8">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Dai</surname> <given-names>F.</given-names>
</name>
<name>
<surname>Zhao</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Sun</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>B.</given-names>
</name>
</person-group> (<year>2022</year>). &#x201c;<article-title>Intelligent audit question answering system based on knowledge graph and semantic similarity</article-title>,&#x201d; in <conf-name>2022 11th International Conference of Information and Communication Technology (ICTech)</conf-name>, <conf-loc>Wuhan, China: IEEE</conf-loc>. <fpage>125</fpage>&#x2013;<lpage>132</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1109/ICTech55460.2022.00033</pub-id>
</citation>
</ref>
<ref id="B9">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Deasy</surname> <given-names>K.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>What we know about the new high seas treaty</article-title>. <source>NPJ Ocean Sustain.</source> <volume>2</volume>, <fpage>7</fpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1038/s44183-023-00013-x</pub-id>
</citation>
</ref>
<ref id="B10">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Devlin</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Chang</surname> <given-names>M.-W.</given-names>
</name>
<name>
<surname>Lee</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Toutanova</surname> <given-names>K.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>BERT: pretraining of deep bidirectional transformers for language understanding</article-title>. <source>arXiv</source>. doi:&#xa0;<pub-id pub-id-type="doi">10.48550/arXiv.1810.04805</pub-id>
</citation>
</ref>
<ref id="B11">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Dong</surname> <given-names>Q. X.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Dai</surname> <given-names>D. M.</given-names>
</name>
<name>
<surname>Zheng</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Wu</surname> <given-names>Z. Y.</given-names>
</name>
<name>
<surname>Chang</surname> <given-names>B. B.</given-names>
</name>
<etal/>
</person-group>. (<year>2023</year>). <article-title>A survey on in-context learning</article-title>. <source>arXiv</source>. doi:&#xa0;<pub-id pub-id-type="doi">10.48550/arXiv.2301.00234</pub-id>
</citation>
</ref>
<ref id="B12">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Eloundou</surname> <given-names>T.</given-names>
</name>
<name>
<surname>Manning</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Mishkin</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Rock</surname> <given-names>D.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>GPTs are GPTs: an early look at the labor market impact potential of large language models</article-title>. <source>arXiv</source>. doi:&#xa0;<pub-id pub-id-type="doi">10.48550/arXiv.2303.10130</pub-id>
</citation>
</ref>
<ref id="B13">
<citation citation-type="web">
<person-group person-group-type="author">
<collab>FAO</collab>
</person-group> (<year>2016</year>) <source>The state of world fisheries and aquaculture 2016. contributing to food security and nutrition for all</source> (<publisher-loc>Rome, Italy</publisher-loc>: <publisher-name>Food and Agriculture Organization of the United Nations</publisher-name>). Available online at: <uri xlink:href="https://www.fao.org/3/i5555e/i5555e.pdf">https://www.fao.org/3/i5555e/i5555e.pdf</uri> (Accessed <access-date>March 4, 2024</access-date>).</citation>
</ref>
<ref id="B14">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hu</surname> <given-names>E. J.</given-names>
</name>
<name>
<surname>Shen</surname> <given-names>Y. L.</given-names>
</name>
<name>
<surname>Wallis</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Allen-Zhu</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>Y. Z.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>S. A.</given-names>
</name>
<etal/>
</person-group>. (<year>2021</year>). <article-title>LoRA: low-rank adaptation of large language models</article-title>. <source>arXiv</source>. doi:&#xa0;<pub-id pub-id-type="doi">10.48550/arXiv.2106.09685</pub-id>
</citation>
</ref>
<ref id="B15">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Huang</surname> <given-names>Q. Z.</given-names>
</name>
<name>
<surname>Tao</surname> <given-names>M. X.</given-names>
</name>
<name>
<surname>An</surname> <given-names>Z. W.</given-names>
</name>
<name>
<surname>Zhang</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Jiang</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Chen</surname> <given-names>Z. B.</given-names>
</name>
<etal/>
</person-group>. (<year>2023</year>). <article-title>Lawyer LLaMA technical report</article-title>. <source>arXiv</source>. doi:&#xa0;<pub-id pub-id-type="doi">10.48550/arXiv.2305.15062</pub-id>
</citation>
</ref>
<ref id="B16">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Humphries</surname> <given-names>F.</given-names>
</name>
<name>
<surname>Harden-Davies</surname> <given-names>H.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Practical policy solutions for the final stage of BBNJ treaty negotiations</article-title>. <source>Mar. Policy.</source> <volume>122</volume>, <elocation-id>104214</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.marpol.2020.104214</pub-id>
</citation>
</ref>
<ref id="B17">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Humphries</surname> <given-names>F.</given-names>
</name>
<name>
<surname>Rabone</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Jaspars</surname> <given-names>M.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Traceability approaches for marine genetic resources under the proposed ocean (BBNJ) treaty</article-title>. <source>Front. Mar. Sci.</source> <volume>8</volume>. doi:&#xa0;<pub-id pub-id-type="doi">10.3389/fmars.2021.661313</pub-id>
</citation>
</ref>
<ref id="B18">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kaplan</surname> <given-names>J.</given-names>
</name>
<name>
<surname>McCandlish</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Henighan</surname> <given-names>T.</given-names>
</name>
<name>
<surname>Brown</surname> <given-names>T. B.</given-names>
</name>
<name>
<surname>Chess</surname> <given-names>B.</given-names>
</name>
<name>
<surname>Child</surname> <given-names>R.</given-names>
</name>
<etal/>
</person-group>. (<year>2020</year>). <article-title>Scaling laws for neural language models</article-title>. <source>arXiv</source>. doi:&#xa0;<pub-id pub-id-type="doi">10.48550/arXiv.2001.08361</pub-id>
</citation>
</ref>
<ref id="B19">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Lau</surname> <given-names>G. T.</given-names>
</name>
<name>
<surname>Law</surname> <given-names>K. H.</given-names>
</name>
<name>
<surname>Wiederhold</surname> <given-names>G.</given-names>
</name>
</person-group> (<year>2005</year>). &#x201c;<article-title>Legal information retrieval and application to e-rulemaking</article-title>,&#x201d; in <source>Proceedings of the 10th International Conference on Artificial Intelligence and Law, ICAIL&#x2019;05</source> (<publisher-name>Association for Computing Machinery</publisher-name>, <publisher-loc>New York, NY</publisher-loc>), <fpage>146</fpage>&#x2013;<lpage>154</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1145/1165485.1165508</pub-id>
</citation>
</ref>
<ref id="B20">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lee</surname> <given-names>S.-H.</given-names>
</name>
<name>
<surname>Choi</surname> <given-names>S.-W.</given-names>
</name>
<name>
<surname>Lee</surname> <given-names>E.-B.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>A question-answering model based on knowledge graphs for the general provisions of equipment purchase orders for steel plants maintenance</article-title>. <source>Electronics</source> <volume>12</volume>, <elocation-id>2504</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.3390/electronics12112504</pub-id>
</citation>
</ref>
<ref id="B21">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname> <given-names>Y. X.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>Z. H.</given-names>
</name>
<name>
<surname>Zhang</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Dan</surname> <given-names>R. L.</given-names>
</name>
<name>
<surname>Jiang</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Zhang</surname> <given-names>Y.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>ChatDoctor: A medical chat model fine-tuned on a large language model meta-AI (LLaMA) using medical domain knowledge</article-title>. <source>Cureus</source> <volume>15</volume>, <elocation-id>e40895</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.7759/cureus.40895</pub-id>
</citation>
</ref>
<ref id="B22">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname> <given-names>J. L.</given-names>
</name>
<name>
<surname>Zhang</surname> <given-names>Z. S.</given-names>
</name>
<name>
<surname>Zhao</surname> <given-names>H.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Self-prompting large language models for zero-shot open-domain QA</article-title>. <source>arXiv</source>. doi:&#xa0;<pub-id pub-id-type="doi">10.48550/arXiv.2212.08635</pub-id>
</citation>
</ref>
<ref id="B23">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Lukichev</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Kryanina</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Bystrova</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Fenogenova</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Tikhonova</surname> <given-names>M.</given-names>
</name>
</person-group> (<year>2023</year>). &#x201c;<article-title>Parameter-efficient tuning of transformer models for Anglicism detection and substitution in Russian</article-title>,&#x201d; in <conf-name>Proceedings of the International Conference &#x201c;Dialogue 2023.&#x201d;</conf-name> Available at: <uri xlink:href="https://www.dialog-21.ru/media/5911/lukichevdplusetal042.pdf">https://www.dialog-21.ru/media/5911/lukichevdplusetal042.pdf</uri> (Accessed April 3, 2024).</citation>
</ref>
<ref id="B24">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ortu&#xf1;o Crespo</surname> <given-names>G.</given-names>
</name>
<name>
<surname>Dunn</surname> <given-names>D. C.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>A review of the impacts of fisheries on open-ocean ecosystems</article-title>. <source>ICES J. Mar. Sci.</source> <volume>74</volume>, <fpage>2283</fpage>&#x2013;<lpage>2297</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1093/icesjms/fsx084</pub-id>
</citation>
</ref>
<ref id="B25">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ouyang</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Wu</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Jiang</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Almeida</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Wainwright</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Mishkin</surname> <given-names>P.</given-names>
</name>
<etal/>
</person-group>. (<year>2022</year>). <article-title>Training language models to follow instructions with human feedback., in Advances in Neural Information Processing Systems. Curran Associates, Inc., 27730&#x2013;27744</article-title>. Available at: <uri xlink:href="https://proceedings.neurips.cc/paper_files/paper/2022/file/b1efde53be364a73914f58805a001731-Paper-Conference.pdf">https://proceedings.neurips.cc/paper_files/paper/2022/file/b1efde53be364a73914f58805a001731-Paper-Conference.pdf</uri> (Accessed April 3, 2024).</citation>
</ref>
<ref id="B26">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Riloff</surname> <given-names>E.</given-names>
</name>
<name>
<surname>Thelen</surname> <given-names>M.</given-names>
</name>
</person-group> (<year>2000</year>). &#x201c;<article-title>A rule-based question answering system for reading comprehension tests</article-title>,&#x201d; in <conf-name>ANLP-NAACL 2000 workshop: reading comprehension tests as evaluation for computer-based language understanding systems</conf-name>, (<publisher-loc>Seattle, WA, USA</publisher-loc>: <publisher-name>ACL</publisher-name>). Vol. <volume>6</volume>. <fpage>13</fpage>&#x2013;<lpage>19</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.3115/1117595.1117598</pub-id>
</citation>
</ref>
<ref id="B27">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Rosenfeld</surname> <given-names>R.</given-names>
</name>
</person-group> (<year>2000</year>). <article-title>Two decades of statistical language modeling: where do we go from here</article-title>? <source>Proc. IEEE</source> <volume>88</volume>, <fpage>1270</fpage>&#x2013;<lpage>1278</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1109/5.880083</pub-id>
</citation>
</ref>
<ref id="B28">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sala</surname> <given-names>E.</given-names>
</name>
<name>
<surname>Mayorga</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Costello</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Kroodsma</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Palomares</surname> <given-names>M. L. D.</given-names>
</name>
<name>
<surname>Pauly</surname> <given-names>D.</given-names>
</name>
<etal/>
</person-group>. (<year>2018</year>). <article-title>The economics of fishing the high seas</article-title>. <source>Sci. Adv.</source> <volume>4</volume>, <elocation-id>eaat2504</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1126/sciadv.aat2504</pub-id>
</citation>
</ref>
<ref id="B29">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Sun</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>S. H.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>Y. K.</given-names>
</name>
<name>
<surname>Feng</surname> <given-names>S. K.</given-names>
</name>
<name>
<surname>Tian</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Wu</surname> <given-names>H.</given-names>
</name>
<etal/>
</person-group>. (<year>2020</year>). &#x201c;<article-title>ERNIE 2.0: a continual pre-training framework for language understanding</article-title>,&#x201d; in <conf-name>Proceedings of the AAAI Conference on Artificial Intelligence</conf-name>, (<publisher-loc>New York, USA</publisher-loc>: <publisher-name>AAAI</publisher-name>). Vol. <volume>34</volume>. <fpage>8968</fpage>&#x2013;<lpage>8975</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1609/aaai.v34i05.6428</pub-id>
</citation>
</ref>
<ref id="B30">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tessnow-von Wysocki</surname> <given-names>I.</given-names>
</name>
<name>
<surname>Vadrot</surname> <given-names>A. B. M.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>The voice of science on marine biodiversity negotiations: a systematic literature review</article-title>. <source>Front. Mar. Sci.</source> <volume>7</volume>. doi:&#xa0;<pub-id pub-id-type="doi">10.3389/fmars.2020.614282</pub-id>
</citation>
</ref>
<ref id="B31">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tiller</surname> <given-names>R.</given-names>
</name>
<name>
<surname>Mendenhall</surname> <given-names>E.</given-names>
</name>
<name>
<surname>De Santo</surname> <given-names>E. D.</given-names>
</name>
<name>
<surname>Nyman</surname> <given-names>E.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Shake it off: negotiations suspended, but hope simmering, after a lack of consensus at the fifth intergovernmental conference on biodiversity beyond national jurisdiction</article-title>. <source>Mar. Policy.</source> <volume>148</volume>, <elocation-id>105457</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.marpol.2022.105457</pub-id>
</citation>
</ref>
<ref id="B32">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Touvron</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Lavril</surname> <given-names>T.</given-names>
</name>
<name>
<surname>Izacard</surname> <given-names>G.</given-names>
</name>
<name>
<surname>Martinet</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Lachaux</surname> <given-names>M.-A.</given-names>
</name>
<name>
<surname>Lacroix</surname> <given-names>T.</given-names>
</name>
<etal/>
</person-group>. (<year>2023</year>a). <article-title>LLaMA: open and efficient foundation language models</article-title>. <source>arXiv</source>. doi:&#xa0;<pub-id pub-id-type="doi">10.48550/arXiv.2302.13971</pub-id>
</citation>
</ref>
<ref id="B33">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Touvron</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Martin</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Stone</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Albert</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Almahairi</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Babaei</surname> <given-names>Y.</given-names>
</name>
<etal/>
</person-group>. (<year>2023</year>b). <article-title>Llama 2: open foundation and fine-tuned chat models</article-title>. <source>arXiv</source>. doi:&#xa0;<pub-id pub-id-type="doi">10.48550/arXiv.2307.09288</pub-id>
</citation>
</ref>
<ref id="B34">
<citation citation-type="web">
<person-group person-group-type="author">
<collab>United Nations</collab>
</person-group> (<year>2023</year>) <source>Agreement under the united nations convention on the law of the sea on the conservation and sustainable use of marine biological diversity of areas beyond national jurisdiction</source>. Available online at: <uri xlink:href="https://undocs.org/en/A/77/L.82">https://undocs.org/en/A/77/L.82</uri> (Accessed <access-date>August 10, 2023</access-date>).</citation>
</ref>
<ref id="B35">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Vaghefi</surname> <given-names>S. A.</given-names>
</name>
<name>
<surname>Stammbach</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Muccione</surname> <given-names>V.</given-names>
</name>
<name>
<surname>Bingler</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Ni</surname> <given-names>J. W.</given-names>
</name>
<name>
<surname>Kraus</surname> <given-names>M.</given-names>
</name>
<etal/>
</person-group>. (<year>2023</year>). <article-title>ChatClimate: Grounding conversational AI in climate science</article-title>. <source>Commun. Earth Environ.</source> <volume>4</volume>, <fpage>480</fpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1038/s43247-023-01084-x</pub-id>
</citation>
</ref>
<ref id="B36">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wang</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Chen</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Zhu</surname> <given-names>W.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>A survey on curriculum learning</article-title>. <source>Proc. IEEE</source> <volume>44</volume>, <fpage>4555</fpage>&#x2013;<lpage>4576</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1109/TPAMI.34</pub-id>
</citation>
</ref>
<ref id="B37">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wang</surname> <given-names>H. C.</given-names>
</name>
<name>
<surname>Liu</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Xi</surname> <given-names>N. W.</given-names>
</name>
<name>
<surname>Qiang</surname> <given-names>Z. W.</given-names>
</name>
<name>
<surname>Zhao</surname> <given-names>S. D.</given-names>
</name>
<name>
<surname>Qin</surname> <given-names>B.</given-names>
</name>
<etal/>
</person-group>. (<year>2023</year>). <article-title>HuaTuo: tuning LLaMA model with Chinese medical knowledge</article-title>. <source>arXiv</source>. doi:&#xa0;<pub-id pub-id-type="doi">10.48550/arXiv.2304.06975</pub-id>
</citation>
</ref>
<ref id="B38">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wang</surname> <given-names>Z. Z.</given-names>
</name>
<name>
<surname>Yang</surname> <given-names>F. K.</given-names>
</name>
<name>
<surname>Zhao</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Zhang</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Garg</surname> <given-names>M.</given-names>
</name>
<etal/>
</person-group>. (<year>2023</year>). <article-title>Empower large language model to perform better on industrial domain-specific question answering</article-title>. <source>arXiv</source>. doi:&#xa0;<pub-id pub-id-type="doi">10.48550/arXiv.2305.11541</pub-id>
</citation>
</ref>
<ref id="B39">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ward</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Melbourne-Thomas</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Pecl</surname> <given-names>G. T.</given-names>
</name>
<name>
<surname>Evans</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Green</surname> <given-names>M.</given-names>
</name>
<name>
<surname>McCormack</surname> <given-names>P. C.</given-names>
</name>
<etal/>
</person-group>. (<year>2022</year>). <article-title>Safeguarding marine life: conservation of biodiversity and ecosystems</article-title>. <source>Rev. Fish Biol. Fisheries.</source> <volume>32</volume>, <fpage>65</fpage>&#x2013;<lpage>100</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1007/s11160-022-09700-3</pub-id>
</citation>
</ref>
<ref id="B40">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wei</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Bosma</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Zhao</surname> <given-names>V. Y.</given-names>
</name>
<name>
<surname>Guu</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Yu</surname> <given-names>A. W.</given-names>
</name>
<name>
<surname>Lester</surname> <given-names>B.</given-names>
</name>
<etal/>
</person-group>. (<year>2021</year>). <article-title>Finetuned language models are zero-shot learners</article-title>. <source>arXiv</source>. doi:&#xa0;<pub-id pub-id-type="doi">10.48550/arXiv.2109.01652</pub-id>
</citation>
</ref>
<ref id="B41">
<citation citation-type="web">
<person-group person-group-type="author">
<collab>Wikipedia</collab>
</person-group> (<year>2023</year>). Available online at: <uri xlink:href="https://en.wikipedia.org/wiki/MainPage">https://en.wikipedia.org/wiki/MainPage</uri> (Accessed <access-date>November 4, 2023</access-date>).</citation>
</ref>
<ref id="B42">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Xiong</surname> <given-names>H. L.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Zhu</surname> <given-names>Y. T.</given-names>
</name>
<name>
<surname>Zhao</surname> <given-names>Z. H.</given-names>
</name>
<name>
<surname>Liu</surname> <given-names>Y. X.</given-names>
</name>
<name>
<surname>Huang</surname> <given-names>L. L.</given-names>
</name>
<etal/>
</person-group>. (<year>2023</year>). <article-title>DoctorGLM: fine-tuning your Chinese doctor is not a herculean task</article-title>. <source>arXiv</source>. doi:&#xa0;<pub-id pub-id-type="doi">10.48550/arXiv.2304.01097</pub-id>
</citation>
</ref>
<ref id="B43">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yang</surname> <given-names>H. Y.</given-names>
</name>
<name>
<surname>Liu</surname> <given-names>X.-Y.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>C. D.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>FinGPT: open-source financial large language models</article-title>. <source>arXiv</source>. doi:&#xa0;<pub-id pub-id-type="doi">10.2139/ssrn.4489826</pub-id>
</citation>
</ref>
<ref id="B44">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zeng</surname> <given-names>A. H.</given-names>
</name>
<name>
<surname>Liu</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Du</surname> <given-names>Z. X.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>Z. H.</given-names>
</name>
<name>
<surname>Lai</surname> <given-names>H. Y.</given-names>
</name>
<name>
<surname>Ding</surname> <given-names>M.</given-names>
</name>
<etal/>
</person-group>. (<year>2022</year>). <article-title>GLM-130B: an open bilingual pre-trained model</article-title>. <source>arXiv</source>. doi:&#xa0;<pub-id pub-id-type="doi">10.48550/arXiv.2210.02414</pub-id>
</citation>
</ref>
<ref id="B45">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhang</surname> <given-names>T. Y.</given-names>
</name>
<name>
<surname>Kishore</surname> <given-names>V.</given-names>
</name>
<name>
<surname>Wu</surname> <given-names>F. L.</given-names>
</name>
<name>
<surname>Weinberger</surname> <given-names>K. Q.</given-names>
</name>
<name>
<surname>Artzi</surname> <given-names>Y.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>BERTScore: evaluating text generation with BERT</article-title>. <source>arXiv</source>. doi:&#xa0;<pub-id pub-id-type="doi">10.48550/arXiv.1904.09675</pub-id>
</citation>
</ref>
<ref id="B46">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhong</surname> <given-names>B.</given-names>
</name>
<name>
<surname>He</surname> <given-names>W.</given-names>
</name>
<name>
<surname>Huang</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Love</surname> <given-names>P. E. D.</given-names>
</name>
<name>
<surname>Tang</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Luo</surname> <given-names>H.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>A building regulation question answering system: a deep learning methodology</article-title>. <source>Adv. Eng. Inform.</source> <volume>46</volume>, <elocation-id>101195</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.aei.2020.101195</pub-id>
</citation>
</ref>
<ref id="B47">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhu</surname> <given-names>F. B.</given-names>
</name>
<name>
<surname>Lei</surname> <given-names>W. Q.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Zheng</surname> <given-names>J. M.</given-names>
</name>
<name>
<surname>Poria</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Chua</surname> <given-names>T.-S.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Retrieving and reading: a comprehensive survey on open-domain question answering</article-title>. <source>arXiv</source>. doi:&#xa0;<pub-id pub-id-type="doi">10.48550/arXiv.2101.00774</pub-id>
</citation>
</ref>
</ref-list>
</back>
</article>