<?xml version="1.0" encoding="UTF-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article article-type="brief-report" dtd-version="2.3" xml:lang="EN" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Mater.</journal-id>
<journal-title>Frontiers in Materials</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Mater.</abbrev-journal-title>
<issn pub-type="epub">2296-8016</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="publisher-id">1364572</article-id>
<article-id pub-id-type="doi">10.3389/fmats.2024.1364572</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Materials</subject>
<subj-group>
<subject>Perspective</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>A framework for computer-aided high performance titanium alloy design based on machine learning</article-title>
<alt-title alt-title-type="left-running-head">An et al.</alt-title>
<alt-title alt-title-type="right-running-head">
<ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3389/fmats.2024.1364572">10.3389/fmats.2024.1364572</ext-link>
</alt-title>
</title-group>
<contrib-group>
<contrib contrib-type="author">
<name>
<surname>An</surname>
<given-names>Suyang</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2616912/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Li</surname>
<given-names>Kun</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<xref ref-type="aff" rid="aff4">
<sup>4</sup>
</xref>
<xref ref-type="corresp" rid="c001">&#x2a;</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1444910/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Zhu</surname>
<given-names>Liang</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff4">
<sup>4</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2631466/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Liang</surname>
<given-names>Haisong</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff4">
<sup>4</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Ma</surname>
<given-names>Ruijin</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff4">
<sup>4</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Liao</surname>
<given-names>Ruobing</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff4">
<sup>4</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Murr</surname>
<given-names>Lawrence E.</given-names>
</name>
<xref ref-type="aff" rid="aff5">
<sup>5</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
</contrib-group>
<aff id="aff1">
<sup>1</sup>
<institution>College of Mechanical and Vehicle Engineering</institution>, <institution>Chongqing University</institution>, <addr-line>Chongqing</addr-line>, <country>China</country>
</aff>
<aff id="aff2">
<sup>2</sup>
<institution>AVIC Guizhou Aircraft Corporation LTD.</institution>, <addr-line>Anshun</addr-line>, <addr-line>Guizhou</addr-line>, <country>China</country>
</aff>
<aff id="aff3">
<sup>3</sup>
<institution>State Key Laboratory of Mechanical Transmission for Advanced Equipment</institution>, <institution>Chongqing University</institution>, <addr-line>Chongqing</addr-line>, <country>China</country>
</aff>
<aff id="aff4">
<sup>4</sup>
<institution>Chongqing Key Laboratory of Metal Additive Manufacturing (3D Printing)</institution>, <institution>Chongqing University</institution>, <addr-line>Chongqing</addr-line>, <country>China</country>
</aff>
<aff id="aff5">
<sup>5</sup>
<institution>W.M. Keck Center for 3D Innovation</institution>, <institution>University of Texas at El Paso</institution>, <addr-line>El Paso</addr-line>, <addr-line>TX</addr-line>, <country>United States</country>
</aff>
<author-notes>
<fn fn-type="edited-by">
<p>
<bold>Edited by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1737784/overview">Shaoping Xiao</ext-link>, The University of Iowa, United States</p>
</fn>
<fn fn-type="edited-by">
<p>
<bold>Reviewed by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1966821/overview">Lei Yang</ext-link>, Wuhan University of Technology, China</p>
<p>
<ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/2067585/overview">Dafan Du</ext-link>, Shanghai Jiao Tong University, China</p>
<p>
<ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/2620713/overview">Ji Yin</ext-link>, China University of Geosciences Wuhan, China</p>
<p>
<ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1168440/overview">Jun Cheng</ext-link>, Northwest Institute For Non-Ferrous Metal Research, China</p>
<p>
<ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/333558/overview">Chao Yang</ext-link>, South China University of Technology, China</p>
</fn>
<corresp id="c001">&#x2a;Correspondence: Kun Li, <email>kun.li@cqu.edu.cn</email>
</corresp>
</author-notes>
<pub-date pub-type="epub">
<day>16</day>
<month>04</month>
<year>2024</year>
</pub-date>
<pub-date pub-type="collection">
<year>2024</year>
</pub-date>
<volume>11</volume>
<elocation-id>1364572</elocation-id>
<history>
<date date-type="received">
<day>02</day>
<month>01</month>
<year>2024</year>
</date>
<date date-type="accepted">
<day>27</day>
<month>03</month>
<year>2024</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#xa9; 2024 An, Li, Zhu, Liang, Ma, Liao and Murr.</copyright-statement>
<copyright-year>2024</copyright-year>
<copyright-holder>An, Li, Zhu, Liang, Ma, Liao and Murr</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<abstract>
<p>Titanium alloy exhibits exceptional performance and a wide range of applications, with the high performance serving as the foundation for the development. However, traditional material design methods encounter numerous calculations and experimental trial-and-error processes, leading to increased costs and decreased efficiency in material design. The data-driven model presents an intriguing alternative to traditional material design methods by offering a novel approach to expedite the materials design process. In this study, a framework for computer-aided design high performance titanium alloys based on machine learning is proposed, which constructs an intelligent search space encompassing various combinations of 18 elements to facilitate alloy design. Firstly, a proprietary dataset was constructed for titanium alloy materials using feature design and a combination of unsupervised and supervised feature engineering methods. Secondly, six machine learning algorithms were employed to establish regression models, and the hyperparameters of each algorithm were optimized to improve model performance. Thirdly, the model was screened using five regression algorithm evaluation methods. The results demonstrated that the selected optimized model achieved an <italic>R</italic>
<sup>2</sup> value of 0.95 on the verification set and 0.93 on the test set, yielding satisfactory outcomes. Finally, a comprehensive model framework along with an intelligent search methodology for designing high-strength titanium alloys has been established. It is believed that this method is also applicable to other properties of titanium alloys and the optimization of other materials.</p>
</abstract>
<kwd-group>
<kwd>titanium alloy</kwd>
<kwd>machine learning</kwd>
<kwd>data-driven</kwd>
<kwd>material design</kwd>
<kwd>feature engineering</kwd>
</kwd-group>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-at-acceptance</meta-name>
<meta-value>Mechanics of Materials</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec id="s1">
<title>1 Introduction</title>
<p>The application of titanium-containing alloys, such as titanium alloy, titanium-niobium alloy, and high entropy alloy, has been extensively observed in the fields of aerospace, navigation, and medicine (<xref ref-type="bibr" rid="B5">Cheng et al., 2021</xref>; <xref ref-type="bibr" rid="B6">Cheng et al., 2022</xref>; <xref ref-type="bibr" rid="B12">Guo et al., 2023</xref>; <xref ref-type="bibr" rid="B20">Liu et al., 2023</xref>; <xref ref-type="bibr" rid="B17">Kang et al., 2024</xref>; <xref ref-type="bibr" rid="B38">Shen et al., 2024</xref>). Among these alloys, titanium alloy stands out due to its exceptional specific strength, corrosion resistance, low-temperature tolerance, high-temperature endurance, and remarkable biocompatibility. In the aerospace field, titanium alloy has been mainly used in aircraft structural components, lips, tubes, fasteners, satellite shells, rocket tubes, and rocket engine shells, etc., especially, the proportion of titanium alloy in the structure of the fifth-generation advanced fighter F-22 in the United States has reached 41% (<xref ref-type="bibr" rid="B1">Boyer, 1995</xref>; <xref ref-type="bibr" rid="B21">Liu et al., 2015</xref>; <xref ref-type="bibr" rid="B22">Liu et al., 2020</xref>). In the field of navigation, titanium alloy has been used in ship structural parts, submarine shells, marine water pipelines, etc. (<xref ref-type="bibr" rid="B4">Chen et al., 2005</xref>; <xref ref-type="bibr" rid="B40">Song et al., 2020</xref>). Moreover, in the medical field, titanium alloy serves as a crucial material for artificial substitutes or implants including joints, craniofacial structures, and dental implants (<xref ref-type="bibr" rid="B13">Hanawa, 2019</xref>; <xref ref-type="bibr" rid="B23">Louren&#xe7;o et al., 2020</xref>; <xref ref-type="bibr" rid="B36">Sarraf et al., 2022</xref>). The demand for new high-performance titanium alloys is increasing due to wide-ranging applications, particularly in the aerospace field where high strength and toughness are emphasized, the marine field where corrosion resistance is prioritized, and the biomedical field where a high elastic modulus is sought after.</p>
<p>There are two approaches in traditional material design, namely, manual design and structural search. Manual design involves the intuitive creation of new materials based on expert knowledge and experience, while structural search entails designing novel materials through structure and calculation methods such as combination experiments, phase diagram calculations (CALPHAD), and density functional theory (DFT) (<xref ref-type="bibr" rid="B15">Ji et al., 2014</xref>; <xref ref-type="bibr" rid="B24">Mao et al., 2017</xref>; <xref ref-type="bibr" rid="B30">Rao et al., 2022a</xref>; <xref ref-type="bibr" rid="B42">Tian et al., 2022</xref>; <xref ref-type="bibr" rid="B41">Song et al., 2023</xref>). However, these traditional methods require extensive calculations and experimental trial-and-error processes with high demands for expertise from material designers, resulting in elevated development costs and low efficiency.</p>
<p>In the era of rapid advances in data-driven and artificial intelligence technologies, computer-aided design (CAD) of novel materials has become feasible (<xref ref-type="bibr" rid="B33">Ren et al., 2018</xref>; <xref ref-type="bibr" rid="B48">Yu et al., 2019</xref>; <xref ref-type="bibr" rid="B7">Deng et al., 2020</xref>; <xref ref-type="bibr" rid="B43">Wahl et al., 2021</xref>; <xref ref-type="bibr" rid="B31">Rao et al., 2022b</xref>; <xref ref-type="bibr" rid="B11">Giles et al., 2022</xref>; <xref ref-type="bibr" rid="B18">Jiang et al., 2022</xref>; <xref ref-type="bibr" rid="B16">Kandavalli et al., 2023</xref>; <xref ref-type="bibr" rid="B19">Li et al., 2023</xref>; <xref ref-type="bibr" rid="B37">Sasidhar et al., 2023</xref>; <xref ref-type="bibr" rid="B45">Wei et al., 2023</xref>), particularly in the realm of high-performance alloys. Lei et al. employed a performance-oriented machine learning design strategy to swiftly discover a new aluminum alloy that exhibits ductility and toughness indexes comparable to the state-of-the-art AA7136 aluminum alloy (<xref ref-type="bibr" rid="B18">Jiang et al., 2022</xref>). <xref ref-type="bibr" rid="B11">Giles et al. (2022)</xref> expedited the exploration of high-entropy alloys with exceptional high-temperature yield strength through an intelligent machine learning model for searching such alloys. <xref ref-type="bibr" rid="B7">Deng et al. (2020)</xref> have promptly identified Cu-Al alloys with tensile strength exceeding 350 MPa by employing six different machine learning algorithms. These breakthroughs challenge traditional design concepts as they eliminate the need for extensive knowledge, experience, calculations, and trial-and-error experiments while effectively enhancing design efficiency and reducing costs. This approach represents a novel alternative to traditional material design methods in CAD for high-performance alloy materials where model performance is paramount. It relies on careful selection and extraction of relevant alloy feature parameters, sample size considerations, and choice of appropriate machine learning algorithms; however, it also confronts several challenges. In terms of feature parameter selection and extraction, it mainly focuses on alloy composition with complex descriptors, such as atomic size mismatch or enthalpy of mixing in the design of high-entropy alloys (<xref ref-type="bibr" rid="B11">Giles et al., 2022</xref>). While these complex descriptors can enhance model performance significantly, their universal applicability remains limited due to computational complexities involved. Furthermore, there still exists a scarcity of samples available for analysis, and even an excellent model achieving <italic>R</italic>
<sup>2</sup> value as high as 0.94 was trained using only 177 samples (<xref ref-type="bibr" rid="B18">Jiang et al., 2022</xref>). In machine learning algorithms, the primary focus lies in algorithm selection, which needs to improve the performance of algorithms for specific application scenarios. In computer-aided design of titanium alloys, it becomes applicable to extract the composition characteristics of alloys and establish models through algorithm screening. However, challenges exist in obtaining descriptor characteristic parameters of titanium alloys, acquiring a larger sample size of titanium alloys, and selecting superior machine learning algorithms specifically tailored for titanium alloys. The research primarily emphasizes fatigue damage analysis and life prediction, low-modulus titanium alloy prediction, manufacturing defect identification, etc. (<xref ref-type="bibr" rid="B46">Wu et al., 2021</xref>; <xref ref-type="bibr" rid="B49">Zhan et al., 2021</xref>; <xref ref-type="bibr" rid="B9">Fotovvati and Chou, 2022</xref>; <xref ref-type="bibr" rid="B47">Wu et al., 2022</xref>; <xref ref-type="bibr" rid="B44">Wang et al., 2023</xref>). This study specifically focuses on the demand for high-strength titanium alloys in the aerospace industry, with limited prior reports on computer-aided such alloys design based on machine learning.</p>
<p>In this study, elemental essential properties is extracted to achieve complex alloy descriptors effectively while simplifying the extraction process and enhancing versatility. Furthermore, the influence of heat treatment is considered on alloy properties by extracting characteristic parameters related to heat treatment. This establishes a novel approach for selecting and extracting characteristic parameters based on alloy elements, essential properties as well as heat treatment systems. Regarding sample size, data sets of significant magnitudes were constructed by comprehensively reviewing a substantial body of literature. In terms of machine learning algorithms, six classical models were adopted, namely, support vector machine, Gaussian process, neural network, CART, boosting tree and random forest regression. These algorithms were further optimized through hyperparameter tuning to enhance model performance. To systematically evaluate the regression model and facilitate model selection, which is conducted using five metrics: root mean squared error (RMSE), mean absolute error (MAE), mean squared error (MSE), coefficient of determination(<italic>R</italic>
<sup>2</sup>) and training time. Consequently, a comprehensive model framework was established. Moreover, an intelligent search space encompassing 18 elements such as Ti, Al, Sn, Mo, V, Mn, Zr, Ni, Si, Nd, B, Cu, Fe, Nb, C, Cr, Y, W was formulated. Accordingly, the second section presents the construction of the dataset, algorithmic model, optimization process, and model evaluation method. The third section provides the results of the model algorithm and selects the optimal model based on thorough evaluation. Furthermore, model verification is conducted using a test set. Finally, a comprehensive framework for the complete model is presented. The fourth section provides a general discussion.</p>
</sec>
<sec sec-type="materials|methods" id="s2">
<title>2 Material and methods</title>
<p>In this study, the framework comprises three sections: feature engineering based on titanium alloy, machine learning algorithm model, and the evaluation and selection of models. Specifically, feature engineering involves the extraction of meaningful features from raw data.</p>
<sec id="s2-1">
<title>2.1 Feature engineering based on titanium alloy</title>
<sec id="s2-1-1">
<title>2.1.1 Collection of original data</title>
<p>Original data consists of two aspects: one aspect is the national standard &#x201c;Designation and composition of titanium and titanium alloys&#x201d; (GB/T 3620.1-2016), which provides information on titanium alloy grades and their chemical composition. The other aspect involves extensive literature, where 66 relevant sources are carefully selected, primarily focusing on forged or rolled bars. The data comprises 60 titanium alloys, encompassing 18 elements, as shown in <xref ref-type="sec" rid="s10">Supplementary Appendix SA1</xref>. To represent these alloys effectively, the proportion of each element is considered as a feature set consisting of 18 dimensions.</p>
<p>Meanwhile, heat treatment has a significant effect on the microstructure and properties of titanium alloys. In this study, two heat treatment systems are retained from selected literature, encompassing parameters such as initial heat treatment temperature and duration, subsequent heat treatment temperature and duration. The employed heat treatment system encompasses solution, aging, and all their possible combinations. Consequently, a total of four distinct dimensions are considered in this analysis. In cases where no heat treatment or only one type of heat treatment is applied, data points without any corresponding treatments are assigned a value of zero.</p>
<p>The original dataset comprised a total of 397 samples across 22 dimensions. In this study, the model predicts the ultimate strength as the performance indicator for titanium alloy, and the corresponding ultimate strength values were collected for each sample.</p>
</sec>
<sec id="s2-1-2">
<title>2.1.2 Process of feature design</title>
<p>The essential properties of the elements in the alloy play a crucial role in determining its performance, thus highlighting their significance. In this study, feature engineering is employed to design the fundamental characteristics of titanium alloy, encompassing parameters such as melting point, density, atomic weight, atomic number, electronegativity, and atomic radius. However, considering that there are 18 elements in the dataset and each element corresponds to six essential characteristics, this leads to a total of 108 features or dimensions, resulting in a dimensional challenge for the dataset. To address this issue effectively while preserving relevant information integrity, a weighted summation method is adopted to reduce these 108 dimensions down to six dimensions. Consequently, each titanium alloy is associated with a set of essential features comprising weighted values for key attributes including melting point (Tm), density (&#x3c1;), atomic weight (u), atomic number (Z), electronegativity (X), and atomic radius (R). These weighted values are calculated using <xref ref-type="disp-formula" rid="e1">Formulas 1</xref>&#x2013;<xref ref-type="disp-formula" rid="e5">5</xref> through <xref ref-type="disp-formula" rid="e6">6</xref>.<disp-formula id="e1">
<mml:math id="m1">
<mml:mrow>
<mml:mtext>Tm</mml:mtext>
<mml:mo>&#x3d;</mml:mo>
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:msubsup>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi mathvariant="normal">i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi mathvariant="normal">n</mml:mi>
</mml:msubsup>
</mml:mstyle>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">T</mml:mi>
<mml:msub>
<mml:mi mathvariant="normal">m</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
<mml:mo>&#xd7;</mml:mo>
<mml:msub>
<mml:mi mathvariant="normal">w</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(1)</label>
</disp-formula>
<disp-formula id="e2">
<mml:math id="m2">
<mml:mrow>
<mml:mi mathvariant="normal">&#x3c1;</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:msubsup>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi mathvariant="normal">i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi mathvariant="normal">n</mml:mi>
</mml:msubsup>
</mml:mstyle>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">&#x3c1;</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
<mml:mo>&#xd7;</mml:mo>
<mml:msub>
<mml:mi mathvariant="normal">w</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(2)</label>
</disp-formula>
<disp-formula id="e3">
<mml:math id="m3">
<mml:mrow>
<mml:mi mathvariant="normal">u</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:msubsup>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi mathvariant="normal">i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi mathvariant="normal">n</mml:mi>
</mml:msubsup>
</mml:mstyle>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">u</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
<mml:mo>&#xd7;</mml:mo>
<mml:msub>
<mml:mi mathvariant="normal">w</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(3)</label>
</disp-formula>
<disp-formula id="e4">
<mml:math id="m4">
<mml:mrow>
<mml:mi mathvariant="normal">Z</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:msubsup>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi mathvariant="normal">i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi mathvariant="normal">n</mml:mi>
</mml:msubsup>
</mml:mstyle>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">Z</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
<mml:mo>&#xd7;</mml:mo>
<mml:msub>
<mml:mi mathvariant="normal">w</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(4)</label>
</disp-formula>
<disp-formula id="e5">
<mml:math id="m5">
<mml:mrow>
<mml:mi mathvariant="normal">X</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:msubsup>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi mathvariant="normal">i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi mathvariant="normal">n</mml:mi>
</mml:msubsup>
</mml:mstyle>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">X</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
<mml:mo>&#xd7;</mml:mo>
<mml:msub>
<mml:mi mathvariant="normal">w</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(5)</label>
</disp-formula>
<disp-formula id="e6">
<mml:math id="m6">
<mml:mrow>
<mml:mi mathvariant="normal">R</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:msubsup>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi mathvariant="normal">i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi mathvariant="normal">n</mml:mi>
</mml:msubsup>
</mml:mstyle>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">r</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
<mml:mo>&#xd7;</mml:mo>
<mml:msub>
<mml:mi mathvariant="normal">w</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(6)</label>
</disp-formula>Where <inline-formula id="inf1">
<mml:math id="m7">
<mml:mrow>
<mml:mi mathvariant="normal">T</mml:mi>
<mml:msub>
<mml:mi mathvariant="normal">m</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> represents the melting point value of the <italic>i</italic>th element, and <inline-formula id="inf2">
<mml:math id="m8">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">w</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> represents the weight of the melting point for that specific element. The value denotes the content proportion of each element in titanium alloy, with a total of 18 elements (<italic>n</italic> &#x3d; 18). Similar formulations are used for others. <xref ref-type="fig" rid="F1">Figure 1</xref> presents the distribution of physical property constants in titanium alloy. The values of constants are presented in <xref ref-type="sec" rid="s10">Supplementary Appendix AS2</xref>.</p>
<fig id="F1" position="float">
<label>FIGURE 1</label>
<caption>
<p>Distribution of physical property constants in titanium alloy: <bold>(A)</bold> Tm; <bold>(B)</bold> <inline-formula id="inf3">
<mml:math id="m9">
<mml:mrow>
<mml:mi mathvariant="normal">&#x3c1;</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>; <bold>(C)</bold> u; <bold>(D)</bold> Z; <bold>(E)</bold> X; <bold>(F)</bold> R.</p>
</caption>
<graphic xlink:href="fmats-11-1364572-g001.tif"/>
</fig>
<p>Through the process of original data collection and feature design, the dataset consists of a total of 28 features, that is, 28-dimensional. The detailed information about these features is presented in <xref ref-type="table" rid="T1">Table 1</xref>.</p>
<table-wrap id="T1" position="float">
<label>TABLE 1</label>
<caption>
<p>Dataset features.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="center">No.</th>
<th align="center">Feature</th>
<th align="center">Significance</th>
<th align="center">No.</th>
<th align="center">Feature</th>
<th align="center">Significance</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="center">1</td>
<td align="center">Tm</td>
<td align="center">Weighted value of the melting point/&#xb0;C</td>
<td align="center">15</td>
<td align="center">Nd</td>
<td align="center">Nd content/wt%</td>
</tr>
<tr>
<td align="center">2</td>
<td align="center">
<inline-formula id="inf4">
<mml:math id="m10">
<mml:mrow>
<mml:mi mathvariant="normal">&#x3c1;</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="center" style="color:#2A2B2E">Weighted value of the density/g/cm&#xb3;</td>
<td align="center">16</td>
<td align="center">B</td>
<td align="center">B content/wt%</td>
</tr>
<tr>
<td align="center">3</td>
<td align="center">u</td>
<td align="center" style="color:#2A2B2E">Weighted value of the atomic weight</td>
<td align="center">17</td>
<td align="center">Cu</td>
<td align="center">Cu content/wt%</td>
</tr>
<tr>
<td align="center">4</td>
<td align="center">Z</td>
<td align="center" style="color:#2A2B2E">Weighted value of the atomic number</td>
<td align="center">18</td>
<td align="center">Fe</td>
<td align="center">Fe content/wt%</td>
</tr>
<tr>
<td align="center">5</td>
<td align="center">X</td>
<td align="center" style="color:#2A2B2E">Weighted value of the electronegativity/Pauling scale</td>
<td align="center">19</td>
<td align="center">Nb</td>
<td align="center">Nb content/wt%</td>
</tr>
<tr>
<td align="center">6</td>
<td align="center">R</td>
<td align="center" style="color:#2A2B2E">Weighted value of the atomic radius/pm</td>
<td align="center">20</td>
<td align="center">C</td>
<td align="center">C content/wt%</td>
</tr>
<tr>
<td align="center">7</td>
<td align="center">Al</td>
<td align="center">Al content/wt%</td>
<td align="center">21</td>
<td align="center">Cr</td>
<td align="center">Cr content/wt%</td>
</tr>
<tr>
<td align="center">8</td>
<td align="center">Sn</td>
<td align="center">Sn content/wt%</td>
<td align="center">22</td>
<td align="center">Y</td>
<td align="center">Y content/wt%</td>
</tr>
<tr>
<td align="center">9</td>
<td align="center">Mo</td>
<td align="center">Mo content/wt%</td>
<td align="center">23</td>
<td align="center">W</td>
<td align="center">W content/wt%</td>
</tr>
<tr>
<td align="center">10</td>
<td align="center">V</td>
<td align="center">V content/wt%</td>
<td align="center">24</td>
<td align="center">Ti</td>
<td align="center">Ti content/wt%</td>
</tr>
<tr>
<td align="center">11</td>
<td align="center">Mn</td>
<td align="center">Mn content/wt%</td>
<td align="center">25</td>
<td align="center">ST</td>
<td align="center">Temperature of heat treatment 1/&#xb0;C</td>
</tr>
<tr>
<td align="center">12</td>
<td align="center">Zr</td>
<td align="center">Zr content/wt%</td>
<td align="center">26</td>
<td align="center">SH</td>
<td align="center">Duration of heat treatment 1/h</td>
</tr>
<tr>
<td align="center">13</td>
<td align="center">Ni</td>
<td align="center">Ni content/wt%</td>
<td align="center">27</td>
<td align="center">AT</td>
<td align="center">Temperature of heat treatment 2/&#xb0;C</td>
</tr>
<tr>
<td align="center">14</td>
<td align="center">Si</td>
<td align="center">Si content/wt%</td>
<td align="center">28</td>
<td align="center">AH</td>
<td align="center">Duration of heat treatment 2/h</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec id="s2-1-3">
<title>2.1.3 Process of feature selection</title>
<p>In this study, feature selection encompasses both supervised and unsupervised analysis. Within the realm of unsupervised analysis, feature correlation analysis is employed to calculate the correlation coefficients between any given features X and Y, as depicted in <xref ref-type="disp-formula" rid="e7">Formula 7</xref>.<disp-formula id="e7">
<mml:math id="m11">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">r</mml:mi>
<mml:mrow>
<mml:mi mathvariant="normal">X</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi mathvariant="normal">Y</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mtext>Cov</mml:mtext>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">X</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi mathvariant="normal">Y</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
<mml:mrow>
<mml:msqrt>
<mml:mrow>
<mml:mi mathvariant="normal">D</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">X</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:msqrt>
<mml:mo>&#xd7;</mml:mo>
<mml:msqrt>
<mml:mrow>
<mml:mi mathvariant="normal">D</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">Y</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:msqrt>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:math>
<label>(7)</label>
</disp-formula>Where Cov (X,Y) denotes the covariance between features X and Y, D(X) represents the variance of feature X, and D(Y) represents the variance of feature Y.</p>
<p>In the supervised analysis, the Minimum Redundancy Maximum Relevance (MRMR) algorithm is employed to compute the mutual information between the feature set and ultimate strength, quantifying feature correlation and ranking them accordingly (<xref ref-type="bibr" rid="B29">Peng et al., 2005</xref>). For comprehensive selection in the final feature choice, both supervised and unsupervised analysis outcomes are utilized with model accuracy as the target.</p>
</sec>
<sec id="s2-1-4">
<title>2.1.4 Standardization of data</title>
<p>Features originate from diverse scales and encompass multiple dimensions in the dataset. To mitigate the influence of dimensionality, data necessitates processing. Data processing techniques are commonly categorized into normalization and standardization. In engineering applications, standardization is typically preferred. Within this dataset, the data without heat treatment is imputed with zeros. When normalization is performed [0,1], the feature with heat treatment would be skewed towards extreme values of 1 and 0, failing to accurately represent the overall distribution of samples in the data set. Henceforth, standardization is selected as a data processing approach, such as employing <xref ref-type="disp-formula" rid="e8">Formula 8</xref>, which centers the dataset around a mean value of 0 while maintaining a normal distribution with a standard deviation of 1.<disp-formula id="e8">
<mml:math id="m12">
<mml:mrow>
<mml:mi mathvariant="normal">Z</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mi mathvariant="normal">C</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mi mathvariant="normal">&#x3bc;</mml:mi>
</mml:mrow>
<mml:mi mathvariant="normal">&#x3c3;</mml:mi>
</mml:mfrac>
</mml:mrow>
</mml:math>
<label>(8)</label>
</disp-formula>
</p>
<p>The feature vector Z is represented as the standardized version of the original arbitrary feature vector C, where &#x3bc; represents the mean value of feature vector C, and &#x3c3; represents its standard deviation.</p>
</sec>
<sec id="s2-1-5">
<title>2.1.5 Partitioning of the dataset</title>
<p>The dataset in this research is divided into three parts: the training set, the validation set, and the test set. The training set is utilized for model training, while the validation set serves the purpose of model evaluation, hyperparameter tuning, and mitigating overfitting risks. On the other hand, the test set is exclusively employed for model evaluation and testing.</p>
</sec>
</sec>
<sec id="s2-2">
<title>2.2 Machine learning algorithm model</title>
<p>The model was constructed using six machine learning algorithms in this study, namely, support vector machine regression, Gaussian process regression, neural network regression, CART tree regression, boosting tree regression, and random forest regression.</p>
<sec id="s2-2-1">
<title>2.2.1 Support vector machine regression</title>
<p>Support vector machine regression (SVR) is a classic machine learning algorithm for regression proposed by Drucker et al., in 1997 (<xref ref-type="bibr" rid="B8">Drucker et al., 1997</xref>), and further developed by Smola et al., in 2004, where the theoretical framework and implementation of SVR was presented (<xref ref-type="bibr" rid="B39">Smola and Sch&#xf6;lkopf, 2004</xref>). Similar to the SVM classification algorithm, the hyperplane used for classification often cannot completely divide the sample space, which can lead to overfitting even if it is possible to achieve complete division. In SVR, an error tolerance interval band &#x3b5; is defined, and samples within this band are considered correct predictions without loss calculation. By partitioning the sample space D &#x3d; {(x<sub>1</sub>,y<sub>1</sub>), (x<sub>2</sub>,y<sub>2</sub>),&#x2026; (x<sub>n</sub>,y<sub>n</sub>)}, SVR aims to find an optimal hyperplane that minimizes the discrepancy between f(x) and y, as shown in <xref ref-type="disp-formula" rid="e9">Formula 9</xref>.<disp-formula id="e9">
<mml:math id="m13">
<mml:mrow>
<mml:mi mathvariant="normal">f</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">x</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x3d;</mml:mo>
<mml:msup>
<mml:mi mathvariant="normal">w</mml:mi>
<mml:mi mathvariant="normal">T</mml:mi>
</mml:msup>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2b;</mml:mo>
<mml:mi mathvariant="normal">b</mml:mi>
</mml:mrow>
</mml:math>
<label>(9)</label>
</disp-formula>Where w represents the normal vector indicating the direction of the obtained hyperplane, and b denotes the displacement term representing the distance from the original point of the hyperplane, SVR computes the loss between f(x) and y outside the interval band. The regression problem is formulated in Eq. <xref ref-type="disp-formula" rid="e10">10</xref>.<disp-formula id="e10">
<mml:math id="m14">
<mml:mrow>
<mml:mtable columnalign="center">
<mml:mtr>
<mml:mtd>
<mml:mi>min</mml:mi>
</mml:mtd>
</mml:mtr>
<mml:mtr>
<mml:mtd>
<mml:mrow>
<mml:mi mathvariant="normal">w</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi mathvariant="normal">b</mml:mi>
</mml:mrow>
</mml:mtd>
</mml:mtr>
</mml:mtable>
<mml:mfrac>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:mfrac>
<mml:msup>
<mml:mrow>
<mml:mfenced open="&#x2016;" close="&#x2016;" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">w</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
<mml:mo>&#x2b;</mml:mo>
<mml:mi mathvariant="normal">C</mml:mi>
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:msubsup>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi mathvariant="normal">i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi mathvariant="normal">n</mml:mi>
</mml:msubsup>
</mml:mstyle>
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">l</mml:mi>
<mml:mi mathvariant="normal">&#x3b5;</mml:mi>
</mml:msub>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">f</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi mathvariant="normal">y</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(10)</label>
</disp-formula>
</p>
<p>The regularization constant C &#x3e; 0 is utilized to balance the minimization of the normal vector with the minimization of the error, while <inline-formula id="inf5">
<mml:math id="m15">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">l</mml:mi>
<mml:mi mathvariant="normal">&#x3b5;</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> denotes the error tolerant interval incorporating &#x3b5; insensitive loss function as defined in <xref ref-type="disp-formula" rid="e11">Formula 11</xref>.<disp-formula id="e11">
<mml:math id="m16">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">l</mml:mi>
<mml:mi mathvariant="normal">&#x3b5;</mml:mi>
</mml:msub>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">z</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x3d;</mml:mo>
<mml:mrow>
<mml:mfenced open="{" close="" separators="|">
<mml:mrow>
<mml:mtable columnalign="left">
<mml:mtr>
<mml:mtd>
<mml:mrow>
<mml:mn>0</mml:mn>
<mml:mo>,</mml:mo>
<mml:mtext>if&#x2009;</mml:mtext>
<mml:mrow>
<mml:mfenced open="|" close="|" separators="|">
<mml:mrow>
<mml:mtext>&#x2009;</mml:mtext>
<mml:mi mathvariant="normal">z</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x2264;</mml:mo>
<mml:mi mathvariant="normal">&#x3b5;</mml:mi>
</mml:mrow>
</mml:mtd>
</mml:mtr>
<mml:mtr>
<mml:mtd>
<mml:mrow>
<mml:mrow>
<mml:mfenced open="|" close="|" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">z</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:mi mathvariant="normal">&#x3b5;</mml:mi>
<mml:mo>,</mml:mo>
<mml:mtext>otherwise</mml:mtext>
</mml:mrow>
</mml:mtd>
</mml:mtr>
</mml:mtable>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(11)</label>
</disp-formula>
</p>
<p>By incorporating relaxation variables, the Lagrange multiplier method, and kernel functions, SVR can be mathematically formulated as Eq. <xref ref-type="disp-formula" rid="e12">12</xref>.<disp-formula id="e12">
<mml:math id="m17">
<mml:mrow>
<mml:mi mathvariant="normal">f</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">x</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x3d;</mml:mo>
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:msubsup>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi mathvariant="normal">i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi mathvariant="normal">n</mml:mi>
</mml:msubsup>
</mml:mstyle>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mover accent="true">
<mml:mi mathvariant="normal">&#x3b1;</mml:mi>
<mml:mo>&#x5e;</mml:mo>
</mml:mover>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi mathvariant="normal">&#x3b1;</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
<mml:mi mathvariant="normal">k</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
<mml:mo>,</mml:mo>
<mml:msub>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mi mathvariant="normal">j</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x2b;</mml:mo>
<mml:mi mathvariant="normal">b</mml:mi>
</mml:mrow>
</mml:math>
<label>(12)</label>
</disp-formula>Where <inline-formula id="inf6">
<mml:math id="m18">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">&#x3b1;</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> represents the Lagrange multiplier, the estimate of <inline-formula id="inf7">
<mml:math id="m19">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">&#x3b1;</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> is denoted as <inline-formula id="inf8">
<mml:math id="m20">
<mml:mrow>
<mml:msub>
<mml:mover accent="true">
<mml:mi mathvariant="normal">&#x3b1;</mml:mi>
<mml:mo>&#x5e;</mml:mo>
</mml:mover>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula>, <inline-formula id="inf9">
<mml:math id="m21">
<mml:mrow>
<mml:mi mathvariant="normal">k</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
<mml:mo>,</mml:mo>
<mml:msub>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mi mathvariant="normal">j</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula> denotes the kernel function, which is a fundamental technique in SVR, and refers to the transformation that maps the indivisible features of the original space to a higher-dimensional divisible space. Since the kernel function implicitly defines the feature space, the specific form of feature mapping remains unknown for a given sample space. Therefore, selecting different kernel functions often yields varying performance outcomes for regression models within specific sample spaces; an inappropriate selection may lead to reduced model performance.</p>
<p>The kernel function is considered as a crucial hyperparameter in this research, and the model is optimized by tuning this hyperparameter. Several commonly used kernel functions are selected, including linear, Gaussian, quadratic, and cubic kernels defined in <xref ref-type="disp-formula" rid="e13">Formulas 13</xref>&#x2013;<xref ref-type="disp-formula" rid="e16">16</xref>. By adjusting these kernel functions, the optimal SVR regression model is identified.<disp-formula id="e13">
<mml:math id="m22">
<mml:mrow>
<mml:mi mathvariant="normal">k</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
<mml:mo>,</mml:mo>
<mml:msub>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mi mathvariant="normal">j</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x3d;</mml:mo>
<mml:msubsup>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
<mml:mi mathvariant="normal">T</mml:mi>
</mml:msubsup>
<mml:msub>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mi mathvariant="normal">j</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
<label>(13)</label>
</disp-formula>
<disp-formula id="e14">
<mml:math id="m23">
<mml:mrow>
<mml:mi mathvariant="normal">k</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
<mml:mo>,</mml:mo>
<mml:msub>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mi mathvariant="normal">j</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x3d;</mml:mo>
<mml:mo>&#x2061;</mml:mo>
<mml:mi>exp</mml:mi>
<mml:mtext>&#x2009;</mml:mtext>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:mfrac>
<mml:msup>
<mml:mrow>
<mml:mfenced open="&#x2016;" close="&#x2016;" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mi mathvariant="normal">j</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
<mml:mrow>
<mml:mn>2</mml:mn>
<mml:msup>
<mml:mi mathvariant="normal">&#x3c3;</mml:mi>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(14)</label>
</disp-formula>
<disp-formula id="e15">
<mml:math id="m24">
<mml:mrow>
<mml:mi mathvariant="normal">k</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
<mml:mo>,</mml:mo>
<mml:msub>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mi mathvariant="normal">j</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x3d;</mml:mo>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msubsup>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
<mml:mi mathvariant="normal">T</mml:mi>
</mml:msubsup>
<mml:msub>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mi mathvariant="normal">j</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:math>
<label>(15)</label>
</disp-formula>
<disp-formula id="e16">
<mml:math id="m25">
<mml:mrow>
<mml:mi mathvariant="normal">k</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
<mml:mo>,</mml:mo>
<mml:msub>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mi mathvariant="normal">j</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x3d;</mml:mo>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msubsup>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
<mml:mi mathvariant="normal">T</mml:mi>
</mml:msubsup>
<mml:msub>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mi mathvariant="normal">j</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mn>3</mml:mn>
</mml:msup>
</mml:mrow>
</mml:math>
<label>(16)</label>
</disp-formula>
</p>
</sec>
<sec id="s2-2-2">
<title>2.2.2 Gaussian process regression</title>
<p>Gaussian process regression (GPR) is a Bayesian non-parametric probabilistic regression model that utilizes a kernel function. The application of Gaussian processes for data fitting was first introduced by O&#x2019;Hagan in 1978 (<xref ref-type="bibr" rid="B28">OHagan, 1978</xref>). <xref ref-type="bibr" rid="B32">Rasmussen et al. (2005)</xref> provided a comprehensive theoretical framework for the Gaussian process regression model in machine learning in 2006. The problem of nonlinear Gaussian regression is formulated using Eq. <xref ref-type="disp-formula" rid="e17">17</xref>.<disp-formula id="e17">
<mml:math id="m26">
<mml:mrow>
<mml:mi mathvariant="normal">y</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mi mathvariant="normal">f</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">x</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x2b;</mml:mo>
<mml:mi mathvariant="normal">&#x3b5;</mml:mi>
</mml:mrow>
</mml:math>
<label>(17)</label>
</disp-formula>Where &#x3b5; represents the additive noise term, &#x3b5; &#x223c; (0, <inline-formula id="inf10">
<mml:math id="m27">
<mml:mrow>
<mml:msup>
<mml:mi mathvariant="normal">&#x3c3;</mml:mi>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:math>
</inline-formula>), where <inline-formula id="inf11">
<mml:math id="m28">
<mml:mrow>
<mml:msup>
<mml:mi mathvariant="normal">&#x3c3;</mml:mi>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:math>
</inline-formula> can be estimated from the sample data. The variable y denotes the predicted value corrupted by noise, while f(x) follows a Gaussian distribution as described in Eq. <xref ref-type="disp-formula" rid="e18">18</xref>.<disp-formula id="e18">
<mml:math id="m29">
<mml:mrow>
<mml:mi mathvariant="normal">f</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">x</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x223c;</mml:mo>
<mml:mtext>GP</mml:mtext>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">m</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">x</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>,</mml:mo>
<mml:mi mathvariant="normal">k</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>,</mml:mo>
<mml:msup>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2032;</mml:mo>
</mml:msup>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(18)</label>
</disp-formula>Where m(x) represents the mean function and takes the value 0. The covariance function k(x,x&#x2019;) captures the relationship between two features x and x&#x2019;. In this regression model, Bayesian and maximum likelihood estimation methods are employed to solve for optimal parameters.</p>
<p>The covariance function, denoted as k(x,x&#x2019;), plays a pivotal role in the model and significantly influences its performance. It is also referred to as the kernel function k(x,x&#x2019;). In this study, the kernel function is treat as a hyperparameter and optimize the model by adjusting five different types of kernels: quadratic rational (<xref ref-type="disp-formula" rid="e19">Formula 19</xref>), square exponent (<xref ref-type="disp-formula" rid="e20">Formula 20</xref>), Matern 5/2 (<xref ref-type="disp-formula" rid="e21">Formula 21</xref>), Matern 3/2 (<xref ref-type="disp-formula" rid="e22">Formula 22</xref>), and exponent (<xref ref-type="disp-formula" rid="e23">Formula 23</xref>).<disp-formula id="e19">
<mml:math id="m30">
<mml:mrow>
<mml:mi mathvariant="normal">k</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>,</mml:mo>
<mml:msup>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2032;</mml:mo>
</mml:msup>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x3d;</mml:mo>
<mml:msubsup>
<mml:mi mathvariant="normal">&#x3c3;</mml:mi>
<mml:mi mathvariant="normal">f</mml:mi>
<mml:mn>2</mml:mn>
</mml:msubsup>
<mml:mtext>&#x2009;</mml:mtext>
<mml:mrow>
<mml:mfenced open="[" close="]" separators="|">
<mml:mrow>
<mml:mn>1</mml:mn>
<mml:mo>&#x2b;</mml:mo>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mfrac>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:msup>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2032;</mml:mo>
</mml:msup>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mi mathvariant="normal">T</mml:mi>
</mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:msup>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2032;</mml:mo>
</mml:msup>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
<mml:mrow>
<mml:mn>2</mml:mn>
<mml:mi mathvariant="normal">&#x3b1;</mml:mi>
<mml:msup>
<mml:mi mathvariant="normal">l</mml:mi>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:mi mathvariant="normal">&#x3b1;</mml:mi>
</mml:mrow>
</mml:msup>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(19)</label>
</disp-formula>
</p>
<p>The symbol &#x3b1; denotes the scaling coefficient, <inline-formula id="inf12">
<mml:math id="m31">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">&#x3c3;</mml:mi>
<mml:mi mathvariant="normal">f</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> represents the standard deviation of the sample, and l signifies the size of the feature length.<disp-formula id="e20">
<mml:math id="m32">
<mml:mrow>
<mml:mi mathvariant="normal">k</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>,</mml:mo>
<mml:msup>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2032;</mml:mo>
</mml:msup>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x3d;</mml:mo>
<mml:msubsup>
<mml:mi mathvariant="normal">&#x3c3;</mml:mi>
<mml:mi mathvariant="normal">f</mml:mi>
<mml:mn>2</mml:mn>
</mml:msubsup>
<mml:mo>&#x2061;</mml:mo>
<mml:mi>exp</mml:mi>
<mml:mtext>&#x2009;</mml:mtext>
<mml:mrow>
<mml:mfenced open="[" close="]" separators="|">
<mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:mfrac>
<mml:mfrac>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:msup>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2032;</mml:mo>
</mml:msup>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mi mathvariant="normal">T</mml:mi>
</mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:msup>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2032;</mml:mo>
</mml:msup>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
<mml:msup>
<mml:mi mathvariant="normal">l</mml:mi>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mfrac>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(20)</label>
</disp-formula>
<disp-formula id="e21">
<mml:math id="m33">
<mml:mtable class="cases" columnalign="left">
<mml:mtr>
<mml:mtd>
<mml:mrow>
<mml:mi mathvariant="normal">k</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>,</mml:mo>
<mml:msup>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2032;</mml:mo>
</mml:msup>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x3d;</mml:mo>
<mml:msubsup>
<mml:mi mathvariant="normal">&#x3c3;</mml:mi>
<mml:mi mathvariant="normal">f</mml:mi>
<mml:mn>2</mml:mn>
</mml:msubsup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mn>1</mml:mn>
<mml:mo>&#x2b;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:msqrt>
<mml:mrow>
<mml:mn>5</mml:mn>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:msup>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2032;</mml:mo>
</mml:msup>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mi mathvariant="normal">T</mml:mi>
</mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:msup>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2032;</mml:mo>
</mml:msup>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:msqrt>
</mml:mrow>
<mml:mrow>
<mml:mi mathvariant="normal">l</mml:mi>
</mml:mrow>
</mml:mfrac>
<mml:mo>&#x2b;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mn>5</mml:mn>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:msup>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2032;</mml:mo>
</mml:msup>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
<mml:mi mathvariant="normal">T</mml:mi>
</mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:msup>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2032;</mml:mo>
</mml:msup>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
<mml:mrow>
<mml:mn>3</mml:mn>
<mml:msup>
<mml:mi mathvariant="normal">l</mml:mi>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:mtd>
</mml:mtr>
<mml:mtr>
<mml:mtd columnalign="center">
<mml:mrow>
<mml:mspace width="4em"/>
<mml:mi>exp</mml:mi>
<mml:mtext>&#x2009;</mml:mtext>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:msqrt>
<mml:mrow>
<mml:mn>5</mml:mn>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:msup>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2032;</mml:mo>
</mml:msup>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mi mathvariant="normal">T</mml:mi>
</mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:msup>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2032;</mml:mo>
</mml:msup>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:msqrt>
</mml:mrow>
<mml:mrow>
<mml:mi mathvariant="normal">l</mml:mi>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mspace width="9em"/>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mo>21</mml:mo>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:mtd>
</mml:mtr>
</mml:mtable>
</mml:math>
</disp-formula>
<disp-formula id="e22">
<mml:math id="m34">
<mml:mrow>
<mml:mi mathvariant="normal">k</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>,</mml:mo>
<mml:msup>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2032;</mml:mo>
</mml:msup>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x3d;</mml:mo>
<mml:msubsup>
<mml:mi mathvariant="normal">&#x3c3;</mml:mi>
<mml:mi mathvariant="normal">f</mml:mi>
<mml:mn>2</mml:mn>
</mml:msubsup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mn>1</mml:mn>
<mml:mo>&#x2b;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:msqrt>
<mml:mrow>
<mml:mn>3</mml:mn>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:msup>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2032;</mml:mo>
</mml:msup>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mi mathvariant="normal">T</mml:mi>
</mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:msup>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2032;</mml:mo>
</mml:msup>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:msqrt>
</mml:mrow>
<mml:mrow>
<mml:mi mathvariant="normal">l</mml:mi>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mi>exp</mml:mi>
<mml:mtext>&#x2009;</mml:mtext>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:msqrt>
<mml:mrow>
<mml:mn>3</mml:mn>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:msup>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2032;</mml:mo>
</mml:msup>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mi mathvariant="normal">T</mml:mi>
</mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:msup>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2032;</mml:mo>
</mml:msup>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:msqrt>
</mml:mrow>
<mml:mrow>
<mml:mi mathvariant="normal">l</mml:mi>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(22)</label>
</disp-formula>
<disp-formula id="e23">
<mml:math id="m35">
<mml:mrow>
<mml:mi mathvariant="normal">k</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>,</mml:mo>
<mml:msup>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2032;</mml:mo>
</mml:msup>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x3d;</mml:mo>
<mml:msubsup>
<mml:mi mathvariant="normal">&#x3c3;</mml:mi>
<mml:mi mathvariant="normal">f</mml:mi>
<mml:mn>2</mml:mn>
</mml:msubsup>
<mml:mo>&#x2061;</mml:mo>
<mml:mi>exp</mml:mi>
<mml:mtext>&#x2009;</mml:mtext>
<mml:mrow>
<mml:mfenced open="[" close="]" separators="|">
<mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:msqrt>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:msup>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2032;</mml:mo>
</mml:msup>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mi mathvariant="normal">T</mml:mi>
</mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:msup>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2032;</mml:mo>
</mml:msup>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:msqrt>
</mml:mrow>
<mml:mrow>
<mml:mi mathvariant="normal">l</mml:mi>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(23)</label>
</disp-formula>
</p>
</sec>
<sec id="s2-2-3">
<title>2.2.3 Neural network regression</title>
<p>The artificial neural network is a network formed by connecting artificial neurons according to a specific topology, which originated from the MP model proposed by McCulloch et al., in 1943 (<xref ref-type="bibr" rid="B26">McCulloch and Pitts, 1943</xref>), as well as the perceptron neural network model introduced by Rosenblatt in 1958 (<xref ref-type="bibr" rid="B34">Rosenblatt, 1958</xref>). These models marked the beginning of the development boom for artificial neural networks. However, it was proven in 1969 that perceptrons were incapable of solving higher-order predicates, such as the XOR problem (<xref ref-type="bibr" rid="B27">Minsky et al., 1969</xref>). Consequently, the field of artificial neural networks experienced a decline until Hopfield&#x2019;s proposal of the Hopfield neural network model in 1982 and Rumelhart et al.&#x2019;s introduction of the BP algorithm in 1986 (<xref ref-type="bibr" rid="B14">Hopfield, 1982</xref>; <xref ref-type="bibr" rid="B25">Mcclelland et al., 1986</xref>; <xref ref-type="bibr" rid="B35">Rumelhart et al., 1986</xref>), which sparked a research upsurge. In this study, a backpropagation (BP) neural network is employed for regression (NNR), where each neuron&#x2019;s threshold was set as &#x3b8; &#x3d; (&#x3b8;&#x2081;, &#x3b8;&#x2082;,&#x2026; &#x3b8;<sub>m</sub>). Here, m represents the number of non-input layer neurons within the neural network structure. The current output Y<sub>j</sub> of neuron j is expressed using <xref ref-type="disp-formula" rid="e24">Formula 24</xref>.<disp-formula id="e24">
<mml:math id="m36">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">Y</mml:mi>
<mml:mi mathvariant="normal">j</mml:mi>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mi mathvariant="normal">f</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:msubsup>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi mathvariant="normal">i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi mathvariant="normal">k</mml:mi>
</mml:msubsup>
</mml:mstyle>
<mml:msub>
<mml:mi mathvariant="normal">w</mml:mi>
<mml:mtext>ij</mml:mtext>
</mml:msub>
<mml:msub>
<mml:mi mathvariant="normal">X</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi mathvariant="normal">&#x3b8;</mml:mi>
<mml:mi mathvariant="normal">j</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(24)</label>
</disp-formula>Where f is the activation function, i&#x3d;(1,2, &#x2026; k), j&#x3d;(1,2, &#x2026; ,m). The value of k corresponds to the number of neurons located above the current neuron j. The weight w<sub>ij</sub> denotes the synaptic connection strength from the <italic>i</italic>th neuron in the preceding layer to the <italic>j</italic>th neuron, X<sub>i</sub> represents the output value of the <italic>i</italic>th neuron in the preceding layer, and &#x3b8;<sub>j</sub> signifies the activation threshold of the current neuron.</p>
<p>The weights and thresholds of the BP neural network are iteratively calculated and updated using the gradient descent strategy, leading to the attainment of optimal model values. In this study, the neural network regression model is empirically defined as depicted in <xref ref-type="fig" rid="F2">Figure 2A</xref>, comprising one input layer, one output layer, and three hidden layers. The input layer consists of 28 neurons, while the output layer comprises a single neuron; each hidden layer encompasses 10 neurons.</p>
<fig id="F2" position="float">
<label>FIGURE 2</label>
<caption>
<p>The neural network regression model structure: <bold>(A)</bold> the neural network regression model based on experience; <bold>(B)</bold> the neural network regression model based on hyperparameter optimization.</p>
</caption>
<graphic xlink:href="fmats-11-1364572-g002.tif"/>
</fig>
<p>The neural network regression model encompasses numerous crucial hyperparameters, which are devised based on empirical knowledge that may not be deemed as ideal or relatively optimal. Hence, the optimal hyperparameters are fine-tuned by adjusting the parameters of the neural network model and iterating through Bayesian optimization to minimize mean squared error. The range for parameter adjustment is presented in <xref ref-type="table" rid="T2">Table 2</xref>. The regularization intensity value of 378 corresponds to the aggregate number of training and validation samples.</p>
<table-wrap id="T2" position="float">
<label>TABLE 2</label>
<caption>
<p>The range of hyperparameter optimization.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="center">Regression model</th>
<th align="center">Hyperparameter</th>
<th align="center">Range</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td rowspan="6" align="center">NNR</td>
<td align="center">Number of hidden layer</td>
<td align="center">{1,2,3}</td>
</tr>
<tr>
<td align="center">Activation function</td>
<td align="center">{Relu, Tanh,None, Sigmoid}</td>
</tr>
<tr>
<td align="center">Number of neurons in hidden layer 1</td>
<td align="center">[1,300]</td>
</tr>
<tr>
<td align="center">Number of neurons in hidden layer 2</td>
<td align="center">[1,300]</td>
</tr>
<tr>
<td align="center">Number of neurons in hidden layer 3</td>
<td align="center">[1,300]</td>
</tr>
<tr>
<td align="center">Intensity of regularization</td>
<td align="center">[1e<sup>&#x2212;5</sup>,1e<sup>5</sup>]/378</td>
</tr>
<tr>
<td rowspan="4" align="center">ER</td>
<td align="center">Ensemble learning method</td>
<td align="center">{Bagging, Boosting}</td>
</tr>
<tr>
<td align="center">Number of learners</td>
<td align="center">[10,500]</td>
</tr>
<tr>
<td align="center">Learning rate</td>
<td align="center">[1e<sup>&#x2212;3</sup>,1]</td>
</tr>
<tr>
<td align="center">Number of sampled predictors</td>
<td align="center">[1,28]</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec id="s2-2-4">
<title>2.2.4 CART regression</title>
<p>The CART algorithm, proposed by Breiman et al., in 1984 (<xref ref-type="bibr" rid="B3">Breiman et al., 1984</xref>), is a classic machine learning algorithm. It comprises two processes: the generation of decision trees and the pruning of decision trees. In the sample feature space, the data is divided into M units (R<sub>1</sub>, R<sub>2</sub>, &#x2026; R<sub>M</sub>), where each unit i has an output value defined as c<sub>i</sub>. The regression tree model is represented by <xref ref-type="disp-formula" rid="e25">Formula 25</xref>.<disp-formula id="e25">
<mml:math id="m37">
<mml:mrow>
<mml:mi mathvariant="normal">f</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">x</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x3d;</mml:mo>
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:msubsup>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi mathvariant="normal">i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi mathvariant="normal">M</mml:mi>
</mml:msubsup>
</mml:mstyle>
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">c</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
<mml:mi mathvariant="normal">I</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>&#x2208;</mml:mo>
<mml:msub>
<mml:mi mathvariant="normal">R</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(25)</label>
</disp-formula>Where I(x&#x2208;R<sub>i</sub>) represents the adaptation function for spatial feature partitioning, which is optimized by minimizing the square error and determining the optimal output value of each partition unit as the average of all observed values. The optimal feature j and segmentation point s on the partitioned feature are solved using <xref ref-type="disp-formula" rid="e26">Formula 26</xref>.<disp-formula id="e26">
<mml:math id="m38">
<mml:mrow>
<mml:munder>
<mml:mi>min</mml:mi>
<mml:mrow>
<mml:mi mathvariant="normal">j</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi mathvariant="normal">s</mml:mi>
</mml:mrow>
</mml:munder>
<mml:mrow>
<mml:mfenced open="[" close="]" separators="|">
<mml:mrow>
<mml:munder>
<mml:mi>min</mml:mi>
<mml:msub>
<mml:mi mathvariant="normal">c</mml:mi>
<mml:mn>1</mml:mn>
</mml:msub>
</mml:munder>
<mml:mstyle displaystyle="true">
<mml:msub>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
<mml:mo>&#x2208;</mml:mo>
<mml:msub>
<mml:mi mathvariant="normal">R</mml:mi>
<mml:mn>1</mml:mn>
</mml:msub>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">j</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi mathvariant="normal">s</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:msub>
</mml:mstyle>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">y</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi mathvariant="normal">c</mml:mi>
<mml:mn>1</mml:mn>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
<mml:mo>&#x2b;</mml:mo>
</mml:mrow>
<mml:munder>
<mml:mi>min</mml:mi>
<mml:msub>
<mml:mi mathvariant="normal">c</mml:mi>
<mml:mn>2</mml:mn>
</mml:msub>
</mml:munder>
<mml:mstyle displaystyle="true">
<mml:msub>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mn>2</mml:mn>
</mml:msub>
<mml:mo>&#x2208;</mml:mo>
<mml:msub>
<mml:mi mathvariant="normal">R</mml:mi>
<mml:mn>2</mml:mn>
</mml:msub>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">j</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi mathvariant="normal">s</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:msub>
</mml:mstyle>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">y</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi mathvariant="normal">c</mml:mi>
<mml:mn>2</mml:mn>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(26)</label>
</disp-formula>
</p>
<p>The optimal value of c<sub>1</sub> in R<sub>1</sub> is indicated by <xref ref-type="disp-formula" rid="e27">Formula 27</xref>, while the optimal value of c<sub>2</sub> in R<sub>2</sub> is denoted by <xref ref-type="disp-formula" rid="e28">Formula 28</xref>.<disp-formula id="e27">
<mml:math id="m39">
<mml:mrow>
<mml:msub>
<mml:mover accent="true">
<mml:mi mathvariant="normal">c</mml:mi>
<mml:mo>&#x5e;</mml:mo>
</mml:mover>
<mml:mn>1</mml:mn>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mtext>ave</mml:mtext>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">y</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
<mml:mo>&#x7c;</mml:mo>
<mml:msub>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
<mml:mo>&#x2208;</mml:mo>
<mml:msub>
<mml:mi mathvariant="normal">R</mml:mi>
<mml:mn>1</mml:mn>
</mml:msub>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">j</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi mathvariant="normal">s</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(27)</label>
</disp-formula>
<disp-formula id="e28">
<mml:math id="m40">
<mml:mrow>
<mml:msub>
<mml:mover accent="true">
<mml:mi mathvariant="normal">c</mml:mi>
<mml:mo>&#x5e;</mml:mo>
</mml:mover>
<mml:mn>2</mml:mn>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mtext>ave</mml:mtext>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">y</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
<mml:mo>&#x7c;</mml:mo>
<mml:msub>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
<mml:mo>&#x2208;</mml:mo>
<mml:msub>
<mml:mi mathvariant="normal">R</mml:mi>
<mml:mn>2</mml:mn>
</mml:msub>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">j</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi mathvariant="normal">s</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(28)</label>
</disp-formula>
</p>
<p>By evaluating the loss function, pruning is performed iteratively from the leaf nodes to the root node of the generated decision tree. The pruned subtree, denoted as {T<sub>0</sub>,T<sub>1</sub>, &#x2026; T<sub>k</sub>}, is determined based on a calculated subtree loss function using <xref ref-type="disp-formula" rid="e29">Formula 29</xref>.<disp-formula id="e29">
<mml:math id="m41">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">C</mml:mi>
<mml:mi mathvariant="normal">&#x3b1;</mml:mi>
</mml:msub>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">T</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x3d;</mml:mo>
<mml:mi mathvariant="normal">C</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">T</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x2b;</mml:mo>
<mml:mi mathvariant="normal">&#x3b1;</mml:mi>
<mml:mrow>
<mml:mfenced open="|" close="|" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">T</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(29)</label>
</disp-formula>
</p>
<p>The cost of pruning, denoted as C(T), is the square error associated with any subtree T, while &#x7c;T&#x7c; represents the complexity of the model in terms of the number of leaf nodes. Here, &#x3b1; is a weight parameter that balances model fitting and complexity, ultimately determining the generalization ability of the model.</p>
<p>In this study, the termination condition of the CART algorithm was defined as the minimum leaf size, which represents the number of samples in a leaf node. Initially, this value is set to 12 based on empirical knowledge for model regression. However, it is important to note that this initial setting may not be optimal or ideal. The hyperparameter for the minimum leaf size of the CART regression model was adjusted accordingly. Bayesian optimization and iterative processes were employed to optimize the hyperparameter based on minimizing mean squared error. Consequently, the minimum leaf size was set to [1,378/2].</p>
</sec>
<sec id="s2-2-5">
<title>2.2.5 Ensemble tree regression</title>
<p>The Boosting tree algorithm, proposed by Friedman et al., in 2000 (<xref ref-type="bibr" rid="B10">Friedman et al., 2000</xref>), is considered one of the most high-performing ensemble learning algorithms. Boosting tree is an ensemble algorithm based on decision trees and utilizes a boosting technique that employs additive models and forward distribution algorithms, following a sequential approach. The weak learners are combined into strong learners by linearly combining basis functions with weights. In this boosting tree regression model, which uses the CART algorithm as the basis function, which is defined by <xref ref-type="disp-formula" rid="e30">Formula 30</xref>.<disp-formula id="e30">
<mml:math id="m42">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">f</mml:mi>
<mml:mi mathvariant="normal">M</mml:mi>
</mml:msub>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">x</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x3d;</mml:mo>
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:msubsup>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi mathvariant="normal">m</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi mathvariant="normal">M</mml:mi>
</mml:msubsup>
</mml:mstyle>
<mml:mrow>
<mml:mi mathvariant="normal">T</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mo>,</mml:mo>
<mml:msub>
<mml:mi mathvariant="normal">&#x398;</mml:mi>
<mml:mi mathvariant="normal">m</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(30)</label>
</disp-formula>Where T(x,&#x398;) is the decision tree generated by CART algorithm. M represents the number of decision trees and &#x398;<sub>m</sub> represents the weight of the <italic>m</italic>th decision tree. The current model of the boosting tree is defined as f<sub>(m-1)</sub>(x), and thus, <xref ref-type="disp-formula" rid="e31">Formula 31</xref> illustrates the weight of the <italic>m</italic>th decision tree.<disp-formula id="e31">
<mml:math id="m43">
<mml:mrow>
<mml:msub>
<mml:mover accent="true">
<mml:mi mathvariant="normal">&#x398;</mml:mi>
<mml:mo>&#x5e;</mml:mo>
</mml:mover>
<mml:mi mathvariant="normal">m</mml:mi>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mo>&#x2061;</mml:mo>
<mml:mi>arg</mml:mi>
<mml:munder>
<mml:mi>min</mml:mi>
<mml:msub>
<mml:mi mathvariant="normal">&#x398;</mml:mi>
<mml:mi mathvariant="normal">m</mml:mi>
</mml:msub>
</mml:munder>
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:msubsup>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi mathvariant="normal">i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi mathvariant="normal">N</mml:mi>
</mml:msubsup>
</mml:mstyle>
<mml:mrow>
<mml:mi mathvariant="normal">L</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">y</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
<mml:mo>,</mml:mo>
<mml:mi mathvariant="normal">&#x3b7;</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">f</mml:mi>
<mml:mrow>
<mml:mi mathvariant="normal">m</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msub>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x2b;</mml:mo>
<mml:mi mathvariant="normal">T</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">x</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
<mml:mo>;</mml:mo>
<mml:msub>
<mml:mi mathvariant="normal">&#x398;</mml:mi>
<mml:mi mathvariant="normal">m</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(31)</label>
</disp-formula>
</p>
<p>The loss function is defined as the mean squared error, and the optimal weight corresponds to the weight that minimizes the mean squared error. Here, &#x3b7; denotes the learning rate. The boosting tree model involves several crucial hyperparameters. Based on empirical knowledge, the minimum leaf node is set to be 8, with a total of 30 learners and a learning rate of 0.1 for boosting tree regression.</p>
<p>Random forest, proposed by Breiman in 2001 (<xref ref-type="bibr" rid="B2">Breiman, 2001</xref>), is a bagging ensemble algorithm based on decision trees. Bagging involves generating multiple decision trees by randomly selecting and replacing feature and sample sets, and the predicted values of all decision trees are averaged during prediction, thereby parallelly combining the basis functions. In this random forest regression model, the CART algorithm is employed as the basis function and consider several important hyperparameters. Based on empirical knowledge, the minimum leaf node hyperparameter is set to 8 and use 30 learners for random forest regression.</p>
<p>However, the empirical values in the regression models of the boosting tree and random forest may not be inherently optimal or ideal. The hyperparameters of the model were adjusted and iterated using Bayesian optimization to optimize the optimal hyperparameters, based on the minimum mean squared error. The range of hyperparameter adjustments is presented in <xref ref-type="table" rid="T2">Table 2</xref>.</p>
</sec>
</sec>
<sec id="s2-3">
<title>2.3 Model evaluation</title>
<p>In this study, five methods were employed to assess the machine learning regression model: root mean squared error (RMSE), mean absolute error (MAE), coefficient of determination (<italic>R</italic>
<sup>2</sup>), mean squared error (MSE), and training time (T).</p>
<sec id="s2-3-1">
<title>2.3.1 Root mean squared error</title>
<p>The root mean squared error is a widely used metric for evaluating regression models, which quantifies the average degree of deviation between predicted and true values as expressed in <xref ref-type="disp-formula" rid="e32">Formula 32</xref>.<disp-formula id="e32">
<mml:math id="m44">
<mml:mrow>
<mml:mtext>RMSE</mml:mtext>
<mml:mo>&#x3d;</mml:mo>
<mml:msqrt>
<mml:mfrac>
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:msubsup>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi mathvariant="normal">i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi mathvariant="normal">n</mml:mi>
</mml:msubsup>
</mml:mstyle>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">y</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mover accent="true">
<mml:mi mathvariant="normal">y</mml:mi>
<mml:mo>&#x5e;</mml:mo>
</mml:mover>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
<mml:mi mathvariant="normal">n</mml:mi>
</mml:mfrac>
</mml:msqrt>
</mml:mrow>
</mml:math>
<label>(32)</label>
</disp-formula>Where n denotes the number of samples, y<sub>i</sub> represents the <italic>i</italic>th true value to be predicted, and <inline-formula id="inf13">
<mml:math id="m45">
<mml:mrow>
<mml:msub>
<mml:mover accent="true">
<mml:mi mathvariant="normal">y</mml:mi>
<mml:mo>&#x5e;</mml:mo>
</mml:mover>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> denotes the <italic>i</italic>th predicted value. The formula reveals that a smaller RMSE value indicates a closer proximity between the predicted and true values, thereby signifying an enhanced fitting degree and performance of the regression model.</p>
</sec>
<sec id="s2-3-2">
<title>2.3.2 Mean absolute error</title>
<p>The mean absolute error is a commonly employed metric for evaluating regression models, quantifying the mean discrepancy between predicted and true values as defined in Eq. <xref ref-type="disp-formula" rid="e33">33</xref>.<disp-formula id="e33">
<mml:math id="m46">
<mml:mrow>
<mml:mtext>MAE</mml:mtext>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:msubsup>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi mathvariant="normal">i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi mathvariant="normal">n</mml:mi>
</mml:msubsup>
</mml:mstyle>
<mml:mrow>
<mml:mfenced open="|" close="|" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">y</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mover accent="true">
<mml:mi mathvariant="normal">y</mml:mi>
<mml:mo>&#x5e;</mml:mo>
</mml:mover>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
<mml:mi mathvariant="normal">n</mml:mi>
</mml:mfrac>
</mml:mrow>
</mml:math>
<label>(33)</label>
</disp-formula>
</p>
<p>The formula reveals that a smaller MAE value corresponds to a narrower discrepancy between the predicted and true values, indicating an enhanced level of fitting and performance for the regression model.</p>
</sec>
<sec id="s2-3-3">
<title>2.3.3 Coefficient of determination</title>
<p>The coefficient of determination, a crucial evaluation index of regression models, quantifies the correlation between dependent and independent variables while representing the explanatory power of the regression model for the dependent variable. The calculation is presented in <xref ref-type="disp-formula" rid="e34">Formula 34</xref>.<disp-formula id="e34">
<mml:math id="m47">
<mml:mrow>
<mml:msup>
<mml:mi mathvariant="normal">R</mml:mi>
<mml:mn>2</mml:mn>
</mml:msup>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
<mml:mo>&#x2212;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mtext>SSE</mml:mtext>
</mml:mrow>
<mml:mrow>
<mml:mtext>SST</mml:mtext>
</mml:mrow>
</mml:mfrac>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
<mml:mo>&#x2212;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:msubsup>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi mathvariant="normal">i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi mathvariant="normal">n</mml:mi>
</mml:msubsup>
</mml:mstyle>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">y</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mover accent="true">
<mml:mi mathvariant="normal">y</mml:mi>
<mml:mo>&#x5e;</mml:mo>
</mml:mover>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:msubsup>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi mathvariant="normal">i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi mathvariant="normal">n</mml:mi>
</mml:msubsup>
</mml:mstyle>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">y</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:mi mathvariant="normal">&#x3bc;</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:math>
<label>(34)</label>
</disp-formula>
</p>
<p>The sum of squares of error (SSE) represents the discrepancy between the predicted values and the true values in the model. On the other hand, the total sum of squares (SST) represents the overall deviation between all predicted true values and their mean value. Here, &#x3bc; denotes the average value of all true values. It is evident from this formula that <italic>R</italic>
<sup>2</sup> typically ranges between 0 and 1, indicating a prediction error lower than that of mean reference. As <italic>R</italic>
<sup>2</sup> approaches 1, it signifies a higher goodness-of-fit for the regression model, implying superior performance. Conversely, when <italic>R</italic>
<sup>2</sup> is negative, it suggests poor predictive ability and greater prediction errors compared to mean reference.</p>
</sec>
<sec id="s2-3-4">
<title>2.3.4 Mean squared error</title>
<p>The mean squared error is a widely used metric for evaluating regression models, quantifying the discrepancy between predicted and true values as computed in <xref ref-type="disp-formula" rid="e35">Formula 35</xref>.<disp-formula id="e35">
<mml:math id="m48">
<mml:mrow>
<mml:mtext>MSE</mml:mtext>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:msubsup>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi mathvariant="normal">i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi mathvariant="normal">n</mml:mi>
</mml:msubsup>
</mml:mstyle>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="normal">y</mml:mi>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mover accent="true">
<mml:mi mathvariant="normal">y</mml:mi>
<mml:mo>&#x5e;</mml:mo>
</mml:mover>
<mml:mi mathvariant="normal">i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
<mml:mi mathvariant="normal">n</mml:mi>
</mml:mfrac>
</mml:mrow>
</mml:math>
<label>(35)</label>
</disp-formula>
</p>
<p>The formula reveals that a smaller MSE value corresponds to a closer proximity between the predicted and true values, indicating an enhanced level of fitting and performance for the regression model.</p>
</sec>
<sec id="s2-3-5">
<title>2.3.5 Training time</title>
<p>The training time reflects the model&#x2019;s complexity and encompasses the time taken from the initiation to completion of training. On a consistent hardware platform, shorter training times indicate higher efficiency in model learning. For models with substantial time requirements, the training time serves as a crucial evaluation metric.</p>
</sec>
</sec>
<sec id="s2-4">
<title>2.4 Model selection</title>
<p>Given this emphasis on model prediction accuracy, the training time serves as a mere reference, while the optimal model is selected based on various hyperparameter optimized models using RMSE, MAE, <italic>R</italic>
<sup>2</sup>, and MSE.</p>
</sec>
</sec>
<sec sec-type="results|discussion" id="s3">
<title>3 Results and discussion</title>
<p>The experiment was conducted using MATLAB software. To enhance the model&#x2019;s generalization ability, the training set and validation set were obtained through a ten-fold cross-validation approach, where 5% of the dataset (19 samples) was randomly selected for testing purposes while the remaining 95% (378 samples) was divided into training and validation sets.</p>
<sec id="s3-1">
<title>3.1 Feature engineering based on titanium alloys</title>
<sec id="s3-1-1">
<title>3.1.1 Feature selection</title>
<p>The unsupervised feature correlation coefficient results are presented in <xref ref-type="fig" rid="F3">Figure 3A</xref>. It is evident that features Z and u, as well as X and Cr, exhibit a relatively high degree of correlation with correlation coefficients exceeding 0.9, which is a correlation threshold of experience.</p>
<fig id="F3" position="float">
<label>FIGURE 3</label>
<caption>
<p>Feature engineering based on titanium alloys: <bold>(A)</bold> heat map of unsupervised feature correlation coefficients; <bold>(B)</bold> ranking of feature correlation under supervision; <bold>(C)</bold> histogram of normalized features.</p>
</caption>
<graphic xlink:href="fmats-11-1364572-g003.tif"/>
</fig>
<p>The ranking results of feature correlation based on the supervised MRMR algorithm are presented in <xref ref-type="fig" rid="F3">Figures 3A, B</xref> comprehensive analysis is conducted by considering the features associated with the supervised comparison. Specifically, Z exhibits a feature correlation score of 0.1395, u has a score of 0.1270, X is assigned a score of 0.1083, and Cr demonstrates a score of 0.1899. These features are all evaluated as significant contributors to the model&#x2019;s performance. In order to construct a high-precision prediction model, all 28 features are retained without any exclusion.</p>
</sec>
<sec id="s3-1-2">
<title>3.1.2 Standardization of data</title>
<p>After standardization, the data for each of the 28 features conforms to a normal distribution with a mean value of 0 and a standard deviation of 1, a random selection of 4 features was made to construct the histogram, as illustrated in <xref ref-type="fig" rid="F3">Figure 3C</xref>.</p>
</sec>
</sec>
<sec id="s3-2">
<title>3.2 Machine learning algorithm model</title>
<sec id="s3-2-1">
<title>3.2.1 Support vector machine regression</title>
<p>In the support vector machine regression model, the hyperparameters were optimized to include linear, Gaussian, quadratic and cubic kernel functions. The resulting performance metrics including RMSE, MAE, <italic>R</italic>
<sup>2</sup> and MSE as well as training time T are presented in <xref ref-type="table" rid="T3">Table 3</xref>.</p>
<table-wrap id="T3" position="float">
<label>TABLE 3</label>
<caption>
<p>Optimization outcomes of hyperparameters for SVR and GPR models.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="center">Regression model</th>
<th align="center">Dataset</th>
<th align="center">Kernel function</th>
<th align="center">RMSE</th>
<th align="center">MAE</th>
<th align="center">
<italic>R</italic>
<sup>2</sup>
</th>
<th align="center">MSE</th>
<th align="center">T(s)</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td rowspan="4" align="center">SVR</td>
<td rowspan="4" align="center">Validation</td>
<td align="center" style="color:#2A2B2E">Linear</td>
<td align="center">169.81</td>
<td align="center">105.82</td>
<td align="center">0.51</td>
<td align="center">28,837</td>
<td align="center">9.9007</td>
</tr>
<tr>
<td align="center" style="color:#2A2B2E">Gaussian</td>
<td align="center">90.253</td>
<td align="center">54.581</td>
<td align="center">0.86</td>
<td align="center">8,145.5</td>
<td align="center">1.1467</td>
</tr>
<tr>
<td align="center" style="color:#2A2B2E">Quadratic</td>
<td align="center">113.8</td>
<td align="center">63.465</td>
<td align="center">0.78</td>
<td align="center">12,951</td>
<td align="center">62.253</td>
</tr>
<tr>
<td align="center" style="color:#2A2B2E">Cubic</td>
<td align="center">2,959.9</td>
<td align="center">503.83</td>
<td align="center">&#x2212;148.39</td>
<td align="center">87,61,100</td>
<td align="center">56.21</td>
</tr>
<tr>
<td rowspan="5" align="center">GPR</td>
<td rowspan="5" align="center">Validation</td>
<td align="center" style="color:#2A2B2E">Rational quadratic</td>
<td align="center">57.338</td>
<td align="center">35.339</td>
<td align="center">0.94</td>
<td align="center">3,287.7</td>
<td align="center">5.0902</td>
</tr>
<tr>
<td align="center" style="color:#2A2B2E">Squared exponential</td>
<td align="center">71.49</td>
<td align="center">40.345</td>
<td align="center">0.91</td>
<td align="center">5,110.9</td>
<td align="center">3.8299</td>
</tr>
<tr>
<td align="center" style="color:#2A2B2E">Matern 5/2</td>
<td align="center">66.104</td>
<td align="center">37.993</td>
<td align="center">0.93</td>
<td align="center">4,369.7</td>
<td align="center">3.8287</td>
</tr>
<tr>
<td align="center" style="color:#2A2B2E">Matern 3/2</td>
<td align="center">62.661</td>
<td align="center">36.644</td>
<td align="center">0.93</td>
<td align="center">3,926.4</td>
<td align="center">3.8182</td>
</tr>
<tr>
<td align="center" style="color:#2A2B2E">Exponential</td>
<td align="center">54.745</td>
<td align="center">33.64</td>
<td align="center">0.95</td>
<td align="center">2,997.1</td>
<td align="center">4.0143</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>The optimal hyperparameter of the model is determined to be the Gaussian kernel function, as evidenced by the smallest values of RMSE, MAE, and MSE, along with an <italic>R</italic>
<sup>2</sup> value of 0.86 (<xref ref-type="table" rid="T3">Table 3</xref>). The model performance is unsatisfactory when the hyperparameter selects the cubic kernel function, as indicated by an <italic>R</italic>
<sup>2</sup> value of &#x2212;148.39, which exceeds the mean reference error and indicates poor prediction accuracy. <xref ref-type="fig" rid="F4">Figure 4A</xref> illustrates the comparison between real and predicted values for this particular response.</p>
<fig id="F4" position="float">
<label>FIGURE 4</label>
<caption>
<p>The response of the true value and predicted value of the model following hyperparameter optimization: <bold>(A)</bold> SVR; <bold>(B)</bold> GPR; <bold>(C)</bold> NNR; <bold>(D)</bold> CART; <bold>(E)</bold> ER.</p>
</caption>
<graphic xlink:href="fmats-11-1364572-g004.tif"/>
</fig>
</sec>
<sec id="s3-2-2">
<title>3.2.2 Gaussian process regression</title>
<p>In the Gaussian process regression model, the hyperparameters of the model were adjusted to incorporate rational quadratic, squared exponential, Matern 5/2, Matern 3/2, and exponential kernel functions. The performance metrics including RMSE, MAE, <italic>R</italic>
<sup>2</sup>, MSE, and training time T were presented in <xref ref-type="table" rid="T3">Table 3</xref>.</p>
<p>The optimal hyperparameter of the model is determined to be the exponential kernel function, as evidenced by the smallest values of RMSE, MAE, and MSE, along with an impressive <italic>R</italic>
<sup>2</sup> value of 0.95 (<xref ref-type="table" rid="T3">Table 3</xref>). The corresponding responses between real and predicted values are visually depicted in <xref ref-type="fig" rid="F4">Figure 4B</xref>.</p>
</sec>
<sec id="s3-2-3">
<title>3.2.3 Neural network regression</title>
<p>The results of the neural network regression, including RMSE, MAE, <italic>R</italic>
<sup>2</sup>, MSE, and training time T, are presented in <xref ref-type="table" rid="T4">Table 4</xref>. Additionally, <xref ref-type="fig" rid="F5">Figure 5A</xref> illustrates the iterative outcomes of the hyperparameter optimization algorithm. It is evident from the figure that convergence occurs during the third iteration and yields optimized hyperparameters as follows: two hidden layers with a Tanh activation function; 257 neurons in hidden layer 1 and 216 neurons in hidden layer 2; Intensity of regularization set at 0.2791. The optimized neural network regression model is depicted in <xref ref-type="fig" rid="F2">Figure 2B</xref>.</p>
<table-wrap id="T4" position="float">
<label>TABLE 4</label>
<caption>
<p>Optimization outcomes of hyperparameters for NNR, CART, Boosting tree, Random forest and Ensemble tree models.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="center">Regression model</th>
<th align="center">Dateset</th>
<th align="center">Model hyperparameter optimization state</th>
<th align="center">RMSE</th>
<th align="center">MAE</th>
<th align="center">
<italic>R</italic>
<sup>2</sup>
</th>
<th align="center">MSE</th>
<th align="center">T(s)</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td rowspan="2" align="center">NNR</td>
<td rowspan="2" align="center">Validation</td>
<td align="center">Unoptimized</td>
<td align="center">142.93</td>
<td align="center">81.85</td>
<td align="center">0.65</td>
<td align="center">20,430</td>
<td align="center">12.814</td>
</tr>
<tr>
<td align="center">Optimized</td>
<td align="center">64.537</td>
<td align="center">37.798</td>
<td align="center">0.93</td>
<td align="center">4,165</td>
<td align="center">781.62</td>
</tr>
<tr>
<td rowspan="2" align="center">CART</td>
<td rowspan="2" align="center">Validation</td>
<td align="center">Unoptimized</td>
<td align="center">97.154</td>
<td align="center">67.808</td>
<td align="center">0.84</td>
<td align="center">9,438.8</td>
<td align="center">2.3318</td>
</tr>
<tr>
<td align="center">Optimized</td>
<td align="center">79.753</td>
<td align="center">50.931</td>
<td align="center">0.89</td>
<td align="center">6,360.6</td>
<td align="center">28.908</td>
</tr>
<tr>
<td align="center">Boosting tree</td>
<td align="center">Validation</td>
<td align="center">Unoptimized</td>
<td align="center">88.138</td>
<td align="center">64.91</td>
<td align="center">0.87</td>
<td align="center">7,768.3</td>
<td align="center">3.7272</td>
</tr>
<tr>
<td align="center">Random forest</td>
<td align="center">Validation</td>
<td align="center">Unoptimized</td>
<td align="center">83.934</td>
<td align="center">57.934</td>
<td align="center">0.88</td>
<td align="center">7,045</td>
<td align="center">8.892</td>
</tr>
<tr>
<td align="center">Ensemble tree</td>
<td align="center">Validation</td>
<td align="center">Optimized</td>
<td align="center">60.859</td>
<td align="center">38.913</td>
<td align="center">0.94</td>
<td align="center">3,703.8</td>
<td align="center">293.09</td>
</tr>
</tbody>
</table>
</table-wrap>
<fig id="F5" position="float">
<label>FIGURE 5</label>
<caption>
<p>Iteration for optimizing model hyperparameters: <bold>(A)</bold> NNR; <bold>(B)</bold> CART; <bold>(C)</bold> ER.</p>
</caption>
<graphic xlink:href="fmats-11-1364572-g005.tif"/>
</fig>
<p>The optimized NNR model has exhibited significant performance enhancement, as evident from the results presented in <xref ref-type="table" rid="T4">Table 4</xref>. Notably, there has been a substantial reduction in the values of RMSE, MAE, and MSE, while <italic>R</italic>
<sup>2</sup> has increased from 0.65 to 0.93. <xref ref-type="fig" rid="F4">Figure 4C</xref> illustrates the comparison between actual and predicted values.</p>
</sec>
<sec id="s3-2-4">
<title>3.2.4 CART regression</title>
<p>The performance evaluation of CART tree regression models, including RMSE, MAE, <italic>R</italic>
<sup>2</sup>, MSE and training time T, is presented in <xref ref-type="table" rid="T4">Table 4</xref>. The iterative results of the hyperparameter optimization algorithm are illustrated in <xref ref-type="fig" rid="F5">Figure 5B</xref>. It can be observed from the figure that the algorithm achieves convergence after three iterations with optimized hyperparameters: a minimum leaf size of 2.</p>
<p>The optimization of the model has led to a noticeable improvement in performance, as evident from the data presented in <xref ref-type="table" rid="T4">Table 4</xref>. Specifically, there has been a reduction in RMSE, MAE, and MSE values, indicating enhanced accuracy. Moreover, the <italic>R</italic>
<sup>2</sup> value has increased significantly from 0.84 to 0.89. <xref ref-type="fig" rid="F4">Figure 4D</xref> illustrates the correlation between actual and predicted values.</p>
</sec>
<sec id="s3-2-5">
<title>3.2.5 Ensemble tree regression</title>
<p>The performance metrics, including RMSE, MAE, <italic>R</italic>
<sup>2</sup>, MSE, and training time T of the boosting tree and random forest regression models in the ensemble tree regression approach are presented in <xref ref-type="table" rid="T4">Table 4</xref>. The iterative results of the hyperparameter optimization algorithm are illustrated in <xref ref-type="fig" rid="F5">Figure 5C</xref>. It can be observed from the figure that the algorithm converges at the 23rd iteration with optimized hyperparameters as follows: Boosting algorithm is employed for ensemble learning method with a total of 497 learners, minimum leaf size set to 7, learning rate set to 0.0675, and number of sampled predictors limited to 5.</p>
<p>The performance of the empirically defined random forest regression model is superior to that of the boosting tree regression model, as evident from <xref ref-type="table" rid="T4">Table 4</xref> prior to hyperparameter optimization. Moreover, the random forest regression model exhibits smaller values for RMSE, MAE, and MSE compared to the boosting tree regression model. Additionally, the <italic>R</italic>
<sup>2</sup> value of the random forest regression model surpasses that of the boosting tree regression model by 0.01. Subsequent optimization resulted in decreased RMSE, MAE, and MSE values along with an increased <italic>R</italic>
<sup>2</sup> value of 0.94, leading to a more refined and optimized model. <xref ref-type="fig" rid="F4">Figure 4E</xref> illustrates the comparison between real and predicted values.</p>
</sec>
</sec>
<sec id="s3-3">
<title>3.3 Model selection</title>
<p>The optimized Gaussian regression model with exponential hyperparameter exhibits superior performance compared to all other models, as evident from the results presented in <xref ref-type="table" rid="T3">Tables 3</xref>, <xref ref-type="table" rid="T4">4</xref>. It achieves a lower RMSE of 54.745, MAE of 33.64, and MSE of 2997.1, outperforming the other optimized models. Additionally, it attains an impressive <italic>R</italic>
<sup>2</sup> value of 0.95 surpasses existing literature benchmarks (refer to <xref ref-type="table" rid="T5">Table 5</xref> for detailed comparisons). This enhancement can be primarily attributed to this meticulous feature engineering and comprehensive model optimization based on titanium alloys.</p>
<table-wrap id="T5" position="float">
<label>TABLE 5</label>
<caption>
<p>The comparison of model <italic>R</italic>
<sup>2</sup>.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="center">Model</th>
<th align="center">
<xref ref-type="bibr" rid="B18">Jiang et al. (2022)</xref>
</th>
<th align="center">
<xref ref-type="bibr" rid="B11">Giles et al. (2022)</xref>
</th>
<th align="center">
<xref ref-type="bibr" rid="B48">Yu et al. (2019)</xref>
</th>
<th align="center">This study</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="center">
<italic>R</italic>
<sup>2</sup>
</td>
<td align="center">0.94</td>
<td align="center">0.895</td>
<td align="center">&#x3c;0.9</td>
<td align="center">0.95</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>In order to validate this model, model verification was conducted on the test set. The RMSE of the model on the test set was 68.53, with an MAE of 42.239, MSE of 4696.3, and <italic>R</italic>
<sup>2</sup> value of 0.93, indicating excellent performance. <xref ref-type="table" rid="T6">Table 6</xref> presents the prediction results obtained from the optimized Gaussian regression model applied to the test set. <xref ref-type="fig" rid="F6">Figure 6A</xref> illustrates both true and predicted responses, while <xref ref-type="fig" rid="F6">Figure 6B</xref> displays the residuals.</p>
<table-wrap id="T6" position="float">
<label>TABLE 6</label>
<caption>
<p>Comparison of predicted and literature values on the test set.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="center">Regression model</th>
<th align="center">Alloy composition</th>
<th align="center">Heat treatment process</th>
<th align="center">Literature value</th>
<th align="center">Predicted value</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td rowspan="19" align="center">GPR</td>
<td align="center">TA10 (Ti0.3Mo0.8Ni)</td>
<td align="center">860&#xb0;C &#xd7; 2 h/WQ</td>
<td align="center">500</td>
<td align="center">522.1</td>
</tr>
<tr>
<td align="center">TA10 (Ti0.3Mo0.8Ni)</td>
<td align="center">700&#xb0;C &#xd7; 0.5 h/AC</td>
<td align="center">557.5</td>
<td align="center">558.7</td>
</tr>
<tr>
<td align="center">TA10 (Ti0.3Mo0.8Ni)</td>
<td align="center">550&#xb0;C &#xd7; 0.5 h/AC</td>
<td align="center">590</td>
<td align="center">546.4</td>
</tr>
<tr>
<td align="center">Ti43 (Ti4Al2.5V1Fe)</td>
<td align="center">850&#xb0;C &#xd7; 1.5 h/AC</td>
<td align="center">873</td>
<td align="center">801</td>
</tr>
<tr>
<td align="center">TA31 (Ti6Al2Zr1Mo3Nb)</td>
<td align="center">980&#xb0;C &#xd7; 1 h/AC &#x2b; 700&#xb0;C &#xd7; 1 h/AC</td>
<td align="center">897</td>
<td align="center">897.6</td>
</tr>
<tr>
<td align="center">Ti-5111 (Ti5Al1Mo1V1Zr1Sn)</td>
<td align="center">1,000&#xb0;C &#xd7; 1 h/AC</td>
<td align="center">940</td>
<td align="center">888.7</td>
</tr>
<tr>
<td align="center">Ti6Al7Nd</td>
<td align="center">985&#xb0;C &#xd7; 1 h/WQ</td>
<td align="center">935</td>
<td align="center">934.9</td>
</tr>
<tr>
<td align="center">TB12 (Ti11Mo5Zr4Sn3Nb)</td>
<td align="center">800&#xb0;C &#xd7; 1 h/WQ</td>
<td align="center">945</td>
<td align="center">977.9</td>
</tr>
<tr>
<td align="center">Ti40 (Ti25V15Cr0.2Si)</td>
<td align="center">850&#xb0;C &#xd7; 1 h/WQ &#x2b; 550&#xb0;C &#xd7; 6 h/AC</td>
<td align="center">970</td>
<td align="center">1,180.2</td>
</tr>
<tr>
<td align="center">TC10 (Ti6Al6V2Sn0.5Fe0.5Cu)</td>
<td align="center">875&#xb0;C &#xd7; 2 h/WC &#x2b; 600&#xb0;C &#xd7; 6 h/AC</td>
<td align="center">1,110</td>
<td align="center">1,046.9</td>
</tr>
<tr>
<td align="center">Ti53311S (Ti5Al3Sn3Zr1Mo1Nb0.3Si)</td>
<td align="center">650&#xb0;C &#xd7; 2 h/AC</td>
<td align="center">1,102</td>
<td align="center">1,123.7</td>
</tr>
<tr>
<td align="center">Ti-62222s (Ti6Al2Sn2Zr2Cr2Mo0.15Si)</td>
<td align="center">750&#xb0;C &#xd7; 1 h/AC</td>
<td align="center">1,109</td>
<td align="center">1,109.1</td>
</tr>
<tr>
<td align="center">TC25 (Ti6.5Al2Zr2Sn2Mo1W0.2Si)</td>
<td align="center">880&#xb0;C &#xd7; 1 h/AC &#x2b; 550&#xb0;C &#xd7; 6 h/AC</td>
<td align="center">1,120</td>
<td align="center">1,144.9</td>
</tr>
<tr>
<td align="center">TC21 (Ti6Al2Zr2Sn2Mo1.5Cr2Nb)</td>
<td align="center">903&#xb0;C &#xd7; 1 h/AC</td>
<td align="center">1,213</td>
<td align="center">1,213.4</td>
</tr>
<tr>
<td align="center">TC9 (Ti6.5Al3.5Mo2.5Sn0.3Si)</td>
<td align="center">970&#xb0;C &#xd7; 1.5 h/AC &#x2b; 530&#xb0;C &#xd7; 6 h/AC</td>
<td align="center">1,257</td>
<td align="center">1,256.4</td>
</tr>
<tr>
<td align="center">TC6 (Ti6Al1.5Cr2.5Mo0.5Fe0.3Si)</td>
<td align="center">870&#xb0;C &#xd7; 1 h/AC &#x2b; 550&#xb0;C &#xd7; 4 h/AC</td>
<td align="center">1,270</td>
<td align="center">1,259.1</td>
</tr>
<tr>
<td align="center">TC18 (Ti5Al5Mo5V1Cr1Fe)</td>
<td align="center">810&#xb0;C &#xd7; 1.5 h/WQ &#x2b; 600&#xb0;C &#xd7; 5 h/AC</td>
<td align="center">1,295</td>
<td align="center">1,303</td>
</tr>
<tr>
<td align="center">TC6 (Ti6Al1.5Cr2.5Mo0.5Fe0.3Si)</td>
<td align="center">900&#xb0;C &#xd7; 0.5 h/AC</td>
<td align="center">1,320</td>
<td align="center">1,202.6</td>
</tr>
<tr>
<td align="center">Ti3.5Al5Mo6V3Cr2Sn0.5Fe</td>
<td align="center">800&#xb0;C &#xd7; 1 h/AC &#x2b; 560&#xb0;C &#xd7; 0.5 h/AC</td>
<td align="center">1,472</td>
<td align="center">1,350.6</td>
</tr>
</tbody>
</table>
</table-wrap>
<fig id="F6" position="float">
<label>FIGURE 6</label>
<caption>
<p>Predict results on the test set: <bold>(A)</bold> response of true value and predicted value; <bold>(B)</bold> residual of true value.</p>
</caption>
<graphic xlink:href="fmats-11-1364572-g006.tif"/>
</fig>
</sec>
<sec id="s3-4">
<title>3.4 Comprehensive model framework</title>
<p>The present study introduces a comprehensive computer-aided design framework for high-performance titanium alloys based on machine learning, as depicted in <xref ref-type="fig" rid="F7">Figure 7</xref>. This framework comprises three main components: feature engineering based on titanium alloys, machine learning algorithm model, and model evaluation and selection. By utilizing the proposed model, it becomes possible to generate any titanium alloy sequence based on the input of titanium and predict its ultimate strength using the established model architecture. Consequently, the output will provide the optimal titanium alloy sequence with superior ultimate strength properties, thereby facilitating the design process of high-performance titanium alloys. Moreover, this model also demonstrates capability in predicting the properties of heat-treated titanium alloys. In theory, the proposed model has unlimited potential to predict ultimate strength properties for all conceivable combinations of titanium based on 18 elements. From a broader perspective, this represents an inexhaustible search for novel materials in the field of titanium alloy design. For the purpose of this illustration the designed Ti6Al4Vx<sub>1</sub>Six<sub>2</sub>Mox<sub>3</sub>Snx<sub>4</sub>Nd series titanium alloy was subjected to computer aided high-performance design using the proposed framework, with x<sub>1</sub>, x<sub>2</sub>, x<sub>3</sub>, and x<sub>4</sub> limited to a value range of [0.1,5] and a step size of 0.1. A total of 6250,000 combinations were generated. Among all combinations, Ti6Al4V0.3Si5Mo2.3Sn0.1Nd exhibited the most optimal performance with an ultimate strength of 1,139.9 MPa.</p>
<fig id="F7" position="float">
<label>FIGURE 7</label>
<caption>
<p>Comprehensive model framework.</p>
</caption>
<graphic xlink:href="fmats-11-1364572-g007.tif"/>
</fig>
</sec>
</sec>
<sec sec-type="conclusion" id="s4">
<title>4 Conclusion</title>
<p>In this study, a computer-aided framework for designing high-performance titanium alloys based on machine learning techniques and an intelligent search space driven by data to facilitate the design process have been proposed. The main results are summarized as follows:<list list-type="simple">
<list-item>
<p>(1) In the feature engineering based on titanium alloy, the data are sourced exclusively from literature, ensuring an open and comprehensive data acquisition process. This approach enables a more universal and accessible dataset. Six essential properties of titanium alloy were meticulously designed to avoid any dimensionality issues. For feature selection, both supervised and unsupervised analysis was conducted, resulting in the establishment of a proprietary dataset for titanium alloy comprising 397 data samples and 28 features.</p>
</list-item>
<list-item>
<p>(2) The machine learning algorithm model incorporates six classical regression algorithms to construct the model, and hyperparameter optimization is employed to enhance its performance.</p>
</list-item>
<list-item>
<p>(3) The model evaluation and selection process involved the utilization of five regression model evaluation methods, ultimately leading to the identification of the optimal Gaussian regression model with an impressive <italic>R</italic>
<sup>2</sup> value of 0.95. This achievement signifies a higher level of technical proficiency. Furthermore, the performance of the model on an independent test set has been validated, which yielded a satisfactory prediction result with an <italic>R</italic>
<sup>2</sup> value of 0.93.</p>
</list-item>
<list-item>
<p>(4) A comprehensive machine learning framework has been proposed, and a model for high-performance titanium alloys has been established. In essence, this model represents an exhaustive intelligent search capable of exploring titanium alloys that incorporate any combination of the remaining 18 elements. Furthermore, the proposed framework is utilized to present a predictive model for a novel titanium alloy, Ti6Al4V0.3Si5Mo2.3Sn0.1Nd, with an ultimate strength of 1,139.9 MPa.</p>
</list-item>
</list>
</p>
<p>In future research, the investigation on laser powder bed fusion additive manufacturing of high performance titanium alloys with the proposed framework was conducted, encompassing the examination of printing process parameters and heat treatment effects on microstructure and mechanical properties.</p>
</sec>
</body>
<back>
<sec sec-type="data-availability" id="s5">
<title>Data availability statement</title>
<p>The raw data supporting the conclusion of this article will be made available by the authors, without undue reservation.</p>
</sec>
<sec id="s6">
<title>Author contributions</title>
<p>SA: Investigation, Methodology, Writing&#x2013;original draft, Writing&#x2013;review and editing. KL: Writing&#x2013;original draft. LZ: Investigation, Writing&#x2013;review and editing. HL: Investigation, Writing&#x2013;review and editing. RM: Investigation, Writing&#x2013;review and editing. RL: Investigation, Methodology, Writing&#x2013;review and editing. LM: Investigation, Methodology, Writing&#x2013;review and editing.</p>
</sec>
<sec sec-type="funding-information" id="s7">
<title>Funding</title>
<p>The author(s) declare that no financial support was received for the research, authorship, and/or publication of this article.</p>
</sec>
<ack>
<p>The authors gratefully acknowledge all the researchers and labs to provide the experimental facilities. KL acknowledges the support from National Natural Science Foundation of China (52201105), Natural Science Foundation of Chongqing (CSTB2022NSCQ-MSX0992), Innovation Support Program for Overseas Returnees in Chongqing (cx2023061), Research, Natural Science Foundation of Sichuan Province (2023NSFSC0407).</p>
</ack>
<sec sec-type="COI-statement" id="s8">
<title>Conflict of interest</title>
<p>Author SA was employed by AVIC Guizhou Aircraft Corporation LTD.</p>
<p>The remaining authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec sec-type="disclaimer" id="s9">
<title>Publisher&#x2019;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<sec id="s10">
<title>Supplementary material</title>
<p>The Supplementary Material for this article can be found online at: <ext-link ext-link-type="uri" xlink:href="https://www.frontiersin.org/articles/10.3389/fmats.2024.1364572/full#supplementary-material">https://www.frontiersin.org/articles/10.3389/fmats.2024.1364572/full&#x23;supplementary-material</ext-link>
</p>
<supplementary-material xlink:href="DataSheet1.pdf" id="SM1" mimetype="application/pdf" xmlns:xlink="http://www.w3.org/1999/xlink"/>
</sec>
<ref-list>
<title>References</title>
<ref id="B1">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Boyer</surname>
<given-names>R. R.</given-names>
</name>
</person-group> (<year>1995</year>). <article-title>Titanium for aerospace: rationale and applications</article-title>. <source>Adv. Perform. Mater.</source> <volume>2</volume>, <fpage>349</fpage>&#x2013;<lpage>368</lpage>. <pub-id pub-id-type="doi">10.1007/bf00705316</pub-id>
</citation>
</ref>
<ref id="B2">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Breiman</surname>
<given-names>L.</given-names>
</name>
</person-group> (<year>2001</year>). <article-title>Random forests</article-title>. <source>Mach. Learn.</source> <volume>45</volume>, <fpage>5</fpage>&#x2013;<lpage>32</lpage>. <pub-id pub-id-type="doi">10.1023/a:1010933404324</pub-id>
</citation>
</ref>
<ref id="B3">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Breiman</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Friedman</surname>
<given-names>J. H.</given-names>
</name>
<name>
<surname>Olshen</surname>
<given-names>R. A.</given-names>
</name>
<name>
<surname>Stone</surname>
<given-names>C. J.</given-names>
</name>
</person-group> (<year>1984</year>). <source>Classification and regression trees</source>. <publisher-loc>Boca Raton, FL, USA</publisher-loc>: <publisher-name>Chapman and Hall/CRC</publisher-name>.</citation>
</ref>
<ref id="B4">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Chen</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Guantao</surname>
<given-names>L.</given-names>
</name>
</person-group> (<year>2005</year>). <article-title>The characteristics and application of titanium alloys in ship</article-title>. <source>Ship Sci. Technol.</source> <volume>27</volume>, <fpage>13</fpage>&#x2013;<lpage>15</lpage>.</citation>
</ref>
<ref id="B5">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Cheng</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Yu</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Du</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>W.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>Influence of isothermal &#x3c9; transitional phase-assisted phase transition from &#x3b2; to &#x3b1; on room-temperature mechanical performance of a meta-stable &#x3b2; titanium alloy Ti&#x2212;10Mo&#x2212;6Zr&#x2212;4Sn&#x2212;3Nb (Ti-B12) for medical application</article-title>. <source>Front. Bioeng. Biotechnol.</source> <volume>8</volume>, <fpage>626665</fpage>. <pub-id pub-id-type="doi">10.3389/fbioe.2020.626665</pub-id>
</citation>
</ref>
<ref id="B6">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Cheng</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Yu</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Gai</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Du</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Dong</surname>
<given-names>F.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). <article-title>Precipitation behavior and microstructural evolution of &#x3b1; phase during hot deformation in a novel &#x3b2;-air-cooled metastable &#x3b2;-type Ti-B12 alloy</article-title>. <source>Metals</source> <volume>12</volume>, <fpage>770</fpage>. <pub-id pub-id-type="doi">10.3390/met12050770</pub-id>
</citation>
</ref>
<ref id="B7">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Deng</surname>
<given-names>Z. H.</given-names>
</name>
<name>
<surname>Yin</surname>
<given-names>H. q.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>G. f.</given-names>
</name>
<name>
<surname>Xu</surname>
<given-names>B.</given-names>
</name>
<etal/>
</person-group> (<year>2020</year>). <article-title>Machine-learning-assisted prediction of the mechanical properties of Cu&#x2013;Al alloy</article-title>. <source>Int. J. Minerals. Metallurgy Mater.</source> <volume>3</volume>, <fpage>362</fpage>&#x2013;<lpage>373</lpage>. <pub-id pub-id-type="doi">10.1007/s12613-019-1894-6</pub-id>
</citation>
</ref>
<ref id="B8">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Drucker</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Burges</surname>
<given-names>C. J. C.</given-names>
</name>
<name>
<surname>Kaufman</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Kaufman</surname>
<given-names>B. L.</given-names>
</name>
<name>
<surname>Smola</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Vapnik</surname>
<given-names>V.</given-names>
</name>
<etal/>
</person-group> (<year>1997</year>). &#x201c;<article-title>Support vector regression machines</article-title>,&#x201d; in <source>Advances in neural information processing systems 9(NIPS)</source> (<publisher-loc>Cambridge, MA, USA</publisher-loc>: <publisher-name>MIT Press</publisher-name>), <fpage>155</fpage>&#x2013;<lpage>161</lpage>.</citation>
</ref>
<ref id="B9">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Fotovvati</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Chou</surname>
<given-names>K.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Build surface study of single-layer raster scanning in selective laser melting: surface roughness prediction using deep learning</article-title>. <source>Manuf. Lett.</source> <volume>33</volume>, <fpage>701</fpage>&#x2013;<lpage>711</lpage>. <pub-id pub-id-type="doi">10.1016/j.mfglet.2022.07.088</pub-id>
</citation>
</ref>
<ref id="B10">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Friedman</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Tibshirani</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Hastie</surname>
<given-names>T.</given-names>
</name>
</person-group> (<year>2000</year>). <article-title>Additive logistic regression: a statistical view of boosting (With discussion and a rejoinder by the authors)</article-title>. <source>Ann. Statistics</source> <volume>28</volume>, <fpage>337</fpage>&#x2013;<lpage>407</lpage>. <pub-id pub-id-type="doi">10.1214/aos/1016120463</pub-id>
</citation>
</ref>
<ref id="B11">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Giles</surname>
<given-names>S. A.</given-names>
</name>
<name>
<surname>Sengupta</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Broderick</surname>
<given-names>S. R.</given-names>
</name>
<name>
<surname>Rajan</surname>
<given-names>K.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Machine-learning-based intelligent framework for discovering refractory high-entropy alloys with improved high-temperature yield strength</article-title>. <source>npj Comput. Mater.</source> <volume>8</volume>, <fpage>235</fpage>. <pub-id pub-id-type="doi">10.1038/s41524-022-00926-0</pub-id>
</citation>
</ref>
<ref id="B12">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Guo</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Shen</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Guan</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Dong</surname>
<given-names>F.</given-names>
</name>
<etal/>
</person-group> (<year>2023</year>). <article-title>Microstructure and mechanical properties of Alx(TiZrTa0.7NbMo) refractory high-entropy alloys</article-title>. <source>J. Alloys Compd.</source> <volume>960</volume>, <fpage>170739</fpage>. <pub-id pub-id-type="doi">10.1016/j.jallcom.2023.170739</pub-id>
</citation>
</ref>
<ref id="B13">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Hanawa</surname>
<given-names>T.</given-names>
</name>
</person-group> (<year>2019</year>). <source>Overview of metals and applications. Metals for biomedical devices</source>. <publisher-loc>New Delhi, India</publisher-loc>: <publisher-name>Woodhead Publishing</publisher-name>, <fpage>3</fpage>&#x2013;<lpage>24</lpage>.</citation>
</ref>
<ref id="B14">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hopfield</surname>
<given-names>J. J.</given-names>
</name>
</person-group> (<year>1982</year>). <article-title>Neural networks and physical systems with emergent collective computational abilities</article-title>. <source>Proc. Natl. Acad. Sci.</source> <volume>79</volume>, <fpage>2554</fpage>&#x2013;<lpage>2558</lpage>. <pub-id pub-id-type="doi">10.1073/pnas.79.8.2554</pub-id>
</citation>
</ref>
<ref id="B15">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ji</surname>
<given-names>Y. Z.</given-names>
</name>
<name>
<surname>Issa</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Heo</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Saal</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Wolverton</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>L. Q.</given-names>
</name>
</person-group> (<year>2014</year>). <article-title>Predicting &#x3b2;&#x2032; precipitate morphology and evolution in Mg&#x2013;RE alloys using a combination of first-principles calculations and phase-field modeling</article-title>. <source>Acta Mater.</source> <volume>76</volume>, <fpage>259</fpage>&#x2013;<lpage>271</lpage>. <pub-id pub-id-type="doi">10.1016/j.actamat.2014.05.002</pub-id>
</citation>
</ref>
<ref id="B18">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Jiang</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Fu</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Shen</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Xie</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Discovery of aluminum alloys with ultra-strength and high-toughness via a property-oriented design strategy</article-title>. <source>J. Mater. Sci. Technol.</source> <volume>3</volume>, <fpage>33</fpage>&#x2013;<lpage>43</lpage>.</citation>
</ref>
<ref id="B16">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kandavalli</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Agarwal</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Poonia</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Kishor</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Ayyagari</surname>
<given-names>K. P. R.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Design of high bulk moduli high entropy alloys using machine learning</article-title>. <source>Sci. Rep.</source> <volume>13</volume>, <fpage>20504</fpage>. <pub-id pub-id-type="doi">10.1038/s41598-023-47181-x</pub-id>
</citation>
</ref>
<ref id="B17">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kang</surname>
<given-names>X. D.</given-names>
</name>
<name>
<surname>Du</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Yue</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>J.</given-names>
</name>
<etal/>
</person-group> (<year>2024</year>). <article-title>Efficient access to ultrafine crystalline metastable-&#x3b2; titanium alloy via dual-phase recrystallization competition</article-title>. <source>J. Mater. Res. Technol.</source> <volume>29</volume>, <fpage>335</fpage>&#x2013;<lpage>343</lpage>. <pub-id pub-id-type="doi">10.1016/j.jmrt.2024.01.101</pub-id>
</citation>
</ref>
<ref id="B19">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Ma</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Qin</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Gong</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Wu</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Wen</surname>
<given-names>P.</given-names>
</name>
<etal/>
</person-group> (<year>2023</year>). <article-title>A review of the multi-dimensional application of machine learning to improve the integrated intelligence of laser powder bed fusion</article-title>. <source>J. Mater. Process. Technol.</source> <volume>318</volume>, <fpage>118032</fpage>. <pub-id pub-id-type="doi">10.1016/j.jmatprotec.2023.118032</pub-id>
</citation>
</ref>
<ref id="B20">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Liu</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Cheng</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Liang</surname>
<given-names>S. X.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>L.</given-names>
</name>
<etal/>
</person-group> (<year>2023</year>). <article-title>Nb-content-dependent passivation behavior of Ti&#x2013;Nb alloys for biomedical applications</article-title>. <source>J. Mater. Res. Technol.</source> <volume>27</volume>, <fpage>7882</fpage>&#x2013;<lpage>7894</lpage>. <pub-id pub-id-type="doi">10.1016/j.jmrt.2023.11.203</pub-id>
</citation>
</ref>
<ref id="B21">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Liu</surname>
<given-names>Q.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Yang</surname>
<given-names>H.</given-names>
</name>
</person-group> (<year>2015</year>). <article-title>Application and development of titanium alloy in aerospace and military hardware</article-title>. <source>J. Iron Steel Res.</source> <volume>27</volume>, <fpage>1</fpage>&#x2013;<lpage>4</lpage>.</citation>
</ref>
<ref id="B22">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Liu</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Song</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Xue</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Ma</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>L.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Application and development of titanium alloy and titanium matrix composites in aerospace field</article-title>. <source>J. Aeronautical Mater.</source> <volume>40</volume>, <fpage>77</fpage>&#x2013;<lpage>94</lpage>.</citation>
</ref>
<ref id="B23">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Louren&#xe7;o</surname>
<given-names>M. L.</given-names>
</name>
<name>
<surname>Cardoso</surname>
<given-names>G. C.</given-names>
</name>
<name>
<surname>Sousa</surname>
<given-names>K. d. S. J.</given-names>
</name>
<name>
<surname>Donato</surname>
<given-names>T. A. G.</given-names>
</name>
<name>
<surname>Pontes</surname>
<given-names>F. M. L.</given-names>
</name>
<name>
<surname>Grandini</surname>
<given-names>C. R.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Development of novel Ti-Mo-Mn alloys for biomedical applications</article-title>. <source>Sci. Rep.</source> <volume>10</volume>, <fpage>6298</fpage>&#x2013;<lpage>8</lpage>. <pub-id pub-id-type="doi">10.1038/s41598-020-62865-4</pub-id>
</citation>
</ref>
<ref id="B24">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Mao</surname>
<given-names>H. H.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>H. L.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>Q.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>TCHEA1: a thermodynamic database not limited for &#x201c;high entropy&#x201d; alloys</article-title>. <source>J. Phase Equilibria Diffusion</source> <volume>4</volume>, <fpage>353</fpage>&#x2013;<lpage>368</lpage>. <pub-id pub-id-type="doi">10.1007/s11669-017-0570-7</pub-id>
</citation>
</ref>
<ref id="B25">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>McClelland</surname>
<given-names>J. L.</given-names>
</name>
<name>
<surname>Rumelhart</surname>
<given-names>D. E.</given-names>
</name>
<collab>the PDP Research Group</collab>
</person-group> (<year>1986</year>). <article-title>Parallel distributed processing: explorations in the microstructure of cognition</article-title>. <source>Psychol. Biol. Models</source> <volume>2</volume>.</citation>
</ref>
<ref id="B26">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>McCulloch</surname>
<given-names>W. S.</given-names>
</name>
<name>
<surname>Pitts</surname>
<given-names>W.</given-names>
</name>
</person-group> (<year>1943</year>). <article-title>A logical calculus of the ideas immanent in nervous activity</article-title>. <source>Bull. Math. Biophysics</source> <volume>5</volume>, <fpage>115</fpage>&#x2013;<lpage>133</lpage>. <pub-id pub-id-type="doi">10.1007/bf02478259</pub-id>
</citation>
</ref>
<ref id="B27">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Minsky</surname>
<given-names>M. L.</given-names>
</name>
<name>
<surname>Papert</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Seymour</surname>
</name>
</person-group> (<year>1969</year>). <source>Perceptrons: an introduction to computational geometry</source>. <publisher-loc>Cambridge, MA, USA</publisher-loc>: <publisher-name>The MIT Press</publisher-name>.</citation>
</ref>
<ref id="B28">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>OHagan</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>1978</year>). <article-title>Curve fitting and optimal design for prediction</article-title>. <source>J. R. Stat. Soc. Ser. B Methodol.</source> <volume>40</volume>, <fpage>1</fpage>&#x2013;<lpage>24</lpage>. <pub-id pub-id-type="doi">10.1111/j.2517-6161.1978.tb01643.x</pub-id>
</citation>
</ref>
<ref id="B29">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Peng</surname>
<given-names>H. C.</given-names>
</name>
<name>
<surname>Fuhui Long</surname>
</name>
<name>
<surname>Ding</surname>
<given-names>C.</given-names>
</name>
</person-group> (<year>2005</year>). <article-title>Feature selection based on mutual information: criteria of max-dependency, max-relevance, and min-redundancy</article-title>. <source>IEEE Trans. Pattern Analysis Mach. Intell.</source> <volume>27</volume>, <fpage>1226</fpage>&#x2013;<lpage>1238</lpage>. <pub-id pub-id-type="doi">10.1109/tpami.2005.159</pub-id>
</citation>
</ref>
<ref id="B30">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Rao</surname>
<given-names>Z. Y.</given-names>
</name>
<name>
<surname>Springer</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Ponge</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>Z.</given-names>
</name>
</person-group> (<year>2022a</year>). <article-title>Combinatorial development of multicomponent invar alloys via rapid alloy prototyping</article-title>. <source>Materialia</source> <volume>21</volume>, <fpage>101326</fpage>. <pub-id pub-id-type="doi">10.1016/j.mtla.2022.101326</pub-id>
</citation>
</ref>
<ref id="B31">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Rao</surname>
<given-names>Z. Y.</given-names>
</name>
<name>
<surname>Tung</surname>
<given-names>P. Y.</given-names>
</name>
<name>
<surname>Xie</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Wei</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Ferrari</surname>
<given-names>A.</given-names>
</name>
<etal/>
</person-group> (<year>2022b</year>). <article-title>Machine learning&#x2013;enabled high-entropy alloy discovery</article-title>. <source>Science</source> <volume>378</volume>, <fpage>78</fpage>&#x2013;<lpage>85</lpage>. <pub-id pub-id-type="doi">10.1126/science.abo4940</pub-id>
</citation>
</ref>
<ref id="B32">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Rasmussen</surname>
<given-names>C. E.</given-names>
</name>
<name>
<surname>Williams</surname>
<given-names>C. K. I.</given-names>
</name>
</person-group> (<year>2005</year>). <source>Gaussian processes for machine learning</source>. <publisher-loc>Cambridge, MA, USA</publisher-loc>: <publisher-name>MIT Press</publisher-name>.</citation>
</ref>
<ref id="B33">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ren</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Ward</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Williams</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Laws</surname>
<given-names>K. J.</given-names>
</name>
<name>
<surname>Wolverton</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Hattrick-Simpers</surname>
<given-names>J.</given-names>
</name>
<etal/>
</person-group> (<year>2018</year>). <article-title>Accelerated discovery of metallic glasses through iteration of machine learning and high-throughput experiments</article-title>. <source>Sci. Adv.</source> <volume>4</volume>, <fpage>eaaq1566</fpage>. <pub-id pub-id-type="doi">10.1126/sciadv.aaq1566</pub-id>
</citation>
</ref>
<ref id="B34">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Rosenblatt</surname>
<given-names>F.</given-names>
</name>
</person-group> (<year>1958</year>). <article-title>The perceptron: a probabilistic model for information storage and organization in the brain</article-title>. <source>Psychol. Rev.</source> <volume>65</volume>, <fpage>386</fpage>&#x2013;<lpage>408</lpage>. <pub-id pub-id-type="doi">10.1037/h0042519</pub-id>
</citation>
</ref>
<ref id="B35">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Rumelhart</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Hinton</surname>
<given-names>G. E.</given-names>
</name>
<name>
<surname>Williams</surname>
<given-names>R. J.</given-names>
</name>
</person-group> (<year>1986</year>). <article-title>Learning representations by back-propagating errors</article-title>. <source>Nature</source> <volume>323</volume>, <fpage>533</fpage>&#x2013;<lpage>536</lpage>. <pub-id pub-id-type="doi">10.1038/323533a0</pub-id>
</citation>
</ref>
<ref id="B36">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sarraf</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Rezvani Ghomi</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Alipour</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Ramakrishna</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Liana Sukiman</surname>
<given-names>N.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>A state-of-the-art review of the fabrication and characteristics of titanium and its alloys for biomedical applications</article-title>. <source>Bio-design Manuf.</source> <volume>5</volume>, <fpage>371</fpage>&#x2013;<lpage>395</lpage>. <pub-id pub-id-type="doi">10.1007/s42242-021-00170-3</pub-id>
</citation>
</ref>
<ref id="B37">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sasidhar</surname>
<given-names>K. N.</given-names>
</name>
<name>
<surname>Siboni</surname>
<given-names>N. H.</given-names>
</name>
<name>
<surname>Mianroodi</surname>
<given-names>J. R.</given-names>
</name>
<name>
<surname>Rohwerder</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Neugebauer</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Raabe</surname>
<given-names>D.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Enhancing corrosion-resistant alloy design through natural language processing and deep learning</article-title>. <source>Sci. Adv.</source> <volume>9</volume>, <fpage>eadg7992</fpage>. <pub-id pub-id-type="doi">10.1126/sciadv.adg7992</pub-id>
</citation>
</ref>
<ref id="B38">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Shen</surname>
<given-names>X. Y.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Guan</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Dong</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Guo</surname>
<given-names>Z.</given-names>
</name>
<etal/>
</person-group> (<year>2024</year>). <article-title>Effect of hydrogen on thermal deformation behavior and microstructure evolution of MoNbHfZrTi refractory high-entropy alloy</article-title>. <source>Intermetallics.</source> <volume>166</volume>, <fpage>108193</fpage>. <pub-id pub-id-type="doi">10.1016/j.intermet.2024.108193</pub-id>
</citation>
</ref>
<ref id="B39">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Smola</surname>
<given-names>A. J.</given-names>
</name>
<name>
<surname>Sch&#xf6;lkopf</surname>
<given-names>B.</given-names>
</name>
</person-group> (<year>2004</year>). <article-title>A tutorial on support vector regression</article-title>. <source>Statistics Comput.</source> <volume>14</volume>, <fpage>199</fpage>&#x2013;<lpage>222</lpage>. <pub-id pub-id-type="doi">10.1023/b:stco.0000035301.49549.88</pub-id>
</citation>
</ref>
<ref id="B40">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Song</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Niu</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Yang</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Research on application technology of titanium alloy in marine pipeline</article-title>. <source>Rare Metal Mater. Eng.</source> <volume>49</volume>, <fpage>1100</fpage>&#x2013;<lpage>1104</lpage>.</citation>
</ref>
<ref id="B41">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Song</surname>
<given-names>X. L.</given-names>
</name>
<name>
<surname>Fu</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>M.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>First-principles study of &#x3b2;&#x2032; phase in Mg-RE alloys</article-title>. <source>Int. J. Mech. Sci.</source> <volume>243</volume>, <fpage>108045</fpage>. <pub-id pub-id-type="doi">10.1016/j.ijmecsci.2022.108045</pub-id>
</citation>
</ref>
<ref id="B42">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tian</surname>
<given-names>X. H.</given-names>
</name>
<name>
<surname>Zhou</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Zhao</surname>
<given-names>Q.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Shi</surname>
<given-names>D.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). <article-title>Screening for shape memory alloys with narrow thermal hysteresis using combined XGBoost and DFT calculation</article-title>. <source>Comput. Mater. Sci.</source> <volume>211</volume>, <fpage>111519</fpage>. <pub-id pub-id-type="doi">10.1016/j.commatsci.2022.111519</pub-id>
</citation>
</ref>
<ref id="B43">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wahl</surname>
<given-names>C. B.</given-names>
</name>
<name>
<surname>Aykol</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Swisher</surname>
<given-names>J. H.</given-names>
</name>
<name>
<surname>Montoya</surname>
<given-names>J. H.</given-names>
</name>
<name>
<surname>Suram</surname>
<given-names>S. K.</given-names>
</name>
<name>
<surname>Mirkin</surname>
<given-names>C. A.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Machine learning&#x2013;accelerated design and synthesis of polyelemental heterostructures</article-title>. <source>Sci. Adv.</source> <volume>7</volume>, <fpage>eabj5505</fpage>. <pub-id pub-id-type="doi">10.1126/sciadv.abj5505</pub-id>
</citation>
</ref>
<ref id="B44">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wang</surname>
<given-names>Y. J.</given-names>
</name>
<name>
<surname>Zhu</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Sha</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Hao</surname>
<given-names>W.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Low cycle fatigue life prediction of titanium alloy using genetic algorithm-optimized BP artificial neural network</article-title>. <source>Int. J. Fatigue</source> <volume>172</volume>, <fpage>107609</fpage>. <pub-id pub-id-type="doi">10.1016/j.ijfatigue.2023.107609</pub-id>
</citation>
</ref>
<ref id="B45">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wei</surname>
<given-names>Q.</given-names>
</name>
<name>
<surname>Cao</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Yuan</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>You</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Yu</surname>
<given-names>S.</given-names>
</name>
<etal/>
</person-group> (<year>2023</year>). <article-title>Divide and conquer: machine learning accelerated design of lead-free solder alloys with high strength and high ductility</article-title>. <source>npj Comput. Mater.</source> <volume>9</volume>, <fpage>201</fpage>. <pub-id pub-id-type="doi">10.1038/s41524-023-01150-0</pub-id>
</citation>
</ref>
<ref id="B46">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wu</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Zhou</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Yang</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Huang</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Ji</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Peng</surname>
<given-names>D.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>An ameliorated deep dense convolutional neural network for accurate recognition of casting defects in X-ray images</article-title>. <source>Knowledge-Based Syst.</source> <volume>226</volume>, <fpage>107096</fpage>. <pub-id pub-id-type="doi">10.1016/j.knosys.2021.107096</pub-id>
</citation>
</ref>
<ref id="B47">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wu</surname>
<given-names>C. T.</given-names>
</name>
<name>
<surname>Lin</surname>
<given-names>P. H.</given-names>
</name>
<name>
<surname>Huang</surname>
<given-names>S. Y.</given-names>
</name>
<name>
<surname>Tseng</surname>
<given-names>Y. J.</given-names>
</name>
<name>
<surname>Chang</surname>
<given-names>H. T.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>S. Y.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). <article-title>Revisiting alloy design of low-modulus biomedical &#x3b2;-Ti alloys using an artificial neural network</article-title>. <source>Materialia</source> <volume>21</volume>, <fpage>101313</fpage>. <pub-id pub-id-type="doi">10.1016/j.mtla.2021.101313</pub-id>
</citation>
</ref>
<ref id="B48">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yu</surname>
<given-names>J. X.</given-names>
</name>
<name>
<surname>Guo</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Han</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Lu</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>Q.</given-names>
</name>
<etal/>
</person-group> (<year>2019</year>). <article-title>A two-stage predicting model for &#x3b3;&#x27; solvus temperature of L1<sub>2</sub>-strengthened Co-base superalloys based on machine learning</article-title>. <source>Intermetallics</source> <volume>110</volume>, <fpage>106466</fpage>. <pub-id pub-id-type="doi">10.1016/j.intermet.2019.04.009</pub-id>
</citation>
</ref>
<ref id="B49">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhan</surname>
<given-names>Z. X.</given-names>
</name>
<name>
<surname>Hu</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Meng</surname>
<given-names>Q.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Data-driven fatigue life prediction in additive manufactured titanium alloy: a damage mechanics based machine learning framework</article-title>. <source>Eng. Fract. Mech.</source> <volume>252</volume>, <fpage>107850</fpage>. <pub-id pub-id-type="doi">10.1016/j.engfracmech.2021.107850</pub-id>
</citation>
</ref>
</ref-list>
</back>
</article>