<?xml version="1.0" encoding="us-ascii"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article article-type="research-article" dtd-version="2.3" xml:lang="EN" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Earth Sci.</journal-id>
<journal-title>Frontiers in Earth Science</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Earth Sci.</abbrev-journal-title>
<issn pub-type="epub">2296-6463</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="publisher-id">1340437</article-id>
<article-id pub-id-type="doi">10.3389/feart.2024.1340437</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Earth Science</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>An interpretable probabilistic prediction algorithm for shield movement performance</article-title>
<alt-title alt-title-type="left-running-head">Zhang et al.</alt-title>
<alt-title alt-title-type="right-running-head">
<ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3389/feart.2024.1340437">10.3389/feart.2024.1340437</ext-link>
</alt-title>
</title-group>
<contrib-group>
<contrib contrib-type="author">
<name>
<surname>Zhang</surname>
<given-names>Yapeng</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2726945/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/resources/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Liu</surname>
<given-names>Long</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Wu</surname>
<given-names>Jian</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Zeng</surname>
<given-names>Shaoxiang</given-names>
</name>
<xref ref-type="aff" rid="aff4">
<sup>4</sup>
</xref>
<xref ref-type="aff" rid="aff5">
<sup>5</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Hu</surname>
<given-names>Jianquan</given-names>
</name>
<xref ref-type="aff" rid="aff6">
<sup>6</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Tao</surname>
<given-names>Yuanqin</given-names>
</name>
<xref ref-type="aff" rid="aff4">
<sup>4</sup>
</xref>
<xref ref-type="aff" rid="aff5">
<sup>5</sup>
</xref>
<xref ref-type="corresp" rid="c001">&#x2a;</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2243518/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/funding-acquisition/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Huang</surname>
<given-names>Yong</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/resources/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Zhou</surname>
<given-names>Xuetao</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Liang</surname>
<given-names>Xu</given-names>
</name>
<xref ref-type="aff" rid="aff6">
<sup>6</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
<role content-type="https://credit.niso.org/contributor-roles/funding-acquisition/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
</contrib-group>
<aff id="aff1">
<sup>1</sup>
<institution>Power China Huadong Engineering Corporation Limited</institution>, <addr-line>Hangzhou</addr-line>, <country>China</country>
</aff>
<aff id="aff2">
<sup>2</sup>
<institution>Zhejiang Engineering Research Center of Green Mine Technology and Intelligent Equipment</institution>, <addr-line>Hangzhou</addr-line>, <country>China</country>
</aff>
<aff id="aff3">
<sup>3</sup>
<institution>Zhejiang Huadong Engineering Construction and Management Corporation Limited</institution>, <addr-line>Hangzhou</addr-line>, <country>China</country>
</aff>
<aff id="aff4">
<sup>4</sup>
<institution>College of Civil Engineering</institution>, <institution>Zhejiang University of Technology</institution>, <addr-line>Hangzhou</addr-line>, <country>China</country>
</aff>
<aff id="aff5">
<sup>5</sup>
<institution>Engineering Research Center of Ministry of Education for Renewable Energy Infrastructure Construction Technology</institution>, <addr-line>Hangzhou</addr-line>, <country>China</country>
</aff>
<aff id="aff6">
<sup>6</sup>
<institution>Hangzhou Urban Infrastructure Construction Management Center</institution>, <addr-line>Hangzhou</addr-line>, <country>China</country>
</aff>
<author-notes>
<fn fn-type="edited-by">
<p>
<bold>Edited by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/87809/overview">Manoj Khandelwal</ext-link>, Federation University Australia, Australia</p>
</fn>
<fn fn-type="edited-by">
<p>
<bold>Reviewed by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/2386656/overview">Muzaffer Can Iban</ext-link>, Mersin University, T&#xfc;rkiye</p>
<p>
<ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1514549/overview">Hossein Moayedi</ext-link>, Southern Illinois University Edwardsville, United States</p>
</fn>
<corresp id="c001">&#x2a;Correspondence: Yuanqin Tao, <email>taoyuanqin@zju.edu.cn</email>
</corresp>
</author-notes>
<pub-date pub-type="epub">
<day>01</day>
<month>07</month>
<year>2024</year>
</pub-date>
<pub-date pub-type="collection">
<year>2024</year>
</pub-date>
<volume>12</volume>
<elocation-id>1340437</elocation-id>
<history>
<date date-type="received">
<day>01</day>
<month>12</month>
<year>2023</year>
</date>
<date date-type="accepted">
<day>03</day>
<month>06</month>
<year>2024</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#xa9; 2024 Zhang, Liu, Wu, Zeng, Hu, Tao, Huang, Zhou and Liang.</copyright-statement>
<copyright-year>2024</copyright-year>
<copyright-holder>Zhang, Liu, Wu, Zeng, Hu, Tao, Huang, Zhou and Liang</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<abstract>
<p>Total thrust and torque are two key indicators of shield movement performance. Most existing data-driven machine learning studies focus on developing more accurate models for predicting total thrust and torque but overlook the interpretability of the models. To address this black-box issue, this study proposes an interpretable probabilistic prediction algorithm for the shield movement performance. The algorithm uses the natural gradient boosting (NGBoost) model to iteratively update the parametric probability distributions (e.g., mean and variance) and achieve probabilistic predictions of the total thrust and torque. The impact of each feature on the prediction values and uncertainty is quantified by extending the importance analysis of a single deterministic predictive value to both the mean and variance. The feature interactions are analyzed and their predictive contributions are quantified by the shapley additive explanations (SHAP) method. The transparency of the NGBoost model is improved through the visualization of the decision-making process. A shield tunneling project in Hangzhou is used to validate the effectiveness of the proposed algorithm. The results indicate that the NGboost model outperforms other five models in terms of accuracy. The prediction results are interpretable, and the interpretable probabilistic model provides decision-makers with a more intuitive and reliable reference.</p>
</abstract>
<kwd-group>
<kwd>shield movement performance</kwd>
<kwd>probabilistic prediction</kwd>
<kwd>model interpretability</kwd>
<kwd>NGBoost</kwd>
<kwd>shap</kwd>
</kwd-group>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-at-acceptance</meta-name>
<meta-value>Geohazards and Georisks</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec id="s1">
<title>1 Introduction</title>
<p>With the progression of urbanization and the increasing demand for underground space, shield tunneling technology has become increasingly important in modern urban infrastructure development (<xref ref-type="bibr" rid="B46">Zhou et al., 2023</xref>). However, poor shield movement performance (total thrust and torque) can lead to inefficient tunneling, excessive ground settlements, tunnel structural damage, and even pose serious threats to surrounding buildings and infrastructures (<xref ref-type="bibr" rid="B4">Chen et al., 2019</xref>). Total thrust refers to the axial force generated by the shield machine to advance forward. It is produced by hydraulic jacks, mechanical rams, or other mechanisms that push against the tunnel face. Torque is the rotational force exerted on cutting tools to break and loosen the soil or rock. Appropriate monitoring, adjustment, and optimization of total thrust and torque are essential for maximizing productivity and reducing risks in shield tunneling projects.</p>
<p>Due to the rapid advancements in big data and computational capabilities, machine learning has received increasing attention in the field of geotechnical engineering (<xref ref-type="bibr" rid="B26">Moayedi et al., 2020</xref>; <xref ref-type="bibr" rid="B43">Zhang et al., 2021</xref>; <xref ref-type="bibr" rid="B2">Baghbani et al., 2022</xref>; <xref ref-type="bibr" rid="B17">Kannangara et al., 2022</xref>; <xref ref-type="bibr" rid="B35">Tao et al., 2022a</xref>; <xref ref-type="bibr" rid="B36">Tao et al., 2022b</xref>; <xref ref-type="bibr" rid="B44">Zhang et al., 2022a</xref>; <xref ref-type="bibr" rid="B30">Phoon and Zhang, 2023</xref>). Machine learning is a powerful tool that can extract nonlinear relationships among features, leading to a better understanding and prediction of geotechnical behaviors. In tunnel constructions, machine learning models have been successfully applied to predict the total thrust and torque (<xref ref-type="bibr" rid="B21">Lin et al., 2022a</xref>; <xref ref-type="bibr" rid="B19">Li et al., 2023a</xref>; <xref ref-type="bibr" rid="B20">Li et al., 2023b</xref>; <xref ref-type="bibr" rid="B42">Yu et al., 2023</xref>). For example, <xref ref-type="bibr" rid="B11">Gao et al. (2019)</xref> employed three recurrent neural network (RNN) models, including basic RNN, long short-term memory (LSTM), and gated recurrent unit (GRU), to predict the total thrust and torque of the shield machine. These models were selected for their inherent ability to process time-series data, effectively capturing dynamic characteristics over time. However, manual parameter tuning of these models is complex and prone to getting stuck in local optima. Based on this, <xref ref-type="bibr" rid="B22">Lin et al. (2022b)</xref> developed a hybrid model combining particle swarm optimization (PSO) and GRU for torque predictions. PSO is a population-based optimization algorithm that simulates the foraging behavior of flocks of birds to select the best hyperparameters automatically. <xref ref-type="bibr" rid="B9">Elbaz et al. (2023)</xref> applied reinforcement learning to optimize the process of PSO hyperparameter tuning, which further improved the accuracy of torque and total thrust predictions. It has been found in practice that, although adding additional optimization layers (such as using reinforcement learning to optimize PSO parameters) can enhance prediction accuracy, the process is complex and time-consuming, and the accuracy improvements are not so significant in most cases. From the perspective of feature engineering, <xref ref-type="bibr" rid="B34">Shi et al. (2021)</xref> used the variational mode decomposition and the empirical wavelet transform to preprocess the raw dataset, which can also improve the prediction accuracy of shield tunneling parameters. The studies above focus on developing innovative algorithms to improve the model accuracy for total thrust and torque, but to some extent overlook the interpretability of the model. Due to the black-box nature of machine learning models, especially deep learning models like LSTM that involve multiple parameters and layers, the prediction results and decision-making processes are difficult to explain. This lack of transparency results in decision-makers having insufficient confidence and a skeptical attitude toward applying these models. <xref ref-type="bibr" rid="B41">Xu et al. (2021)</xref> compared the performance of various machine learning models for predicting the total thrust and torque, demonstrating that random forest, a tree-based model, offered the best balance between model accuracy and computation time while also being more interpretable compared to other models. Given its performance and greater transparency relative to complex models like LSTM, tree-based models like random forests provide a viable option for interpretable total thrust and torque predictions.</p>
<p>A series of interpretable methods, such as shapley additive explanations (SHAP) (<xref ref-type="bibr" rid="B24">Lundberg and Lee, 2017</xref>) and causal artificial intelligence methods (<xref ref-type="bibr" rid="B18">Kuang et al., 2020</xref>; <xref ref-type="bibr" rid="B39">Wang et al., 2023</xref>), have been proposed to make complex machine learning models more transparent and understandable (<xref ref-type="bibr" rid="B47">Zhou et al., 2021</xref>; <xref ref-type="bibr" rid="B15">Iban, 2022</xref>; <xref ref-type="bibr" rid="B40">Wen et al., 2023</xref>; <xref ref-type="bibr" rid="B6">Das et al., 2024</xref>). For example, <xref ref-type="bibr" rid="B16">Iban and Bilgilioglu (2023)</xref> employed the local explainable artificial intelligence method of SHAP to attain the contribution of each factor to the avalanches. Based on SHAP values, the most critical factors triggering avalanches were identified. <xref ref-type="bibr" rid="B33">Scavuzzo et al. (2022)</xref> used SHAP with extreme gradient boosting (XGBoost) for geospatial health prediction, aiming to understand the impact of each input feature on the predicted outcomes. Similarly, <xref ref-type="bibr" rid="B28">Parsa et al. (2020)</xref> used XGBoost to predict real-time traffic accidents and employed SHAP to analyze the factors leading to risk. Although SHAP has been successfully applied in various fields, its effectiveness in providing interpretability for thrust and torque predictions remains unclear. While SHAP can explain how features affect predictions, it still lacks transparency in clarifying the decision-making process. Moreover, it is commonly acknowledged that there is inevitable uncertainty in geotechnical engineering (<xref ref-type="bibr" rid="B29">Phoon and Kulhawy, 1999</xref>; <xref ref-type="bibr" rid="B12">Gu et al., 2023</xref>; <xref ref-type="bibr" rid="B37">Tao et al., 2023</xref>; <xref ref-type="bibr" rid="B38">Tao, et al., 2024</xref>). However, the existing studies on shield tunneling often overlook such uncertainty, which may result in potential safety risks of the project.</p>
<p>This paper thoroughly compares and analyzes multiple machine learning models used in current shield tunneling prediction research, summarizing the findings as follows: First, although time series models like LSTM can effectively utilize lagged time series information to predict shield movement performance, tuning their parameters is notably complex and time-consuming. Second, the multi-layered and opaque structure of LSTM models complicates the interpretation of their decision-making processes, adversely affecting model transparency and explainability. Moreover, while the SHAP method is a mature interpretability technique that has proven effective in other domains, its effectiveness for predicting total thrust and torque in shield tunneling applications has not yet been validated. Crucially, SHAP struggles to assess the importance of variance, which limits its utility in conducting uncertainty analysis for risk assessments. To address these limitations, this paper proposes an interpretable probabilistic prediction algorithm for predicting the movement performance of shield machines (i.e., total thrust and torque). By constructing a natural gradient boosting (NGBoost) probabilistic prediction model, the approach not only provides predictions for total thrust and torque but also offers prediction uncertainty. This is essential for comprehensive risk assessment, as improper loads can lead to equipment failures, ground collapse, or tunnel instability. Unlike standard tree model importance analyses, which typically evaluate only the mean predictions, this study extends the importance analysis to include both the mean and variance of the model. This allows for a more comprehensive understanding of how each feature impacts the overall prediction outcomes and associated uncertainties. The model can quantitatively assess how geological conditions and shield tunneling parameters specifically affect predictions of total thrust and torque. The interactions between features are analyzed, and their contributions to the predictions are quantified using the SHAP method. The predictive mechanism of the model is explained from both global and local perspectives, enhancing the transparency and explainability of decision-making. Furthermore, the transparency of the model is further enhanced by visualizing the decision path. This visualization not only helps operators better understand how the model predicts total thrust and torque but can also show how the model responds to varying geological conditions, thus optimizing decision-making and preventing potential risks.</p>
</sec>
<sec sec-type="methods" id="s2">
<title>2 Methodology</title>
<p>This section is composed of three main parts: first, the principle of the NGBoost model is introduced to provide a theoretical basis for the subsequent model interpretability analysis; second, the interpretability of model predictions is explored; finally, an overview of the proposed interpretable probabilistic algorithm is given.</p>
<sec id="s2-1">
<title>2.1 Principle of the NGBoost model</title>
<p>Natural gradient boosting (NGBoost) (<xref ref-type="bibr" rid="B8">Duan et al., 2020</xref>) is an advanced model designed for probabilistic forecasting. <xref ref-type="fig" rid="F1">Figure 1</xref> shows the structure of the NGBoost model. Unlike gradient boosting and XGBoost methods (<xref ref-type="bibr" rid="B10">Friedman, 2001</xref>; <xref ref-type="bibr" rid="B5">Chen and Guestrin, 2016</xref>) that focus on point predictions, NGBoost emphasizes the entire probability distribution, making it possible to estimate prediction uncertainty. In addition, an iterative method is employed to improve model predictions based on the errors in the previous steps. This approach consists of several key steps:</p>
<fig id="F1" position="float">
<label>FIGURE 1</label>
<caption>
<p>Structure of the NGBoost model.</p>
</caption>
<graphic xlink:href="feart-12-1340437-g001.tif"/>
</fig>
<p>The natural gradient (<xref ref-type="bibr" rid="B1">Amari, 1998</xref>) is computed firstly. In a standard gradient descent, the geometric structure of the parameter space is not considered, resulting in potential instability and inefficiency during optimization. By considering the local curvature, the natural gradient can alleviate these problems. Specifically, the Fisher information matrix measures the local curvature of the parameter space. It shows how changes in pairs of parameters affect the output distribution of the model. For a parameterized distribution model, the Fisher information matrix is defined in Eq. <xref ref-type="disp-formula" rid="e1">(1)</xref>:<disp-formula id="e1">
<mml:math id="m1">
<mml:mrow>
<mml:mi>I</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi>&#x3b8;</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x3d;</mml:mo>
<mml:mi>E</mml:mi>
<mml:mrow>
<mml:mfenced open="[" close="]" separators="|">
<mml:mrow>
<mml:msub>
<mml:mo>&#x2207;</mml:mo>
<mml:mi>&#x3b8;</mml:mi>
</mml:msub>
<mml:mo>&#x2061;</mml:mo>
<mml:mi>log</mml:mi>
<mml:mo>&#x2061;</mml:mo>
<mml:mi>p</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi>y</mml:mi>
<mml:mo>&#x7c;</mml:mo>
<mml:mi>&#x3b8;</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:msub>
<mml:mo>&#x2207;</mml:mo>
<mml:mi>&#x3b8;</mml:mi>
</mml:msub>
<mml:mo>&#x2061;</mml:mo>
<mml:mi>log</mml:mi>
<mml:mo>&#x2061;</mml:mo>
<mml:mi>p</mml:mi>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi>y</mml:mi>
<mml:mo>&#x7c;</mml:mo>
<mml:mi>&#x3b8;</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mi>T</mml:mi>
</mml:msup>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(1)</label>
</disp-formula>where <italic>&#x3b8;</italic> represents the parameter of the model; <italic>I</italic>(<italic>&#x3b8;</italic>) represents the Fisher information matrix; <italic>E</italic> denotes the expectation, which is an average of all possible outcomes; &#x2207;<sub>
<italic>&#x3b8;</italic>
</sub> is the gradient; log <italic>p</italic> (<italic>y</italic> &#x7c; <italic>&#x3b8;</italic>) is the log-likelihood of the outcome <italic>y</italic> given the parameter <italic>&#x3b8;</italic>. The natural gradient &#x2207;<sub>natural</sub> is then computed by using the Fisher information matrix to correct the standard gradient, as shown in Eq. <xref ref-type="disp-formula" rid="e2">(2)</xref>:<disp-formula id="e2">
<mml:math id="m2">
<mml:mrow>
<mml:msub>
<mml:mo>&#x2207;</mml:mo>
<mml:mtext>natural</mml:mtext>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mi>I</mml:mi>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi>&#x3b8;</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msup>
<mml:msub>
<mml:mo>&#x2207;</mml:mo>
<mml:mi>&#x3b8;</mml:mi>
</mml:msub>
<mml:mi>L</mml:mi>
</mml:mrow>
</mml:math>
<label>(2)</label>
</disp-formula>where <italic>L</italic> is the loss function, and <italic>I</italic>(<italic>&#x3b8;</italic>)<sup>&#x2212;1</sup> represents the inverse of the Fisher information matrix <italic>I</italic>(<italic>&#x3b8;</italic>). The natural gradient gives a direction to the model optimization so that the prediction errors are minimized more efficiently. This adjustment ensures that parameter updates are in line with the geometry of the parameter space, leading to faster convergence rates and improved robustness.</p>
<p>After calculating the natural gradient, the model residuals can then be determined, which represent the prediction errors in the current iteration. Specifically, <italic>y</italic>
<sub>actual</sub> represents the actual measurement, and <inline-formula id="inf1">
<mml:math id="m3">
<mml:mrow>
<mml:msubsup>
<mml:mi>y</mml:mi>
<mml:mtext>pred</mml:mtext>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:msubsup>
</mml:mrow>
</mml:math>
</inline-formula> is the predicted value in the first iteration. Then, the residual in the first iteration is given by <inline-formula id="inf2">
<mml:math id="m4">
<mml:mrow>
<mml:msub>
<mml:mi>r</mml:mi>
<mml:mn>1</mml:mn>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:msub>
<mml:mi>y</mml:mi>
<mml:mtext>actual</mml:mtext>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msubsup>
<mml:mi>y</mml:mi>
<mml:mtext>pred</mml:mtext>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:msubsup>
</mml:mrow>
</mml:math>
</inline-formula>. It is worth noting that the residuals in NGBoost include both the mean <italic>&#x3bc;</italic> and standard deviation <italic>&#x3c3;</italic> of the predicted distribution. Within the framework of NGBoost, &#x201c;base learner fitting&#x201d; is a crucial step in the iterative training loop. This step involves selecting and training a base learner to fit the residual <italic>r</italic>
<sub>1</sub> of the current model. In this study, the decision tree is chosen as the base learner due to its simple structure and high interpretability. In each iteration, the model takes the residual <italic>r</italic>
<sub>1</sub> as the input and makes predictions through a series of decision nodes. For example, as shown by the red path in <xref ref-type="fig" rid="F1">Figure 1</xref>, the decision tree processes the residuals by examining &#x201c;Feature 2&#x201d;. The data point moves to the left when the value of &#x201c;Feature 2&#x201d; &#x2264; 168 and &#x2264;76 and then goes to the right when the value &#x3e;30. The data point finally reaches the leaf node, which provides the predicted residual <inline-formula id="inf3">
<mml:math id="m5">
<mml:mrow>
<mml:msubsup>
<mml:mi>r</mml:mi>
<mml:mrow>
<mml:mtext>&#x2002;</mml:mtext>
<mml:mtext>pred</mml:mtext>
</mml:mrow>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:msubsup>
</mml:mrow>
</mml:math>
</inline-formula>.</p>
<p>Once the predicted residual <inline-formula id="inf4">
<mml:math id="m6">
<mml:mrow>
<mml:msubsup>
<mml:mi>r</mml:mi>
<mml:mrow>
<mml:mtext>pred</mml:mtext>
</mml:mrow>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:msubsup>
</mml:mrow>
</mml:math>
</inline-formula> is obtained from the leaf node, the first prediction <inline-formula id="inf5">
<mml:math id="m7">
<mml:mrow>
<mml:msubsup>
<mml:mi>y</mml:mi>
<mml:mtext>pred</mml:mtext>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:msubsup>
</mml:mrow>
</mml:math>
</inline-formula> can be adjusted accordingly. Specifically, <inline-formula id="inf6">
<mml:math id="m8">
<mml:mrow>
<mml:msubsup>
<mml:mi>y</mml:mi>
<mml:mtext>pred</mml:mtext>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:msubsup>
</mml:mrow>
</mml:math>
</inline-formula> is updated by adding the predicted residual <inline-formula id="inf7">
<mml:math id="m9">
<mml:mrow>
<mml:msubsup>
<mml:mi>r</mml:mi>
<mml:mrow>
<mml:mtext>pred</mml:mtext>
</mml:mrow>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:msubsup>
</mml:mrow>
</mml:math>
</inline-formula> (<inline-formula id="inf8">
<mml:math id="m10">
<mml:mrow>
<mml:msubsup>
<mml:mi>y</mml:mi>
<mml:mtext>pred</mml:mtext>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:msubsup>
<mml:mo>&#x3d;</mml:mo>
<mml:msubsup>
<mml:mi>y</mml:mi>
<mml:mtext>pred</mml:mtext>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:msubsup>
<mml:mo>&#x2b;</mml:mo>
<mml:msubsup>
<mml:mi>r</mml:mi>
<mml:mrow>
<mml:mtext>&#x2002;</mml:mtext>
<mml:mtext>pred</mml:mtext>
</mml:mrow>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:msubsup>
</mml:mrow>
</mml:math>
</inline-formula>). After <italic>n</italic> iterations, the model reaches the preset maximum iteration number, and the model training stops. The final prediction accumulates results from all decision trees throughout the whole iterations (<inline-formula id="inf9">
<mml:math id="m11">
<mml:mrow>
<mml:msubsup>
<mml:mi>y</mml:mi>
<mml:mtext>pred</mml:mtext>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi>n</mml:mi>
<mml:mo>&#x2b;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:msubsup>
<mml:mo>&#x3d;</mml:mo>
<mml:msubsup>
<mml:mi>y</mml:mi>
<mml:mtext>pred</mml:mtext>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:msubsup>
<mml:mo>&#x2b;</mml:mo>
<mml:mstyle displaystyle="true">
<mml:mover>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>n</mml:mi>
<mml:mo>&#x2b;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:mover>
</mml:mstyle>
<mml:msubsup>
<mml:mi>r</mml:mi>
<mml:mrow>
<mml:mtext>pred</mml:mtext>
</mml:mrow>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:msubsup>
</mml:mrow>
</mml:math>
</inline-formula>).</p>
</sec>
<sec id="s2-2">
<title>2.2 Interpretability methods for NGBoost model</title>
<p>Most existing machine learning models are seen as black boxes, with unclear principles and decision processes (<xref ref-type="bibr" rid="B13">Guidotti et al., 2018</xref>). To address this issue, this study interprets the model from two main perspectives: global interpretability and local interpretability. For local interpretability, the decision tree-based NGBoost model can visualize the decision-making process to track and explain how an individual prediction is made through the branches of the trees. In addition, the SHAP method (<xref ref-type="bibr" rid="B24">Lundberg and Lee, 2017</xref>) is employed to explain the interactions between multiple features and their contributions to the target individual prediction. SHAP utilizes the principles of Shapley values from cooperative game theory to quantify the importance of each feature within the NGBoost model. Within cooperative game theory, Shapley values offer a way to fairly share game rewards, ensuring participants receive compensation in line with their overall contribution to the total gain. When applied in machine learning, SHAP assesses how much each feature contributes to the model prediction, specifically how each feature impacts the final prediction result. Through this process, SHAP uncovers the importance and function of each feature within the internal decision-making mechanism of the model, significantly improving model transparency.</p>
<p>For a given feature <italic>x</italic>
<sub>
<italic>j</italic>
</sub>, its Shapley value is computed as shown in Eq. <xref ref-type="disp-formula" rid="e3">(3)</xref>:<disp-formula id="e3">
<mml:math id="m12">
<mml:mrow>
<mml:msub>
<mml:mi>&#x3d5;</mml:mi>
<mml:mi>j</mml:mi>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mstyle displaystyle="true">
<mml:munder>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mo>&#x2286;</mml:mo>
<mml:mi>N</mml:mi>
<mml:mo>\</mml:mo>
<mml:mrow>
<mml:mfenced open="{" close="}" separators="|">
<mml:mrow>
<mml:mi>j</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:munder>
</mml:mstyle>
<mml:mfrac>
<mml:mrow>
<mml:mi>n</mml:mi>
<mml:mo>!</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mrow>
<mml:mfenced open="|" close="|" separators="|">
<mml:mrow>
<mml:mi>S</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>!</mml:mo>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi>n</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mrow>
<mml:mfenced open="|" close="|" separators="|">
<mml:mrow>
<mml:mi>S</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>!</mml:mo>
</mml:mrow>
</mml:mfrac>
<mml:mrow>
<mml:mfenced open="[" close="]" separators="|">
<mml:mrow>
<mml:mi>f</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mo>&#x222a;</mml:mo>
<mml:mrow>
<mml:mfenced open="{" close="}" separators="|">
<mml:mrow>
<mml:mi>j</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:mi>f</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi>S</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(3)</label>
</disp-formula>where <italic>N</italic> denotes a set of all features; <italic>S</italic> is a subset excluding <italic>x</italic>
<sub>
<italic>j</italic>
</sub>; <italic>n</italic> is the total number of features in the set <italic>N</italic>; &#x7c;<italic>S</italic>&#x7c; is the size (number of elements) of the subset <italic>S</italic>; <italic>&#x3d5;</italic>
<sub>
<italic>j</italic>
</sub> represents the Shapley value of feature; <italic>N</italic>\{<italic>j</italic>} means the set of all features except <italic>x</italic>
<sub>
<italic>j</italic>
</sub>; and <inline-formula id="inf10">
<mml:math id="m13">
<mml:mrow>
<mml:mi>f</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mo>&#x222a;</mml:mo>
<mml:mrow>
<mml:mfenced open="{" close="}" separators="|">
<mml:mrow>
<mml:mi>j</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula> represents the output of a function when the feature subset <italic>S</italic> and the feature <italic>x</italic>
<sub>
<italic>j</italic>
</sub> are used together.</p>
<p>For global interpretability, SHAP summary plots are used to help understand the overall trends and patterns in model predictions by assessing the average contribution of features across all data points, revealing which features have the greatest impact on the model prediction. In addition, the decision tree-based NGBoost model can also show the global importance analysis of each feature on the target prediction, including the mean and the standard deviation of the prediction, aiming to provide a global analysis method based on the model itself for assessing prediction accuracy and uncertainty. The importance of each feature is evaluated using Eq. <xref ref-type="disp-formula" rid="e4">(4)</xref>:<disp-formula id="e4">
<mml:math id="m14">
<mml:mrow>
<mml:mi>F</mml:mi>
<mml:mi>I</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mi>j</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x3d;</mml:mo>
<mml:mstyle displaystyle="true">
<mml:munder>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x2208;</mml:mo>
<mml:mtext>nodes</mml:mtext>
</mml:mrow>
</mml:munder>
</mml:mstyle>
<mml:mi>I</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#xd7;</mml:mo>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mtext>Error</mml:mtext>
<mml:mtext>parent</mml:mtext>
</mml:msub>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mtext>Error</mml:mtext>
<mml:mtext>child</mml:mtext>
</mml:msub>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(4)</label>
</disp-formula>where <italic>FI</italic>(<italic>x</italic>
<sub>
<italic>j</italic>
</sub>) denotes the importance of feature <italic>x</italic>
<sub>
<italic>j</italic>
</sub>; <italic>I</italic>(<italic>i</italic>) is an indicator function, which is equal to one if the node <italic>i</italic> splits and 0 otherwise; <inline-formula id="inf11">
<mml:math id="m15">
<mml:mrow>
<mml:msub>
<mml:mtext>Error</mml:mtext>
<mml:mtext>parent</mml:mtext>
</mml:msub>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula> and <inline-formula id="inf12">
<mml:math id="m16">
<mml:mrow>
<mml:msub>
<mml:mtext>Error</mml:mtext>
<mml:mtext>child</mml:mtext>
</mml:msub>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula> represent the error of the parent node before splitting and the error of the child node after splitting, respectively.</p>
</sec>
<sec id="s2-3">
<title>2.3 Flowchart of the interpretable probabilistic prediction algorithm</title>
<p>The proposed algorithm for shield movement performance has five main steps: data acquisition, data preprocessing, model construction, model evaluation, and model interpretability, as shown in <xref ref-type="fig" rid="F2">Figure 2</xref>. For the tunnel boring excavation, data are acquired from the installed sensors on the shield machine equipment. Typical parameters include the total thrust, torque, advance rate, <italic>etc</italic>. These data are sequentially stored in a secured database, either locally or in the cloud. In the data preprocessing step, the downtime data and outliers are first removed. The reserved dataset is then divided into a training set and a test set at a ratio of 8:2. In the model construction step, the NGBoost model is constructed using a decision tree as a base learner. Decision trees are chosen because they are highly interpretable (<xref ref-type="bibr" rid="B32">Quinlan, 1986</xref>). The model undergoes an iterative training loop that consists of natural gradient calculation for optimization, fitting of base learners, and adaptive model updating. Specifically, the NGBoost model uses the natural gradient descent to adjust the model parameters. The base decision tree fits the residuals from the previous iteration and serves as the current base learner. A natural gradient is then computed from the loss function MSE (mean squared error) and model parameters, guiding an adaptive parameter update. This process is repeated until a predefined maximum iteration number is met. Next, the trained NGBoost model can provide probabilistic prediction results due to its inherent capability to estimate the distribution of the target variable. Optimal hyperparameters and the superior model are determined through comparative analysis. Finally, based on the prediction results, the importance and influence of each feature on the target variable are analyzed. By visualizing the decision-making process, the transparency of the model can be improved.</p>
<fig id="F2" position="float">
<label>FIGURE 2</label>
<caption>
<p>Flowchart of shield movement performance predictions.</p>
</caption>
<graphic xlink:href="feart-12-1340437-g002.tif"/>
</fig>
</sec>
</sec>
<sec id="s3">
<title>3 Case study</title>
<sec id="s3-1">
<title>3.1 Data acquisition</title>
<p>A shield tunneling project in Hangzhou is used to illustrate the proposed method. The dataset has real-time tunneling information from 100 rings, all gathered by the monitoring system. Specifically, this tunneling dataset is collected from seven subsystems: the propulsion system, cylinder stroke system, mud and water silo pressure system, air cushion silo pressure system, cutter system, grouting system, and shield attitude system, totaling 50 tunneling parameters. Among these parameters, the total thrust and torque directly reflect the shield movement performance. The values of total thrust and torque at the next moment are chosen as the output for the prediction model (<xref ref-type="bibr" rid="B27">Ntoutsi et al., 2020</xref>). Advance rate is the distance that the shield machine tunnels within a unit of time. Penetration is an important indicator for assessing the propulsive capacity of a shield machine. A higher penetration means the shield can traverse hard geological conditions more easily. Ring number is usually used to label segments of the tunnel, and it is an indicator of construction progress. The aforementioned three parameters (advance rate, penetration, and ring number) are closely related to the shield movement performance. Therefore, they are selected as input parameters for the predictive model. In addition, as this task is a time-series prediction, the future states are influenced by the past states. Therefore, the torque and total thrust at the current moment are also selected as the input parameters for the model. The monitoring system automatically logs data every 10 s, resulting in approximately 350,000 data samples.</p>
</sec>
<sec id="s3-2">
<title>3.2 Data preprocessing</title>
<p>Data preprocessing is crucial to ensure data quality and model accuracy (<xref ref-type="bibr" rid="B23">Liu et al., 2019</xref>; <xref ref-type="bibr" rid="B45">Zhang et al., 2022b</xref>). The data preprocessing in this study includes removing the downtime data and the outliers.</p>
<p>In practical engineering projects, there is a significant amount of downtime data for shield machines due to cutterhead replacement, electrical circuit inspections, lubrication system maintenance, and so on. <xref ref-type="fig" rid="F3">Figure 3</xref> shows the change in advance rate during the shield tunneling process. A complete shield tunneling operational cycle, as illustrated in the upper subfigure, consists of a rising phase, a stable phase, and a declining phase. When the advance rate &#x3d; 0, it indicates that the shield machine has stopped tunneling. In this case, the downtime data cannot provide practical value for predicting the shield movement performance, therefore, they are removed, and 66,224 data samples are then reserved.</p>
<fig id="F3" position="float">
<label>FIGURE 3</label>
<caption>
<p>Advance rate during the shield tunneling process.</p>
</caption>
<graphic xlink:href="feart-12-1340437-g003.tif"/>
</fig>
<p>The raw dataset may include outliers caused by factors such as sensor errors or human operational mistakes, which could negatively affect the training of the model. To solve this problem, this paper firstly divides the 66,224 samples into a training set and a test set at a ratio of 8:2. Boxplot analysis has been proved as an effective method for identifying and addressing outliers in shield tunneling, as shown in <xref ref-type="bibr" rid="B14">Hou et al. (2022)</xref>, <xref ref-type="bibr" rid="B25">Ma et al. (2024)</xref>, and <xref ref-type="bibr" rid="B3">Chen et al. (2024)</xref>. Given its proven efficacy, this technique is utilized to discern and remove outliers from the training set. Specifically, the boxplot analysis is used to identify and eliminate the outliers in the training set to improve the data quality. <xref ref-type="fig" rid="F4">Figure 4</xref> illustrates the outlier removal for the total thrust and torque. In the boxplot, the central horizontal line rep-resents the median, while the bottom and top of the box denote the first and third quartiles, respectively. The upper and lower edges of the boxplot represent the maxi-mum and minimum values of the data. The scattered points represent the outliers. There are 1782 outliers for total thrust and 1,508 outliers for torque in the training set. To improve data quality, records including these outliers are removed (<xref ref-type="bibr" rid="B31">Qin et al., 2023</xref>).</p>
<fig id="F4" position="float">
<label>FIGURE 4</label>
<caption>
<p>Outliers for the total thrust and torque (training set).</p>
</caption>
<graphic xlink:href="feart-12-1340437-g004.tif"/>
</fig>
</sec>
<sec id="s3-3">
<title>3.3 Model construction</title>
<p>Based on the preprocessed dataset, a shield movement performance prediction model is developed. The hyperparameters of the NGBoost model are determined through random search, and the last 20% of the training set is used as a validation set for model evaluation. The optimal hyperparameters are summarized in <xref ref-type="table" rid="T1">Table 1</xref>. The &#x201c;n_estimators&#x201d; parameter is set to 300 iterations. The normal distribution is selected as the target distribution for the NGBoost model. This is because when the number of independent random variables is large enough, their sum tends to approximate a normal distribution. In regression problems, the target variable can be considered as the sum of multiple factors (<xref ref-type="bibr" rid="B7">Devore, 2011</xref>). To determine the learning rate and the base learner combination, a systematic comparison of model performance is then conducted. Mean absolute error (MAE), R-squared (<italic>R</italic>
<sup>2</sup>), and Root mean square error (RMSE) shown in Eqs <xref ref-type="disp-formula" rid="e5">5</xref>&#x2013;<xref ref-type="disp-formula" rid="e7">7</xref> are used as the evaluation metrics.<disp-formula id="e5">
<mml:math id="m17">
<mml:mrow>
<mml:msup>
<mml:mi mathvariant="normal">R</mml:mi>
<mml:mn>2</mml:mn>
</mml:msup>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
<mml:mo>&#x2212;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:mo>&#x2211;</mml:mo>
</mml:mstyle>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi>y</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mstyle displaystyle="true">
<mml:mover accent="true">
<mml:mi>y</mml:mi>
<mml:mo>&#x5e;</mml:mo>
</mml:mover>
</mml:mstyle>
<mml:mi>i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:mo>&#x2211;</mml:mo>
</mml:mstyle>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi>y</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mstyle displaystyle="true">
<mml:mover accent="true">
<mml:mi>y</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:mstyle>
<mml:mi>i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:math>
<label>(5)</label>
</disp-formula>
<disp-formula id="e6">
<mml:math id="m18">
<mml:mrow>
<mml:mtext>MAE</mml:mtext>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:mfrac>
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:munderover>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi>n</mml:mi>
</mml:munderover>
</mml:mstyle>
<mml:mrow>
<mml:mfenced open="|" close="|" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi>y</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mstyle displaystyle="true">
<mml:mover accent="true">
<mml:mi>y</mml:mi>
<mml:mo>&#x5e;</mml:mo>
</mml:mover>
</mml:mstyle>
<mml:mi>i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(6)</label>
</disp-formula>
<disp-formula id="e7">
<mml:math id="m19">
<mml:mrow>
<mml:mtext>RMSE</mml:mtext>
<mml:mo>&#x3d;</mml:mo>
<mml:msqrt>
<mml:mrow>
<mml:mfrac>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:mfrac>
<mml:mstyle displaystyle="true">
<mml:munderover>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi>n</mml:mi>
</mml:munderover>
</mml:mstyle>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi>y</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mstyle displaystyle="true">
<mml:mover accent="true">
<mml:mi>y</mml:mi>
<mml:mo>&#x5e;</mml:mo>
</mml:mover>
</mml:mstyle>
<mml:mi>i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:msqrt>
</mml:mrow>
</mml:math>
<label>(7)</label>
</disp-formula>where <italic>n</italic> represents the number of measurements, <italic>y</italic>
<sub>
<italic>i</italic>
</sub> is the actual value for the <italic>i</italic>th measurement. <inline-formula id="inf13">
<mml:math id="m20">
<mml:mrow>
<mml:msub>
<mml:mstyle displaystyle="true">
<mml:mover accent="true">
<mml:mi>y</mml:mi>
<mml:mo>&#x5e;</mml:mo>
</mml:mover>
</mml:mstyle>
<mml:mi>i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> is the prediction for the <italic>i</italic>th measurement; <inline-formula id="inf14">
<mml:math id="m21">
<mml:mrow>
<mml:msub>
<mml:mstyle displaystyle="true">
<mml:mover accent="true">
<mml:mi>y</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:mstyle>
<mml:mi>i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> is the mean of the measurements.</p>
<table-wrap id="T1" position="float">
<label>TABLE 1</label>
<caption>
<p>Optimal hyperparameter settings for the NGBoost model.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="center">Hyperparameter</th>
<th align="center">Description</th>
<th align="center">Value</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="center">n_estimators</td>
<td align="center">Number of boosting iterations to be run</td>
<td align="center">300</td>
</tr>
<tr>
<td align="center">Learning rate</td>
<td align="center">A parameter controlling the step size during model updates. The value bound is (0,1)</td>
<td align="center">0.001</td>
</tr>
<tr>
<td align="center">Base learner</td>
<td align="center">The individual learning algorithm or model used in ensemble methods</td>
<td align="center">Decision tree</td>
</tr>
<tr>
<td align="center">Distribution</td>
<td align="center">Probability distribution of the targets</td>
<td align="center">Normal</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>The selection of base learner and learning rate has significant impact on the model performance. <xref ref-type="fig" rid="F5">Figure 5</xref> shows the model performance for the total thrust and torque. The decision tree is a simple and effective model for making decisions based on previous data, and random forest is an ensemble method that aggregates multiple decision trees. The two models both show steady performances across different learning rates. Extreme gradient boosting (XGBoost) improves predictions by adding weak learners iteratively, but it does not perform well in this case. Specifically, for the total thrust shown in <xref ref-type="fig" rid="F5">Figure 5A</xref>, the decision tree with a learning rate of 0.001 achieves the highest <italic>R</italic>
<sup>2</sup> value of 0.948. In contrast, XGBoost with a learning rate of 0.1 results in the lowest <italic>R</italic>
<sup>2</sup> of 0.840. For the torque shown in <xref ref-type="fig" rid="F5">Figure 5B</xref>, the decision tree with a learning rate of 0.001 leads to the best <italic>R</italic>
<sup>2</sup> value of 0.933. In contrast, XGBoost with a learning rate of 0.05 has the lowest <italic>R</italic>
<sup>2</sup> of 0.852. Therefore, the base learner is selected as the decision tree, with a learning rate of 0.001.</p>
<fig id="F5" position="float">
<label>FIGURE 5</label>
<caption>
<p>Model performance using various learning rates and base learner combinations for <bold>(A)</bold> Total thrust; and <bold>(B)</bold> Torque.</p>
</caption>
<graphic xlink:href="feart-12-1340437-g005.tif"/>
</fig>
</sec>
<sec id="s3-4">
<title>3.4 Model evaluation</title>
<p>Probabilistic prediction is helpful for risk management and decision-making as it provides decision makers with a prediction interval besides a single prediction value. <xref ref-type="fig" rid="F6">Figure 6</xref> displays the prediction results of the NGBoost model. The 95% confidence interval (CI) is plotted, which means that it includes the actual dataset parameter in about 95 out of 100 cases when sampled repeatedly. Overall, the measurements predominantly reside within the 95% confidence interval. This indicates that the model has a high level of predictive accuracy. Specifically, between the 3000th and 7500th sample points, the CI is considerably wider, which indicates a larger prediction uncertainty in this region. According to the geological survey report, this region is transitioning from a uniform layer to an upper-soft and lower-hard composite stratum, leading to heightened prediction challenges. In this complex geological situation, it is particularly important to control the operational parameters of the shield machine with greater precision and caution. In practice, the shield control system can automatically adjust the advance rate and rotation speed of the shield machine based on the confidence interval to adapt to the current geological conditions. When the model provides significant uncertainty in complex strata, such as the area between the 3000th and 7500th sample points, the system may suggest reducing the advance rate and increasing torque to ensure that the cutting head can effectively break through hard strata without causing mechanical overload.</p>
<fig id="F6" position="float">
<label>FIGURE 6</label>
<caption>
<p>Probabilistic prediction results of the NGBoost model for <bold>(A)</bold> total thrust and <bold>(B)</bold> torque.</p>
</caption>
<graphic xlink:href="feart-12-1340437-g006.tif"/>
</fig>
<p>In addition, the width of the confidence interval is also suddenly increased at the area marked by the red box, where no anomalies are identified by the geological survey. The possible reason is that the geological survey report is not detailed enough to reflect the local variations in strata. The scatter plot shows the comparison between the model predictions and the actual monitored values. Most of the points are densely clustered around the diagonal line. This indicates that the model can predict the total thrust force and torque accurately. However, <xref ref-type="fig" rid="F6">Figure 6</xref> also displays several noticeable outliers, where the model overestimates the monitored values. Specifically, these outliers mainly appear during the descending phase of the shield tunneling operational cycle. The appearance of outliers during the descending phase can be attributed to various factors, such as geological shifts, equipment malfunctions, or regular inspections of the shield machine. These factors introduce significant randomness to abrupt changes, which makes the prediction more challenging.</p>
<p>The performance of the NGBoost model is compared to five other machine learning models. The hyperparameter settings for the five models are determined through random search, as detailed in <xref ref-type="table" rid="T2">Table 2</xref>. Three popular decision-tree-based models are used for comparison, specifically, random forest consists of a collection of decision trees. Each tree is trained on a random subset of the data and gives its own prediction. The random forest model then averages these results to produce a final outcome. Extreme gradient boosting (XGBoost) is a decision tree-based ensemble model that uses a gradient boosting framework. It iteratively adds new trees and corrects errors made by previously trained trees. Light gradient boosting machine (LightGBM) also follows the gradient boosting framework. Instead of simply growing trees level by level, it prioritizes splitting the most optimal leaf nodes. In addition, to further demonstrate the effectiveness of the established NGBoost model, two commonly used machine learning models for the shield tunneling problem, namely extreme learning machine (ELM) and K-nearest neighbors (KNN), are also included for comparison. ELM is a feedforward neural network with a single layer of hidden nodes, where the weights are randomly assigned and do not require iterative tuning. KNN considers the &#x201c;k&#x201d; nearest data points to estimate the value of a data point based on the average of neighbors. <xref ref-type="table" rid="T3">Table 3</xref> summarizes the performance of all the models for predicting the total thrust. The LightGBM and NGBoost models show superior accuracy, with an <italic>R</italic>
<sup>2</sup> value of 0.9477. The ELM model has the worst performance, with an <italic>R</italic>
<sup>2</sup> of 0.3633. In terms of MAE, NGBoost has the lowest MAE of 361.60 kN, which is less than 1% of the total thrust during the stable phase, indicating that the average deviation of its predictions from the measurements is small. <xref ref-type="table" rid="T4">Table 4</xref> shows the model performance for predicting the torque. The NGBoost model performs the best with an <italic>R</italic>
<sup>2</sup> value of 0.9329 and an MAE of 432.22 kN m. For both torque and total thrust, the NGBoost model performs best in term of the <italic>R</italic>
<sup>2</sup>, MAE, and RMSE. The predictive accuracy of the LightGBM model is close to that of the NGBoost. However, the NGBoost model can provide probabilistic predictions for the shield movement performance. Based on these probabilistic results, decision-makers can better understand the uncertainty of predictions and make more reasonable adjustments accordingly.</p>
<table-wrap id="T2" position="float">
<label>TABLE 2</label>
<caption>
<p>Hyperparameter settings of the five machine learning models.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="left"/>
<th align="center">Hyperparameter</th>
<th align="center">Description</th>
<th align="center">Value</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td rowspan="3" align="center">Random forest</td>
<td align="center">num_models</td>
<td align="center">The number of trees in the forest</td>
<td align="center">100</td>
</tr>
<tr>
<td align="center">max_depth</td>
<td align="center">The maximum depth of the tree</td>
<td align="center">4</td>
</tr>
<tr>
<td align="center">min_samples_leaf</td>
<td align="center">The minimum number of samples required to be at a leaf node</td>
<td align="center">1</td>
</tr>
<tr>
<td align="center">ELM</td>
<td align="center">hidden_units</td>
<td align="center">The number of computational units in the hidden layer</td>
<td align="center">128</td>
</tr>
<tr>
<td rowspan="3" align="center">XGBoost</td>
<td align="center">n_estimators</td>
<td align="center">Number of boosting iterations to be run</td>
<td align="center">100</td>
</tr>
<tr>
<td align="center">learning_rate</td>
<td align="center">A parameter controlling the step size during model updates. Range is (0,1]</td>
<td align="center">0.03</td>
</tr>
<tr>
<td align="center">max_depth</td>
<td align="center">The maximum depth of the tree</td>
<td align="center">6</td>
</tr>
<tr>
<td align="center">KNN</td>
<td align="center">n_neighbors</td>
<td align="center">Number of neighboring points used for averaging</td>
<td align="center">5</td>
</tr>
<tr>
<td rowspan="4" align="center">LightGBM</td>
<td align="center">n_estimators</td>
<td align="center">Number of boosting iterations to be run</td>
<td align="center">100</td>
</tr>
<tr>
<td align="center">learning_rate</td>
<td align="center">A parameter controlling the step size during model updates. Range is (0,1)</td>
<td align="center">0.01</td>
</tr>
<tr>
<td align="center">max_depth</td>
<td align="center">The maximum depth of the tree</td>
<td align="center">4</td>
</tr>
<tr>
<td align="center">num_leaves</td>
<td align="center">The maximum number of leaves in one tree</td>
<td align="center">31</td>
</tr>
</tbody>
</table>
</table-wrap>
<table-wrap id="T3" position="float">
<label>TABLE 3</label>
<caption>
<p>Performance of different machine learning models for predicting the total thrust.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="center">Model</th>
<th align="center">Random forest</th>
<th align="center">ELM</th>
<th align="center">XGBoost</th>
<th align="center">KNN</th>
<th align="center">LightGBM</th>
<th align="center">NGBoost</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="center">R<sup>2</sup>
</td>
<td align="center">0.8752</td>
<td align="center">0.3633</td>
<td align="center">0.6468</td>
<td align="center">0.6288</td>
<td align="center">0.9477</td>
<td align="center">0.9477</td>
</tr>
<tr>
<td align="center">MAE</td>
<td align="center">599.94</td>
<td align="center">1,490.91</td>
<td align="center">1,010.14</td>
<td align="center">1740.69</td>
<td align="center">370.47</td>
<td align="center">361.60</td>
</tr>
</tbody>
</table>
</table-wrap>
<table-wrap id="T4" position="float">
<label>TABLE 4</label>
<caption>
<p>Performance of different machine learning models for predicting the torque.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="center">Model</th>
<th align="center">Random forest</th>
<th align="center">ELM</th>
<th align="center">XGBoost</th>
<th align="center">KNN</th>
<th align="center">LightGBM</th>
<th align="center">NGBoost</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="center">R<sup>2</sup>
</td>
<td align="center">0.8394</td>
<td align="center">0.4531</td>
<td align="center">0.5995</td>
<td align="center">0.8488</td>
<td align="center">0.9229</td>
<td align="center">0.9329</td>
</tr>
<tr>
<td align="center">MAE</td>
<td align="center">613.37</td>
<td align="center">1,043.54</td>
<td align="center">1,049</td>
<td align="center">795.54</td>
<td align="center">432.22</td>
<td align="center">367.84</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec id="s3-5">
<title>3.5 Model interpretability</title>
<p>
<xref ref-type="sec" rid="s3-4">Sections 3.4</xref>, <xref ref-type="sec" rid="s3-5">3.5</xref> show the accuracy of the decision tree-based NGBoost model. In practical engineering applications, an interpretable model is preferable because it improves transparency and acceptance of the model. The paper visualizes the decision pathways of NGBoost and utilizes the SHAP technique to explore the influence of different features on the prediction of shield movement performance.</p>
<p>Each decision tree path can translate into clear rules to improve model interpretability and transparency of its decisions. This study combines this highly interpretable feature with the NGBoost model to explain the decision path changes during each iteration of NGBoost. <xref ref-type="fig" rid="F7">Figure 7</xref> shows the decision path of the 1st and 300th iterations in the NGBoost model for predicting the total thrust. Node 0 is the root node of the decision tree; &#x201c;Samples&#x201d; denotes the number of samples under the path; and &#x201c;Correction value&#x201d; denotes an adjustment to the model prediction. The decision path directs samples from the root node to the leaf nodes based on their feature values and node rules. Samples go left if they meet the rule; otherwise, they go right. For example, during the 300th iteration as shown in <xref ref-type="fig" rid="F7">Figure 7B</xref>, given the input conditions of penetration &#x3d; 16 mm/r, advance rate &#x3d; 25 mm/min, and total thrust &#x3d; 5,000 kN, the path proceeds left to Node 1 as the penetration 16 mm/r &#x2264; 17 mm/r meets the condition at Node 0. At Node 1, since the advance rate 25 mm/min &#x2264;28 mm/min is met, the path continues left to Node 2. However, failing to meet the condition of advance rate &#x2264;12 mm/min at Node two (i.e., 25 &#x2265; 12), it finally reaches Node 4 with a correction value of &#x2212;1.7 kN. Note that the correction from the 300th iteration is just one example. The actual prediction is based on the cumulative sum of the corrections from 300 iterations.</p>
<fig id="F7" position="float">
<label>FIGURE 7</label>
<caption>
<p>Decision path of the NGBoost model for predicting the total thrust at the next moment: <bold>(A)</bold> first iteration; <bold>(B)</bold> 300th iteration.</p>
</caption>
<graphic xlink:href="feart-12-1340437-g007.tif"/>
</fig>
<p>In the first iteration of the NGBoost model, only the feature &#x201c;total thrust&#x201d; is used as the criterion for decision-making. This indicates that in this task of predicting the total thrust, the total thrust of the current moment is the most crucial feature for forecasting the total thrust of the next moment. The large correction values in the first iteration indicate that the model initially tends to make a wide range of parameter adjustments to move quickly towards the target value. In contrast, in the 300th iteration of NGBoost, more features are used for parameter tuning. The corrections are small and are mainly used to fine-tune the model. These results indicate that in the 300th iteration, the model has gradually stabilized and attempts to reduce the prediction error using as many features as possible.</p>
<p>The importance of input features for total thrust and torque at the next moment is shown in <xref ref-type="fig" rid="F8">Figure 8</xref>. Overall, &#x201c;total thrust&#x201d; and &#x201c;torque&#x201d; show the highest levels of importance when predicting their respective future values, and the value of importance is significantly higher than all other input features. The torque and total thrust significantly affect the standard deviation of their respective predicted outcomes (see <xref ref-type="fig" rid="F8">Figures 8A, B</xref>). This is because the torque mainly overcomes the friction between the cutter and the soil, and the total thrust pushes the shield machine forward. Excessive thrust with insufficient torque may cause the cutting tools to slip or jam, while insufficient thrust can lead to slow progress or machine stalling. Therefore, an increase in one parameter leads to a higher demand for the other, resulting in increased uncertainty in the predictive results. The ring number in shield tunneling significantly affects the standard deviation of the predicted values for total thrust and torque. The possible reason might be the variation in geological layers traversed by shield machine at different ring numbers, which also leads to increased predictive uncertainty.</p>
<fig id="F8" position="float">
<label>FIGURE 8</label>
<caption>
<p>Importance of input features for predicting the <bold>(A)</bold> total thrust and <bold>(B)</bold> torque.</p>
</caption>
<graphic xlink:href="feart-12-1340437-g008.tif"/>
</fig>
<p>A SHAP summary plot shown in <xref ref-type="fig" rid="F9">Figure 9</xref> is used to illustrate how the different features of shield tunneling affect model predictions. The SHAP values on the horizontal axis show how much each feature contributes to the prediction, and the features are listed on the vertical axis in order of their importance, from most to least. The color indicates the magnitude of the feature values, with blue colors for lower values and green colors for higher values. Overall, the two subfigures show distinct orders of impact for different prediction targets. For the total thrust at the next moment, the order of impact is total thrust &#x3e; penetration &#x3e; ring number &#x3e; advance rate &#x3e; torque. For the torque at the next moment, the order of impact is torque &#x3e; penetration &#x3e; ring number &#x3e; advance rate &#x3e; total thrust. Specifically, <xref ref-type="fig" rid="F9">Figure 9A</xref> explains that the total thrust at the next moment is mostly influenced by its previous value, highlighting the significance of past conditions on future predictions. Other factors like penetration rate, advance rate, and torque have a more minor effect, and they usually reduce the total thrust at the next moment. Although the impact of the ring number is similarly small, higher ring numbers tend to reduce the total thrust for the next ring more significantly. <xref ref-type="fig" rid="F9">Figure 9B</xref> discusses the torque required at the next moment and shows that, except for a significant impact of torque, all other features have minor impacts and no apparent trends.</p>
<fig id="F9" position="float">
<label>FIGURE 9</label>
<caption>
<p>SHAP summary plot of <bold>(A)</bold> total thrust and <bold>(B)</bold> torque.</p>
</caption>
<graphic xlink:href="feart-12-1340437-g009.tif"/>
</fig>
<p>Shield tunneling is a highly intricate process involving multiple variables and complex interactions. The contribution of a single feature to the tunneling performance is not isolated but depends on the other features. Partial dependence plots for double features are employed to examine the influence of these feature on the predictive accuracy, which illustrate how two features interact to affect the predictions of the NGBoost model. The SHAP value quantifies the contribution of each feature to the predictions made by the NGBoost model. The larger the value, the greater the contribution. The horizonal axis represents the value of the most important feature, and the vertical axis represents the SHAP values contributed by the main feature. The color intensity of each point represents the value of the second important feature, with darker colors indicating larger values. <xref ref-type="fig" rid="F10">Figure 10A</xref> shows the interplay between the advance rate and the total thrust and how they affect the torque at the next moment. When the advance rate ranges from 0 mm/min to 1.3 mm/min, it makes a negative contribution to the predicted value of torque at the next moment. The negative contribution increases as the advance rate is gradually decreased. This phenomenon indicates that the shield machine faces less resistance at a lower advance rate, leading to a lower torque. When the range of the advance rate is between 1.3 mm/min and 3.2 mm/min, a total thrust of 70,000 kN to 75,000 kN causes no increase in the torque. Similarly, <xref ref-type="fig" rid="F10">Figure 10B</xref> shows the interplay between the penetration and the torque and how they affect the torque at the next moment. When the value of penetration is less than 1 mm/r, it results in a negative contribution. However, when it exceeds 1 mm/r, the contribution becomes positive, peaking at around a SHAP value of 3. <xref ref-type="fig" rid="F10">Figure 10C</xref> shows that as the total thrust increases, its negative contribution to the torque at the next moment decreases to about 0. This means that the total thrust can be adjusted to reduce the torque, but not to increase the torque. <xref ref-type="fig" rid="F10">Figure 10D</xref> demonstrates that the SHAP values remain unchanged with the changes in the ring number. The ring number has little effect on the subsequent torque. Based on these interpretable results, the proposed algorithm can adjust various tunneling parameters more precisely, thereby significantly enhancing shield movement performance. The partial dependence plots and SHAP value analysis provide a scientific basis for this predictive adjustment. For example, these analyses reveal that the relationship between increasing the advance rate, torque and the total thrust at the next moment is not simply linear but exhibits complex non-linear dynamics, with significant boosts in some ranges and reductions in others. This not only helps maintain the shield machine in optimal working condition but also effectively minimizes downtime caused by inappropriate adjustments, thereby improving the overall construction efficiency and safety of the project.</p>
<fig id="F10" position="float">
<label>FIGURE 10</label>
<caption>
<p>Interaction of two features on the next-moment torque: <bold>(A)</bold> advance rate vs. total thrust; <bold>(B)</bold> penetration vs. torque; <bold>(C)</bold> total thrust vs. torque; and <bold>(D)</bold> ring number vs. advance rate.</p>
</caption>
<graphic xlink:href="feart-12-1340437-g010.tif"/>
</fig>
<p>
<xref ref-type="fig" rid="F11">Figure 11</xref> shows the partial dependence plots for the total thrust at the next moment. Overall, the trends observed in <xref ref-type="fig" rid="F11">Figure 11</xref> are similar to those for the next-moment torque shown in <xref ref-type="fig" rid="F10">Figure 10</xref>. Specifically, <xref ref-type="fig" rid="F11">Figure 11C</xref> demonstrates that as the penetration increases to about 1.7 mm/r, its negative impact on the next-moment thrust decreases to 0. As the penetration increases further, its impact suddenly changes, dropping to about &#x2212;0.7 mm/r and continuing to decline to &#x2212;3 mm/r. This pattern of change is different from that for the next-moment torque shown in <xref ref-type="fig" rid="F10">Figure 10C</xref>. In <xref ref-type="fig" rid="F10">Figure 10C</xref>, as the penetration increases, the negative impact on the torque changes from negative to positive, eventually increasing to about &#x2b;2 mm/r. This difference may be due to a change in the penetration mechanism from cutting to squeezing upon reaching a certain penetration. After the transition, the shield tunneling needs to overcome greater resistance, thus negatively impacting the total thrust.</p>
<fig id="F11" position="float">
<label>FIGURE 11</label>
<caption>
<p>Interaction of two features on the next-moment total thrust: <bold>(A)</bold> advance rate vs. torque; <bold>(B)</bold> penetration vs. total thrust; <bold>(C)</bold> torque vs. total thrust; and <bold>(D)</bold> ring number vs. advance rate.</p>
</caption>
<graphic xlink:href="feart-12-1340437-g011.tif"/>
</fig>
</sec>
</sec>
<sec sec-type="conclusion" id="s4">
<title>4 Conclusion</title>
<p>In this study, an interpretable probabilistic prediction algorithm is proposed for predicting the shield movement performance (i.e., total thrust and torque). Specifically, the natural gradient boosting (NGBoost) model is used to iteratively update the parametric probability distributions and achieve probabilistic predictions. The impact of each feature on the prediction values and uncertainty is quantified by extending the importance analysis of a single predictive value to the mean and variance. The feature interactions are analyzed and their contributions to the predictions are quantified by the SHAP method. The transparency of the NGBoost model is improved through the visualization of the decision-making process. The main conclusions are as follows.<list list-type="simple">
<list-item>
<p>(1) The NGBoost model can provide a probabilistic prediction of the total thrust and torque. The results show that the model produces a wider confidence interval in the complex geological conditions, guiding decision-makers to adjust the tunneling parameters with more caution.</p>
</list-item>
<list-item>
<p>(2) In terms of model interpretability, the total thrust and torque rank the most important features for predicting their respective future values. During the decision visualization process, the NGBoost model mainly uses the most important features for making decisions at the first iteration and makes significant corrections to the predicted outcomes. However, as the number of iterations increases, the model begins to incorporate additional features and fine-tunes the outcomes. In addition, partial dependence plots for double features show that the interaction between different features significantly affects the torque at the next moment. Based on these interpretable results, the proposed algorithm can acquire more information to help decision-making and adjust tunneling parameters more accurately, leading to improved shield movement performance.</p>
</list-item>
<list-item>
<p>(3) The NGBoost model outperforms the random forest, ELM, XGBoost, KNN, and LightGBM in terms of prediction accuracy. Optimal performance is achieved using the decision trees as the base learner and a learning rate of 0.001. These indicate that the NGBoost model, under these hyperparameter settings, is most suitable for predicting the torque and total thrust for the studied project.</p>
</list-item>
</list>
</p>
</sec>
</body>
<back>
<sec sec-type="data-availability" id="s5">
<title>Data availability statement</title>
<p>The original contributions presented in the study are included in the article/Supplementary material, further inquiries can be directed to the corresponding authors.</p>
</sec>
<sec id="s6">
<title>Author contributions</title>
<p>YZ: Conceptualization, Data curation, Resources, Writing&#x2013;review and editing. LL: Conceptualization, Data curation, Supervision, Writing&#x2013;review and editing. JW: Methodology, Writing&#x2013;review and editing. SZ: Software, Writing&#x2013;original draft. JH: Project administration, Writing&#x2013;review and editing. YT: Funding acquisition, Validation, Visualization, Writing&#x2013;review and editing. YH: Resources, Writing&#x2013;review and editing. XZ: Project administration, Validation, Writing&#x2013;review and editing. XL: Project administration, Funding acquisition, Writing&#x2013;review and editing.</p>
</sec>
<sec sec-type="funding-information" id="s7">
<title>Funding</title>
<p>The author(s) declare financial support was received for the research, authorship, and/or publication of this article. This study is supported by the Postdoctoral Research Foundation of Zhejiang Province (Grant No. ZJ2023057) and the Construction Research Project of Zhejiang Province (2023K155).</p>
</sec>
<sec sec-type="COI-statement" id="s8">
<title>Conflict of interest</title>
<p>Authors YZ, LL, JW, YH, and XZ were employed by Power China Huadong Engineering Corporation Limited. Authors LL, JW, YH, and XZ were employed by Zhejiang Huadong Engineering Construction and Management Corporation limited.</p>
<p>The remaining authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec sec-type="disclaimer" id="s9">
<title>Publisher&#x2019;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<ref-list>
<title>References</title>
<ref id="B1">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Amari</surname>
<given-names>S. I.</given-names>
</name>
</person-group> (<year>1998</year>). <article-title>Natural gradient works efficiently in learning</article-title>. <source>Neural Comput.</source> <volume>10</volume> (<issue>2</issue>), <fpage>251</fpage>&#x2013;<lpage>276</lpage>. <pub-id pub-id-type="doi">10.1162/089976698300017746</pub-id>
</citation>
</ref>
<ref id="B2">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Baghbani</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Choudhury</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Costa</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Costa</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Reiner</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Application of artificial intelligence in geotechnical engineering: a state-of-the-art review</article-title>. <source>Earth-Science Rev.</source> <volume>228</volume>, <fpage>103991</fpage>. <pub-id pub-id-type="doi">10.1016/j.earscirev.2022.103991</pub-id>
</citation>
</ref>
<ref id="B3">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Chen</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Tian</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Zhou</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Gong</surname>
<given-names>Q.</given-names>
</name>
<name>
<surname>Di</surname>
<given-names>H.</given-names>
</name>
</person-group> (<year>2024</year>). <article-title>Attitude deviation prediction of shield tunneling machine using time-aware LSTM networks</article-title>. <source>Transp. Geotech.</source> <volume>45</volume>, <fpage>101195</fpage>. <pub-id pub-id-type="doi">10.1016/j.trgeo.2024.101195</pub-id>
</citation>
</ref>
<ref id="B4">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Chen</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Wu</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Zhong</surname>
<given-names>Z.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Prediction of shield tunneling-induced ground settlement using machine learning techniques</article-title>. <source>Front. Struct. Civ. Eng.</source> <volume>13</volume>, <fpage>1363</fpage>&#x2013;<lpage>1378</lpage>. <pub-id pub-id-type="doi">10.1007/s11709-019-0561-3</pub-id>
</citation>
</ref>
<ref id="B5">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Chen</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Guestrin</surname>
<given-names>C.</given-names>
</name>
</person-group> (<year>2016</year>). &#x201c;<article-title>Xgboost: a scalable tree boosting system</article-title>,&#x201d; in <conf-name>Proceedings of the 22nd ACM SIGKDD international conference on knowledge discovery and data mining</conf-name>, <conf-loc>San Francisco, CA</conf-loc>, <conf-date>August 13&#x2013;17, 2016</conf-date>, <fpage>785</fpage>&#x2013;<lpage>794</lpage>.</citation>
</ref>
<ref id="B6">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Das</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Kashem</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Hasan</surname>
<given-names>I.</given-names>
</name>
<name>
<surname>Islam</surname>
<given-names>M.</given-names>
</name>
</person-group> (<year>2024</year>). <article-title>A comparative study of machine learning models for construction costs prediction with natural gradient boosting algorithm and SHAP analysis</article-title>. <source>Asian J. Civ. Eng.</source> <volume>25</volume>, <fpage>3301</fpage>&#x2013;<lpage>3316</lpage>. <pub-id pub-id-type="doi">10.1007/s42107-023-00980-z</pub-id>
</citation>
</ref>
<ref id="B7">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Devore</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2011</year>). <source>Probability and statistics for engineering and the sciences</source>. <publisher-loc>Boston, MA</publisher-loc>: <publisher-name>Cengage Learning</publisher-name>.</citation>
</ref>
<ref id="B8">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Duan</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Anand</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Ding</surname>
<given-names>D. Y.</given-names>
</name>
<name>
<surname>Thai</surname>
<given-names>K. K.</given-names>
</name>
<name>
<surname>Basu</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Ng</surname>
<given-names>A.</given-names>
</name>
<etal/>
</person-group> (<year>2020</year>). &#x201c;<article-title>Ngboost: natural gradient boosting for probabilistic prediction</article-title>,&#x201d; in <conf-name>Proceedings of the international conference on machine learning</conf-name>, <conf-loc>(Vienna: The International Machine Learning Society)</conf-loc>, <fpage>2690</fpage>&#x2013;<lpage>2700</lpage>.</citation>
</ref>
<ref id="B9">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Elbaz</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Zhou</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Shen</surname>
<given-names>S. L.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Deep reinforcement learning approach to optimize the driving performance of shield tunnelling machines</article-title>. <source>Tunn. Undergr. Space Technol.</source> <volume>136</volume>, <fpage>105104</fpage>. <pub-id pub-id-type="doi">10.1016/j.tust.2023.105104</pub-id>
</citation>
</ref>
<ref id="B10">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Friedman</surname>
<given-names>J. H.</given-names>
</name>
</person-group> (<year>2001</year>). <article-title>Greedy function approximation: a gradient boosting machine</article-title>. <source>Ann. Statistics</source> <volume>29</volume>, <fpage>1189</fpage>&#x2013;<lpage>1232</lpage>. <pub-id pub-id-type="doi">10.1214/aos/1013203451</pub-id>
</citation>
</ref>
<ref id="B11">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Gao</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Shi</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Song</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>H.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Recurrent neural networks for real-time prediction of TBM operating parameters</article-title>. <source>Automation Constr.</source> <volume>98</volume>, <fpage>225</fpage>&#x2013;<lpage>235</lpage>. <pub-id pub-id-type="doi">10.1016/j.autcon.2018.11.013</pub-id>
</citation>
</ref>
<ref id="B12">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Gu</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Ou</surname>
<given-names>Q.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>W. G.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>G. H.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Reliability assessment of rainfall-induced slope stability using Chebyshev-Galerkin-KL expansion and Bayesian approach</article-title>. <source>Can. Geotechnical J.</source> <volume>60</volume>, <fpage>1909</fpage>&#x2013;<lpage>1922</lpage>. <pub-id pub-id-type="doi">10.1139/cgj-2022-0671</pub-id>
</citation>
</ref>
<ref id="B13">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Guidotti</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Monreale</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Ruggieri</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Turini</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Giannotti</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Pedreschi</surname>
<given-names>D.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>A survey of methods for explaining black box models</article-title>. <source>ACM Comput. Surv. (CSUR)</source> <volume>51</volume> (<issue>5</issue>), <fpage>1</fpage>&#x2013;<lpage>42</lpage>. <pub-id pub-id-type="doi">10.1145/3236009</pub-id>
</citation>
</ref>
<ref id="B14">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hou</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Yang</surname>
<given-names>Q.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Real-time prediction of rock mass classification based on TBM operation big data and stacking technique of ensemble learning</article-title>. <source>J. Rock Mech. Geotechnical Eng.</source> <volume>14</volume> (<issue>1</issue>), <fpage>123</fpage>&#x2013;<lpage>143</lpage>. <pub-id pub-id-type="doi">10.1016/j.jrmge.2021.05.004</pub-id>
</citation>
</ref>
<ref id="B15">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Iban</surname>
<given-names>M. C.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>An explainable model for the mass appraisal of residences: the application of tree-based Machine Learning algorithms and interpretation of value determinants</article-title>. <source>Habitat Int.</source> <volume>128</volume>, <fpage>102660</fpage>. <pub-id pub-id-type="doi">10.1016/j.habitatint.2022.102660</pub-id>
</citation>
</ref>
<ref id="B16">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Iban</surname>
<given-names>M. C.</given-names>
</name>
<name>
<surname>Bilgilioglu</surname>
<given-names>S. S.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Snow avalanche susceptibility mapping using novel tree-based machine learning algorithms (XGBoost, NGBoost, and LightGBM) with eXplainable Artificial Intelligence (XAI) approach. Stochastic Environmental Research and Risk Assessment</article-title>. <source>Stoch. Environ. Res. Risk Assess.</source> <volume>37</volume> (<issue>6</issue>), <fpage>2243</fpage>&#x2013;<lpage>2270</lpage>. <pub-id pub-id-type="doi">10.1007/s00477-023-02392-6</pub-id>
</citation>
</ref>
<ref id="B17">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kannangara</surname>
<given-names>K. K. P. M.</given-names>
</name>
<name>
<surname>Zhou</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Ding</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Hong</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Hong</surname>
<given-names>Z.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Investigation of feature contribution to shield tunneling-induced settlement using Shapley additive explanations method</article-title>. <source>J. Rock Mech. Geotechnical Eng.</source> <volume>14</volume> (<issue>4</issue>), <fpage>1052</fpage>&#x2013;<lpage>1063</lpage>. <pub-id pub-id-type="doi">10.1016/j.jrmge.2022.01.002</pub-id>
</citation>
</ref>
<ref id="B18">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kuang</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Geng</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Xu</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Liao</surname>
<given-names>B.</given-names>
</name>
<etal/>
</person-group> (<year>2020</year>). <article-title>Causal inference</article-title>. <source>Causal Inference. Eng.</source> <volume>6</volume>, <fpage>253</fpage>&#x2013;<lpage>263</lpage>. <pub-id pub-id-type="doi">10.1016/j.eng.2019.08.016</pub-id>
</citation>
</ref>
<ref id="B19">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname>
<given-names>J. B.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>Z. Y.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Jing</surname>
<given-names>L. J.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>Y. P.</given-names>
</name>
<name>
<surname>Xiao</surname>
<given-names>H. H.</given-names>
</name>
<etal/>
</person-group> (<year>2023a</year>). <article-title>Feedback on a shared big dataset for intelligent TBM Part I: feature extraction and machine learning methods</article-title>. <source>Undergr. Space</source> <volume>11</volume>, <fpage>1</fpage>&#x2013;<lpage>25</lpage>. <pub-id pub-id-type="doi">10.1016/j.undsp.2023.01.001</pub-id>
</citation>
</ref>
<ref id="B20">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname>
<given-names>J. B.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>Z. Y.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Jing</surname>
<given-names>L. J.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>Y. P.</given-names>
</name>
<name>
<surname>Xiao</surname>
<given-names>H. H.</given-names>
</name>
<etal/>
</person-group> (<year>2023b</year>). <article-title>Feedback on a shared big dataset for intelligent TBM, part II: application and forward look</article-title>. <source>Undergr. Space</source> <volume>11</volume>, <fpage>26</fpage>&#x2013;<lpage>45</lpage>. <pub-id pub-id-type="doi">10.1016/j.undsp.2023.01.002</pub-id>
</citation>
</ref>
<ref id="B21">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lin</surname>
<given-names>S. S.</given-names>
</name>
<name>
<surname>Shen</surname>
<given-names>S. L.</given-names>
</name>
<name>
<surname>Zhou</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2022a</year>). <article-title>Real-time analysis and prediction of shield cutterhead torque using optimized gated recurrent unit neural network</article-title>. <source>J. Rock Mech. Geotechnical Eng.</source> <volume>14</volume> (<issue>4</issue>), <fpage>1232</fpage>&#x2013;<lpage>1240</lpage>. <pub-id pub-id-type="doi">10.1016/j.jrmge.2022.06.006</pub-id>
</citation>
</ref>
<ref id="B22">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lin</surname>
<given-names>S. S.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Zhou</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Shen</surname>
<given-names>S. L.</given-names>
</name>
<name>
<surname>Shen</surname>
<given-names>S. L.</given-names>
</name>
</person-group> (<year>2022b</year>). <article-title>Time-series prediction of shield movement performance during tunneling based on hybrid model</article-title>. <source>Tunn. Undergr. Space Technol.</source> <volume>119</volume>, <fpage>104245</fpage>. <pub-id pub-id-type="doi">10.1016/j.tust.2021.104245</pub-id>
</citation>
</ref>
<ref id="B23">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Liu</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Guan</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Xu</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Guo</surname>
<given-names>X.</given-names>
</name>
<etal/>
</person-group> (<year>2019</year>). <article-title>Improved support vector regression models for predicting rock mass parameters using tunnel boring machine driving data</article-title>. <source>Tunn. Undergr. Space Technol.</source> <volume>91</volume>, <fpage>102958</fpage>. <pub-id pub-id-type="doi">10.1016/j.tust.2019.04.014</pub-id>
</citation>
</ref>
<ref id="B24">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lundberg</surname>
<given-names>S. M.</given-names>
</name>
<name>
<surname>Lee</surname>
<given-names>S.-I.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>A unified approach to interpreting model predictions</article-title>. <source>Adv. Neural Inf. Process. Syst.</source> <volume>30</volume>, <fpage>4765</fpage>&#x2013;<lpage>4774</lpage>.</citation>
</ref>
<ref id="B25">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ma</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Ma</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Yan</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Faradonbeh</surname>
<given-names>R. S.</given-names>
</name>
<name>
<surname>Long</surname>
<given-names>H.</given-names>
</name>
<etal/>
</person-group> (<year>2024</year>). <article-title>Real-time classification model for tunnel surrounding rocks based on high-resolution neural network and structure&#x2013;optimizer hyperparameter optimization</article-title>. <source>Comput. Geotechnics</source> <volume>168</volume>, <fpage>106155</fpage>. <pub-id pub-id-type="doi">10.1016/j.compgeo.2024.106155</pub-id>
</citation>
</ref>
<ref id="B26">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Moayedi</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Mosallanezhad</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Rashid</surname>
<given-names>A. S. A.</given-names>
</name>
<name>
<surname>Jusoh</surname>
<given-names>W. A. W.</given-names>
</name>
<name>
<surname>Muazu</surname>
<given-names>M.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>A systematic review and meta-analysis of artificial neural network application in geotechnical engineering: theory and applications</article-title>. <source>Neural Comput. Appl.</source> <volume>32</volume>, <fpage>495</fpage>&#x2013;<lpage>518</lpage>. <pub-id pub-id-type="doi">10.1007/s00521-019-04109-9</pub-id>
</citation>
</ref>
<ref id="B27">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ntoutsi</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Fafalios</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Gadiraju</surname>
<given-names>U.</given-names>
</name>
<name>
<surname>Iosifidis</surname>
<given-names>V.</given-names>
</name>
<name>
<surname>Nejdl</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Vidal</surname>
<given-names>M.</given-names>
</name>
<etal/>
</person-group> (<year>2020</year>). <article-title>Bias in data-driven artificial intelligence systems&#x2014;an introductory survey</article-title>. <source>WIREs Data Min. Knowl. Discov.</source> <volume>10</volume>, <fpage>3</fpage>. <pub-id pub-id-type="doi">10.1002/widm.1356</pub-id>
</citation>
</ref>
<ref id="B28">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Parsa</surname>
<given-names>A. B.</given-names>
</name>
<name>
<surname>Movahedi</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Taghipour</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Derrible</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Mohammadian</surname>
<given-names>A. K.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Toward safer highways, application of XGBoost and SHAP for real-time accident detection and feature analysis</article-title>. <source>Accid. Analysis Prev.</source> <volume>136</volume>, <fpage>105405</fpage>. <pub-id pub-id-type="doi">10.1016/j.aap.2019.105405</pub-id>
</citation>
</ref>
<ref id="B29">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Phoon</surname>
<given-names>K. K.</given-names>
</name>
<name>
<surname>Kulhawy</surname>
<given-names>F. H.</given-names>
</name>
</person-group> (<year>1999</year>). <article-title>Characterization of geotechnical variability</article-title>. <source>Can. Geotechnical J.</source> <volume>36</volume> (<issue>4</issue>), <fpage>612</fpage>&#x2013;<lpage>624</lpage>. <pub-id pub-id-type="doi">10.1139/cgj-36-4-612</pub-id>
</citation>
</ref>
<ref id="B30">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Phoon</surname>
<given-names>K. K.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>W.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Future of machine learning in geotechnics</article-title>. <source>Georisk Assess. Manag. Risk Eng. Syst. Geohazards</source> <volume>17</volume> (<issue>1</issue>), <fpage>7</fpage>&#x2013;<lpage>22</lpage>. <pub-id pub-id-type="doi">10.1080/17499518.2022.2087884</pub-id>
</citation>
</ref>
<ref id="B31">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Qin</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Wu</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Huang</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Tao</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>C.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>A novel LSTM-autoencoder and enhanced transformer-based detection method for shield machine cutterhead clogging</article-title>. <source>Sci. China Technol. Sci.</source> <volume>66</volume> (<issue>2</issue>), <fpage>512</fpage>&#x2013;<lpage>527</lpage>. <pub-id pub-id-type="doi">10.1007/s11431-022-2218-9</pub-id>
</citation>
</ref>
<ref id="B32">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Quinlan</surname>
<given-names>J. R.</given-names>
</name>
</person-group> (<year>1986</year>). <article-title>Induction of decision trees</article-title>. <source>Mach. Learn.</source> <volume>1</volume>, <fpage>81</fpage>&#x2013;<lpage>106</lpage>. <pub-id pub-id-type="doi">10.1007/bf00116251</pub-id>
</citation>
</ref>
<ref id="B33">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Scavuzzo</surname>
<given-names>C. M.</given-names>
</name>
<name>
<surname>Scavuzzo</surname>
<given-names>J. M.</given-names>
</name>
<name>
<surname>Campero</surname>
<given-names>M. N.</given-names>
</name>
<name>
<surname>Anegagrie</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Aramendia</surname>
<given-names>A. A.</given-names>
</name>
<name>
<surname>Benito</surname>
<given-names>A.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). <article-title>Feature importance: opening a soil-transmitted helminth machine learning model via SHAP</article-title>. <source>Infect. Dis. Model</source> <volume>7</volume> (<issue>1</issue>), <fpage>262</fpage>&#x2013;<lpage>276</lpage>. <pub-id pub-id-type="doi">10.1016/j.idm.2022.01.004</pub-id>
</citation>
</ref>
<ref id="B34">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Shi</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Qin</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Tao</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>C.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>A VMD-EWT-LSTM-based multi-step prediction approach for shield tunneling machine cutterhead torque</article-title>. <source>Knowl. Based Syst.</source> <volume>228</volume>, <fpage>107213</fpage>. <pub-id pub-id-type="doi">10.1016/j.knosys.2021.107213</pub-id>
</citation>
</ref>
<ref id="B35">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tao</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>He</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Cai</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Cai</surname>
<given-names>Y.</given-names>
</name>
</person-group> (<year>2022a</year>). <article-title>Multi-objective optimization-based prediction of excavation-induced tunnel displacement</article-title>. <source>Undergr. Space</source> <volume>7</volume> (<issue>5</issue>), <fpage>735</fpage>&#x2013;<lpage>747</lpage>. <pub-id pub-id-type="doi">10.1016/j.undsp.2021.12.005</pub-id>
</citation>
</ref>
<ref id="B36">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tao</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Cai</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Cai</surname>
<given-names>Y.</given-names>
</name>
</person-group> (<year>2022b</year>). <article-title>Predictions of deep excavation responses considering model uncertainty: integrating BiLSTM neural networks with Bayesian updating</article-title>. <source>Int. J. Geomechanics</source> <volume>22</volume> (<issue>1</issue>), <fpage>04021250</fpage>. <pub-id pub-id-type="doi">10.1061/(asce)gm.1943-5622.0002245</pub-id>
</citation>
</ref>
<ref id="B37">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tao</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Phoon</surname>
<given-names>K. K.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Cai</surname>
<given-names>Y.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Hierarchical Bayesian model for predicting small-strain stiffness of sand</article-title>. <source>Can. Geotechnical J.</source> <volume>61</volume>, <fpage>668</fpage>&#x2013;<lpage>683</lpage>. <comment>online</comment>. <pub-id pub-id-type="doi">10.1139/cgj-2022-0598</pub-id>
</citation>
</ref>
<ref id="B38">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tao</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Phoon</surname>
<given-names>K. K.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Ching</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2024</year>). <article-title>Variance reduction function for a potential inclined slip line in a spatially variable soil</article-title>. <source>Struct. Saf.</source> <volume>106</volume>, <fpage>102395</fpage>. <pub-id pub-id-type="doi">10.1016/j.strusafe.2023.102395</pub-id>
</citation>
</ref>
<ref id="B39">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wang</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Fu</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Fu</surname>
<given-names>X.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Time series prediction of tunnel boring machine (TBM) performance during excavation using causal explainable artificial intelligence (CX-AI)</article-title>. <source>Automation Constr.</source> <volume>147</volume>, <fpage>104730</fpage>. <pub-id pub-id-type="doi">10.1016/j.autcon.2022.104730</pub-id>
</citation>
</ref>
<ref id="B40">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wen</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Lin</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Huang</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>Z.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Lithium battery health state assessment based on vehicle-to-grid (V2G) real-world data and natural gradient boosting model</article-title>. <source>Energy</source> <volume>284</volume>, <fpage>129246</fpage>. <pub-id pub-id-type="doi">10.1016/j.energy.2023.129246</pub-id>
</citation>
</ref>
<ref id="B41">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Xu</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Prediction of tunnel boring machine operating parameters using various machine learning algorithms</article-title>. <source>Tunn. Undergr. Space Technol.</source> <volume>109</volume>, <fpage>103699</fpage>. <pub-id pub-id-type="doi">10.1016/j.tust.2020.103699</pub-id>
</citation>
</ref>
<ref id="B42">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yu</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Qin</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Tao</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>Q.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>A multi-channel decoupled deep neural network for tunnel boring machine torque and thrust prediction</article-title>. <source>Tunn. Undergr. Space Technol.</source> <volume>133</volume>, <fpage>104949</fpage>. <pub-id pub-id-type="doi">10.1016/j.tust.2022.104949</pub-id>
</citation>
</ref>
<ref id="B43">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhang</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Ding</surname>
<given-names>X.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Application of deep learning algorithms in geotechnical engineering: a short critical review</article-title>. <source>Artif. Intell. Rev.</source> <volume>54</volume>, <fpage>5633</fpage>&#x2013;<lpage>5673</lpage>. <pub-id pub-id-type="doi">10.1007/s10462-021-09967-1</pub-id>
</citation>
</ref>
<ref id="B44">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhang</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Yin</surname>
<given-names>Z. Y.</given-names>
</name>
<name>
<surname>Jin</surname>
<given-names>Y. F.</given-names>
</name>
<name>
<surname>Jin</surname>
<given-names>Y. F.</given-names>
</name>
</person-group> (<year>2022a</year>). <article-title>Machine learning-based modelling of soil properties for geotechnical design: review, tool development and comparison</article-title>. <source>Archives Comput. Methods Eng.</source> <volume>29</volume>, <fpage>1229</fpage>&#x2013;<lpage>1245</lpage>. <pub-id pub-id-type="doi">10.1007/s11831-021-09615-5</pub-id>
</citation>
</ref>
<ref id="B45">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhang</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Safdar</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Xie</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Sage</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Zhao</surname>
<given-names>Y. F.</given-names>
</name>
</person-group> (<year>2022b</year>). <article-title>A systematic review on data of additive manufacturing for machine learning applications: the data quality, type, preprocessing, and management</article-title>. <source>J. Intelligent Manuf.</source> <volume>34</volume>, <fpage>3305</fpage>&#x2013;<lpage>3340</lpage>. <pub-id pub-id-type="doi">10.1007/s10845-022-02017-9</pub-id>
</citation>
</ref>
<ref id="B46">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhou</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Gao</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>E. J.</given-names>
</name>
<name>
<surname>Ding</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Qin</surname>
<given-names>W.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Deep learning technologies for shield tunneling: challenges and opportunities</article-title>. <source>Automation Constr.</source> <volume>154</volume>, <fpage>104982</fpage>. <pub-id pub-id-type="doi">10.1016/j.autcon.2023.104982</pub-id>
</citation>
</ref>
<ref id="B47">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhou</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Qiu</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Khandelwal</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Zhu</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>X.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Developing a hybrid model of Jaya algorithm-based extreme gradient boosting machine to estimate blast-induced ground vibrations</article-title>. <source>Int. J. Rock Mech. Min. Sci.</source> <volume>145</volume>, <fpage>104856</fpage>. <pub-id pub-id-type="doi">10.1016/j.ijrmms.2021.104856</pub-id>
</citation>
</ref>
</ref-list>
</back>
</article>