<?xml version="1.0" encoding="UTF-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article article-type="research-article" dtd-version="2.3" xml:lang="EN" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Earth Sci.</journal-id>
<journal-title>Frontiers in Earth Science</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Earth Sci.</abbrev-journal-title>
<issn pub-type="epub">2296-6463</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="publisher-id">1364515</article-id>
<article-id pub-id-type="doi">10.3389/feart.2024.1364515</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Earth Science</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Prediction of permeability in a tight sandstone reservoir using a gated network stacking model driven by data and physical models</article-title>
<alt-title alt-title-type="left-running-head">Shi et al.</alt-title>
<alt-title alt-title-type="right-running-head">
<ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3389/feart.2024.1364515">10.3389/feart.2024.1364515</ext-link>
</alt-title>
</title-group>
<contrib-group>
<contrib contrib-type="author">
<name>
<surname>Shi</surname>
<given-names>Pengyu</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2617946/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Shi</surname>
<given-names>Pengda</given-names>
</name>
<xref ref-type="aff" rid="aff4">
<sup>4</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Bie</surname>
<given-names>Kang</given-names>
</name>
<xref ref-type="aff" rid="aff5">
<sup>5</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Han</surname>
<given-names>Chuang</given-names>
</name>
<xref ref-type="aff" rid="aff5">
<sup>5</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Ni</surname>
<given-names>Xiaowei</given-names>
</name>
<xref ref-type="aff" rid="aff6">
<sup>6</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Mao</surname>
<given-names>Zhiqiang</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<xref ref-type="corresp" rid="c001">&#x2a;</xref>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Zhao</surname>
<given-names>Peiqiang</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<xref ref-type="corresp" rid="c001">&#x2a;</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1820443/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/funding-acquisition/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
</contrib>
</contrib-group>
<aff id="aff1">
<sup>1</sup>
<institution>College of Geophysics</institution>, <institution>China University of Petroleum</institution>, <addr-line>Beijing</addr-line>, <country>China</country>
</aff>
<aff id="aff2">
<sup>2</sup>
<institution>National Key Laboratory of Petroleum Resources and Engineering</institution>, <addr-line>Beijing</addr-line>, <country>China</country>
</aff>
<aff id="aff3">
<sup>3</sup>
<institution>Beijing Key Laboratory of Earth Prospecting and Information Technology</institution>, <addr-line>Beijing</addr-line>, <country>China</country>
</aff>
<aff id="aff4">
<sup>4</sup>
<institution>School of Software Engineering</institution>, <institution>Chengdu University of Information Technology</institution>, <addr-line>Chengdu</addr-line>, <country>China</country>
</aff>
<aff id="aff5">
<sup>5</sup>
<institution>Exploration and Production Research Institute</institution>, <institution>PetroChina Tarim Oilfield Company</institution>, <addr-line>Korla</addr-line>, <country>China</country>
</aff>
<aff id="aff6">
<sup>6</sup>
<institution>Oil &#x26; Gas Field Productivity Construction Department</institution>, <institution>PetroChina Tarim Oilfield Company</institution>, <addr-line>Korla</addr-line>, <country>China</country>
</aff>
<author-notes>
<fn fn-type="edited-by">
<p>
<bold>Edited by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1670302/overview">Xixin Wang</ext-link>, Yangtze University, China</p>
</fn>
<fn fn-type="edited-by">
<p>
<bold>Reviewed by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/2623566/overview">Zhishui Liu</ext-link>, Chang&#x2019;an University, China</p>
<p>
<ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/2623541/overview">Yuhang Gyo</ext-link>, Jilin University, China</p>
</fn>
<corresp id="c001">&#x2a;Correspondence: Zhiqiang Mao, <email>maozq@cup.edu.cn</email>; Peiqiang Zhao, <email>pqzhao@cup.edu.cn</email>
</corresp>
</author-notes>
<pub-date pub-type="epub">
<day>15</day>
<month>02</month>
<year>2024</year>
</pub-date>
<pub-date pub-type="collection">
<year>2024</year>
</pub-date>
<volume>12</volume>
<elocation-id>1364515</elocation-id>
<history>
<date date-type="received">
<day>02</day>
<month>01</month>
<year>2024</year>
</date>
<date date-type="accepted">
<day>26</day>
<month>01</month>
<year>2024</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#xa9; 2024 Shi, Shi, Bie, Han, Ni, Mao and Zhao.</copyright-statement>
<copyright-year>2024</copyright-year>
<copyright-holder>Shi, Shi, Bie, Han, Ni, Mao and Zhao</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<abstract>
<p>
<bold>Introduction:</bold> Permeability is one of the most important parameters for reservoir evaluation. It is commonly measured in laboratories using underground core samples. However, it cannot describe the entire reservoir because of the limited number of cores. Therefore, petrophysicists use well logs to establish empirical equations to estimate permeability. This method has been widely used in conventional sandstone reservoirs, but it is not applicable to tight sandstone reservoirs with low porosity, extremely low permeability, and complex pore structures.</p>
<p>
<bold>Methods:</bold> Machine learning models can identify potential relationships between input features and sample labels, making them a good choice for establishing permeability prediction models. A stacking model is an ensemble learning method that aims to train a meta-learner to learn an optimal combination of expert models. However, the meta-learner does not evaluate or control the experts, making it difficult to interpret the contribution of each model. In this study, we design a gate network stacking (GNS) model, which is an algorithm that combines data and model-driven methods. First, an input log combination is selected for each expert model to ensure the best performance of the expert model and selfoptimization of the hyperparameters. Petrophysical constraints are then added to the inputs of the expert model and meta-learner, and weights are dynamically assigned to the output of the expert model. Finally, the overall performance of the model is evaluated iteratively to enhance its interpretability and robustness.</p> <p>
<bold>Results and discussion:</bold> The GNS model is then used to predict the permeability of a tight sandstone reservoir in the Jurassic Ahe Formation in the Tarim Basin. The case study shows that the permeability predicted by the GNS model is more accurate than that of other ensemble models. This study provides a new approach for predicting the parameters of tight sandstone reservoirs.</p>
</abstract>
<kwd-group>
<kwd>machine learning</kwd>
<kwd>ensemble model</kwd>
<kwd>gate network</kwd>
<kwd>tight sandstone reservoir</kwd>
<kwd>permeability prediction</kwd>
</kwd-group>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-at-acceptance</meta-name>
<meta-value>Geochemistry</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec sec-type="intro" id="s1">
<title>1 Introduction</title>
<p>The absolute permeability, k, is a measure of the ability of a porous medium to pass through a certain fluid in the presence of one or more fluid phases. The accurate prediction of permeability plays an important role in the evaluation of reservoir quality, numerical simulations of reservoirs, and estimation of geological reserves. Rock permeability is usually obtained through laboratory core measurements (<xref ref-type="bibr" rid="B31">Wu, 2004</xref>). However, this method is time-consuming and costly, and it has problems such as non-random sampling locations and limited quantity, making it impossible to characterize the permeability characteristics of the entire reservoir.</p>
<p>Scholars have performed extensive research on the problem of permeability predictions. The Kozeny-Carman equation (KC equation) based on the tube-like model applies the porosity and Kozeny constant to calculate permeability (<xref ref-type="bibr" rid="B14">Kozeny, 1927</xref>; <xref ref-type="bibr" rid="B5">Carman, 1937</xref>). However, the formation is not a uniform porous medium; thus, the prediction results are unreliable (<xref ref-type="bibr" rid="B24">Paterson, 1983</xref>; <xref ref-type="bibr" rid="B20">Mauran et al., 2001</xref>). Timur established a relationship using the porosity, bulk volume irreducible (BVI), and free fluid index (FFI) to predict permeability, but the predicted value was sensitive to the value of BVI, and it generally tended to overestimate permeability (<xref ref-type="bibr" rid="B29">Timur, 1968</xref>). Coates and Dumanoir derived a new free fluid model that ensured zero permeability at zero porosity when the irreducible water saturation was 100%. However, this model was only valid for intergranular pores and was not applicable for tight sandstones (<xref ref-type="bibr" rid="B7">Coates and Dumanoir, 1973</xref>). Ahmed considered the influence of minerals based on the Timur model and introduced a dual-mineral diffusion water model to estimate the permeability values. However, the irreducible water saturation is related to the shale volume and particle size; therefore, in complex heterogeneous or fractured reservoirs, the results are unreliable (<xref ref-type="bibr" rid="B2">Ahmed et al., 1991</xref>). The above research shows that it is very difficult to establish a comprehensive permeability prediction equation for highly heterogeneous reservoirs, and it is necessary to find a nonlinear method that uses well log curves for prediction.</p>
<p>In some cases, the models used in petrophysical calculations are nonlinear and cannot be explained by empirical theory. Machine learning is a data-driven approach that can provide alternative models in the absence of deterministic physical models; therefore, machine learning is a suitable technique for building regression models (<xref ref-type="bibr" rid="B21">Mohaghegh and Ameri, 1995</xref>; <xref ref-type="bibr" rid="B26">Saemi et al., 2007</xref>; <xref ref-type="bibr" rid="B1">Ahmadi et al., 2013</xref>; <xref ref-type="bibr" rid="B27">Saljooghi and Hezarkhani, 2014</xref>). Rogers applied a backpropagation neural network (BPNN) to predict the permeability and input the porosity log (<xref ref-type="bibr" rid="B25">Rogers et al., 1995</xref>). Jamialahmadi predicted the permeability of typical Iranian oil fields based on radial basis function (RBF) neural networks (<xref ref-type="bibr" rid="B11">Jamialahmadi and Javadpour, 2000</xref>), and Nazari applied support vector regression (SVR) to extract data into hyperplane dimensions to avoid overfitting (<xref ref-type="bibr" rid="B22">Nazari et al., 2011</xref>). Al&#x2013;Anazi applied the gamma-ray log (GR), density (DEN), neutron (CN) and compressional slow-ness (DT) to predict permeability; the results showed that SVR was better than artificial neural networks (ANN) (<xref ref-type="bibr" rid="B3">Al-Anazi and Gates, 2012</xref>). The above research shows that machine learning methods have significant advantages over empirical models. However, overtraining often occurs in the process of using these models for prediction, resulting in individual models that are not robust and have poor generalization. Zhang constructed a visual prediction model and concluded that ResNet learned the best-fitting nonlinear porosity&#x2013;permeability relationship from the input feature (<xref ref-type="bibr" rid="B32">Zhang et al., 2021</xref>).</p>
<p>To improve the problem of limited performance of a single model, ensemble learning has been developed; this method combines separate machine learning models through different combination strategies to improve the performance of the combined model (<xref ref-type="bibr" rid="B23">Nilsson, 1965</xref>; <xref ref-type="bibr" rid="B9">Friedman, 2001</xref>; <xref ref-type="bibr" rid="B28">Sammut and Webb, 2011</xref>; <xref ref-type="bibr" rid="B34">Zhang et al., 2022</xref>; <xref ref-type="bibr" rid="B12">Kalule et al., 2023</xref>). Chen and Lin used a committee machine with empirical formulas (CMEF) model to predict permeability using a collection of empirical formulas as experts. The results showed that the proposed model was more accurate than any single empirical formula (<xref ref-type="bibr" rid="B6">Chen and Lin, 2006</xref>). Sadegh built a supervised committee machine neural network (SCMNN), and each estimator of the SCMNN was the combination of two simple networks and one gating network; the prediction results were in good agreement with the core experimental results (<xref ref-type="bibr" rid="B13">Karimpouli et al., 2010</xref>). Zhu built a hybrid intelligent algorithm that combined the AdaBoost algorithm, adaptive rain forest optimization algorithm, and improved Back Propagation Neural Network (BPNN) with Nuclear Magnetic Resonance (NMR) logging data to improve and reconstruct the original artificial intelligence algorithm (<xref ref-type="bibr" rid="B36">Zhu et al., 2017</xref>). Zhang established a regression model based on the fusional temporal convolutional network (FTCN) method that can accurately predict the formation properties when there are major changes (<xref ref-type="bibr" rid="B33">Zhang et al., 2022</xref>). Bai established an RCM model to predict reservoir parameters, and the prediction results of the integrated model were more accurate than those of individual expert models (<xref ref-type="bibr" rid="B4">Bai et al., 2020</xref>). Morteza combined the social ski-driver (SSD) algorithm with a multilayer perception (MLP) neural network and presented a new hybrid algorithm to predict the value of rock permeability. The results indicated that the hybrid models can predict rock permeability with excellent accuracy (<xref ref-type="bibr" rid="B19">Matinkia et al., 2023</xref>).</p>
<p>However, in current ensemble learning training, the input data are mostly well log measured <italic>in situ</italic>, resulting in input features that often lack petrophysical constraints. In log interpretation, the classic petrophysical model is applicable to sandstones with medium and high porosity. However, for tight sandstones, the porosity is low and the pore structure is complex; therefore, the porosity&#x2013;permeability relationship cannot be explained by a simple linear model. Therefore, we aim to build an improved stacking model in this study.</p>
<p>In this study, we improve a classic stacking model called the gate network stacking (GNS) model. In the methodology section, we introduce the structure of the GNS model, which includes four parts: the input layer, expert layer, meta-learner, and gate network. In the case study section, we apply the GNS model to predict the permeability of tight sandstones in the Jurassic formation of the Tarim Basin and compare its performance with other ensemble learning models. In the discussion section, we discuss the advantages of the GNS model over other models and present the conclusions of the study.</p>
</sec>
<sec sec-type="methods" id="s2">
<title>2 Methodology</title>
<sec id="s2-1">
<title>2.1 Input layer</title>
<sec id="s2-1-1">
<title>2.1.1 Data processing</title>
<p>The input layer of the model is the data import port, which is responsible for data standardization and preprocessing. The input data for the gated network stacking model are a variety of well log curves, and the units and orders of magnitude of these log curves are quite different. Data normalization converts raw data into dimension-less and order-of-magnitude standardized values that can be compared between different input indicators. We use Z-value standardization for data preprocessing of the conventional well-logging curves. This method converts the data into a normal distribution with a mean of zero and a standard deviation of one without changing the distribution characteristics. To ensure that the model exhibits the best possible performance, the value of the output variable (core permeability) is considered as a logarithm. The Z-value normalization formula is shown in Eqs <xref ref-type="disp-formula" rid="e1">1</xref>&#x2013;<xref ref-type="disp-formula" rid="e3">3</xref>.<disp-formula id="e1">
<mml:math id="m1">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="bold-italic">Z</mml:mi>
<mml:mi mathvariant="bold-italic">i</mml:mi>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="bold-italic">x</mml:mi>
<mml:mi mathvariant="bold-italic">i</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:mi mathvariant="bold-italic">&#x3bc;</mml:mi>
</mml:mrow>
<mml:mi mathvariant="bold-italic">&#x3c3;</mml:mi>
</mml:mfrac>
</mml:mrow>
</mml:math>
<label>(1)</label>
</disp-formula>
<disp-formula id="e2">
<mml:math id="m2">
<mml:mrow>
<mml:mi mathvariant="bold-italic">&#x3bc;</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mn mathvariant="bold">1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi mathvariant="bold-italic">n</mml:mi>
</mml:mrow>
</mml:mfrac>
<mml:mstyle displaystyle="true">
<mml:munderover>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi mathvariant="bold-italic">i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn mathvariant="bold">1</mml:mn>
</mml:mrow>
<mml:mi mathvariant="bold-italic">n</mml:mi>
</mml:munderover>
</mml:mstyle>
<mml:msub>
<mml:mi mathvariant="bold-italic">x</mml:mi>
<mml:mi mathvariant="bold-italic">i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
<label>(2)</label>
</disp-formula>
<disp-formula id="e3">
<mml:math id="m3">
<mml:mrow>
<mml:mi mathvariant="bold-italic">&#x3c3;</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:msqrt>
<mml:mrow>
<mml:mfrac>
<mml:mrow>
<mml:mn mathvariant="bold">1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi mathvariant="bold-italic">n</mml:mi>
</mml:mrow>
</mml:mfrac>
<mml:mstyle displaystyle="true">
<mml:munderover>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi mathvariant="bold-italic">i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn mathvariant="bold">1</mml:mn>
</mml:mrow>
<mml:mi mathvariant="bold-italic">n</mml:mi>
</mml:munderover>
</mml:mstyle>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="bold-italic">x</mml:mi>
<mml:mi mathvariant="bold-italic">i</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:mi mathvariant="bold-italic">&#x3bc;</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mn mathvariant="bold">2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:msqrt>
</mml:mrow>
</mml:math>
<label>(3)</label>
</disp-formula>where <inline-formula id="inf1">
<mml:math id="m4">
<mml:mrow>
<mml:msub>
<mml:mi>Z</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> is the standardized result of the i-th sample, <inline-formula id="inf2">
<mml:math id="m5">
<mml:mrow>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> is the i-th sample data, <inline-formula id="inf3">
<mml:math id="m6">
<mml:mrow>
<mml:mi>&#x3bc;</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> is the mean calculated for all samples, <inline-formula id="inf4">
<mml:math id="m7">
<mml:mrow>
<mml:mi>&#x3c3;</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> is the variance calculated for all samples, and <inline-formula id="inf5">
<mml:math id="m8">
<mml:mrow>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> is the total number of samples in the dataset.</p>
<p>After data preprocessing, based on the game theory proposed by Shapley (1988), we used Shapley values to analyze the contribution of features and evaluate the importance of features in the ensemble learning model. The Shapley value is a method from cooperative game theory used to fairly distribute the total gains or losses among players based on their individual contributions to the collective effort.</p>
<p>The Shapley regression value represents the feature importance of a linear model in the presence of multicollinearity (<xref ref-type="bibr" rid="B15">Lundberg and Lee, 2016</xref>). It assigns an important value to each feature, indicating the impact of including the feature in the model prediction. SHAP (SHapley Additive exPlanations) explains the output of a model by attributing the importance of each feature to the prediction, which can be used to identify a new class of additive feature importance measures and show that there is a unique solution and a desirable set of properties in this class (<xref ref-type="bibr" rid="B18">Lundberg and Lee, 2017</xref>; <xref ref-type="bibr" rid="B17">Lundberg et al., 2018</xref>; <xref ref-type="bibr" rid="B16">Lundberg et al., 2020</xref>). The SHAP value is a unified measure of feature importance that combines these conditional expectation functions with the classic Shapley value from game theory to attribute <inline-formula id="inf6">
<mml:math id="m9">
<mml:mrow>
<mml:msub>
<mml:mi>&#x3d5;</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> values to each feature, as shown in Eqs <xref ref-type="disp-formula" rid="e4">4</xref>, <xref ref-type="disp-formula" rid="e5">5</xref>.<disp-formula id="e4">
<mml:math id="m10">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="bold-italic">f</mml:mi>
<mml:mi mathvariant="bold-italic">x</mml:mi>
</mml:msub>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="bold-italic">S</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x3d;</mml:mo>
<mml:mi mathvariant="bold-italic">E</mml:mi>
<mml:mrow>
<mml:mfenced open="[" close="]" separators="|">
<mml:mrow>
<mml:mi mathvariant="bold-italic">f</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="bold-italic">x</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x7c;</mml:mo>
<mml:msub>
<mml:mi mathvariant="bold-italic">x</mml:mi>
<mml:mi mathvariant="bold-italic">S</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(4)</label>
</disp-formula>
<disp-formula id="e5">
<mml:math id="m11">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="bold-italic">&#x3d5;</mml:mi>
<mml:mi mathvariant="bold-italic">i</mml:mi>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:munder>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi mathvariant="bold-italic">S</mml:mi>
<mml:mo>&#x2286;</mml:mo>
<mml:mi mathvariant="bold-italic">N</mml:mi>
<mml:mo>\</mml:mo>
<mml:mrow>
<mml:mfenced open="{" close="}" separators="|">
<mml:mrow>
<mml:mi mathvariant="bold-italic">i</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:munder>
</mml:mstyle>
<mml:mrow>
<mml:mfrac>
<mml:mrow>
<mml:mrow>
<mml:mfenced open="|" close="|" separators="|">
<mml:mrow>
<mml:mi mathvariant="bold-italic">S</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>!</mml:mo>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mrow>
<mml:mfenced open="|" close="|" separators="|">
<mml:mrow>
<mml:mi mathvariant="bold-italic">M</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:mrow>
<mml:mfenced open="|" close="|" separators="|">
<mml:mrow>
<mml:mi mathvariant="bold-italic">S</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:mn mathvariant="bold">1</mml:mn>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>!</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi mathvariant="bold-italic">M</mml:mi>
<mml:mo>!</mml:mo>
</mml:mrow>
</mml:mfrac>
<mml:mrow>
<mml:mfenced open="[" close="]" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="bold-italic">f</mml:mi>
<mml:mi mathvariant="bold-italic">x</mml:mi>
</mml:msub>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="bold-italic">S</mml:mi>
<mml:mo>&#x222a;</mml:mo>
<mml:mrow>
<mml:mfenced open="{" close="}" separators="|">
<mml:mrow>
<mml:mi mathvariant="bold-italic">i</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi mathvariant="bold-italic">f</mml:mi>
<mml:mi mathvariant="bold-italic">x</mml:mi>
</mml:msub>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="bold-italic">S</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(5)</label>
</disp-formula>where <inline-formula id="inf7">
<mml:math id="m12">
<mml:mrow>
<mml:mi>S</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> is the subset and <inline-formula id="inf8">
<mml:math id="m13">
<mml:mrow>
<mml:mi>E</mml:mi>
<mml:mrow>
<mml:mfenced open="[" close="]" separators="|">
<mml:mrow>
<mml:mi>f</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi>x</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x7c;</mml:mo>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mi>S</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula> is the expected value of the function conditioned on a subset S of the input features. <inline-formula id="inf9">
<mml:math id="m14">
<mml:mrow>
<mml:msub>
<mml:mi>&#x3d5;</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> is the Shapley value. N is the set of all input features. <inline-formula id="inf10">
<mml:math id="m15">
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mo>&#x2286;</mml:mo>
<mml:mi>N</mml:mi>
<mml:mo>\</mml:mo>
<mml:mrow>
<mml:mfenced open="{" close="}" separators="|">
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula> comprises all the possible subsets with feature i. <inline-formula id="inf11">
<mml:math id="m16">
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mo>&#x222a;</mml:mo>
<mml:mrow>
<mml:mfenced open="{" close="}" separators="|">
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula> is a subset S with added i, and <inline-formula id="inf12">
<mml:math id="m17">
<mml:mrow>
<mml:mfenced open="|" close="|" separators="|">
<mml:mrow>
<mml:mi>S</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:math>
</inline-formula> is the size of the subset before the ith feature is added to the interaction. <inline-formula id="inf13">
<mml:math id="m18">
<mml:mrow>
<mml:mfrac>
<mml:mrow>
<mml:mrow>
<mml:mfenced open="|" close="|" separators="|">
<mml:mrow>
<mml:mi>S</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>!</mml:mo>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mrow>
<mml:mfenced open="|" close="|" separators="|">
<mml:mrow>
<mml:mi>M</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:mrow>
<mml:mfenced open="|" close="|" separators="|">
<mml:mrow>
<mml:mi>S</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>!</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>M</mml:mi>
<mml:mo>!</mml:mo>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:math>
</inline-formula> is the weight of the combination, and <inline-formula id="inf14">
<mml:math id="m19">
<mml:mrow>
<mml:msub>
<mml:mi>f</mml:mi>
<mml:mi>x</mml:mi>
</mml:msub>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mo>&#x222a;</mml:mo>
<mml:mrow>
<mml:mfenced open="{" close="}" separators="|">
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>f</mml:mi>
<mml:mi>x</mml:mi>
</mml:msub>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi>S</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula> is the marginal contribution.</p>
</sec>
<sec id="s2-1-2">
<title>2.1.2 Petrophysical models</title>
<p>In the process of intelligent well-logging interpretation, simply using a data-driven strategy without geological constraints can easily lead to training model prediction results that are significantly different from the objective understanding. Even if the evaluation index of the trained model is good, its prediction results for unknown samples are not convincing. Therefore, we add the results of the petrophysical model as the input.</p>
<p>The volume of clay from the gamma ray log, Vsh, is as follows:<disp-formula id="e6">
<mml:math id="m20">
<mml:mrow>
<mml:mi mathvariant="bold-italic">S</mml:mi>
<mml:mi mathvariant="bold-italic">H</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mi mathvariant="bold-italic">G</mml:mi>
<mml:mi mathvariant="bold-italic">R</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi mathvariant="bold-italic">G</mml:mi>
<mml:mi mathvariant="bold-italic">R</mml:mi>
</mml:mrow>
<mml:mi mathvariant="bold-italic">min</mml:mi>
</mml:msub>
</mml:mrow>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi mathvariant="bold-italic">G</mml:mi>
<mml:mi mathvariant="bold-italic">R</mml:mi>
</mml:mrow>
<mml:mi mathvariant="bold-italic">max</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi mathvariant="bold-italic">G</mml:mi>
<mml:mi mathvariant="bold-italic">R</mml:mi>
</mml:mrow>
<mml:mi mathvariant="bold-italic">min</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:math>
<label>(6)</label>
</disp-formula>
<disp-formula id="e7">
<mml:math id="m21">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="bold-italic">V</mml:mi>
<mml:mrow>
<mml:mi mathvariant="bold-italic">s</mml:mi>
<mml:mi mathvariant="bold-italic">h</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:msup>
<mml:mn mathvariant="bold">2</mml:mn>
<mml:mrow>
<mml:mi mathvariant="bold-italic">G</mml:mi>
<mml:mi mathvariant="bold-italic">C</mml:mi>
<mml:mi mathvariant="bold-italic">U</mml:mi>
<mml:mi mathvariant="bold-italic">R</mml:mi>
<mml:mo>&#x2a;</mml:mo>
<mml:mi mathvariant="bold-italic">S</mml:mi>
<mml:mi mathvariant="bold-italic">H</mml:mi>
</mml:mrow>
</mml:msup>
<mml:mo>&#x2212;</mml:mo>
<mml:mn mathvariant="bold">1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:msup>
<mml:mn mathvariant="bold">2</mml:mn>
<mml:mrow>
<mml:mi mathvariant="bold-italic">G</mml:mi>
<mml:mi mathvariant="bold-italic">C</mml:mi>
<mml:mi mathvariant="bold-italic">U</mml:mi>
<mml:mi mathvariant="bold-italic">R</mml:mi>
</mml:mrow>
</mml:msup>
<mml:mo>&#x2212;</mml:mo>
<mml:mn mathvariant="bold">1</mml:mn>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:math>
<label>(7)</label>
</disp-formula>where <inline-formula id="inf15">
<mml:math id="m22">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>G</mml:mi>
<mml:mi>R</mml:mi>
</mml:mrow>
<mml:mi mathvariant="italic">min</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> is the minimum value of GR, which is about approximately 40 API; <inline-formula id="inf16">
<mml:math id="m23">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>G</mml:mi>
<mml:mi>R</mml:mi>
</mml:mrow>
<mml:mi mathvariant="italic">max</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> is the maximum value of GR, which is about approximately 140 API; <inline-formula id="inf17">
<mml:math id="m24">
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mi>H</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> is the relative value of the clay volume, and GCUR is an empirical coefficient that is approximately 2.</p>
<p>For the volume model of shaley sandstone, the density porosity from the density log, <inline-formula id="inf18">
<mml:math id="m25">
<mml:mrow>
<mml:msub>
<mml:mi>&#x3d5;</mml:mi>
<mml:mi>D</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula>, is as follows:<disp-formula id="e8">
<mml:math id="m26">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="bold-italic">&#x3d5;</mml:mi>
<mml:mi mathvariant="bold-italic">D</mml:mi>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mi mathvariant="bold-italic">&#x3c1;</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi mathvariant="bold-italic">&#x3c1;</mml:mi>
<mml:mrow>
<mml:mi mathvariant="bold-italic">m</mml:mi>
<mml:mi mathvariant="bold-italic">a</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="bold-italic">&#x3c1;</mml:mi>
<mml:mrow>
<mml:mi mathvariant="bold-italic">m</mml:mi>
<mml:mi mathvariant="bold-italic">f</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi mathvariant="bold-italic">&#x3c1;</mml:mi>
<mml:mrow>
<mml:mi mathvariant="bold-italic">m</mml:mi>
<mml:mi mathvariant="bold-italic">a</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfrac>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi mathvariant="bold-italic">V</mml:mi>
<mml:mrow>
<mml:mi mathvariant="bold-italic">s</mml:mi>
<mml:mi mathvariant="bold-italic">h</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mfrac>
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="bold-italic">&#x3c1;</mml:mi>
<mml:mrow>
<mml:mi mathvariant="bold-italic">s</mml:mi>
<mml:mi mathvariant="bold-italic">h</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi mathvariant="bold-italic">&#x3c1;</mml:mi>
<mml:mrow>
<mml:mi mathvariant="bold-italic">m</mml:mi>
<mml:mi mathvariant="bold-italic">a</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="bold-italic">&#x3c1;</mml:mi>
<mml:mrow>
<mml:mi mathvariant="bold-italic">m</mml:mi>
<mml:mi mathvariant="bold-italic">f</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi mathvariant="bold-italic">&#x3c1;</mml:mi>
<mml:mrow>
<mml:mi mathvariant="bold-italic">m</mml:mi>
<mml:mi mathvariant="bold-italic">a</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:math>
<label>(8)</label>
</disp-formula>where <inline-formula id="inf19">
<mml:math id="m27">
<mml:mrow>
<mml:msub>
<mml:mi>&#x3c1;</mml:mi>
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mi>a</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> is the matrix density of tight sandstone (density of quartz), which is about approximately 2.65 g/cm<sup>3</sup>; <inline-formula id="inf20">
<mml:math id="m28">
<mml:mrow>
<mml:msub>
<mml:mi>&#x3c1;</mml:mi>
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mi>f</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> is the density of mud filtrate, which is approximately 1.0 g/cm<sup>3</sup>; and <inline-formula id="inf21">
<mml:math id="m29">
<mml:mrow>
<mml:msub>
<mml:mi>&#x3c1;</mml:mi>
<mml:mrow>
<mml:mi>s</mml:mi>
<mml:mi>h</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> is the density of shale, which is approximately 2.3 g/cm<sup>3</sup>.</p>
<p>In addition, the difference between the neutron porosity and density porosity, ND, is as follows: <disp-formula id="e9">
<mml:math id="m30">
<mml:mrow>
<mml:mi mathvariant="bold-italic">N</mml:mi>
<mml:mi mathvariant="bold-italic">D</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="bold-italic">&#x3d5;</mml:mi>
<mml:mi mathvariant="bold-italic">N</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi mathvariant="bold-italic">&#x3d5;</mml:mi>
<mml:mrow>
<mml:mi mathvariant="bold-italic">N</mml:mi>
<mml:mtext>ma</mml:mtext>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mi mathvariant="bold-italic">&#x3c1;</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi mathvariant="bold-italic">&#x3c1;</mml:mi>
<mml:mrow>
<mml:mi mathvariant="bold-italic">m</mml:mi>
<mml:mi mathvariant="bold-italic">a</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="bold-italic">&#x3c1;</mml:mi>
<mml:mrow>
<mml:mi mathvariant="bold-italic">m</mml:mi>
<mml:mi mathvariant="bold-italic">f</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi mathvariant="bold-italic">&#x3c1;</mml:mi>
<mml:mrow>
<mml:mi mathvariant="bold-italic">m</mml:mi>
<mml:mi mathvariant="bold-italic">a</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:math>
<label>(9)</label>
</disp-formula>where <inline-formula id="inf22">
<mml:math id="m31">
<mml:mrow>
<mml:msub>
<mml:mi>&#x3d5;</mml:mi>
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mtext>ma</mml:mtext>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> is the matrix neutron porosity of sandstone, i.e., the neutron porosity of quartz, which is about &#x2212;0.02, fraction.</p>
</sec>
</sec>
<sec id="s2-2">
<title>2.2 Expert layer</title>
<p>The expert layer in this study has the same structure as the expert layer of the stacking model (<xref ref-type="bibr" rid="B30">Wolpert, 1992</xref>). It is an intelligent system composed of multiple algorithms (experts), and all experts handle the same tasks. During the model training process, the data set is divided into k folds, and all of the experts output simulation prediction results through K simulation training and prediction processes as the input data set of the meta-learner. Finally, all data sets are used to complete the training of the expert. In the actual prediction process, the test set is input to each expert, and their prediction results are input to the meta-learner for combination.</p>
<p>Permeability prediction is an important aspect of formation evaluation. In the stacking model, using strong learners as experts can improve the accuracy and robustness of the model. Therefore, we selected five experts to construct the expert layer: multilayer perceptron (MLP), support vector regression (SVR), ElasticNet (EN), light gradient boosting machine (LightGBM), and category boosting (CatBoost). The characteristics and advantages of the five experts are discussed in detail below:</p>
<p>An MLP is a feed-forward neural network composed of multiple layers of nodes and neurons. It performs nonlinear mapping through activation functions to complete classification and regression tasks. The error between the predicted and actual results is then measured based on the loss function, and a backpropagation algorithm is applied to update the weights and biases of the neural network. The MLP can adapt to different datasets and tasks, has good capabilities for solving complex function-fitting problems, and has certain generalization capabilities.</p>
<p>SVR is a regression algorithm based on a support vector machine (SVM). It maps the original features to a high-dimensional feature space, transforms the regression problem into a convex optimization problem by determining the optimal hyperplane, and solves the optimization problem to determine the best hyperplane. SVR has strong nonlinear modeling capabilities, can solve complex regression tasks involving nonlinear relationships and high-dimensional data, is robust to noise and outliers in the training data, and can effectively avoid overfitting.</p>
<p>ElasticNet regression is an extended form of linear regression. It uses both L1 and L2 regularization terms in the objective function. The L1 regularization term judges the importance of features through the size of the coefficient, such that for unimportant features the coefficient tends to zero, which is suitable for dealing with problems with redundant features. The L2 regularization term controls the complexity of the model and reduces the impact of the correlation between features on the model. The model is more stable when highly correlated features are present. ElasticNet regression has good robustness when dealing with regression problems with redundant features or fewer samples, and it is relatively insensitive to noise and outliers in the data.</p>
<p>LightGBM iteratively trains multiple weak learners using the gradient boosting algorithm and optimizes and adjusts the newly generated weak learners according to the objective function to improve the accuracy of the model. It retains samples with larger gradients to accelerate the training process and uses a leaf-wise growth algorithm with depth restrictions to prevent overfitting. LightGBM sorts features according to their importance and performs feature selection based on thresholds, thereby improving the generalization ability and interpretability of the model. Further, LightGBM is suitable for processing large-scale data sets with uneven feature distributions.</p>
<p>CatBoost is a gradient boosting algorithm that is specially designed to effectively handle categorical features. The impact of most noise and outliers is diluted in the entire tree structure, thereby reducing their impact on individual nodes. CatBoost dynamically adjusts the learning rate according to the distribution of data and the complexity of the model and has strong accuracy and robustness when processing datasets with noise and outliers.</p>
</sec>
<sec id="s2-3">
<title>2.3 Meta-learner</title>
<p>In the classic stacking model, the meta-learner is between the expert and output layers. Its main function is to combine the prediction results of different experts in a manner that minimizes errors, as shown in <xref ref-type="fig" rid="F1">Figure 1</xref>. By generalizing the outputs of multiple expert models, the prediction accuracy of the overall model is improved, and the final prediction result is output. When choosing a meta-learner, the capacity and complexity should be considered. Stronger meta-learners may overfit the model, whereas weaker meta-learners may fail to capture complex relationships. Because the meta-learner only learns the combination of expert models, only simple machine learning models (weak learners) can be used; when dealing with high-dimensional feature sets output by multiple expert models, the fitting effect is often poor.</p>
<fig id="F1" position="float">
<label>FIGURE 1</label>
<caption>
<p>Stacking model architecture consisting of an input layer, experts, prediction results, meta-learner, and output layer.</p>
</caption>
<graphic xlink:href="feart-12-1364515-g001.tif"/>
</fig>
<p>In this study, the AdaBoost model (<xref ref-type="bibr" rid="B8">Freund and Schapire, 1997</xref>) is used as the meta-learner, and the weak learner in the model is a linear regression model. During the training process for each learner, the sample weight is adjusted according to the error rate of the previous round, and greater weights are assigned to the misclassified samples to iteratively learn and correct the errors, thereby improving the accuracy of the meta-learner model. The GNS model proposed in this study uses a gate network to assign different weights to the input data of the meta-learner and adds physical model constraints to allow the meta-learner to better find the best combination of expert models while complying with the physical model constraints.</p>
</sec>
<sec id="s2-4">
<title>2.4 Gate network</title>
<p>A gate network is a neural network structure that is used to control the flow and processing of information in a model. By adding a gate network mechanism, the input features can be selectively filtered, amplified, or suppressed, thus forming a constraint function for the model. In the classic stacking model, differences in expert performance cause the relationship between the features and learning objectives to become more complex. In addition, the impact of the added experts on the overall model cannot be measured. Therefore, this study introduces gate networks to solve the aforementioned problems, adjust the model architecture, and improve the accuracy and interpretability of the overall model.</p>
<p>In the GNS model, the role of the gate network includes three main aspects.<list list-type="simple">
<list-item>
<p>a. In the input layer, the gate network extracts the input petrophysical model and adds physical constraints to the input data of all expert models and meta-learners to improve the prediction performance and stability of the model and achieve model driving.</p>
</list-item>
<list-item>
<p>b. In the k-fold cross-validation of the expert model, the weight calculated according to the mean squared error (MSE) of each expert model is saved to the gate network, and the features corresponding to each expert model are weighted according to the weight combination saved by the gate network, which reduces the complexity of the meta-learner learning process and improves the robustness. The expert weights of the storage entry network are shown in Eq. <xref ref-type="disp-formula" rid="e10">10</xref>, and the expert prediction results calculated based on the combination of weights are shown in Eq. <xref ref-type="disp-formula" rid="e11">11</xref>.</p>
</list-item>
</list>
<disp-formula id="e10">
<mml:math id="m32">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="bold-italic">w</mml:mi>
<mml:mrow>
<mml:mi mathvariant="bold-italic">t</mml:mi>
<mml:mo>_</mml:mo>
<mml:mi mathvariant="bold-italic">G</mml:mi>
<mml:mi mathvariant="bold-italic">N</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mn mathvariant="bold">1</mml:mn>
<mml:mrow>
<mml:mi mathvariant="bold-italic">M</mml:mi>
<mml:mi mathvariant="bold-italic">S</mml:mi>
<mml:msub>
<mml:mi mathvariant="bold-italic">E</mml:mi>
<mml:mi mathvariant="bold-italic">t</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfrac>
<mml:mo>/</mml:mo>
<mml:mstyle displaystyle="true">
<mml:munderover>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi mathvariant="bold-italic">i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn mathvariant="bold">1</mml:mn>
</mml:mrow>
<mml:mi mathvariant="bold-italic">T</mml:mi>
</mml:munderover>
</mml:mstyle>
<mml:mfrac>
<mml:mn mathvariant="bold">1</mml:mn>
<mml:mrow>
<mml:mi mathvariant="bold-italic">M</mml:mi>
<mml:mi mathvariant="bold-italic">S</mml:mi>
<mml:msub>
<mml:mi mathvariant="bold-italic">E</mml:mi>
<mml:mi mathvariant="bold-italic">i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:math>
<label>(10)</label>
</disp-formula>where T is the total number of expert models, and <inline-formula id="inf23">
<mml:math id="m33">
<mml:mrow>
<mml:mi>M</mml:mi>
<mml:mi>S</mml:mi>
<mml:msub>
<mml:mi>E</mml:mi>
<mml:mi>t</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> is the MSE of the t-th expert.<disp-formula id="e11">
<mml:math id="m34">
<mml:mrow>
<mml:msub>
<mml:mi mathvariant="bold-italic">H</mml:mi>
<mml:mi mathvariant="bold-italic">t</mml:mi>
</mml:msub>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="bold-italic">x</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x3d;</mml:mo>
<mml:msub>
<mml:mi mathvariant="bold-italic">w</mml:mi>
<mml:mrow>
<mml:mi mathvariant="bold-italic">t</mml:mi>
<mml:mo>_</mml:mo>
<mml:mi mathvariant="bold-italic">G</mml:mi>
<mml:mi mathvariant="bold-italic">N</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x22c5;</mml:mo>
<mml:msub>
<mml:mi mathvariant="bold-italic">h</mml:mi>
<mml:mi mathvariant="bold-italic">t</mml:mi>
</mml:msub>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi mathvariant="bold-italic">x</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(11)</label>
</disp-formula>where <inline-formula id="inf24">
<mml:math id="m35">
<mml:mrow>
<mml:msub>
<mml:mi>h</mml:mi>
<mml:mi>t</mml:mi>
</mml:msub>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi>x</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula> is the prediction result of the t-th expert and <inline-formula id="inf25">
<mml:math id="m36">
<mml:mrow>
<mml:msub>
<mml:mi>H</mml:mi>
<mml:mi>t</mml:mi>
</mml:msub>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi>x</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula> is the prediction result of the t-th expert after assigning weights through the gate network.<list list-type="simple">
<list-item>
<p>c. When the meta learner outputs results, the accuracy indicators of the expert model combination are evaluated through a gate network. Experts who have a positive impact on the accuracy indicators are retained, abandoning experts who have a negative impact on the accuracy indicators are abandoned to ensure that the overall performance of the model is improved.</p>
</list-item>
</list>
</p>
</sec>
<sec id="s2-5">
<title>2.5 GNS model architecture and workflow</title>
<p>For the formation permeability prediction problem, the aforementioned components are used in this study to form a GNS model. First, a dataset is created and labeled. The logging dataset includes the natural gamma-ray log (GR), spectral gamma-ray log (K-TH-U), compressional slow-ness log (DT), neutron log (CN), and density log (DEN). Lab-measured core permeability data points are used as labels. In addition to the logging data mentioned above, <inline-formula id="inf26">
<mml:math id="m37">
<mml:mrow>
<mml:msub>
<mml:mi>V</mml:mi>
<mml:mrow>
<mml:mi>s</mml:mi>
<mml:mi>h</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula>, <inline-formula id="inf27">
<mml:math id="m38">
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mi>D</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, and <inline-formula id="inf28">
<mml:math id="m39">
<mml:mrow>
<mml:msub>
<mml:mi>&#x3d5;</mml:mi>
<mml:mi>D</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula>, calculated using Eqs <xref ref-type="disp-formula" rid="e6">6</xref>&#x2013;<xref ref-type="disp-formula" rid="e9">9</xref>, are also input into the GNS, called petrophysical constraints. Finally, a petrophysical and data-driven intelligent model is developed.</p>
<p>In the input layer, the optimal combination of the corresponding well log curves is selected for each expert model based on the SHAP value, and the petrophysical constraints are input into the gate network. In the expert layer, petrophysical constraints are added to the input data of the expert model through the gate network, and the hyperparameters of the model are self-optimized. The Bayesian optimization method is used in this study to optimize the parameters; it estimates the posterior distribution of the objective function by constructing a Gaussian process (GP) model and determines the hyperparameter value for the next sampling, thereby quickly finding the global optimal solution. Cross-validation partitioning the dataset into multiple subsets, training the model on some of these subsets, and evaluating its performance on the remaining subsets. After model optimization is completed, each expert model is simulated and predicted using a five-fold cross-validation method. The performance of the expert model is evaluated according to the MSE of the model and stored in the gate network. Each expert model is then trained using all of the datasets.</p>
<p>According to the ranking of expert performance indicators, the expert model simulation prediction results are added to the meta-learner and given weights to form the input set of the meta-learner. At the same time, the gate network adds petrophysical constraints.</p>
<p>Finally, the meta-learner is trained. We evaluate the output results, retain the expert models that positively impact the overall model performance, eliminate the expert models that negatively impact the overall model, and finally output the prediction results of the best model combination.</p>
<p>Because the expert model adds dynamic weight constraints, it is more robust and interpretable than the classic stacking algorithm. <xref ref-type="fig" rid="F2">Figure 2</xref> shows the GNS workflow based on petrophysical and data-driven methods.</p>
<fig id="F2" position="float">
<label>FIGURE 2</label>
<caption>
<p>GNS model workflow consisting of five parts: the input layer, expert layer, meta-learner, gate network, and output layer.</p>
</caption>
<graphic xlink:href="feart-12-1364515-g002.tif"/>
</fig>
</sec>
<sec id="s2-6">
<title>2.6 Performance evaluation</title>
<p>To verify the reliability and accuracy of the proposed GNS model, three statistical parameters are introduced as performance evaluation indicators. The equations used to calculate each parameter are listed in <xref ref-type="table" rid="T1">Table 1</xref>.</p>
<table-wrap id="T1" position="float">
<label>TABLE 1</label>
<caption>
<p>Parameters for assessment to evaluate the model performance.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="center">Accuracy measure</th>
<th align="center">Mathematical exp</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="center">Coefficient of determination (R<sup>2</sup>)</td>
<td align="center">
<inline-formula id="inf29">
<mml:math id="m40">
<mml:mrow>
<mml:msup>
<mml:mi>R</mml:mi>
<mml:mn>2</mml:mn>
</mml:msup>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
<mml:mo>&#x2212;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:munderover>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi>N</mml:mi>
</mml:munderover>
</mml:mstyle>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi>y</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mover accent="true">
<mml:mi>y</mml:mi>
<mml:mo>&#x5e;</mml:mo>
</mml:mover>
<mml:mi>i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:munderover>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi>N</mml:mi>
</mml:munderover>
</mml:mstyle>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi>y</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mover accent="true">
<mml:mi>y</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
<mml:mi>i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
</tr>
<tr>
<td align="center">Mean square error (MSE)</td>
<td align="center">
<inline-formula id="inf30">
<mml:math id="m41">
<mml:mrow>
<mml:mi>M</mml:mi>
<mml:mi>S</mml:mi>
<mml:mi>E</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>N</mml:mi>
</mml:mrow>
</mml:mfrac>
<mml:mstyle displaystyle="true">
<mml:munderover>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi>N</mml:mi>
</mml:munderover>
</mml:mstyle>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi>y</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:mover accent="true">
<mml:msub>
<mml:mi>y</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
<mml:mo>&#x5e;</mml:mo>
</mml:mover>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
</tr>
<tr>
<td align="center">Maximum absolute error (MAE)</td>
<td align="center">
<inline-formula id="inf31">
<mml:math id="m42">
<mml:mrow>
<mml:mi>M</mml:mi>
<mml:mi>A</mml:mi>
<mml:mi>E</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>N</mml:mi>
</mml:mrow>
</mml:mfrac>
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:munderover>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi>N</mml:mi>
</mml:munderover>
</mml:mstyle>
<mml:mrow>
<mml:mfenced open="|" close="|" separators="|">
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi>y</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:mover accent="true">
<mml:msub>
<mml:mi>y</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
<mml:mo>&#x5e;</mml:mo>
</mml:mover>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
</sec>
<sec id="s3">
<title>3 Case study</title>
<sec id="s3-1">
<title>3.1 Geological background</title>
<p>The Dibei tight gas reservoir in the Kuqa Depression of the Tarim Basin has become a key exploration area because of its large gas reserve potential. The main production and storage unit is the Jurassic Ahe formation. The lithology is mainly coarse sandstone, and the sedimentary type is a braided river delta plain channel. The rock type in the sedimentation is mainly lithic sandstone, followed by feldspathic lithic sandstone, with a quartz content of more than 60%. Feldspar, clay minerals, calcite, and dolomite are all developed. The Dibei gas reservoir has low porosity, a complex pore throat structure, and a wide range of permeability changes. Multiple factors jointly control the permeability of the reservoir, and the predictive effect of the empirical model is poor. It is necessary to apply a model that can identify nonlinear characteristics between well logs and core data. Therefore, we apply the GNS model to predict the permeability.</p>
</sec>
<sec id="s3-2">
<title>3.2 Data acquisition</title>
<p>The stability and accuracy of the model depend on the reliability of the training data. In this study, 1,088 samples were collected from five wells in the Dibei area. The locations of the core wells are shown in <xref ref-type="fig" rid="F3">Figure 3</xref>. The input variables include seven conventional well logs (GR, K, TH, U, DT, DEN, and CN) and three petrophysical constraints (ND, <inline-formula id="inf32">
<mml:math id="m43">
<mml:mrow>
<mml:msub>
<mml:mi>V</mml:mi>
<mml:mrow>
<mml:mi>s</mml:mi>
<mml:mi>h</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> , and <inline-formula id="inf33">
<mml:math id="m44">
<mml:mrow>
<mml:msub>
<mml:mi>&#x3d5;</mml:mi>
<mml:mi>D</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula>); the permeability value is the desired output. To ensure that the GNS model has the best possible performance, the training data were normalized. Based on the data normalization, we calculated the SHAP value of each input feature, as shown in <xref ref-type="fig" rid="F4">Figure 4</xref>, which reflects the contribution of each feature to the model prediction, further indicating the relative importance of each feature to the prediction results.</p>
<fig id="F3" position="float">
<label>FIGURE 3</label>
<caption>
<p>The Dibei gas reservoir is located in the northern part of the Kuqa Depression, as indicated by the red box in the thumbnail; the red dot indicates the location of the coring well.</p>
</caption>
<graphic xlink:href="feart-12-1364515-g003.tif"/>
</fig>
<fig id="F4" position="float">
<label>FIGURE 4</label>
<caption>
<p>SHAP contribution of well logs and changes in the MSE of well log combinations in each expert model.</p>
</caption>
<graphic xlink:href="feart-12-1364515-g004.tif"/>
</fig>
<p>The bar chart in <xref ref-type="fig" rid="F4">Figure 4</xref> shows that the contributions of the SHAP values in descending order are DEN, GR, TH, CN, K, DT, and U. Five different expert models were trained based on the SHAP value contribution of the well log, and the MSE of the training results is shown in the line chart in <xref ref-type="fig" rid="F4">Figure 4</xref>. The MSE reduction is defined as a log that has a positive contribution to the model, while conversely, when it has a negative contribution to the model, the positive contribution log is retained as the input of the expert model, and different log combinations are input for different expert models. Together with the permeability label, these constitute the training set, as summarized in <xref ref-type="table" rid="T2">Table 2</xref>.</p>
<table-wrap id="T2" position="float">
<label>TABLE 2</label>
<caption>
<p>Well log input combinations for different experts.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="center">Model</th>
<th align="center">Well logs</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="center">SVR</td>
<td align="center">DEN, GR, TH, CN, DT</td>
</tr>
<tr>
<td align="center">MLP</td>
<td align="center">DEN, TH, K, DT, U</td>
</tr>
<tr>
<td align="center">EN</td>
<td align="center">DEN, GR, TH</td>
</tr>
<tr>
<td align="center">LightGBM</td>
<td align="center">DEN, GR, TH, CN, K, DT, U</td>
</tr>
<tr>
<td align="center">CatBoost</td>
<td align="center">DEN, GR, TH, CN, K, DT, U</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec id="s3-3">
<title>3.3 Model training and performance evaluation</title>
<p>The feature combinations in <xref ref-type="table" rid="T2">Table 2</xref> are input into the expert model for training, and the Bayesian optimization method is used to find the optimal hyperparameter values such that the model can achieve the best performance on the target task. A performance comparison of the expert model using optimal hyperparameter values and the expert model using default parameters is summarized in <xref ref-type="table" rid="T3">Table 3</xref>.</p>
<table-wrap id="T3" position="float">
<label>TABLE 3</label>
<caption>
<p>Expert performance comparison when applying optimized and default hyperparameters.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="center"/>
<th align="center">SVR</th>
<th align="center">MLP</th>
<th align="center">EN</th>
<th align="center">LightGBM</th>
<th align="center">CatBoost</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="center">MSE of experts using default hyperparameter values</td>
<td align="center">0.38</td>
<td align="center">0.39</td>
<td align="center">0.44</td>
<td align="center">0.21</td>
<td align="center">0.23</td>
</tr>
<tr>
<td align="center">MSE of experts using optimal hyperparameter values</td>
<td align="center">0.34</td>
<td align="center">0.36</td>
<td align="center">0.39</td>
<td align="center">0.18</td>
<td align="center">0.21</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>The petrophysical constraints stored in the gate network are added to the feature combination of each expert to form the input dataset for each expert. Targeting the core permeability, multiple random sampling and no-replacement iterative predictions are conducted on experts to obtain a set of simulation prediction results, record the expert model MSE, save it to the gate network, and use all datasets to train each expert model. The expert combination is formed iteratively based on the MSE of the expert model (from small to large), the weight of each expert in the combination is calculated (Eq. <xref ref-type="disp-formula" rid="e8">8</xref>), and the corresponding simulation prediction results are weighted (Eq. <xref ref-type="disp-formula" rid="e9">9</xref>). Next, the weighted prediction results and petrophysical constraints are input into the meta-learner for training to determine the optimal expert model combination strategy. To demonstrate the effectiveness and advantages of the GNS algorithm, a classic stacking model is used for comparison. To ensure consistency, the same experts and meta-learners are used in the models. A performance comparison of the two models is shown in <xref ref-type="fig" rid="F5">Figure 5</xref>.</p>
<fig id="F5" position="float">
<label>FIGURE 5</label>
<caption>
<p>MSE of the expert model and changes in the MSE when iteratively inputting the ensemble learning model.</p>
</caption>
<graphic xlink:href="feart-12-1364515-g005.tif"/>
</fig>
<p>It can be seen that the addition of the gate network clearly improves the overall model. In contrast, the classic stacking model lacks petrophysical constraints and expert evaluation mechanisms and cannot optimize over-fitting, coupling correlation, and other problems existing in the integration process of experts; as a result, the overall performance of the stacking model is inferior to that of the GNS. For example, in the fourth iteration shown in <xref ref-type="fig" rid="F5">Figure 5</xref>, the MLP model may have problems with dataset misfit or poor coupling with other expert results, resulting in negative improvements to the overall model when an expert is added to the stacking model. In the GNS algorithm, the addition of this expert is judged to be a negative improvement. Therefore, the prediction results of this expert are eliminated from the overall model, ultimately maintaining the quality of the overall expert group.</p>
<p>The above analysis shows that the performance of the GNS model is better than that of the stacking model, indicating that the improvement in the gate network is effective. To discuss the superiority of the model architecture, this study conducts a horizontal comparison with other ensemble learning models. We selected three representative integration strategies: two heterogeneous integration models (RCM and Voting) and a bagging integration model constructed using LightGBM. Three evaluation indicators&#x2014;MSE, MAE, and R2&#x2014;are used to evaluate the performance of the model. The results are summarized in <xref ref-type="table" rid="T4">Table 4</xref>. The GNS and RCM models have better prediction performance, and the GNS model achieves the best prediction results.</p>
<table-wrap id="T4" position="float">
<label>TABLE 4</label>
<caption>
<p>Performance evaluation parameters of different ensemble learning models.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="center">Model</th>
<th align="center">MSE</th>
<th align="center">MAE</th>
<th align="center">R<sup>2</sup>
</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="center">GNS</td>
<td align="center">0.1364</td>
<td align="center">0.2802</td>
<td align="center">0.7291</td>
</tr>
<tr>
<td align="center">RCM</td>
<td align="center">0.1762</td>
<td align="center">0.2941</td>
<td align="center">0.6124</td>
</tr>
<tr>
<td align="center">Bagging</td>
<td align="center">0.2230</td>
<td align="center">0.3467</td>
<td align="center">0.5566</td>
</tr>
<tr>
<td align="center">Voting</td>
<td align="center">0.2764</td>
<td align="center">0.4005</td>
<td align="center">0.4512</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>
<xref ref-type="fig" rid="F6">Figure 6A</xref> shows the error between the core permeability measured in the laboratory and that predicted using the GNS model. To compare the prediction effect more intuitively, the permeability is plotted on a cross diagram (<xref ref-type="fig" rid="F6">Figure 6B</xref>), where the 45&#xb0; diagonal line represents a perfect match between the predicted and true penetration rates. The prediction error of the model conforms to a normal distribution; therefore, the accuracy of the model can be judged based on the variance and MSE (<xref ref-type="bibr" rid="B10">Helle et al., 2001</xref>; <xref ref-type="bibr" rid="B35">Zhong et al., 2019</xref>). <xref ref-type="fig" rid="F7">Figures 7</xref>&#x2013;<xref ref-type="fig" rid="F9">9</xref> show the prediction results of the RCM, bagging, and voting models, respectively.</p>
<fig id="F6" position="float">
<label>FIGURE 6</label>
<caption>
<p>
<bold>(A)</bold> Histogram of the error between the core measured permeability value and the permeability predicted by the GNS model. <bold>(B)</bold> Cross plot of the core measured permeability value and permeability predicted by the GNS model.</p>
</caption>
<graphic xlink:href="feart-12-1364515-g006.tif"/>
</fig>
<fig id="F7" position="float">
<label>FIGURE 7</label>
<caption>
<p>
<bold>(A)</bold> Histogram of the error between the core measured permeability value and the permeability predicted by the RCM model. <bold>(B)</bold> Cross plot of the core measured permeability value and permeability predicted by the RCM model.</p>
</caption>
<graphic xlink:href="feart-12-1364515-g007.tif"/>
</fig>
<fig id="F8" position="float">
<label>FIGURE 8</label>
<caption>
<p>
<bold>(A)</bold> Histogram of the error between the core measured permeability value and the permeability predicted by the bagging model. <bold>(B)</bold> Cross plot of the core measured permeability value and permeability predicted by the bagging model.</p>
</caption>
<graphic xlink:href="feart-12-1364515-g008.tif"/>
</fig>
<fig id="F9" position="float">
<label>FIGURE 9</label>
<caption>
<p>
<bold>(A)</bold> Histogram of the error between the core measured permeability value and the permeability predicted by the voting model. <bold>(B)</bold> Cross plot of the core measured permeability value and permeability predicted by the voting model.</p>
</caption>
<graphic xlink:href="feart-12-1364515-g009.tif"/>
</fig>
<p>The GNS model has a variance of 0.0931 and MSE of 0.1364. It has the highest degree of fit with the core permeability values. The error does not exceed one order of magnitude for any of the core samples, showing the best performance (<xref ref-type="fig" rid="F6">Figure 6</xref>). The variance of the RCM model is 0.1109 and the MSE is 0.1762. The prediction results are slightly lower in the high-permeability interval, and the performance is slightly weaker than that of the GNS (<xref ref-type="fig" rid="F7">Figure 7</xref>). The variance of the bagging model is 0.1272, and the MSE is 0.2230. Some sample points exhibit large errors at different permeability intervals. This may be because random sampling will lead to a loss of some useful information (<xref ref-type="fig" rid="F8">Figure 8</xref>). Finally, the voting model has the largest variance and MSE, and the predicted permeability has a larger error than the core permeability (<xref ref-type="fig" rid="F9">Figure 9</xref>). These studies demonstrate that the GNS model is consistent with the measured permeability of the cores.</p>
</sec>
<sec id="s3-4">
<title>3.4 Prediction results</title>
<p>Well N4 is a core well in the Dibei gas reservoir in the Tarim Basin. The target reservoir is the Jurassic Ahe Formation. Logging data includes natural gamma ray log (GR), porosity logs (DEN-CN-DT), and natural gamma spectrum logs (K-TH-U). In this study, we selected the above seven well logs and three petrophysical constraints (VSH, PHID, and ND) as inputs, used the core permeability measured in the laboratory as label data, and entered it into the GNS model for prediction. For comparison, the results of several other ensemble learning models were obtained. <xref ref-type="fig" rid="F10">Figure 10</xref> shows the prediction results for each model in Well N4. The first track is the natural gamma curve, the second track is the natural gamma spectrum curve, the third track is the porosity log curve, and the fourth track is the petrophysical constraints. Tracks 6, 7, 8, and 9 are the comparisons between the prediction results of the GNS model, RCM model, bagging model, and voting model and the true value of the core permeability, respectively. The prediction results of the GNS model exhibit the best agreement with the core measurement results. The RCM and bagging models have certain errors at lower permeabilities. The voting model exhibits the largest difference from the core measurement results. The case study shows that the GNS model has advantages in predicting permeability. Importantly, we applied the gate network to integrate the petrophysical constraints with the model, rather than simply taking the petrophysical constraints as inputs. In addition, we controlled the weight of the model through the gate network, avoiding the influence of weak experts. Thus, the proposed model is more flexible than previously established empirical equations and other machine learning models.</p>
<fig id="F10" position="float">
<label>FIGURE 10</label>
<caption>
<p>Permeability prediction results for Well N4 in the Dibei Gas Reservoir, Tarim Basin.</p>
</caption>
<graphic xlink:href="feart-12-1364515-g010.tif"/>
</fig>
</sec>
</sec>
<sec sec-type="discussion" id="s4">
<title>4 Discussion</title>
<p>In recent years, machine learning has promoted the intelligent development of log interpretation. In the process of supervised learning, machine learning automatically adjusts the internal structure to satisfy the mapping relationship between the feature data and label data and establishes a prediction model with certain generalization capabilities. However, data-driven learning is the core of current machine learning. Most studies have only considered the use of indicators such as accuracy and mean square error to evaluate the performance of machine learning models, ignoring the impact of physical constraints on the overall model. Therefore, even a model with superior performance is not convincing in its predictions for unknown samples. It is clear that a single data-driven method cannot meet the requirements of log interpretation tasks. The combination of model- and data-driven methods is a development trend for the future application of machine learning models in the field of log interpretation. This approach considers model performance and interpretability. While ensuring the accuracy of the model, it also ensures that it conforms to the physical laws of the logging interpretation process.</p>
<p>The stacking model typically consists of an expert model layer and a meta-learner, which has better performance than a single model in actual tasks. However, the meta-learner only learns the combination of experts and does not perform any evaluation or control of the experts, resulting in poor interpretability and difficulty in explaining the contribution of each model. The RCM model relies on the performance of experts, and some weak models may have a greater impact on the results and a weak ability to solve the problem of data imbalance. The bagging model aims to reduce the variance. When the data distribution was uneven, each basic model learned different features from different data subsets. In addition, if the basic model itself has overfitting or underfitting problems, bagging cannot solve them. The voting algorithm relies on the performance of experts. Because the weights are the same, the information obtained by each model may not be fully utilized. The voting model is sensitive to noise. In this study, it is also found that for sample sets with large differences in permeability, noise causes the model to obtain poor results. The GNS model controls the model through the gate network mechanism, adds petrophysical constraints to the inputs of the expert model and meta-learner, dynamically assigns weights to the output of the expert model, and finally iteratively evaluates the overall performance of the model. Connecting each training step of the model and adding constraints uniformly improves the shortcomings of the traditional model and significantly enhances the interpretability and credibility of the overall model.</p>
<p>However, petrophysical constraints are calculated from well logs, and there is still a certain similarity in their features. During the model training process, feature redundancy is possible, and the information dimension may not be sufficiently rich. Therefore, the choice of petrophysical constraints is important. In the future research, while applying petrophysical constraints, more logging information can be used as the input or constraint of the model, such as imaging logging to achieve better results.</p>
</sec>
<sec sec-type="conclusion" id="s5">
<title>5 Conclusion</title>
<p>This paper proposes a workflow for reservoir permeability prediction based on the GNS model. The SHAP value is used to analyze the contribution of the feature curve, select the optimal feature combination to build the traditional stacking model, and use the gate network to control and dynamically optimize the overall model. Finally, the results of other models are compared to verify the superior performance of the GNS model. The following conclusions can be drawn from this study.<list list-type="simple">
<list-item>
<p>1) The selection of a well log is crucial for the prediction performance of intelligent algorithms. Intelligent algorithms based on appropriate features can output more accurate results. SHAP contribution analysis is an effective means of evaluating the importance of well logs.</p>
</list-item>
<list-item>
<p>2) Based on the classical stacking model, the GNS model applies gate networks to control and dynamically optimize the overall model. GNS is an algorithm driven by both data and models. It retains the advantages of the stacking model&#x2019;s strong expressive ability and significantly improves the interpretability of the algorithm. Compared with single data-driven algorithms, it is more reliable and superior in practical logging task applications.</p>
</list-item>
<list-item>
<p>3) The GNS model can effectively solve the nonlinear prediction problem of high-dimensional features. In the tight sandstone reservoir permeability prediction example, the results of the GNS model are closer to the laboratory-measured permeability than the results of the RCM, voting, and bagging models, thus verifying that the proposed model is an accurate method.</p>
</list-item>
</list>
</p>
</sec>
</body>
<back>
<sec sec-type="data-availability" id="s6">
<title>Data availability statement</title>
<p>The data analyzed in this study is subject to the following licenses/restrictions: exploratory research. Requests to access these datasets should be directed to PS, <email>shipy0410@126.com</email>.</p>
</sec>
<sec id="s7">
<title>Author contributions</title>
<p>PS: Conceptualization, Investigation, Methodology, Software, Validation, Writing&#x2013;original draft, Writing&#x2013;review and editing, Data curation, Formal Analysis. PS: Data curation, Methodology, Software, Validation, Writing&#x2013;original draft, Writing&#x2013;review and editing. KB: Validation, Writing&#x2013;review and editing. CH: Validation, Writing&#x2013;review and editing. XN: Validation, Writing&#x2013;review and editing. ZM: Supervision, Validation, Writing&#x2013;review and editing. PZ: Funding acquisition, Validation, Writing&#x2013;review and editing, Supervision.</p>
</sec>
<sec sec-type="funding-information" id="s8">
<title>Funding</title>
<p>The author(s) declare financial support was received for the research, authorship, and/or publication of this article. The Science Foundation of China University of Petroleum, Beijing (2462020BJRC001).</p>
</sec>
<sec sec-type="COI-statement" id="s9">
<title>Conflict of interest</title>
<p>Authors KB, CH and XN were employed by PetroChina Tarim Oilfield Company.</p>
<p>The remaining authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec sec-type="disclaimer" id="s10">
<title>Publisher&#x2019;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<ref-list>
<title>References</title>
<ref id="B1">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ahmadi</surname>
<given-names>M. A.</given-names>
</name>
<name>
<surname>Ebadi</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Shokrollahi</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Majidi</surname>
<given-names>S. M. J.</given-names>
</name>
</person-group> (<year>2013</year>). <article-title>Evolving artificial neural network and imperialist competitive algorithm for prediction oil flow rate of the reservoir</article-title>. <source>Appl. Soft Comput.</source> <volume>13</volume> (<issue>2</issue>), <fpage>1085</fpage>&#x2013;<lpage>1098</lpage>. <pub-id pub-id-type="doi">10.1016/j.asoc.2012.10.009</pub-id>
</citation>
</ref>
<ref id="B2">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ahmed</surname>
<given-names>U.</given-names>
</name>
<name>
<surname>Crary</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Coates</surname>
<given-names>G.</given-names>
</name>
</person-group> (<year>1991</year>). <article-title>Permeability estimation: the various sources and their interrelationships</article-title>. <source>J. Pet. Technol.</source> <volume>43</volume> (<issue>5</issue>), <fpage>578</fpage>&#x2013;<lpage>587</lpage>. <pub-id pub-id-type="doi">10.2118/19604-PA</pub-id>
</citation>
</ref>
<ref id="B3">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Al-Anazi</surname>
<given-names>A. F.</given-names>
</name>
<name>
<surname>Gates</surname>
<given-names>I. D.</given-names>
</name>
</person-group> (<year>2012</year>). <article-title>Support vector regression to predict porosity and permeability: effect of sample size</article-title>. <source>Comput. Geosci.</source> <volume>39</volume>, <fpage>64</fpage>&#x2013;<lpage>76</lpage>. <pub-id pub-id-type="doi">10.1016/j.cageo.2011.06.011</pub-id>
</citation>
</ref>
<ref id="B4">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Bai</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Tan</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Shi</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>G.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Regression committee machine and petrophysical model jointly driven parameters prediction from wireline logs in tight sandstone reservoirs</article-title>. <source>IEEE Trans. Geosci. Remote Sens.</source> <volume>60</volume>, <fpage>1</fpage>&#x2013;<lpage>9</lpage>. <pub-id pub-id-type="doi">10.1109/TGRS.2020.3041366</pub-id>
</citation>
</ref>
<ref id="B5">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Carman</surname>
<given-names>P. C.</given-names>
</name>
</person-group> (<year>1937</year>). <article-title>Fluid flow through a granular bed</article-title>. <source>Trans. Instit. Chem. Eng.</source> <volume>15</volume>, <fpage>150</fpage>&#x2013;<lpage>156</lpage>. <pub-id pub-id-type="doi">10.1016/S0263-8762(97)80003-2</pub-id>
</citation>
</ref>
<ref id="B6">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Chen</surname>
<given-names>C. H.</given-names>
</name>
<name>
<surname>Lin</surname>
<given-names>Z. S.</given-names>
</name>
</person-group> (<year>2006</year>). <article-title>A committee machine with empirical formulas for permeability prediction</article-title>. <source>Comput. Geosci.</source> <volume>32</volume> (<issue>4</issue>), <fpage>485</fpage>&#x2013;<lpage>496</lpage>. <pub-id pub-id-type="doi">10.1016/j.cageo.2005.08.003</pub-id>
</citation>
</ref>
<ref id="B7">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Coates</surname>
<given-names>G. R.</given-names>
</name>
<name>
<surname>Dumanoir</surname>
<given-names>J. L.</given-names>
</name>
</person-group> (<year>1973</year>). &#x201c;<article-title>A new approach to improved log-derived permeability</article-title>,&#x201d; in <source>SPWLA annual logging symposium</source> (<publisher-loc>Lafayette, Louisiana</publisher-loc>: <publisher-name>SPWLA</publisher-name>).</citation>
</ref>
<ref id="B8">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Freund</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Schapire</surname>
<given-names>R. E.</given-names>
</name>
</person-group> (<year>1997</year>). <article-title>A decision-theoretic generalization of on-line learning and an application to boosting</article-title>. <source>J. Comput. Syst. Sci.</source> <volume>55</volume> (<issue>1</issue>), <fpage>119</fpage>&#x2013;<lpage>139</lpage>. <pub-id pub-id-type="doi">10.1006/jcss.1997.1504</pub-id>
</citation>
</ref>
<ref id="B9">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Friedman</surname>
<given-names>J. H.</given-names>
</name>
</person-group> (<year>2001</year>). <article-title>Greedy function approximation: a gradient boosting machine</article-title>. <source>Ann. Stat.</source> <volume>29</volume> (<issue>5</issue>), <fpage>1189</fpage>&#x2013;<lpage>1232</lpage>. <pub-id pub-id-type="doi">10.1214/aos/1013203451</pub-id>
</citation>
</ref>
<ref id="B10">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Helle</surname>
<given-names>H. B.</given-names>
</name>
<name>
<surname>Bhatt</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Ursin</surname>
<given-names>B.</given-names>
</name>
</person-group> (<year>2001</year>). <article-title>Porosity and permeability prediction from wireline logs using artificial neural networks: a North Sea case study</article-title>. <source>Geophys. Prospect.</source> <volume>49</volume> (<issue>4</issue>), <fpage>431</fpage>&#x2013;<lpage>444</lpage>. <pub-id pub-id-type="doi">10.1046/j.1365-2478.2001.00271.x</pub-id>
</citation>
</ref>
<ref id="B11">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Jamialahmadi</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Javadpour</surname>
<given-names>F.</given-names>
</name>
</person-group> (<year>2000</year>). <article-title>Relationship of permeability, porosity and depth using an artificial neural network</article-title>. <source>J. Pet. Sci. Eng.</source> <volume>26</volume> (<issue>1-4</issue>), <fpage>235</fpage>&#x2013;<lpage>239</lpage>. <pub-id pub-id-type="doi">10.1016/S0920-4105(00)00037-1</pub-id>
</citation>
</ref>
<ref id="B12">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kalule</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Abderrahmane</surname>
<given-names>H. A.</given-names>
</name>
<name>
<surname>Alameri</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Sassi</surname>
<given-names>M.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Stacked ensemble machine learning for porosity and absolute permeability prediction of carbonate rock plugs</article-title>. <source>Sci. Rep.</source> <volume>13</volume> (<issue>1</issue>), <fpage>9855</fpage>. <pub-id pub-id-type="doi">10.1038/s41598-023-36096-2</pub-id>
</citation>
</ref>
<ref id="B13">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Karimpouli</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Fathianpour</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Roohi</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2010</year>). <article-title>A new approach to improve neural networks&#x27; algorithm in permeability prediction of petroleum reservoirs using supervised committee machine neural network (SCMNN)</article-title>. <source>J. Pet. Sci. Eng.</source> <volume>73</volume> (<issue>3-4</issue>), <fpage>227</fpage>&#x2013;<lpage>232</lpage>. <pub-id pub-id-type="doi">10.1016/j.petrol.2010.07.003</pub-id>
</citation>
</ref>
<ref id="B14">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kozeny</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>1927</year>). <article-title>Uber kapillare leitung der Wasser in boden</article-title>. <source>R. Acad. Sci.</source> <volume>136</volume>, <fpage>271</fpage>&#x2013;<lpage>306</lpage>.</citation>
</ref>
<ref id="B15">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lundberg</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Lee</surname>
<given-names>S. I.</given-names>
</name>
</person-group> (<year>2016</year>). <article-title>An unexpected unity among methods for interpreting model predictions</article-title>. <comment>arXiv Preprint</comment>. <pub-id pub-id-type="doi">10.48550/arXiv.1611.07478</pub-id>
</citation>
</ref>
<ref id="B16">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lundberg</surname>
<given-names>S. M.</given-names>
</name>
<name>
<surname>Erion</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Degrave</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Prutkin</surname>
<given-names>J. M.</given-names>
</name>
<name>
<surname>Nair</surname>
<given-names>B.</given-names>
</name>
<etal/>
</person-group> (<year>2020</year>). <article-title>From local explanations to global understanding with explainable AI for trees</article-title>. <source>Nat. Mach. Intell.</source> <volume>2</volume> (<issue>1</issue>), <fpage>56</fpage>&#x2013;<lpage>67</lpage>. <pub-id pub-id-type="doi">10.1038/s42256-019-0138-9</pub-id>
</citation>
</ref>
<ref id="B17">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lundberg</surname>
<given-names>S. M.</given-names>
</name>
<name>
<surname>Erion</surname>
<given-names>G. G.</given-names>
</name>
<name>
<surname>Lee</surname>
<given-names>S. I.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>Consistent individualized feature attribution for tree ensembles</article-title>. <comment>arXiv Preprint</comment>. <pub-id pub-id-type="doi">10.48550/arXiv.1802.03888</pub-id>
</citation>
</ref>
<ref id="B18">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Lundberg</surname>
<given-names>S. M.</given-names>
</name>
<name>
<surname>Lee</surname>
<given-names>S. I.</given-names>
</name>
</person-group> (<year>2017</year>). &#x201c;<article-title>A unified approach to interpreting model predictions</article-title>,&#x201d; in <source>NIPS&#x27;17: proceedings of the 31st international conference on neural information processing systems (ACM).</source> <comment>Available at: <ext-link ext-link-type="uri" xlink:href="https://www.aap.org/sections/perinatal/NCE08/TheRole2.pdf">https://www.aap.org/sections/perinatal/NCE08/TheRole2.pdf</ext-link>
</comment>
</citation>
</ref>
<ref id="B19">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Matinkia</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Hashami</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Mehrad</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Hajsaeedi</surname>
<given-names>M. R.</given-names>
</name>
<name>
<surname>Velayati</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Prediction of permeability from well logs using a new hybrid machine learning algorithm</article-title>. <source>Petroleum</source> <volume>9</volume> (<issue>1</issue>), <fpage>108</fpage>&#x2013;<lpage>123</lpage>. <pub-id pub-id-type="doi">10.1016/j.petlm.2022.03.003</pub-id>
</citation>
</ref>
<ref id="B20">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Mauran</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Rigaud</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Coudevylle</surname>
<given-names>O.</given-names>
</name>
</person-group> (<year>2001</year>). <article-title>Application of the carman&#x2013;kozeny correlation to a high&#x2010;porosity and anisotropic consolidated medium: the compressed expanded natural graphite</article-title>. <source>Transp. Porous Med.</source> <volume>43</volume> (<issue>2</issue>), <fpage>355</fpage>&#x2013;<lpage>376</lpage>. <pub-id pub-id-type="doi">10.1023/a:1010735118136</pub-id>
</citation>
</ref>
<ref id="B21">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Mohaghegh</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Ameri</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>1995</year>). Artificial neural network as a valuable tool for petroleum engineers. <publisher-name>Paper SPE, 29220</publisher-name>, <fpage>1</fpage>&#x2013;<lpage>6</lpage>.</citation>
</ref>
<ref id="B22">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Nazari</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Kuzma</surname>
<given-names>H. A.</given-names>
</name>
<name>
<surname>Rector</surname>
<given-names>J. W.</given-names>
</name>
</person-group> (<year>2011</year>). &#x201c;<article-title>Predicting permeability from well log data and core measurements using support vector machines</article-title>,&#x201d; in <source>SEG technical program expanded abstracts 2011</source>. Editor <person-group person-group-type="editor">
<name>
<surname>Fomel</surname>
<given-names>S.</given-names>
</name>
</person-group> (<publisher-loc>Tulsa, Okla</publisher-loc>: <publisher-name>Society of Exploration Geophysicists</publisher-name>). <pub-id pub-id-type="doi">10.1190/1.3627601</pub-id>
</citation>
</ref>
<ref id="B23">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Nilsson</surname>
<given-names>N. J.</given-names>
</name>
</person-group> (<year>1965</year>). <source>Learning machines</source>. <publisher-loc>New York</publisher-loc>: <publisher-name>McGraw-Hill</publisher-name>.</citation>
</ref>
<ref id="B24">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Paterson</surname>
<given-names>M.</given-names>
</name>
</person-group> (<year>1983</year>). <article-title>The equivalent channel model for permeability and resistivity in fluid-saturated rock&#x2014;a re-appraisal</article-title>. <source>Mech. Mat.</source> <volume>2</volume> (<issue>4</issue>), <fpage>345</fpage>&#x2013;<lpage>352</lpage>. <pub-id pub-id-type="doi">10.1016/0167-6636(83)90025-X</pub-id>
</citation>
</ref>
<ref id="B25">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Rogers</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Kopaska-Merkel</surname>
<given-names>D. T.</given-names>
</name>
<name>
<surname>Fang</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>1995</year>). <article-title>Predicting permeability from porosity using artificial neural networks</article-title>. <source>AAPG Bull.</source> <volume>79</volume> (<issue>12</issue>), <fpage>1786</fpage>&#x2013;<lpage>1797</lpage>. <pub-id pub-id-type="doi">10.1306/7834DEFE-1721-11D7-8645000102C1865D</pub-id>
</citation>
</ref>
<ref id="B26">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Saemi</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Ahmadi</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Varjani</surname>
<given-names>A. Y.</given-names>
</name>
</person-group> (<year>2007</year>). <article-title>Design of neural networks using genetic algorithm for the permeability estimation of the reservoir</article-title>. <source>J. Pet. Sci. Eng.</source> <volume>59</volume> (<issue>1-2</issue>), <fpage>97</fpage>&#x2013;<lpage>105</lpage>. <pub-id pub-id-type="doi">10.1016/j.petrol.2007.03.007</pub-id>
</citation>
</ref>
<ref id="B27">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Saljooghi</surname>
<given-names>B. S.</given-names>
</name>
<name>
<surname>Hezarkhani</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2014</year>). <article-title>Comparison of WAVENET and ANN for predicting the porosity obtained from well log data</article-title>. <source>J. Pet. Sci. Eng.</source> <volume>123</volume>, <fpage>172</fpage>&#x2013;<lpage>182</lpage>. <pub-id pub-id-type="doi">10.1016/j.petrol.2014.08.025</pub-id>
</citation>
</ref>
<ref id="B28">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Sammut</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Webb</surname>
<given-names>G. I.</given-names>
</name>
</person-group> (<year>2011</year>). <source>Encyclopedia of machine learning</source>. <publisher-loc>Berlin</publisher-loc>: <publisher-name>Springer Science &#x26; Business Media</publisher-name>.</citation>
</ref>
<ref id="B29">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Timur</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>1968</year>). &#x201c;<article-title>An investigation of permeability, porosity, and residual water saturation relationships</article-title>,&#x201d; in <source>SPWLA 9th annual logging symposium</source> (<publisher-loc>New Orleans, Louisiana</publisher-loc>: <publisher-name>SPWLA</publisher-name>).</citation>
</ref>
<ref id="B30">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wolpert</surname>
<given-names>D. H.</given-names>
</name>
</person-group> (<year>1992</year>). <article-title>Stacked generalization</article-title>. <source>Neural Netw.</source> <volume>5</volume> (<issue>2</issue>), <fpage>241</fpage>&#x2013;<lpage>259</lpage>. <pub-id pub-id-type="doi">10.1016/S0893-6080(05)80023-1</pub-id>
</citation>
</ref>
<ref id="B31">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Wu</surname>
<given-names>T.</given-names>
</name>
</person-group> (<year>2004</year>). <article-title>Permeability prediction and drainage capillary pressure simulation in sandstone reservoirs</article-title>. <comment>Doctoral thesis</comment>. <publisher-loc>Texas</publisher-loc>: <publisher-name>Texas A&#x26;M University</publisher-name>.</citation>
</ref>
<ref id="B32">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhang</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Mohaghegh</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Lin</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Pei</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Pattern visualization and understanding of machine learning models for permeability prediction in tight sandstone reservoirs</article-title>. <source>J. Pet. Sci. Eng.</source> <volume>200</volume>, <fpage>108142</fpage>. <pub-id pub-id-type="doi">10.1016/j.petrol.2020.108142</pub-id>
</citation>
</ref>
<ref id="B33">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhang</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Fu</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Lv</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Shi</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Yu</surname>
<given-names>H.</given-names>
</name>
<etal/>
</person-group> (<year>2022a</year>). <article-title>FTCN: a reservoir parameter prediction method based on a fusional temporal convolutional network</article-title>. <source>Energies</source> <volume>15</volume> (<issue>15</issue>), <fpage>5680</fpage>. <pub-id pub-id-type="doi">10.3390/en15155680</pub-id>
</citation>
</ref>
<ref id="B34">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhang</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Gu</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Gao</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Zhou</surname>
<given-names>L.</given-names>
</name>
</person-group> (<year>2022b</year>). <article-title>Petrophysical regression regarding porosity, permeability, and water saturation driven by logging-based ensemble and transfer learnings: a case study of sandy-mud reservoirs</article-title>. <source>Geofluids</source> <volume>2022</volume>, <fpage>1</fpage>&#x2013;<lpage>31</lpage>. <pub-id pub-id-type="doi">10.1155/2022/9443955</pub-id>
</citation>
</ref>
<ref id="B35">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhong</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Carr</surname>
<given-names>T. R.</given-names>
</name>
<name>
<surname>Wu</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>G.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Application of a convolutional neural network in permeability prediction: a case study in the Jacksonburg-Stringtown oil field, West Virginia, USA</article-title>. <source>Geophysics</source> <volume>84</volume> (<issue>6</issue>), <fpage>B363</fpage>&#x2013;<lpage>B373</lpage>. <pub-id pub-id-type="doi">10.1190/geo2018-0588.1</pub-id>
</citation>
</ref>
<ref id="B36">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhu</surname>
<given-names>L. Q.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Wei</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>C. M.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>Permeability prediction of the tight sandstone reservoirs using hybrid intelligent algorithm and nuclear magnetic resonance logging data</article-title>. <source>Arab. J. Sci. Eng.</source> <volume>42</volume> (<issue>4</issue>), <fpage>1643</fpage>&#x2013;<lpage>1654</lpage>. <pub-id pub-id-type="doi">10.1007/s13369-016-2365-2</pub-id>
</citation>
</ref>
</ref-list>
</back>
</article>