<?xml version="1.0" encoding="UTF-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article article-type="research-article" dtd-version="2.3" xml:lang="EN" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Mater.</journal-id>
<journal-title>Frontiers in Materials</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Mater.</abbrev-journal-title>
<issn pub-type="epub">2296-8016</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="publisher-id">1128954</article-id>
<article-id pub-id-type="doi">10.3389/fmats.2023.1128954</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Materials</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Convolution, aggregation and attention based deep neural networks for accelerating simulations in mechanics</article-title>
<alt-title alt-title-type="left-running-head">Deshpande et&#xa0;al.</alt-title>
<alt-title alt-title-type="right-running-head">
<ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3389/fmats.2023.1128954">10.3389/fmats.2023.1128954</ext-link>
</alt-title>
</title-group>
<contrib-group>
<contrib contrib-type="author">
<name>
<surname>Deshpande</surname>
<given-names>Saurabh</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1919088/overview"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Sosa</surname>
<given-names>Ra&#xfa;l I.</given-names>
</name>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2164848/overview"/>
<xref ref-type="fn" rid="fn1">
<sup>&#x2020;</sup>
</xref>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Bordas</surname>
<given-names>St&#xe9;phane P. A.</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="corresp" rid="c001">&#x2a;</xref>
<uri xlink:href="https://loop.frontiersin.org/people/72792/overview"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Lengiewicz</surname>
<given-names>Jakub</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1280188/overview"/>
</contrib>
</contrib-group>
<aff id="aff1">
<sup>1</sup>
<institution>Department of Engineering</institution>, <institution>Faculty of Science</institution>, <institution>Technology and Medicine</institution>, <institution>University of Luxembourg</institution>, <addr-line>Belval</addr-line>, <country>Luxembourg</country>
</aff>
<aff id="aff2">
<sup>2</sup>
<institution>Department of Physics and Materials Science</institution>, <institution>Faculty of Science</institution>, <institution>Technology and Medicine</institution>, <institution>University of Luxembourg</institution>, <addr-line>Belval</addr-line>, <country>Luxembourg</country>
</aff>
<aff id="aff3">
<sup>3</sup>
<institution>Institute of Fundamental Technological Research</institution>, <institution>Polish Academy of Sciences</institution>, <addr-line>Warsaw</addr-line>, <addr-line>Masovian</addr-line>, <country>Poland</country>
</aff>
<author-notes>
<fn fn-type="edited-by">
<p>
<bold>Edited by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1835649/overview">Tae Yeon Kim</ext-link>, Khalifa University, United Arab Emirates</p>
</fn>
<fn fn-type="edited-by">
<p>
<bold>Reviewed by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/2001146/overview">Jesus Martinez-Frutos</ext-link>, Polytechnic University of Cartagena, Spain</p>
<p>
<ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/2134323/overview">Xingfei Wei</ext-link>, Johns Hopkins University, United States</p>
</fn>
<corresp id="c001">&#x2a;Correspondence: St&#xe9;phane P.A. Bordas, <email>stephane.bordas@alum.northwestern.edu</email>
</corresp>
<fn fn-type="equal" id="fn1">
<label>
<sup>&#x2020;</sup>
</label>
<p>ORCID: Ra&#xfa;l I. Sosa, <ext-link ext-link-type="uri" xlink:href="http://orcid.org/0000-0003-4116-2973">orcid.org/0000-0003-4116-2973</ext-link>
</p>
</fn>
<fn fn-type="other">
<p>This article was submitted to Computational Materials Science, a section of the journal Frontiers in Materials</p>
</fn>
</author-notes>
<pub-date pub-type="epub">
<day>24</day>
<month>03</month>
<year>2023</year>
</pub-date>
<pub-date pub-type="collection">
<year>2023</year>
</pub-date>
<volume>10</volume>
<elocation-id>1128954</elocation-id>
<history>
<date date-type="received">
<day>21</day>
<month>12</month>
<year>2022</year>
</date>
<date date-type="accepted">
<day>06</day>
<month>03</month>
<year>2023</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#xa9; 2023 Deshpande, Sosa, Bordas and Lengiewicz.</copyright-statement>
<copyright-year>2023</copyright-year>
<copyright-holder>Deshpande, Sosa, Bordas and Lengiewicz</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<abstract>
<p>Deep learning surrogate models are being increasingly used in accelerating scientific simulations as a replacement for costly conventional numerical techniques. However, their use remains a significant challenge when dealing with real-world complex examples. In this work, we demonstrate three types of neural network architectures for efficient learning of highly non-linear deformations of solid bodies. The first two architectures are based on the recently proposed CNN U-NET and MAgNET (graph U-NET) frameworks which have shown promising performance for learning on mesh-based data. The third architecture is Perceiver IO, a very recent architecture that belongs to the family of attention-based neural networks&#x2013;a class that has revolutionised diverse engineering fields and is still unexplored in computational mechanics. We study and compare the performance of all three networks on two benchmark examples, and show their capabilities to accurately predict the non-linear mechanical responses of soft bodies.</p>
</abstract>
<kwd-group>
<kwd>surrogate modeling</kwd>
<kwd>deep learning-artificial neural network</kwd>
<kwd>CNN U-NET</kwd>
<kwd>graph U-net</kwd>
<kwd>perceiver IO</kwd>
<kwd>finite element method</kwd>
</kwd-group>
<contract-num rid="cn001">764644 800150</contract-num>
<contract-num rid="cn002">QuaC C20/MS/14782078</contract-num>
<contract-num rid="cn003">811099</contract-num>
<contract-sponsor id="cn001">H2020 Marie Sk&#x142;odowska-Curie Actions<named-content content-type="fundref-id">10.13039/100010665</named-content>
</contract-sponsor>
<contract-sponsor id="cn002">Fonds National de La Recherche Luxembourg<named-content content-type="fundref-id">10.13039/501100001866</named-content>
</contract-sponsor>
<contract-sponsor id="cn003">Horizon 2020 Framework Programme<named-content content-type="fundref-id">10.13039/100010661</named-content>
</contract-sponsor>
</article-meta>
</front>
<body>
<sec id="s1">
<title>1 Introduction</title>
<p>The ability to make fast or real-time predictions of the response of physical systems is essential for a variety of engineering applications. Notable examples of this can be found in the field of robotics (<xref ref-type="bibr" rid="B52">Rus and Tolley, 2015</xref>; <xref ref-type="bibr" rid="B15">Choi&#xa0;et&#xa0;al., 2021</xref>) and medical simulations (<xref ref-type="bibr" rid="B17">Cotin&#xa0;et&#xa0;al., 1999</xref>; <xref ref-type="bibr" rid="B18">Courtecuisse&#xa0;et&#xa0;al., 2014</xref>; <xref ref-type="bibr" rid="B11">Bui&#xa0;et&#xa0;al., 2018</xref>; <xref ref-type="bibr" rid="B42">Mazier&#xa0;et&#xa0;al., 2021</xref>), which have the potential to advance personalized medicine and improve computer-assisted and robotic surgery (<xref ref-type="bibr" rid="B14">Chen&#xa0;et&#xa0;al., 2020</xref>; <xref ref-type="bibr" rid="B20">Dennler&#xa0;et&#xa0;al., 2021</xref>). In computational physics and chemistry, fast and accurate predictions are fundamental for studying complex systems, such as those arising in biology and materials science (<xref ref-type="bibr" rid="B27">Friesner, 2005</xref>), or in drug discovery (<xref ref-type="bibr" rid="B19">De&#xa0;Vivo&#xa0;et&#xa0;al., 2016</xref>). In many cases, the necessary accuracy of these predictions requires complex models that can be expressed through partial differential equations and solved numerically using methods such as the finite element method (FEM) at continuum scales or specialized <italic>ab initio</italic> approaches at the atomic or quantum scales. However, these high-fidelity computational models are often too slow for real-time or practical purposes, and therefore approximate or surrogate models must be developed to achieve the necessary speed-ups.</p>
<p>At the same time, the 21st century has seen an explosion of measurement data, much of which is available as public datasets in various scientific domains, including structural mechanics, material science, and meteorology (<xref ref-type="bibr" rid="B68">Zakutayev&#xa0;et&#xa0;al., 2018</xref>; <xref ref-type="bibr" rid="B25">Elouneg&#xa0;et&#xa0;al., 2022</xref>; <xref ref-type="bibr" rid="B28">Gholamalizadeh&#xa0;et&#xa0;al., 2022</xref>). The availability of this data, combined with the rapid growth in computational resources, has led to the increasing importance of machine learning (ML) techniques (<xref ref-type="bibr" rid="B12">Butler&#xa0;et&#xa0;al., 2018</xref>; <xref ref-type="bibr" rid="B7">Bock&#xa0;et&#xa0;al., 2019</xref>; <xref ref-type="bibr" rid="B54">Schleder&#xa0;et&#xa0;al., 2019</xref>) for solving forward and inverse engineering problems. This includes surrogate and data-driven approaches that aim to enable modeling (<xref ref-type="bibr" rid="B5">Barrios and Romero, 2019</xref>), accelerate computationally costly direct numerical simulations (<xref ref-type="bibr" rid="B51">Rupp&#xa0;et&#xa0;al., 2012</xref>; <xref ref-type="bibr" rid="B66">Wirtz&#xa0;et&#xa0;al., 2015</xref>; <xref ref-type="bibr" rid="B13">Capuano and Rimoli, 2019</xref>; <xref ref-type="bibr" rid="B65">Weerasuriya&#xa0;et&#xa0;al., 2021</xref>), and even discover new material laws (<xref ref-type="bibr" rid="B39">Liu&#xa0;et&#xa0;al., 2017</xref>; <xref ref-type="bibr" rid="B26">Flaschel&#xa0;et&#xa0;al., 2021</xref>). The increasing use of ML in engineering and other fields has also spurred the development of various methods and algorithms for improving the accuracy and efficiency of these techniques.</p>
<p>Within the class of machine learning methods for surrogate and data-driven modeling, deep learning (DL) approaches have seen great success due to their ability to efficiently extract complex relationships present in the underlying data. DL models have been successfully employed for a range of tasks in diverse fields such as computational physics and chemistry, material science, computational mechanics, computer vision, natural language processing, and many others (<xref ref-type="bibr" rid="B47">Oishi and Yagawa, 2017</xref>; <xref ref-type="bibr" rid="B56">Sch&#xfc;tt&#xa0;et&#xa0;al., 2017</xref>; <xref ref-type="bibr" rid="B32">Jha&#xa0;et&#xa0;al., 2018</xref>; <xref ref-type="bibr" rid="B64">Voulodimos&#xa0;et&#xa0;al., 2018</xref>; <xref ref-type="bibr" rid="B55">Schmidt&#xa0;et&#xa0;al., 2019</xref>; <xref ref-type="bibr" rid="B9">Brown&#xa0;et&#xa0;al., 2020</xref>; <xref ref-type="bibr" rid="B16">Choudhary&#xa0;et&#xa0;al., 2022</xref>). In computational chemistry, machine learning force fields (MLFFs), see (<xref ref-type="bibr" rid="B59">Unke&#xa0;et&#xa0;al., 2021</xref>), have seen great success in recent years for accelerating costly <italic>ab initio</italic> simulations. For instance, Deep Tensor Neural Network (DTNN) (<xref ref-type="bibr" rid="B57">Sch&#xfc;tt&#xa0;et&#xa0;al., 2017</xref>) and SchNet (<xref ref-type="bibr" rid="B56">Sch&#xfc;tt&#xa0;et&#xa0;al., 2017</xref>) models have been shown to accurately predict forces in a variety of molecules and could be used in applications such as protein folding and material design. Similarly, computational mechanics has witnessed an increasing use of DL surrogate models as a replacement for costly direct numerical simulations (<xref ref-type="bibr" rid="B2">Abueidda&#xa0;et&#xa0;al., 2021</xref>; <xref ref-type="bibr" rid="B45">Mianroodi&#xa0;et&#xa0;al., 2021</xref>). What is common to all the above-mentioned cases is that deep learning techniques rely on deep artificial neural networks (deep ANNs, or DNNs), which must be trained on a sufficiently large amount of data. While this training process is computationally costly, once trained, the predictions of DL models are extremely efficient.</p>
<p>Obtaining necessary amount of training data is often difficult when it originates from physical experiments. This can be due to multiple factors, such as high costs, risks and difficulties associated with the experiments, or data privacy clauses. There are two possible approaches to deal with the scarcity of experimental data. The first approach relies on enhancing the DL model with the information on underlying physics&#x2014;an approach popularly termed as Physics Informed Neural Networks (PINN) (<xref ref-type="bibr" rid="B43">McFall and Mahan, 2009</xref>; <xref ref-type="bibr" rid="B41">Mao&#xa0;et&#xa0;al., 2020</xref>; <xref ref-type="bibr" rid="B53">Samaniego&#xa0;et&#xa0;al., 2020</xref>; <xref ref-type="bibr" rid="B46">Odot&#xa0;et&#xa0;al., 2022</xref>). The second approach includes the underlying physics implicitly, through high-fidelity simulations done <italic>in silico</italic> to provide the necessary amount of synthetically generated data, which has shown to be useful in various applications (<xref ref-type="bibr" rid="B38">Le&#xa0;et&#xa0;al., 2017</xref>; <xref ref-type="bibr" rid="B3">Aydin&#xa0;et&#xa0;al., 2019</xref>; <xref ref-type="bibr" rid="B50">Pfeiffer&#xa0;et&#xa0;al., 2019</xref>; <xref ref-type="bibr" rid="B62">Vijayaraghavan&#xa0;et&#xa0;al., 2021</xref>; <xref ref-type="bibr" rid="B33">Kim&#xa0;et&#xa0;al., 2022</xref>). In this work, we will follow the latter approach and will focus on DL surrogate models that are trained on synthetically generated data from finite element simulations in non-linear elasticity.</p>
<p>One of the most important aspects that will be studied in this work is the architecture of deep neural networks. The majority of DL approaches that are present in the literature are based on fully connected networks, which can be inefficient and prone to overfitting when applied to high-dimensional inputs. If such large inputs are structured, they fall under the umbrella of geometric deep learning (GDL) (<xref ref-type="bibr" rid="B8">Bronstein&#xa0;et&#xa0;al., 2021</xref>), a concept that has gained increasing interest in recent years. In this work, we will compare three architectures that can efficiently handle high-dimensional structured inputs: convolutional neural networks (CNNs), graph neural networks (GNNs), and attention-based networks.</p>
<p>Convolutional neural networks (CNNs) are known to outperform traditional fully-connected ANNs, and this has been demonstrated in various domains, including physics-based simulations (<xref ref-type="bibr" rid="B29">Guo&#xa0;et&#xa0;al., 2016</xref>; <xref ref-type="bibr" rid="B22">Deshpande&#xa0;et&#xa0;al., 2022a</xref>; <xref ref-type="bibr" rid="B37">Krokos&#xa0;et&#xa0;al., 2022b</xref>; <xref ref-type="bibr" rid="B24">El&#xa0;Haber&#xa0;et&#xa0;al., 2022</xref>). CNNs work on the principle of parameter sharing and local convolution operations, which enables efficient training on large inputs. Their disadvantage is that the inputs/outputs of CNNs are restricted to grid inputs, such as images, videos, or structured FE meshes. However, CNNs have found their successors, the graph neural networks (GNNs), that can work with any structure of inputs/outputs.</p>
<p>Graph-based approaches leverage the topological information of the input to perform local operations in the respective neighborhood only, and can learn efficiently on generally structured data. Recently, GDL methods have shown promising performance for their applications as well in the field of mechanics, (<xref ref-type="bibr" rid="B6">Battaglia&#xa0;et&#xa0;al., 2018</xref>; <xref ref-type="bibr" rid="B63">Vlassis&#xa0;et&#xa0;al., 2020</xref>; <xref ref-type="bibr" rid="B49">Pfaff&#xa0;et&#xa0;al., 2021</xref>; <xref ref-type="bibr" rid="B36">Krokos&#xa0;et&#xa0;al., 2022a</xref>; <xref ref-type="bibr" rid="B58">Str&#xf6;nisch&#xa0;et&#xa0;al., 2022</xref>). More recently (<xref ref-type="bibr" rid="B21">Deshpande&#xa0;et&#xa0;al., 2022b</xref>), proposed MAgNET, a novel graph U-Net framework for efficiently learning on mesh-based data. In this work we utilise it to accurately predict non-linear deformations of solids.</p>
<p>Attention-based approaches, similar to human cognitive attention, work by allowing the DL model to focus on certain parts of the input data that are relevant to the task at hand. This is done through a fully trainable process that, without the need to introduce topological information or enforce structural restrictions, allows the neural network to extract dependencies from throughout the whole input domain. This type of approach has led to significant strides in a wide range of areas, starting from computer vision (<xref ref-type="bibr" rid="B67">Xu&#xa0;et&#xa0;al., 2015</xref>) to natural language processing (<xref ref-type="bibr" rid="B23">Devlin&#xa0;et&#xa0;al., 2018</xref>; <xref ref-type="bibr" rid="B4">Baevski&#xa0;et&#xa0;al., 2020</xref>), as well as becoming the basic building block of the Transformer architecture (<xref ref-type="bibr" rid="B61">Vaswani&#xa0;et&#xa0;al., 2017</xref>). Recently the Perceiver IO (<xref ref-type="bibr" rid="B31">Jaegle&#xa0;et&#xa0;al., 2022</xref>), a new type of architecture that builds upon Transformers, has been proposed as a general-purpose model that can handle data from arbitrary settings. Since Perceiver IO has been shown to achieve several state-of-the-art results without the need for problem-specific architecture engineering, we will compare its performance on non-linear deformation prediction of solids based on mesh data against the previously discussed models.</p>
<p>To summarise, in this work we will compare three DNN architectures: two architectures presented in our earlier works, i.e., CNN U-Net framework (<xref ref-type="bibr" rid="B22">Deshpande&#xa0;et&#xa0;al., 2022a</xref>), and MAgNET framework (<xref ref-type="bibr" rid="B21">Deshpande&#xa0;et&#xa0;al., 2022b</xref>), as well as the attention-based architecture, Perceiver IO (<xref ref-type="bibr" rid="B31">Jaegle&#xa0;et&#xa0;al., 2022</xref>), which has not been explored for its applications in mechanics yet. We show the capabilities of three frameworks by learning on non-linear FEM datasets and by cross-comparing their performance. In <xref ref-type="sec" rid="s2">Section&#xa0;2</xref>, we will introduce the three DNN architectures, in <xref ref-type="sec" rid="s3">Section&#xa0;3</xref>, we will study their performance, and in <xref ref-type="sec" rid="s4">Section&#xa0;4</xref> we will summarize the results and discuss future directions.</p>
</sec>
<sec sec-type="materials|methods" id="s2">
<title>2 Materials and methods</title>
<p>As previously mentioned in the introduction, in this paper we propose three types of deep neural network (DNN) frameworks that can be used as surrogate models to replace computationally expensive non-linear FEM solvers. The proposed DNN frameworks are trained on force-displacement FEM datasets that are given in the mesh format. Once trained, these surrogate DNN models are able to quickly and accurately simulate the mechanical responses of bodies subjected to external forces. The outline of study pursued in this paper is shown in <xref ref-type="fig" rid="F1">Figure&#xa0;1</xref>.</p>
<fig id="F1" position="float">
<label>FIGURE 1</label>
<caption>
<p>Outline of the neural network surrogate frameworks for predicting body deformations. (Left) Training datasets for structured and arbitrary mesh cases are generated by using a non-linear FEM solver. (Middle) Proposed neural network frameworks are trained on these datasets. For structured mesh case, all NN frameworks are used while for arbitrary unstructured meshes only MAgNET and Perceiver IO networks are used. (Right) Trained networks are then used as surrogate models to predict the deformation of bodies under unseen forces.</p>
</caption>
<graphic xlink:href="fmats-10-1128954-g001.tif"/>
</fig>
<p>Any input mesh can be categorised into either a structured mesh or an arbitrary unstructured mesh. In this work we introduce three types of DNN network architectures. The CNN U-Net network can only be (straightforwardly) used for structured meshes, while the MAgNET and Perceiver IO networks are more general and are capable of handling arbitrary mesh inputs. All these frameworks are discussed in detail in the following subsections.</p>
<sec id="s2-1">
<title>2.1 DNN frameworks for predicting mechanical deformations</title>
<p>Below we introduce three different types of neural network architectures which can efficiently predict non-linear deformations of bodies subjected to external traction and body forces. All the proposed DNN frameworks directly operate on the finite element mesh data thereby making it very convenient to be used as surrogate models in place of conventional FEM solver. The first two i.e., CNN U-Net and MAgNET belong to the family of U-Net architecture while Perceiver IO is based on Transformer-type attention, see <xref ref-type="table" rid="T1">Table&#xa0;1</xref>.</p>
<table-wrap id="T1" position="float">
<label>TABLE 1</label>
<caption>
<p>Properties of deep neural network architectures studied in this work.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="center">Framework</th>
<th align="center">Type</th>
<th align="center">Supported mesh</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="center">CNN</td>
<td align="center">U-Net (convolution operation)</td>
<td align="center">structured</td>
</tr>
<tr>
<td align="center">MAgNET</td>
<td align="center">Graph U-NET (MAg operation)</td>
<td align="center">arbitrary</td>
</tr>
<tr>
<td align="center">Perceiver IO</td>
<td align="center">Transformer (attention mechanism)</td>
<td align="center">arbitrary</td>
</tr>
</tbody>
</table>
</table-wrap>
<sec id="s2-1-1">
<title>2.1.1 CNN U-Net</title>
<p>CNNs were originally proposed for performing classification and regression tasks on image, video like data but lately are even being used for generic inputs such as mesh data which is crucial to many scientific applications. In particular, U-Net like architectures have shown great potential in learning on large-scale inputs and lately have been successfully used for simulating mechanical responses of materials as well (<xref ref-type="bibr" rid="B45">Mianroodi&#xa0;et&#xa0;al., 2021</xref>; <xref ref-type="bibr" rid="B22">Deshpande&#xa0;et&#xa0;al., 2022a</xref>). The name U-Net comes from the particular U-shaped architecture which involves a series of convolutional and pooling operations. Convolutional layers are responsible for non-linear transformations whereas pooling enables learning through low-fidelity representation thus making the network capable of learning on high-dimensional inputs. Experiments presented in this work are carried out by using the CNN U-Net framework proposed by (<xref ref-type="bibr" rid="B22">Deshpande&#xa0;et&#xa0;al., 2022a</xref>), see <xref ref-type="fig" rid="F2">Figure 2</xref>.</p>
<fig id="F2" position="float">
<label>FIGURE 2</label>
<caption>
<p>Schematic of CNN architecture used for generic structured 2D mesh inputs.</p>
</caption>
<graphic xlink:href="fmats-10-1128954-g002.tif"/>
</fig>
<p>One major limitation of the CNN is that it cannot straightforwardly accommodate unstructured mesh inputs. To overcome this issue, the simplest approach embeds a structured grid on unstructured meshes with a naive mapping between unstructured and grid node values (<xref ref-type="bibr" rid="B44">Mendizabal&#xa0;et&#xa0;al., 2019</xref>). While more sophisticated approaches are proposed to make unstructured meshes compatible to be used with CNN framework (<xref ref-type="bibr" rid="B10">Brunet&#xa0;et&#xa0;al., 2019</xref>). However, they perform poorly on complicated geometries and come with an associated preprocessing cost; they are not considered in the scope of this work.</p>
</sec>
<sec id="s2-1-2">
<title>2.1.2 MAgNET</title>
<p>In an attempt to generalise CNN to arbitrary unstructured meshes, very recently (<xref ref-type="bibr" rid="B21">Deshpande&#xa0;et&#xa0;al., 2022b</xref>) proposed the MAgNET framework, see <xref ref-type="fig" rid="F3">Figure 3</xref>. MAgNET architecture belongs to the family of graph U-Net architectures and it is proposed for efficient learning on mesh structured data. MAgNET directly accepts arbitrary mesh inputs (such as forces/stresses/displacements of nodes in the mesh) values thus making it very convenient to be used with existing numerical solvers.</p>
<fig id="F3" position="float">
<label>FIGURE 3</label>
<caption>
<p>Schematic of the MAgNET architecture used in this work. It takes external forces on arbitrary mesh as an input to gives mesh displacements as output.</p>
</caption>
<graphic xlink:href="fmats-10-1128954-g003.tif"/>
</fig>
<p>MAgNET relies on the so-called MAg layer (Multichannel Aggregation layer) which is capable of learning non-linear transformations between input and output data existing in the mesh format. MAg extends the concept of local operations in convolution layers to arbitrary mesh inputs by performing aggregation of nodal feature values in the respective neighborhood nodes only. It leverages the topology of inputs and performs learnable local aggregations with heterogeneous window sizes as opposed to the fixed-size window in the case of CNN. While its graph pooling/unpooling layers enable efficient learning on large-dimension inputs through reduced graph representation.</p>
</sec>
<sec id="s2-1-3">
<title>2.1.3 Perceiver IO</title>
<p>The Perceiver IO architecture (<xref ref-type="bibr" rid="B31">Jaegle&#xa0;et&#xa0;al., 2022</xref>), was developed with the goal of achieving a DL scheme that can easily integrate and transform arbitrary information for arbitrary tasks. This architecture employs an attention encoder that maps inputs from a wide range of modalities to a fixed-size latent space using cross-attention, this latent space is then further processed using self-attention as an usual Transformer and decoded into the output domain <italic>via</italic> cross-attention, see <xref ref-type="fig" rid="F4">Figure 4</xref>. This process allows the network to scale to large and multi-modal data since it decouples the bulk of the network&#x2019;s processing from the size and modality-specific details of the input. In this work, we will leverage this property and use Perceiver IO to learn non-linear deformations on unstructured meshes without adding any information or restrictions about how to treat the underlying data structure. During training, Perceiver IO automatically learns the important dependencies that exist in the input domain, composed of arbitrarily unstructured mesh data, and transforms them into the corresponding output which consists of displacement data.</p>
<fig id="F4" position="float">
<label>FIGURE 4</label>
<caption>
<p>Schematic of the Perceiver IO architecture, (<xref ref-type="bibr" rid="B31">Jaegle&#xa0;et&#xa0;al., 2022</xref>), used for external forces on arbitrary mesh as inputs and mesh displacement as outputs.</p>
</caption>
<graphic xlink:href="fmats-10-1128954-g004.tif"/>
</fig>
</sec>
</sec>
<sec id="s2-2">
<title>2.2 Input/output and training of DNN surrogate models</title>
<p>As motivated in <xref ref-type="sec" rid="s2">Section&#xa0;2</xref>, the proposed DNN frameworks are trained on the force-displacement datasets. Let us denote the neural network in consideration as <italic>h</italic>, it is parameterised by trainable parameters, <bold>
<italic>&#x3b8;</italic>
</bold>. In all the cases, <italic>h</italic> accepts external forces, <bold>f</bold>, on all degrees of freedom (dofs) of mesh as the input. And as an output it predicts displacement vector, <bold>u</bold>, of all dofs (same size as the input), i.e., <italic>h</italic>&#xa0;:&#xa0;<bold>
<italic>f</italic>
</bold> &#x2192; <bold>
<italic>u</italic>
</bold>.</p>
<p>In the case of CNN and MagNET, forces associated with X, Y, Z degrees of freedom are kept in different channels. For instance, in the case of CNN U-Net, X, Y directional forces on a 2D quad mesh with <italic>n</italic>
<sub>
<italic>x</italic>
</sub> &#xd7; <italic>n</italic>
<sub>
<italic>y</italic>
</sub> nodes are fed to the network as a 2 &#xd7; <italic>n</italic>
<sub>
<italic>x</italic>
</sub> &#xd7; <italic>n</italic>
<sub>
<italic>y</italic>
</sub> tensor. For MAgNET, they are provided as a single dimensional tensor of shape 2 &#x22c5; <italic>n</italic>
<sub>
<italic>x</italic>
</sub> &#x22c5; <italic>n</italic>
<sub>
<italic>y</italic>
</sub>, with X and Y directional forces concatenated together (each representing a different channel). On the other hand, Perceiver IO inputs are first kept as a single array of size 2 &#x22c5; <italic>n</italic>
<sub>
<italic>x</italic>
</sub> &#x22c5; <italic>n</italic>
<sub>
<italic>y</italic>
</sub> to which we apply 1 &#xd7; 1 convolution kernels to project it to a tensor of shape 2 &#x22c5; <italic>n</italic>
<sub>
<italic>x</italic>
</sub> &#x22c5; <italic>n</italic>
<sub>
<italic>y</italic>
</sub> &#xd7; 256, adding 256 channels. Finally, to this channel dimension we concatenate trainable 1D positional embeddings, thus leaving us with a tensor of shape 2 &#x22c5; <italic>n</italic>
<sub>
<italic>x</italic>
</sub> &#x22c5; <italic>n</italic>
<sub>
<italic>y</italic>
</sub> &#xd7; 512 that we will use as an input for the network.</p>
<p>Now, for a given training dataset <inline-formula id="inf1">
<mml:math id="m1">
<mml:mfenced open="{" close="}">
<mml:mrow>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi mathvariant="bold">f</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msub>
<mml:mo>,</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi mathvariant="bold">u</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfenced>
<mml:mo>,</mml:mo>
<mml:mo>&#x2026;</mml:mo>
<mml:mo>,</mml:mo>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi mathvariant="bold">f</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>N</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>,</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi mathvariant="bold">u</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>N</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mfenced>
</mml:math>
</inline-formula>, <italic>h</italic> is trained by minimizing mean squared error between true and predicted values as to get optimised parameters <bold>
<italic>&#x3b8;</italic>
</bold>&#x2a; as:<disp-formula id="e1">
<mml:math id="m2">
<mml:msup>
<mml:mrow>
<mml:mi mathvariant="bold-italic">&#x3b8;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mo>&#x2a;</mml:mo>
</mml:mrow>
</mml:msup>
<mml:mo>&#x3d;</mml:mo>
<mml:mtext>arg</mml:mtext>
<mml:munder>
<mml:mrow>
<mml:mtext>min</mml:mtext>
</mml:mrow>
<mml:mrow>
<mml:mi mathvariant="bold-italic">&#x3b8;</mml:mi>
</mml:mrow>
</mml:munder>
<mml:mspace width="0.3333em" class="nbsp"/>
<mml:mfrac>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>N</mml:mi>
</mml:mrow>
</mml:mfrac>
<mml:munderover accentunder="false" accent="true">
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>N</mml:mi>
</mml:mrow>
</mml:munderover>
<mml:msubsup>
<mml:mrow>
<mml:mo>&#x2225;</mml:mo>
<mml:mi>h</mml:mi>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi mathvariant="bold">f</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>,</mml:mo>
<mml:mi mathvariant="bold-italic">&#x3b8;</mml:mi>
</mml:mrow>
</mml:mfenced>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi mathvariant="bold">u</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2225;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:msubsup>
</mml:math>
<label>(1)</label>
</disp-formula>
</p>
<p>Performance of <italic>h</italic> (i.e., the neural network in consideration) over respective test dataset <inline-formula id="inf2">
<mml:math id="m3">
<mml:mfenced open="{" close="}">
<mml:mrow>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi mathvariant="bold">f</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msub>
<mml:mo>,</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi mathvariant="bold">u</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfenced>
<mml:mo>,</mml:mo>
<mml:mo>&#x2026;</mml:mo>
<mml:mo>,</mml:mo>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi mathvariant="bold">f</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>M</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>,</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi mathvariant="bold">u</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>M</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mfenced>
</mml:math>
</inline-formula> is measured in terms of mean absolute error which is computed for each example (<italic>e</italic>
<sub>
<italic>m</italic>
</sub>) and for the entire test set <inline-formula id="inf3">
<mml:math id="m4">
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mover accent="true">
<mml:mrow>
<mml:mi>e</mml:mi>
</mml:mrow>
<mml:mo>&#x304;</mml:mo>
</mml:mover>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:math>
</inline-formula> as follows:<disp-formula id="e2">
<mml:math id="m5">
<mml:mtable class="aligned">
<mml:mtr>
<mml:mtd columnalign="right">
<mml:msub>
<mml:mrow>
<mml:mi>e</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mtd>
<mml:mtd columnalign="left">
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi mathvariant="script">F</mml:mi>
</mml:mrow>
</mml:mfrac>
<mml:munderover accentunder="false" accent="true">
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi mathvariant="script">F</mml:mi>
</mml:mrow>
</mml:munderover>
<mml:mo stretchy="false">&#x7c;</mml:mo>
<mml:mi>h</mml:mi>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi mathvariant="bold">f</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msup>
<mml:mo>&#x2212;</mml:mo>
<mml:msubsup>
<mml:mrow>
<mml:mi mathvariant="bold">u</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msubsup>
<mml:mo stretchy="false">&#x7c;</mml:mo>
<mml:mo>,</mml:mo>
</mml:mtd>
</mml:mtr>
<mml:mtr>
<mml:mtd columnalign="right">
<mml:mrow>
<mml:mover accent="true">
<mml:mrow>
<mml:mi>e</mml:mi>
</mml:mrow>
<mml:mo>&#x304;</mml:mo>
</mml:mover>
</mml:mrow>
</mml:mtd>
<mml:mtd columnalign="left">
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>M</mml:mi>
</mml:mrow>
</mml:mfrac>
<mml:munderover accentunder="false" accent="true">
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>M</mml:mi>
</mml:mrow>
</mml:munderover>
<mml:msub>
<mml:mrow>
<mml:mi>e</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>,</mml:mo>
<mml:mspace width="2em"/>
</mml:mtd>
</mml:mtr>
<mml:mtr>
<mml:mtd columnalign="right">
<mml:mi>&#x3c3;</mml:mi>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:mi>e</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mtd>
<mml:mtd columnalign="left">
<mml:mo>&#x3d;</mml:mo>
<mml:msqrt>
<mml:mrow>
<mml:mfrac>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>M</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:mfrac>
<mml:munderover accentunder="false" accent="true">
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>M</mml:mi>
</mml:mrow>
</mml:munderover>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>e</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:mrow>
<mml:mover accent="true">
<mml:mrow>
<mml:mi>e</mml:mi>
</mml:mrow>
<mml:mo>&#x304;</mml:mo>
</mml:mover>
</mml:mrow>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:msup>
</mml:mrow>
</mml:msqrt>
<mml:mo>,</mml:mo>
</mml:mtd>
</mml:mtr>
</mml:mtable>
</mml:math>
<label>(2)</label>
</disp-formula>where <inline-formula id="inf4">
<mml:math id="m6">
<mml:mi mathvariant="script">F</mml:mi>
</mml:math>
</inline-formula> stands for the number of dofs of the mesh. The maximum error over the entire test dataset is defined as<disp-formula id="e3">
<mml:math id="m7">
<mml:msub>
<mml:mrow>
<mml:mi>e</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mtext>max</mml:mtext>
</mml:mrow>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:munder>
<mml:mrow>
<mml:mtext>max</mml:mtext>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:munder>
<mml:mo stretchy="false">&#x7c;</mml:mo>
<mml:mi>h</mml:mi>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi mathvariant="bold">f</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msup>
<mml:mo>&#x2212;</mml:mo>
<mml:msubsup>
<mml:mrow>
<mml:mi mathvariant="bold">u</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msubsup>
<mml:mo stretchy="false">&#x7c;</mml:mo>
<mml:mo>.</mml:mo>
</mml:math>
<label>(3)</label>
</disp-formula>
</p>
</sec>
</sec>
<sec sec-type="results" id="s3">
<title>3 Results</title>
<p>We validate proposed neural network frameworks on two examples, representing a 2D and a 3D problem respectively. For the 2D example, a structured mesh is considered so that all frameworks including CNN can be applied to it. Whereas for the 3D example an arbitrary unstructured mesh is considered.</p>
<sec id="s3-1">
<title>3.1 Generation of hyperelastic FEM training data</title>
<p>As motivated in the methodology section, training datasets of non-linear displacement solutions are generated by applying random traction and body forces on the given discretisation. The number of cases generated randomly (dataset size) has been chosen large enough to generalise well to unseen arbitrary forces. The dataset is split into the training (95%) and testing (5%) part, and the pairs of input force and output displacement solutions from the training dataset are then fed to train different types of neural networks. The proposed neural network frameworks are validated on two examples, both following Neo-Hookean hyperelastic law. To avoid the divergence of the non-linear FEM solver, both traction and body forces are applied in incremental load steps. All the computations are performed using AceFEM framework (<xref ref-type="bibr" rid="B35">Korelc, 2002</xref>).</p>
<p>For the 2D case, a rectangular domain made of soft material and discretized by 8 &#xd7; 32 mesh with 217 quad elements is considered. It is constrained at 4 corner nodes as shown in <xref ref-type="fig" rid="F5">Figure&#xa0;5A</xref>. It is subjected to traction forces of random magnitude, location, and direction in the region prescribed by the pink line. Body forces are ignored in this case. <xref ref-type="table" rid="T2">Table&#xa0;2</xref> provides detailed information about datasets including the material properties and external force ranges used for the generation of 2D and 3D datasets. For the 2D case, a lower range of <italic>Y</italic>-direction force density is chosen since it does not contribute much to generating large deformation solutions.</p>
<fig id="F5" position="float">
<label>FIGURE 5</label>
<caption>
<p>Schematics of dataset generation for <bold>(A)</bold> 2D example subjected to external traction forces. <bold>(B)</bold> 3D example subjected to external body forces. External tractions and body forces are indicated with pink arrows.</p>
</caption>
<graphic xlink:href="fmats-10-1128954-g005.tif"/>
</fig>
<table-wrap id="T2" position="float">
<label>TABLE 2</label>
<caption>
<p>Desciption of FEM datasets.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="left">Case</th>
<th align="center">Material properties (<italic>E</italic> [Pa], <italic>&#x3bd;</italic>)</th>
<th align="center">N. of FEM DOFs <inline-formula id="inf5">
<mml:math id="m8">
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mi mathvariant="script">F</mml:mi>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:math>
</inline-formula>
</th>
<th align="left">External traction/body force density range</th>
<th align="center">Dataset size (train &#x2b; test)</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="left">2D domain</td>
<td align="center">100, 0.3</td>
<td align="center">512</td>
<td align="left">
<italic>f</italic>
<sub>
<italic>x</italic>
</sub> &#x3d; &#x2212;24 to 24&#xa0;N/m, <italic>f</italic>
<sub>
<italic>y</italic>
</sub> &#x3d; &#x2212;8 to 8&#xa0;N/m</td>
<td align="center">7,124 &#x2b; 372</td>
</tr>
<tr>
<td align="left">3D elephant</td>
<td align="center">3 &#xd7; 10<sup>6</sup>, 0.4</td>
<td align="center">5,835</td>
<td align="left">
<italic>b</italic>
<sub>
<italic>x</italic>
</sub>, <italic>b</italic>
<sub>
<italic>z</italic>
</sub> &#x3d; &#x2212;0.35 to 0.35&#xa0;N/kg, <italic>b</italic>
<sub>
<italic>y</italic>
</sub> &#x3d; 0&#xa0;N/kg</td>
<td align="center">7,600 &#x2b; 400</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>In the 3D case, a continuum toy elephant model is discretised with 6,627 tetrahedron elements. It is subjected to fixed boundary conditions by constraining nodes on the bottom region of the legs. Body forces in random magnitude and direction are applied in the transverse directions (see <xref ref-type="fig" rid="F5">Figure&#xa0;5B</xref>) to generate datasets of force-displacement pairs. External tractions are ignored in this case.</p>
</sec>
<sec id="s3-2">
<title>3.2 Implementation details</title>
<p>As introduced in the methodology section, CNN and MAgNET frameworks belong to the family of U-Net architectures, while Perceiver IO leverages Transformer-style attention. It has to be noted that all the frameworks are robust, and do not need fine hyperparameter tuning.</p>
<p>2D case: For the CNN U-Net, a 4-level architecture with 3 max-pooling/upsampling operations is used. At each level, two convolution layers with 3 &#xd7; 3 filters are applied with 64, 128, 256, and 256 channels at respective levels. In the case of MAgNET, a 5-level graph U-Net architecture with 4 graph pooling/unpooling operations is used. At each level 2 MAg layers (with <italic>A</italic>
<sup>2</sup> adjacency, refer (<xref ref-type="bibr" rid="B21">Deshpande&#xa0;et&#xa0;al., 2022b</xref>)) are applied with 8, 16, 16, 32, and 32 channels at respective levels. For both 2D and 3D case MAgNET architectures, the seed of the first graph pooling operation is chosen by grid search, while seeds for other pooling layers are kept constant. This is done to have the maximum possible coarsened representation of the lowest-level graph. This ensures propagating boundary condition information with a minimum number of MAg operations at the lowest level. CNN U-Net is trained with a batch size of 16 for 32,000 epochs and MAgNET is trained with a batch size of 4 for 10,000 epochs.</p>
<p>In the case of Perceiver IO we defined 512 inputs (standing for dofs of the example) with a total embedding size of 512 following the procedure detailed in <xref ref-type="sec" rid="s2-2">Section&#xa0;2.2</xref>. We also selected a total of 128 latent arrays of dimension 210 for performing cross-attention in the encoder, self attention in latent space and inputs for the decoder. For the decoder&#x2019;s output query array we used an index dimension of 512, which defines the size of the outputs, and a channel dimension of 210 equal to the dimension size of the latents. We used a total of 3 blocks for the latent array processing, with 2 self-attention layers per block and 2 self-attention heads per layer. Both the encoder and decoder worked with 2 cross-attention heads each. The selection of these hyper-parameters was determined in a coarse exploratory fashion, with the goal of reducing the number of network parameters while maintaining its performance. Perceiver IO is trained with a batch size of 16 for 264,140 epochs.</p>
<p>3D case: For the MAgNET, 7-level graph U-Net architecture is used with 6 graph pooling/unpooling operations. Again at each level, 2 MAg layers (with <italic>A</italic>
<sup>2</sup> adjacency) are applied with 6, 6, 6, 12, 12, 24, 24 channels at respective levels. The complex topology of this particular mesh demands more number graph pooling layers, this ensures propagating boundary condition information with a minimum number of MAg operations at the lowest level. On the other hand, the only change with respect to the 2D case for Perceiver IO is an increase of the input and output dimension from 512 to 5,835 in both the encoder and decoder. MAgNET architecture for this case is trained for 1200 epochs with a batch size of 4, whereas Perceiver IO is trained for 32,580 epochs with a batch size of 16.</p>
<p>CNN U-Net and MAgNET networks are trained using the Adam optimizer (<xref ref-type="bibr" rid="B34">Kingma and Ba, 2014</xref>), whereas Perceiver IO is trained using AdamW optimizer (<xref ref-type="bibr" rid="B40">Loshchilov and Hutter, 2017</xref>) as implemented in the original paper. CNN and MAgNET are implemented using TensorFlow (<xref ref-type="bibr" rid="B1">Abadi&#xa0;et&#xa0;al., 2015</xref>), while Perceiver IO is implemented using PyTorch (<xref ref-type="bibr" rid="B48">Paszke&#xa0;et&#xa0;al., 2019</xref>). All the implementations in this work are performed using HPC facilities of the University of Luxembourg (<xref ref-type="bibr" rid="B60">Varrette&#xa0;et&#xa0;al., 2014</xref>).</p>
</sec>
<sec id="s3-3">
<title>3.3 Performance on unseen examples</title>
<p>Proposed DNN frameworks are trained and tested on the datasets generated as illustrated in <xref ref-type="sec" rid="s3-1">Section&#xa0;3.1</xref>. The maximum nodal displacement for the 2D case is 0.35&#xa0;m and for the 3D case it is 140.04&#xa0;m, i.e., we compute displacements of all the nodes for every single example, and then choose particular examples for which the maximum nodal displacement is observed. <xref ref-type="table" rid="T3">Table&#xa0;3</xref> summarises the performance of neural networks on the two test datasets. It shows that all three networks are capable of predicting mechanical deformation responses with a very low error.</p>
<table-wrap id="T3" position="float">
<label>TABLE 3</label>
<caption>
<p>Error metrics over the test set using the proposed NN frameworks. <italic>M</italic> stands for the number of test examples, and <inline-formula id="inf6">
<mml:math id="m9">
<mml:mrow>
<mml:mover accent="true">
<mml:mrow>
<mml:mi>e</mml:mi>
</mml:mrow>
<mml:mo>&#x304;</mml:mo>
</mml:mover>
</mml:mrow>
</mml:math>
</inline-formula>, <italic>&#x3c3;</italic>(<italic>e</italic>), <italic>e</italic>
<sub>max</sub> are error metrics defined in <xref ref-type="sec" rid="s2-2">Section&#xa0;2.2</xref>.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="left">Example</th>
<th align="center">Framework</th>
<th align="center">
<italic>M</italic>
</th>
<th align="center">
<inline-formula id="inf7">
<mml:math id="m10">
<mml:mrow>
<mml:mover accent="true">
<mml:mrow>
<mml:mi>e</mml:mi>
</mml:mrow>
<mml:mo>&#x304;</mml:mo>
</mml:mover>
</mml:mrow>
</mml:math>
</inline-formula> [m]</th>
<th align="center">
<italic>&#x3c3;</italic>(<italic>e</italic>)&#xa0;[m]</th>
<th align="center">
<italic>e</italic> (E)<sub>max</sub>&#xa0;[m]</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="left">2D</td>
<td align="center">CNN</td>
<td align="center">372</td>
<td align="center">0.06 E-3</td>
<td align="center">0.2&#x2013;4</td>
<td align="center">0.001</td>
</tr>
<tr>
<td align="left"/>
<td align="center">MAgNET</td>
<td align="left"/>
<td align="center">0.17 E-3</td>
<td align="center">0.7&#x2013;4</td>
<td align="center">0.021</td>
</tr>
<tr>
<td align="left"/>
<td align="center">Perceiver IO</td>
<td align="left"/>
<td align="center">0.02 E-3</td>
<td align="center">0.1&#x2013;4</td>
<td align="center">0.001</td>
</tr>
<tr>
<td align="left">3D</td>
<td align="center">MAgNET</td>
<td align="center">400</td>
<td align="center">8.92 E-3</td>
<td align="center">1.9&#x2013;3</td>
<td align="center">0.307</td>
</tr>
<tr>
<td align="left"/>
<td align="center">Perceiver IO</td>
<td align="left"/>
<td align="center">2.60 E-3</td>
<td align="center">1.1&#x2013;3</td>
<td align="center">0.098</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>To compare, we observed that CNN U-NET and Perceiver IO gave lower error metrics than MAgNET for the 2D case, which has relatively low dimensional input. As the size and complexity of the mesh increased in the 3D case, both MAgNET and Perceiver IO performed well, with Perceiver IO giving slightly better error metrics.</p>
</sec>
<sec id="s3-4">
<title>3.4 Training and inference of DNN frameworks</title>
<p>First, we compare the training convergence of proposed DNN frameworks by comparing the mean square loss plots for both 2D and 3D cases. <xref ref-type="fig" rid="F6">Figure&#xa0;6A</xref> shows that for the relatively smaller dimension inputs as in the 2D case, both CNN and Perceiver IO observed to learn more efficiently when compared to MAgNET. <xref ref-type="fig" rid="F6">Figure&#xa0;6B</xref> shows that both MAgNET and perceiver are able to learn efficiently on the complex mesh data as observed in the 3D case. However, MAgNET could learn more quickly than Perceiver IO, also MAgNET can learn efficiently even with the increased mesh complexity and input size. In case of Perceiver IO, as the size and complexity of mesh further increases, it becomes less and less robust and fails to learn efficiently. We observed that Perceiver IO failed to learn on the data with input dimension higher than 10<sup>4</sup>.</p>
<fig id="F6" position="float">
<label>FIGURE 6</label>
<caption>
<p>Training convergence for the proposed neural network frameworks for the <bold>(A)</bold> 2D case <bold>(B)</bold> 3D case.</p>
</caption>
<graphic xlink:href="fmats-10-1128954-g006.tif"/>
</fig>
<p>A possible interpretation of this behavior is that in the case of MAgNET topological information is externally provided through the adjacency matrix. Hence MAgNET can efficiently learn by leveraging inter-dependencies between different nodal feature values in the data. On the other hand Perceiver IO implicitly learns the nodal data dependencies and as the size of the input data increases, the task to find these inter-dependencies gets more difficult. Evidence of this behavior can be seen in <xref ref-type="fig" rid="F6">Figure&#xa0;6</xref>, where it is clear that Perceiver IO is optimizing through a much more complex objective function with a higher density of local minimas.</p>
<p>Once trained, proposed DNN frameworks are fast at the inference stage while predicting unseen examples. <xref ref-type="table" rid="T4">Table&#xa0;4</xref> provides training and inference time (for a single test example) for all three frameworks, for comparison FEM solution time is also provided. In particular, Perceiver IO takes a much longer time during the training phase but is extremely fast at the inference stage. It could make predictions on both small scale (2D) and large scale (3D) inputs in almost similar time. It has to be noted that to ensure the convergence of the iterative solver, the non-linear FEM problem is solved with incremental load steps. Hence the solution time for FEM increases with the magnitude of external force. Whereas trained DNN frameworks take almost similar time at the inference stage irrespective of external force magnitudes.</p>
<table-wrap id="T4" position="float">
<label>TABLE 4</label>
<caption>
<p>Comparison of training and inference times for all the three networks implemented in this work.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="left">Example</th>
<th align="center">Framework</th>
<th align="center">N. Parameters (&#xd7; E6)</th>
<th align="center">Training time (hrs)</th>
<th align="center">Inference time (s)</th>
<th align="center">FEM solver time (s)</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="left">2D</td>
<td align="center">CNN U-Net</td>
<td align="center">4.8</td>
<td align="center">18</td>
<td align="center">0.021</td>
<td align="center">0.6</td>
</tr>
<tr>
<td align="left"/>
<td align="center">MAgNET</td>
<td align="center">4.5</td>
<td align="center">132</td>
<td align="center">0.040</td>
<td align="left"/>
</tr>
<tr>
<td align="left"/>
<td align="center">Perceiver IO</td>
<td align="center">1.9</td>
<td align="center">521</td>
<td align="center">0.006</td>
<td align="left"/>
</tr>
<tr>
<td align="left">3D</td>
<td align="center">MAgNET</td>
<td align="center">33.9</td>
<td align="center">161</td>
<td align="center">0.217</td>
<td align="center">2.5</td>
</tr>
<tr>
<td align="left"/>
<td align="center">Perceiver IO</td>
<td align="center">4.4</td>
<td align="center">312</td>
<td align="center">0.006</td>
<td align="left"/>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec id="s3-5">
<title>3.5 Qualitative analysis of individual examples</title>
<p>In this section, we analyze deformations of individual examples by giving a qualitative comparison of predictions obtained using different networks. In particular, we analyze test examples with the maximum nodal displacement for 2D as well as the 3D case. In both cases, we plot nodal error contours standing for the absolute difference between the DNN prediction and the true FEM solution.</p>
<sec id="s3-5-1">
<title>3.5.1 2D case</title>
<p>The analyzed example in the <xref ref-type="fig" rid="F7">Figure&#xa0;7</xref> stands for the maximum nodal displacement example in the 2D test dataset. The node indicated by the green dot has the maximum nodal displacement of 0.35&#xa0;m. While the pink arrows represent corresponding nodal forces for the line density force applied on those four nodes.</p>
<fig id="F7" position="float">
<label>FIGURE 7</label>
<caption>
<p>Prediction error for different neural network frameworks when compared to true FEM solution, plotted on the deformed mesh (obtained using the same framework). Force with line density of (&#x2212;21.6645, &#x2212;2.99384) N is applied as shown with pink arrows. True displacement of the green node is 0.35&#xa0;m. Nodal error contours obtained <bold>(A)</bold> using CNN U-Net <bold>(B)</bold> using MAgNET <bold>(C)</bold> using Perceiver IO.</p>
</caption>
<graphic xlink:href="fmats-10-1128954-g007.tif"/>
</fig>
<p>
<xref ref-type="fig" rid="F7">Figure&#xa0;7</xref> shows absolute error counters of nodal displacements predicted using proposed DNN frameworks when compared with the FEM solution. All the proposed neural network frameworks can accurately predict the deformed mesh. Percentage prediction error (when compared to the true FEM solution) for the green node is 0.03%, 0.42%, 0.06% using CNN, MAgNET, and Perceiver IO network respectively. We observed that for the small-scale structured inputs, both Perceiver IO and CNN U-Net could make better predictions when compared to MAgNET. In case of Perceiver IO, advantage likely comes from the network leveraging its capability of learning long-range correlations more accurately. This stems from the fact that Perceiver&#x2019;s inputs are not constrained by any topological assumption.</p>
</sec>
<sec id="s3-5-2">
<title>3.5.2 3D case</title>
<p>The 3D case is aiming to demonstrate the performance for unstructured meshes, that are commonly used when dealing with real world geometries. Tetrahedron discretisation of the continuous 3D domain is irregular and the mesh topology is complex. The fixed nodes are topologically far from the tip of the elephant trunk, thus making it challenging to communicate the boundary condition information. We show that both MAgNET and Perceiver IO can learn on such complex real-world examples efficiently.</p>
<p>Again we consider the test example with maximum nodal displacement. <xref ref-type="fig" rid="F8">Figure&#xa0;8</xref> shows that both MAgNET (<xref ref-type="fig" rid="F8">Figures&#xa0;8A,&#xa0;C,&#xa0;E</xref>) and Perceiver IO (<xref ref-type="fig" rid="F8">Figures&#xa0;8B,&#xa0;D,&#xa0;F</xref>) solutions are able to predict non-linear mesh deformations accurately. We further analyse the absolute nodal error counters for both predictions by plotting them on the deformed meshes predicted using respective DNN frameworks. Both front (<xref ref-type="fig" rid="F8">Figures&#xa0;8E,&#xa0;F</xref>) and side views (<xref ref-type="fig" rid="F8">Figures&#xa0;8C,&#xa0;D</xref>) obtained using MAgNET and Perceiver IO solutions respectively indicate low prediction errors for both frameworks. The green node (at the tip of the ear) shown in (<xref ref-type="fig" rid="F8">Figures&#xa0;8E,&#xa0;F</xref>) has a maximum nodal displacement of 140.04&#xa0;m for this example. The percentage prediction error when compared to the true FEM solution for this green node is 0.03%, and 0.02% using MAgNET and Perceiver IO networks respectively. Both MAgNET and Perceiver IO are observed to make efficient predictions, with Perceiver IO giving relatively low nodal errors for this demanding case. Also, owing to the lesser number of trainable parameters, Perceiver IO is much faster at the inference stage.</p>
<fig id="F8" position="float">
<label>FIGURE 8</label>
<caption>
<p>Deformation of elephant mesh subjected to external body force density (0.34, 0.0, 0.35) N/kg. First column represents MAgNET solutions while the second column represents Perceiver IO solutions. <bold>(A,B)</bold> Deformed meshes using MAgNET (dark blue) and Perceiver IO (sky blue) respectively, for comparison FEM mesh is presented in red. The rest position is indicated with gray mesh. <bold>(C,D)</bold> Side view of nodal error contours when compared to the FEM solution, plotted on the deformed meshes for MAgNET and Perceiver IO solution respectively. <bold>(E,F)</bold> Front view of nodal error contours for MAgNET and Perceiver IO respectively. The true displacement of the green node is 140.04&#xa0;m.</p>
</caption>
<graphic xlink:href="fmats-10-1128954-g008.tif"/>
</fig>
</sec>
</sec>
</sec>
<sec sec-type="conclusion" id="s4">
<title>4 Conclusion</title>
<p>In this work, we demonstrated the capabilities of three promising deep neural network (DNN) frameworks for accurate and fast predictions of non-linear deformations of solid bodies. We compared their performance on two benchmark examples, in which data was generated by the finite element method. Although we only tested the frameworks for the Noe-Hoohean material model, they are compatible with more general hyperelastic models, such as Mooney&#x2013;Rivlin or Ogden models. As such, they promise to be used as surrogate models for non-linear computational models in mechanics.</p>
<p>The comparison included two very recent DNN frameworks, MAgNET and Perceiver IO, that are naturally able to work with arbitrarily structured data at inputs/outputs, including complex finite element meshes that originate from real-world applications. The third compared framework, CNN U-Net, could only operate on grid inputs/outputs, and we suggested possible remedies to extend it to work with arbitrary unstructured meshes. When looking at prediction capabilities, especially interesting are the capabilities of the Perceiver IO network, which demonstrated to give better predictions with a lesser number of parameters, as compared to MAgNET and CNN U-Net. Additionally, the use of Perceiver IO creates a direct link to rapidly advancing research in ML and AI communities, which promises further advancements.</p>
<p>MAgNET and Perceiver IO are designed to be flexible in terms of the input and output structures, allowing them to potentially be applied to a wide range of problems. One possible application for these types of neural networks is in <italic>ab initio</italic> multi-scale modeling, which is also pursued in our team, see <xref ref-type="bibr" rid="B30">Hauseux&#xa0;et&#xa0;al. (2020)</xref>. These methods could be used to accelerate computationally expensive accurate simulations of large atomic systems by helping to connect atomic-level simulations with the macroscopic continuum description of materials. As such, these neural networks could lead to significant strides in the field of materials science.</p>
<p>One of the first possible future extensions of the presented frameworks would be to incorporate the physics-informed neural network paradigm. This can be easily achieved by incorporating relevant physical laws in the optimization objective of the training procedure. Such extension can further increase the accuracy of predictions and accelerate the training procedure. Another possible extension is to consider a much wider class of phenomena and models, including buckling instabilities and more general history/time-dependent phenomena (visco-elasticity, dynamics, plasticity, etc.), which would allow tackling more challenging problems in solid mechanics, see e.g., (<xref ref-type="bibr" rid="B62">Vijayaraghavan&#xa0;et&#xa0;al., 2021</xref>). Going beyond mechanics, these approaches can be also adopted for a much wider range of engineering and scientific applications.</p>
</sec>
</body>
<back>
<sec sec-type="data-availability" id="s5">
<title>Data availability statement</title>
<p>The datasets presented in this study can be found in online repositories. The names of the repository/repositories and accession number(s) can be found below: <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.5281/zenodo.7585319">https://doi.org/10.5281/zenodo.7585319</ext-link>, <ext-link ext-link-type="uri" xlink:href="https://github.com/saurabhdeshpande93/convolution-aggregation-attention">https://github.com/saurabhdeshpande93/convolution-aggregation-attention</ext-link>.</p>
</sec>
<sec id="s6">
<title>Author contributions</title>
<p>SD conceptualized and created the core structure of the manuscript with the help of JL. SD&#xa0;and RS carried out all the simulations presented in this work. SD&#xa0;wrote the main part and assembled all of the manuscript, RS contributed to multiple sections. JL and SB revised the text in detail. All authors read, discussed, and approved the final version.</p>
</sec>
<sec id="s7">
<title>Funding</title>
<p>
<inline-graphic xlink:href="fmats-10-1128954-fx1.tif"/>This project has received funding from the European Union&#x2019;s Horizon 2020 research and innovation programme under the Marie Sklodowska-Curie grant agreement No. 764644. Jakub Lengiewicz would like to acknowledge the support from EU Horizon 2020 Marie Sklodowska Curie Individual Fellowship <italic>MOrPhEM</italic> under Grant 800150. St&#xe9;phane Bordas, Jakub Lengiewicz and Ra&#xfa;l I. Sosa are grateful for the support of the Fonds National de la Recherche Luxembourg FNR grant QuaC C20/MS/14782078. St&#xe9;phane Bordas received funding from the European Union&#x2019;s Horizon 2020 research and innovation programme under grant agreement No 811099 TWINNING Project DRIVEN for the University of Luxembourg.</p>
</sec>
<sec sec-type="COI-statement" id="s8">
<title>Conflict of interest</title>
<p>The authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec sec-type="disclaimer" id="s10">
<title>Publisher&#x2019;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<sec id="s9">
<title>Author disclaimer</title>
<p>This paper only contains the author&#x2019;s views and the Research Executive Agency and the Commission are not responsible for any use that may be made of the information it contains.</p>
</sec>
<ref-list>
<title>References</title>
<ref id="B1">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Abadi</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Agarwal</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Barham</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Brevdo</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Citro</surname>
<given-names>C.</given-names>
</name>
<etal/>
</person-group> (<year>2015</year>). <source>TensorFlow: Large-scale machine learning on heterogeneous systems Software available from tensorflow.org</source>.</citation>
</ref>
<ref id="B2">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Abueidda</surname>
<given-names>D. W.</given-names>
</name>
<name>
<surname>Koric</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Sobh</surname>
<given-names>N. A.</given-names>
</name>
<name>
<surname>Sehitoglu</surname>
<given-names>H.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Deep learning for plasticity and thermo-viscoplasticity</article-title>. <source>Int. J. Plasticity</source> <volume>136</volume>, <fpage>102852</fpage>. <pub-id pub-id-type="doi">10.1016/j.ijplas.2020.102852</pub-id>
</citation>
</ref>
<ref id="B3">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Aydin</surname>
<given-names>R. C.</given-names>
</name>
<name>
<surname>Braeu</surname>
<given-names>F. A.</given-names>
</name>
<name>
<surname>Cyron</surname>
<given-names>C. J.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>General multi-fidelity framework for training artificial neural networks with computational models</article-title>. <source>Front. Mater.</source> <volume>6</volume>, <fpage>61</fpage>. <pub-id pub-id-type="doi">10.3389/fmats.2019.00061</pub-id>
</citation>
</ref>
<ref id="B4">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Baevski</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Zhou</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Mohamed</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Auli</surname>
<given-names>M.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>wav2vec 2.0: A framework for self-supervised learning of speech representations</article-title>. <source>Adv. Neural Inf. Process. Syst.</source> <volume>33</volume>, <fpage>12449</fpage>&#x2013;<lpage>12460</lpage>.</citation>
</ref>
<ref id="B5">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Barrios</surname>
<given-names>J. M.</given-names>
</name>
<name>
<surname>Romero</surname>
<given-names>P. E.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Decision tree methods for predicting surface roughness in fused deposition modeling parts</article-title>. <source>Materials</source> <volume>12</volume>, <fpage>2574</fpage>. <pub-id pub-id-type="doi">10.3390/ma12162574</pub-id>
</citation>
</ref>
<ref id="B6">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Battaglia</surname>
<given-names>P. W.</given-names>
</name>
<name>
<surname>Hamrick</surname>
<given-names>J. B.</given-names>
</name>
<name>
<surname>Bapst</surname>
<given-names>V.</given-names>
</name>
<name>
<surname>Sanchez-Gonzalez</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Zambaldi</surname>
<given-names>V.</given-names>
</name>
<name>
<surname>Malinowski</surname>
<given-names>M.</given-names>
</name>
<etal/>
</person-group> (<year>2018</year>). <source>Relational inductive biases, deep learning, and graph networks</source>. <comment>arXiv preprint arXiv:1806.01261</comment>.</citation>
</ref>
<ref id="B7">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Bock</surname>
<given-names>F. E.</given-names>
</name>
<name>
<surname>Aydin</surname>
<given-names>R. C.</given-names>
</name>
<name>
<surname>Cyron</surname>
<given-names>C. J.</given-names>
</name>
<name>
<surname>Huber</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Kalidindi</surname>
<given-names>S. R.</given-names>
</name>
<name>
<surname>Klusemann</surname>
<given-names>B.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>A review of the application of machine learning and data mining approaches in continuum materials mechanics</article-title>. <source>Front. Mater.</source> <volume>6</volume>, <fpage>110</fpage>, <pub-id pub-id-type="doi">10.3389/fmats.2019.00110</pub-id>
</citation>
</ref>
<ref id="B8">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Bronstein</surname>
<given-names>M. M.</given-names>
</name>
<name>
<surname>Bruna</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Cohen</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Veli&#x10d;kovi&#x107;</surname>
<given-names>P.</given-names>
</name>
</person-group> (<year>2021</year>). <source>Geometric deep learning: Grids, groups, graphs, geodesics, and gauges</source>. <comment>arXiv preprint arXiv:2104.13478</comment>.</citation>
</ref>
<ref id="B9">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Brown</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Mann</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Ryder</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Subbiah</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Kaplan</surname>
<given-names>J. D.</given-names>
</name>
<name>
<surname>Dhariwal</surname>
<given-names>P.</given-names>
</name>
<etal/>
</person-group> (<year>2020</year>). <article-title>Language models are few-shot learners</article-title>. <source>Adv. neural Inf. Process. Syst.</source> <volume>33</volume>, <fpage>1877</fpage>&#x2013;<lpage>1901</lpage>.</citation>
</ref>
<ref id="B10">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Brunet</surname>
<given-names>J.-N.</given-names>
</name>
<name>
<surname>Mendizabal</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Petit</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Golse</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Vibert</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Cotin</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2019</year>). &#x201c;<article-title>Physics-based deep neural network for augmented reality during liver surgery</article-title>,&#x201d; in <source>Medical image computing and computer assisted intervention &#x2013; miccai 2019</source>. Editors <person-group person-group-type="editor">
<name>
<surname>Shen</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Peters</surname>
<given-names>T. M.</given-names>
</name>
<name>
<surname>Staib</surname>
<given-names>L. H.</given-names>
</name>
<name>
<surname>Essert</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Zhou</surname>
<given-names>S.</given-names>
</name>
<etal/>
</person-group> (<publisher-loc>Cham</publisher-loc>: <publisher-name>Springer International Publishing</publisher-name>), <fpage>137</fpage>&#x2013;<lpage>145</lpage>.</citation>
</ref>
<ref id="B11">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Bui</surname>
<given-names>H. P.</given-names>
</name>
<name>
<surname>Tomar</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Courtecuisse</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Cotin</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Bordas</surname>
<given-names>S. P. A.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>Real-time error control for surgical simulation</article-title>. <source>IEEE Trans. Biomed. Eng.</source> <volume>65</volume>, <fpage>596</fpage>&#x2013;<lpage>607</lpage>. <pub-id pub-id-type="doi">10.1109/TBME.2017.2695587</pub-id>
</citation>
</ref>
<ref id="B12">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Butler</surname>
<given-names>K. T.</given-names>
</name>
<name>
<surname>Davies</surname>
<given-names>D. W.</given-names>
</name>
<name>
<surname>Cartwright</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Isayev</surname>
<given-names>O.</given-names>
</name>
<name>
<surname>Walsh</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>Machine learning for molecular and materials science</article-title>. <source>Nature</source> <volume>559</volume>, <fpage>547</fpage>&#x2013;<lpage>555</lpage>. <pub-id pub-id-type="doi">10.1038/s41586-018-0337-2</pub-id>
</citation>
</ref>
<ref id="B13">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Capuano</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Rimoli</surname>
<given-names>J. J.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Smart finite elements: A novel machine learning application</article-title>. <source>Comput. Methods Appl. Mech. Eng.</source> <volume>345</volume>, <fpage>363</fpage>&#x2013;<lpage>381</lpage>. <pub-id pub-id-type="doi">10.1016/j.cma.2018.10.046</pub-id>
</citation>
</ref>
<ref id="B14">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Chen</surname>
<given-names>A. I.</given-names>
</name>
<name>
<surname>Balter</surname>
<given-names>M. L.</given-names>
</name>
<name>
<surname>Maguire</surname>
<given-names>T. J.</given-names>
</name>
<name>
<surname>Yarmush</surname>
<given-names>M. L.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Deep learning robotic guidance for autonomous vascular access</article-title>. <source>Nat. Mach. Intell.</source> <volume>2</volume>, <fpage>104</fpage>&#x2013;<lpage>115</lpage>. <pub-id pub-id-type="doi">10.1038/s42256-020-0148-7</pub-id>
</citation>
</ref>
<ref id="B15">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Choi</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Crump</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Duriez</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Elmquist</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Hager</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Han</surname>
<given-names>D.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>On the use of simulation in robotics: Opportunities, challenges, and suggestions for moving forward</article-title>. <source>Proc. Natl. Acad. Sci.</source> <volume>118</volume>, <fpage>e1907856118</fpage>. <pub-id pub-id-type="doi">10.1073/pnas.1907856118</pub-id>
</citation>
</ref>
<ref id="B16">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Choudhary</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>DeCost</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Jain</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Tavazza</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Cohn</surname>
<given-names>R.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). <article-title>Recent advances and applications of deep learning methods in materials science</article-title>. <source>npj Comput. Mater.</source> <volume>8</volume>, <fpage>59</fpage>&#x2013;<lpage>26</lpage>. <pub-id pub-id-type="doi">10.1038/s41524-022-00734-6</pub-id>
</citation>
</ref>
<ref id="B17">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Cotin</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Delingette</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Ayache</surname>
<given-names>N.</given-names>
</name>
</person-group> (<year>1999</year>). <article-title>Real-time elastic deformations of soft tissues for surgery simulation</article-title>. <source>IEEE Trans. Vis. Comput. Graph.</source> <volume>5</volume>, <fpage>62</fpage>&#x2013;<lpage>73</lpage>. <pub-id pub-id-type="doi">10.1109/2945.764872</pub-id>
</citation>
</ref>
<ref id="B18">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Courtecuisse</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Allard</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Kerfriden</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Bordas</surname>
<given-names>S. P.</given-names>
</name>
<name>
<surname>Cotin</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Duriez</surname>
<given-names>C.</given-names>
</name>
</person-group> (<year>2014</year>). <article-title>Real-time simulation of contact and cutting of heterogeneous soft-tissues</article-title>. <source>Med. image Anal.</source> <volume>18</volume>, <fpage>394</fpage>&#x2013;<lpage>410</lpage>. <pub-id pub-id-type="doi">10.1016/j.media.2013.11.001</pub-id>
</citation>
</ref>
<ref id="B19">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>De Vivo</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Masetti</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Bottegoni</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Cavalli</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2016</year>). <article-title>Role of molecular dynamics and related methods in drug discovery</article-title>. <source>J. Med. Chem.</source> <volume>59</volume>, <fpage>4035</fpage>&#x2013;<lpage>4061</lpage>. <pub-id pub-id-type="doi">10.1021/acs.jmedchem.5b01684</pub-id>
</citation>
</ref>
<ref id="B20">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Dennler</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Bauer</surname>
<given-names>D. E.</given-names>
</name>
<name>
<surname>Scheibler</surname>
<given-names>A.-G.</given-names>
</name>
<name>
<surname>Spirig</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>G&#xf6;tschi</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>F&#xfc;rnstahl</surname>
<given-names>P.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>Augmented reality in the operating room: A clinical feasibility study</article-title>. <source>BMC Musculoskelet. Disord.</source> <volume>22</volume>, <fpage>451</fpage>. <pub-id pub-id-type="doi">10.1186/s12891-021-04339-w</pub-id>
</citation>
</ref>
<ref id="B21">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Deshpande</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Bordas</surname>
<given-names>S. P. A.</given-names>
</name>
<name>
<surname>Lengiewicz</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2022b</year>). <article-title>MAgNET: A graph U-net architecture for mesh-based simulations</article-title>. <source>arXiv</source>. <pub-id pub-id-type="doi">10.48550/ARXIV.2211.00713</pub-id>
</citation>
</ref>
<ref id="B22">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Deshpande</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Lengiewicz</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Bordas</surname>
<given-names>S. P.</given-names>
</name>
</person-group> (<year>2022a</year>). <article-title>Probabilistic deep learning for real-time large deformation simulations</article-title>. <source>Comput. Methods Appl. Mech. Eng.</source> <volume>398</volume>, <fpage>115307</fpage>. <pub-id pub-id-type="doi">10.1016/j.cma.2022.115307</pub-id>
</citation>
</ref>
<ref id="B23">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Devlin</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Chang</surname>
<given-names>M.-W.</given-names>
</name>
<name>
<surname>Lee</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Toutanova</surname>
<given-names>K.</given-names>
</name>
</person-group> (<year>2018</year>). <source>Bert: Pre-training of deep bidirectional transformers for language understanding</source>. <comment>arXiv preprint arXiv:1810.04805</comment>.</citation>
</ref>
<ref id="B24">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>El Haber</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Viquerat</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Larcher</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Ryckelynck</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Alves</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Patil</surname>
<given-names>A.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). <article-title>Deep learning model to assist multiphysics conjugate problems</article-title>. <source>Phys. Fluids</source> <volume>34</volume>, <fpage>015131</fpage>. <pub-id pub-id-type="doi">10.1063/5.0077723</pub-id>
</citation>
</ref>
<ref id="B25">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Elouneg</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Bertin</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Lucot</surname>
<given-names>Q.</given-names>
</name>
<name>
<surname>Tissot</surname>
<given-names>V.</given-names>
</name>
<name>
<surname>Jacquet</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Chambert</surname>
<given-names>J.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). <article-title>
<italic>In vivo</italic> skin anisotropy dataset from annular suction test</article-title>. <source>Data Brief</source> <volume>40</volume>, <fpage>107835</fpage>. <pub-id pub-id-type="doi">10.1016/j.dib.2022.107835</pub-id>
</citation>
</ref>
<ref id="B26">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Flaschel</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Kumar</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>De Lorenzis</surname>
<given-names>L.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Unsupervised discovery of interpretable hyperelastic constitutive laws</article-title>. <source>Comput. Methods Appl. Mech. Eng.</source> <volume>381</volume>, <fpage>113852</fpage>. <pub-id pub-id-type="doi">10.1016/j.cma.2021.113852</pub-id>
</citation>
</ref>
<ref id="B27">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Friesner</surname>
<given-names>R. A.</given-names>
</name>
</person-group> (<year>2005</year>). <article-title>
<italic>Ab initio</italic> quantum chemistry: Methodology and applications</article-title>. <source>Proc. Natl. Acad. Sci.</source> <volume>102</volume>, <fpage>6648</fpage>&#x2013;<lpage>6653</lpage>. <pub-id pub-id-type="doi">10.1073/pnas.0408036102</pub-id>
</citation>
</ref>
<ref id="B28">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Gholamalizadeh</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Moshfeghifar</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Ferguson</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Schneider</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Panozzo</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Darkner</surname>
<given-names>S.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). <article-title>Open-full-jaw: An open-access dataset and pipeline for finite element models of human jaw</article-title>. <source>Comput. Methods Programs Biomed.</source> <volume>224</volume>, <fpage>107009</fpage>. <pub-id pub-id-type="doi">10.1016/j.cmpb.2022.107009</pub-id>
</citation>
</ref>
<ref id="B29">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Guo</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Iorio</surname>
<given-names>F.</given-names>
</name>
</person-group> (<year>2016</year>). &#x201c;<article-title>Convolutional neural networks for steady flow approximation</article-title>,&#x201d; in <source>Proceedings of the 22nd ACM SIGKDD international conference on knowledge discovery and data mining</source> (<publisher-loc>New York, NY, USA</publisher-loc>: <publisher-name>Association for Computing Machinery</publisher-name>), <fpage>481</fpage>&#x2013;<lpage>490</lpage>. <pub-id pub-id-type="doi">10.1145/2939672.2939738</pub-id>
</citation>
</ref>
<ref id="B30">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hauseux</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Nguyen</surname>
<given-names>T.-T.</given-names>
</name>
<name>
<surname>Ambrosetti</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Ruiz</surname>
<given-names>K. S.</given-names>
</name>
<name>
<surname>Bordas</surname>
<given-names>S. P. A.</given-names>
</name>
<name>
<surname>Tkatchenko</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>From quantum to continuum mechanics in the delamination of atomically-thin layers from substrates</article-title>. <source>Nat. Commun.</source> <volume>11</volume>, <fpage>1651</fpage>. <pub-id pub-id-type="doi">10.1038/s41467-020-15480-w</pub-id>
</citation>
</ref>
<ref id="B31">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Jaegle</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Borgeaud</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Alayrac</surname>
<given-names>J.-B.</given-names>
</name>
<name>
<surname>Doersch</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Ionescu</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Ding</surname>
<given-names>D.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). &#x201c;<article-title>Perceiver IO: A general architecture for structured inputs and outputs</article-title>,&#x201d; in <source>International conference on learning representations</source>.</citation>
</ref>
<ref id="B32">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Jha</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Ward</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Paul</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Liao</surname>
<given-names>W.-k.</given-names>
</name>
<name>
<surname>Choudhary</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Wolverton</surname>
<given-names>C.</given-names>
</name>
<etal/>
</person-group> (<year>2018</year>). <article-title>Elemnet: Deep learning the chemistry of materials from only elemental composition</article-title>. <source>Sci. Rep.</source> <volume>8</volume>, <fpage>1</fpage>&#x2013;<lpage>13</lpage>. <pub-id pub-id-type="doi">10.1038/s41598-018-35934-y</pub-id>
</citation>
</ref>
<ref id="B33">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Kim</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Mishra</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Jin</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Panda</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Kuehne</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Karlinsky</surname>
<given-names>L.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). &#x201c;<article-title>How transferable are video representations based on synthetic data?</article-title>,&#x201d; in <source>Thirty-sixth conference on neural information processing systems datasets and benchmarks track</source>.</citation>
</ref>
<ref id="B34">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kingma</surname>
<given-names>D. P.</given-names>
</name>
<name>
<surname>Ba</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2014</year>). <article-title>Adam: A method for stochastic optimization</article-title>. <comment>
<italic>arXiv</italic>
</comment>. <pub-id pub-id-type="doi">10.48550/ARXIV.1412.6980</pub-id>
</citation>
</ref>
<ref id="B35">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Korelc</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2002</year>). <article-title>Multi-language and multi-environment generation of nonlinear finite element codes</article-title>. <source>Eng. Comput.</source> <volume>18</volume>, <fpage>312</fpage>&#x2013;<lpage>327</lpage>. <pub-id pub-id-type="doi">10.1007/s003660200028</pub-id>
</citation>
</ref>
<ref id="B36">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Krokos</surname>
<given-names>V.</given-names>
</name>
<name>
<surname>Bordas</surname>
<given-names>S. P. A.</given-names>
</name>
<name>
<surname>Kerfriden</surname>
<given-names>P.</given-names>
</name>
</person-group> (<year>2022a</year>). <source>A graph-based probabilistic geometric deep learning framework with online physics-based corrections to predict the criticality of defects in porous materials</source>. <pub-id pub-id-type="doi">10.48550/ARXIV.2205.06562</pub-id>
</citation>
</ref>
<ref id="B37">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Krokos</surname>
<given-names>V.</given-names>
</name>
<name>
<surname>Bui Xuan</surname>
<given-names>V.</given-names>
</name>
<name>
<surname>Bordas</surname>
<given-names>S. P. A.</given-names>
</name>
<name>
<surname>Young</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Kerfriden</surname>
<given-names>P.</given-names>
</name>
</person-group> (<year>2022b</year>). <article-title>A bayesian multiscale cnn framework to predict local stress fields in structures with microscale features</article-title>. <source>Comput. Mech.</source> <volume>69</volume>, <fpage>733</fpage>&#x2013;<lpage>766</lpage>. <pub-id pub-id-type="doi">10.1007/s00466-021-02112-3</pub-id>
</citation>
</ref>
<ref id="B38">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Le</surname>
<given-names>T. A.</given-names>
</name>
<name>
<surname>Baydin</surname>
<given-names>A. G.</given-names>
</name>
<name>
<surname>Zinkov</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Wood</surname>
<given-names>F.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>Using synthetic data to train neural networks is model-based reasoning</article-title>. In <source>2017 international joint conference on neural networks (IJCNN)</source> (<publisher-name>IEEE</publisher-name>), <fpage>3514</fpage>&#x2013;<lpage>3521</lpage>.</citation>
</ref>
<ref id="B39">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Liu</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Zhao</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Ju</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Shi</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>Materials discovery and design using machine learning</article-title>. <source>J. Materiomics</source> <volume>3</volume>, <fpage>159</fpage>&#x2013;<lpage>177</lpage>. <pub-id pub-id-type="doi">10.1016/j.jmat.2017.08.002</pub-id>
</citation>
</ref>
<ref id="B40">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Loshchilov</surname>
<given-names>I.</given-names>
</name>
<name>
<surname>Hutter</surname>
<given-names>F.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>Decoupled weight decay regularization</article-title>. <source>arXiv</source>. <pub-id pub-id-type="doi">10.48550/ARXIV.1711.05101</pub-id>
</citation>
</ref>
<ref id="B41">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Mao</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Jagtap</surname>
<given-names>A. D.</given-names>
</name>
<name>
<surname>Karniadakis</surname>
<given-names>G. E.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Physics-informed neural networks for high-speed flows</article-title>. <source>Comput. Methods Appl. Mech. Eng.</source> <volume>360</volume>, <fpage>112789</fpage>. <pub-id pub-id-type="doi">10.1016/j.cma.2019.112789</pub-id>
</citation>
</ref>
<ref id="B42">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Mazier</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Ribes</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Gilles</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Bordas</surname>
<given-names>S. P.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>A rigged model of the breast for preoperative surgical planning</article-title>. <source>J. Biomechanics</source> <volume>128</volume>, <fpage>110645</fpage>. <pub-id pub-id-type="doi">10.1016/j.jbiomech.2021.110645</pub-id>
</citation>
</ref>
<ref id="B43">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>McFall</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Mahan</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2009</year>). <article-title>Artificial neural network method for solution of boundary value problems with exact satisfaction of arbitrary boundary conditions</article-title>. <source>IEEE Trans. neural Netw.</source> <volume>20</volume>, <fpage>1221</fpage>. <lpage>1233</lpage>. <pub-id pub-id-type="doi">10.1109/tnn.2009.2020735</pub-id>
</citation>
</ref>
<ref id="B44">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Mendizabal</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>M&#xe1;rquez-Neila</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Cotin</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Simulation of hyperelastic materials in real-time using deep learning</article-title>. <source>Med. Image Anal.</source> <volume>59</volume>, <fpage>101569</fpage>. <pub-id pub-id-type="doi">10.1016/j.media.2019.101569</pub-id>
</citation>
</ref>
<ref id="B45">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Mianroodi</surname>
<given-names>J. R.</given-names>
</name>
<name>
<surname>H Siboni</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Raabe</surname>
<given-names>D.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Teaching solid mechanics to artificial intelligence&#x2014;A fast solver for heterogeneous materials</article-title>. <source>Npj Comput. Mater.</source> <volume>7</volume>, <fpage>99</fpage>&#x2013;<lpage>10</lpage>. <pub-id pub-id-type="doi">10.1038/s41524-021-00571-z</pub-id>
</citation>
</ref>
<ref id="B46">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Odot</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Haferssas</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Cotin</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Deepphysics: A physics aware deep learning framework for real-time simulation</article-title>. <source>Int. J. Numer. Methods Eng.</source> <volume>123</volume>, <fpage>2381</fpage>&#x2013;<lpage>2398</lpage>. <pub-id pub-id-type="doi">10.1002/nme.6943</pub-id>
</citation>
</ref>
<ref id="B47">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Oishi</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Yagawa</surname>
<given-names>G.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>Computational mechanics enhanced by deep learning</article-title>. <source>Comput. Methods Appl. Mech. Eng.</source> <volume>327</volume>, <fpage>327</fpage>&#x2013;<lpage>351</lpage>. <pub-id pub-id-type="doi">10.1016/j.cma.2017.08.040</pub-id>
</citation>
</ref>
<ref id="B48">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Paszke</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Gross</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Massa</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Lerer</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Bradbury</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Chanan</surname>
<given-names>G.</given-names>
</name>
<etal/>
</person-group> (<year>2019</year>). &#x201c;<article-title>Pytorch: An imperative style, high-performance deep learning library</article-title>,&#x201d; in <source>Advances in neural information processing systems</source> (<publisher-name>Curran Associates, Inc.</publisher-name>), <volume>32</volume>, <fpage>8024</fpage>&#x2013;<lpage>8035</lpage>.</citation>
</ref>
<ref id="B49">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Pfaff</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Fortunato</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Gonzalez</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Battaglia</surname>
<given-names>P.</given-names>
</name>
</person-group> (<year>2021</year>). &#x201c;<article-title>Learning mesh-based simulation with graph networks</article-title>,&#x201d; in <source>International conference on learning representations</source>.</citation>
</ref>
<ref id="B50">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Pfeiffer</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Riediger</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Weitz</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Speidel</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Learning soft tissue behavior of organs for surgical navigation with convolutional neural networks</article-title>. <source>Int. J. Comput. Assisted Radiology Surg.</source> <volume>14</volume>, <fpage>1147</fpage>&#x2013;<lpage>1155</lpage>. <pub-id pub-id-type="doi">10.1007/s11548-019-01965-7</pub-id>
</citation>
</ref>
<ref id="B51">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Rupp</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Tkatchenko</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>M&#xfc;ller</surname>
<given-names>K.-R.</given-names>
</name>
<name>
<surname>von Lilienfeld</surname>
<given-names>O. A.</given-names>
</name>
</person-group> (<year>2012</year>). <article-title>Fast and accurate modeling of molecular atomization energies with machine learning</article-title>. <source>Phys. Rev. Lett.</source> <volume>108</volume>, <fpage>058301</fpage>. <pub-id pub-id-type="doi">10.1103/PhysRevLett.108.058301</pub-id>
</citation>
</ref>
<ref id="B52">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Rus</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Tolley</surname>
<given-names>M. T.</given-names>
</name>
</person-group> (<year>2015</year>). <article-title>Design, fabrication and control of soft robots</article-title>. <source>Nature</source> <volume>521</volume>, <fpage>467</fpage>&#x2013;<lpage>475</lpage>. <pub-id pub-id-type="doi">10.1038/nature14543</pub-id>
</citation>
</ref>
<ref id="B53">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Samaniego</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Anitescu</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Goswami</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Nguyen-Thanh</surname>
<given-names>V.</given-names>
</name>
<name>
<surname>Guo</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Hamdia</surname>
<given-names>K.</given-names>
</name>
<etal/>
</person-group> (<year>2020</year>). <article-title>An energy approach to the solution of partial differential equations in computational mechanics via machine learning: Concepts, implementation and applications</article-title>. <source>Comput. Methods Appl. Mech. Eng.</source> <volume>362</volume>, <fpage>112790</fpage>. <pub-id pub-id-type="doi">10.1016/j.cma.2019.112790</pub-id>
</citation>
</ref>
<ref id="B54">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Schleder</surname>
<given-names>G. R.</given-names>
</name>
<name>
<surname>Padilha</surname>
<given-names>A. C.</given-names>
</name>
<name>
<surname>Acosta</surname>
<given-names>C. M.</given-names>
</name>
<name>
<surname>Costa</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Fazzio</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>From DFT to machine learning: Recent approaches to materials science&#x2013;a review</article-title>. <source>J. Phys. Mater.</source> <volume>2</volume>, <fpage>032001</fpage>. <pub-id pub-id-type="doi">10.1088/2515-7639/ab084b</pub-id>
</citation>
</ref>
<ref id="B55">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Schmidt</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Marques</surname>
<given-names>M. R.</given-names>
</name>
<name>
<surname>Botti</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Marques</surname>
<given-names>M. A.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Recent advances and applications of machine learning in solid-state materials science</article-title>. <source>npj Comput. Mater.</source> <volume>5</volume>, <fpage>83</fpage>&#x2013;<lpage>36</lpage>. <pub-id pub-id-type="doi">10.1038/s41524-019-0221-0</pub-id>
</citation>
</ref>
<ref id="B56">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Sch&#xfc;tt</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Kindermans</surname>
<given-names>P.-J.</given-names>
</name>
<name>
<surname>Sauceda Felix</surname>
<given-names>H. E.</given-names>
</name>
<name>
<surname>Chmiela</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Tkatchenko</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>M&#xfc;ller</surname>
<given-names>K.-R.</given-names>
</name>
</person-group> (<year>2017</year>). &#x201c;<article-title>Schnet: A continuous-filter convolutional neural network for modeling quantum interactions</article-title>,&#x201d; in <source>Advances in neural information processing systems</source>. Editors <person-group person-group-type="editor">
<name>
<surname>Guyon</surname>
<given-names>I.</given-names>
</name>
<name>
<surname>Luxburg</surname>
<given-names>U. V.</given-names>
</name>
<name>
<surname>Bengio</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Wallach</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Fergus</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Vishwanathan</surname>
<given-names>S.</given-names>
</name>
<etal/>
</person-group> (<publisher-name>Curran Associates, Inc.</publisher-name>), <volume>30</volume>.</citation>
</ref>
<ref id="B57">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sch&#xfc;tt</surname>
<given-names>K. T.</given-names>
</name>
<name>
<surname>Arbabzadah</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Chmiela</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>M&#xfc;ller</surname>
<given-names>K. R.</given-names>
</name>
<name>
<surname>Tkatchenko</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>Quantum-chemical insights from deep tensor neural networks</article-title>. <source>Nat. Commun.</source> <volume>8</volume>, <fpage>13890</fpage>&#x2013;<lpage>13898</lpage>. <pub-id pub-id-type="doi">10.1038/ncomms13890</pub-id>
</citation>
</ref>
<ref id="B58">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Str&#xf6;nisch</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Meyer</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Lehmann</surname>
<given-names>C.</given-names>
</name>
</person-group> (<year>2022</year>). &#x201c;<article-title>Flow field prediction on large variable sized 2d point clouds with graph convolution</article-title>,&#x201d; in <source>Proceedings of the platform for advanced scientific computing conference</source> (<publisher-loc>New York, NY, USA</publisher-loc>: <publisher-name>Association for Computing Machinery</publisher-name>). <comment>PASC &#x2019;22</comment>. <pub-id pub-id-type="doi">10.1145/3539781.3539789</pub-id>
</citation>
</ref>
<ref id="B59">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Unke</surname>
<given-names>O. T.</given-names>
</name>
<name>
<surname>Chmiela</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Sauceda</surname>
<given-names>H. E.</given-names>
</name>
<name>
<surname>Gastegger</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Poltavsky</surname>
<given-names>I.</given-names>
</name>
<name>
<surname>Sch&#xfc;tt</surname>
<given-names>K. T.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>Machine learning force fields</article-title>. <source>Chem. Rev.</source> <volume>121</volume>, <fpage>10142</fpage>&#x2013;<lpage>10186</lpage>. <pub-id pub-id-type="doi">10.1021/acs.chemrev.0c01111</pub-id>
</citation>
</ref>
<ref id="B60">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Varrette</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Bouvry</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Cartiaux</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Georgatos</surname>
<given-names>F.</given-names>
</name>
</person-group> (<year>2014</year>). <article-title>Management of an academic hpc cluster: The ul experience</article-title>. <pub-id pub-id-type="doi">10.1109/HPCSim.2014.6903792</pub-id>
</citation>
</ref>
<ref id="B61">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Vaswani</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Shazeer</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Parmar</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Uszkoreit</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Jones</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Gomez</surname>
<given-names>A. N.</given-names>
</name>
<etal/>
</person-group> (<year>2017</year>). <article-title>Attention is all you need</article-title>. <source>Adv. neural Inf. Process. Syst.</source> <volume>30</volume>.</citation>
</ref>
<ref id="B62">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Vijayaraghavan</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Wu</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Noels</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Bordas</surname>
<given-names>S. P. A.</given-names>
</name>
<name>
<surname>Natarajan</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Beex</surname>
<given-names>L. A. A.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Neural-network acceleration of projection-based model-order-reduction for finite plasticity: Application to RVEs</article-title>. <source>arXiv</source>. <pub-id pub-id-type="doi">10.48550/ARXIV.2109.07747</pub-id>
</citation>
</ref>
<ref id="B63">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Vlassis</surname>
<given-names>N. N.</given-names>
</name>
<name>
<surname>Ma</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>W.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Geometric deep learning for computational mechanics part i: Anisotropic hyperelasticity</article-title>. <source>Comput. Methods Appl. Mech. Eng.</source> <volume>371</volume>, <fpage>113299</fpage>. <pub-id pub-id-type="doi">10.1016/j.cma.2020.113299</pub-id>
</citation>
</ref>
<ref id="B64">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Voulodimos</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Doulamis</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Doulamis</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Protopapadakis</surname>
<given-names>E.</given-names>
</name>
</person-group> (<year>2018</year>). <source>Deep learning for computer vision: A brief review</source>. <publisher-name>Computational intelligence and neuroscience</publisher-name>, <fpage>2018</fpage>.</citation>
</ref>
<ref id="B65">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Weerasuriya</surname>
<given-names>A. U.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Lu</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Tse</surname>
<given-names>K. T.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>C.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>A Gaussian process-based emulator for modeling pedestrian-level wind field</article-title>. <source>Build. Environ.</source> <volume>188</volume>, <fpage>107500</fpage>. <pub-id pub-id-type="doi">10.1016/j.buildenv.2020.107500</pub-id>
</citation>
</ref>
<ref id="B66">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wirtz</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Karajan</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Haasdonk</surname>
<given-names>B.</given-names>
</name>
</person-group> (<year>2015</year>). <article-title>Surrogate modeling of multiscale models using kernel methods</article-title>. <source>Int. J. Numer. Methods Eng.</source> <volume>101</volume>, <fpage>1</fpage>&#x2013;<lpage>28</lpage>. <pub-id pub-id-type="doi">10.1002/nme.4767</pub-id>
</citation>
</ref>
<ref id="B67">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Xu</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Ba</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Kiros</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Cho</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Courville</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Salakhudinov</surname>
<given-names>R.</given-names>
</name>
<etal/>
</person-group> (<year>2015</year>). &#x201c;<article-title>Show, attend and tell: Neural image caption generation with visual attention</article-title>,&#x201d; in <source>International conference on machine learning</source> (<publisher-name>PMLR</publisher-name>), <fpage>2048</fpage>&#x2013;<lpage>2057</lpage>.</citation>
</ref>
<ref id="B68">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zakutayev</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Wunder</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Schwarting</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Perkins</surname>
<given-names>J. D.</given-names>
</name>
<name>
<surname>White</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Munch</surname>
<given-names>K.</given-names>
</name>
<etal/>
</person-group> (<year>2018</year>). <article-title>An open experimental database for exploring inorganic materials</article-title>. <source>Sci. data</source> <volume>5</volume>, <fpage>1</fpage>&#x2013;<lpage>12</lpage>. <pub-id pub-id-type="doi">10.1038/sdata.2018.53</pub-id>
</citation>
</ref>
</ref-list>
</back>
</article>