<?xml version="1.0" encoding="UTF-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Archiving and Interchange DTD v2.3 20070202//EN" "archivearticle.dtd">
<article article-type="data-paper" dtd-version="2.3" xml:lang="EN" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Robot. AI</journal-id>
<journal-title>Frontiers in Robotics and AI</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Robot. AI</abbrev-journal-title>
<issn pub-type="epub">2296-9144</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="publisher-id">1384575</article-id>
<article-id pub-id-type="doi">10.3389/frobt.2024.1384575</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Robotics and AI</subject>
<subj-group>
<subject>Data Report</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>L-AVATeD: The lidar and visual walking terrain dataset</article-title>
<alt-title alt-title-type="left-running-head">Whipps et al.</alt-title>
<alt-title alt-title-type="right-running-head">
<ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3389/frobt.2024.1384575">10.3389/frobt.2024.1384575</ext-link>
</alt-title>
</title-group>
<contrib-group>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Whipps</surname>
<given-names>David</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="corresp" rid="c001">&#x2a;</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2618176/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Ippersiel</surname>
<given-names>Patrick</given-names>
</name>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2859541/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Dixon</surname>
<given-names>Philippe C.</given-names>
</name>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<xref ref-type="aff" rid="aff4">
<sup>4</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2859614/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/funding-acquisition/"/>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
<role content-type="https://credit.niso.org/contributor-roles/resources/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
</contrib-group>
<aff id="aff1">
<sup>1</sup>
<institution>D&#xe9;partement d&#x2019;informatique et de recherche op&#xe9;rationnelle</institution>, <institution>Universit&#xe9; de Montr&#xe9;al</institution>, <addr-line>Montr&#xe9;al</addr-line>, <addr-line>QC</addr-line>, <country>Canada</country>
</aff>
<aff id="aff2">
<sup>2</sup>
<institution>Mila, the Quebec Artificial Intelligence Institute</institution>, <addr-line>Montr&#xe9;al</addr-line>, <addr-line>QC</addr-line>, <country>Canada</country>
</aff>
<aff id="aff3">
<sup>3</sup>
<institution>School of Kinesiology and Physical Activity Sciences</institution>, <institution>Universit&#xe9; de Montr&#xe9;al</institution>, <addr-line>Montr&#xe9;al</addr-line>, <addr-line>QC</addr-line>, <country>Canada</country>
</aff>
<aff id="aff4">
<sup>4</sup>
<institution>Department of Kinesiology and Physical Education</institution>, <institution>McGill University</institution>, <addr-line>Montr&#xe9;al</addr-line>, <addr-line>QC</addr-line>, <country>Canada</country>
</aff>
<author-notes>
<fn fn-type="edited-by">
<p>
<bold>Edited by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1210346/overview">Luis Paya</ext-link>, Miguel Hern&#xe1;ndez University of Elche, Spain</p>
</fn>
<fn fn-type="edited-by">
<p>
<bold>Reviewed by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1419380/overview">Swarn Singh Rathour</ext-link>, Hitachi, Japan</p>
<p>
<ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1640296/overview">Alwin Poulose</ext-link>, Indian Institute of Science Education and Research, Thiruvananthapuram, India</p>
</fn>
<corresp id="c001">&#x2a;Correspondence: David Whipps, <email>david.whipps@umontreal.ca</email>
</corresp>
</author-notes>
<pub-date pub-type="epub">
<day>04</day>
<month>12</month>
<year>2024</year>
</pub-date>
<pub-date pub-type="collection">
<year>2024</year>
</pub-date>
<volume>11</volume>
<elocation-id>1384575</elocation-id>
<history>
<date date-type="received">
<day>09</day>
<month>02</month>
<year>2024</year>
</date>
<date date-type="accepted">
<day>18</day>
<month>11</month>
<year>2024</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#xa9; 2024 Whipps, Ippersiel and Dixon.</copyright-statement>
<copyright-year>2024</copyright-year>
<copyright-holder>Whipps, Ippersiel and Dixon</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<kwd-group>
<kwd>artificial intelligence</kwd>
<kwd>environment recognition</kwd>
<kwd>computer vision</kwd>
<kwd>wearable technology</kwd>
<kwd>assistive robotics</kwd>
<kwd>gait</kwd>
<kwd>3D</kwd>
<kwd>lidar</kwd>
</kwd-group>
<contract-num rid="cn001">RRGPIN-2022-04217</contract-num>
<contract-num rid="cn002">Junior 1</contract-num>
<contract-sponsor id="cn001">Natural Sciences and Engineering Research Council of Canada<named-content content-type="fundref-id">10.13039/501100000038</named-content>
</contract-sponsor>
<contract-sponsor id="cn002">Fonds de Recherche du Qu&#xe9;bec - Sant&#xe9;<named-content content-type="fundref-id">10.13039/501100000156</named-content>
</contract-sponsor>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-at-acceptance</meta-name>
<meta-value>Robot Vision and Artificial Perception</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec id="s1">
<title>1 Introduction</title>
<p>Worldwide, a significant number of individuals depend on assistive robotic technologies, such as sophisticated lower-limb prostheses and comprehensive exoskeletons to address challenges related to gait and mobility (<xref ref-type="bibr" rid="B13">Laschowski et al., 2020</xref>; <xref ref-type="bibr" rid="B3">Hu et al., 2018</xref>). These innovative devices significantly improve the quality of life for people with various disabilities, enabling them to navigate real-world environments more effectively and with greater independence. Despite significant advancements in recent decades, these systems often struggle with prompt and accurate responses to changes in the local environment <xref ref-type="bibr" rid="B3">Hu et al. (2018)</xref>. Typically, users either control these systems directly or use semi-autonomous modes, frequently needing to manually switch locomotion modes to adapt to different terrains (<xref ref-type="bibr" rid="B14">Laschowski et al., 2022</xref>). This requirement imposes an additional cognitive burden, potentially leading to distraction or injury. Recent research highlights the potential for substantial enhancements in exoskeleton control systems by transitioning away from user-initiated to automated locomotion mode changes (such as shifting from walking on flat surfaces to climbing stairs) to more autonomous control systems. A critical aspect of such systems is their ability to precisely classify the local environment <xref ref-type="bibr" rid="B13">Laschowski et al. (2020)</xref>).</p>
<p>Conventionally, the study and diagnosis of gait conditions is conducted in laboratory settings, where clinicians can control numerous variables and observe patients on a known walking surface. While this ensures accuracy, it limits the evaluation of gait patterns in varied natural environments and walking terrains <xref ref-type="bibr" rid="B1">Adamczyk et al. (2023)</xref>.</p>
<p>Emerging studies, (e.g., <xref ref-type="bibr" rid="B5">Dixon et al., 2018</xref>) have demonstrated that gait is materially affected by the walking surface. Specifically, walking surfaces have been shown to impact lower-extremity muscle activity (<xref ref-type="bibr" rid="B22">Voloshina and Ferris, 2015</xref>), joint kinematics (<xref ref-type="bibr" rid="B6">Dixon et al., 2021</xref>), inter-joint coordination and variability (<xref ref-type="bibr" rid="B9">Ippersiel et al., 2021</xref>; <xref ref-type="bibr" rid="B10">Ippersiel et al., 2022</xref>), and joint kinetics (<xref ref-type="bibr" rid="B15">Lay et al., 2006</xref>). Freeing researchers from the constraints imposed by the laboratory requires moving away from measurement systems designed to operate in a fixed environment, towards more portable hardware solutions that can function in ecological contexts. One issue that remains, however, is the ability of portable systems to accurately classify walking surfaces.</p>
<p>Portable gait analysis systems with multiple participant-mounted inertial measurement units (IMUs) can themselves be used to classify walking terrain. <xref ref-type="bibr" rid="B21">Shah et al. (2022)</xref> used six sensors to distinguish nine surface types in a group of 30 young healthy adults using a machine learning algorithm. It remains unclear however if terrain classification accuracy from IMU data would be affected by patient pathology. That is, algorithms based on IMU data alone may incorrectly assess patient pathology or severity of impairment due to particularities of terrain on which walking is performed. The data IMUs produce is valuable, but data collection outside of the lab could cause them to suffer from exactly the problem they are trying to solve; namely, the lack of a known, controlled walking surface. While these mobile systems capture patient gait characteristics, they must also leverage additional sensor capabilities to simultaneously deliver an accurate, ground-truth classification of the local walking terrain.</p>
<p>Advancements in Deep Learning, particularly in Convolutional Neural Network (CNN) architectures, have shown remarkable success in image classification <xref ref-type="bibr" rid="B12">Krizhevsky et al. (2012)</xref>. Various studies (<xref ref-type="bibr" rid="B14">Laschowski et al., 2022</xref>; <xref ref-type="bibr" rid="B4">Diaz et al., 2018</xref>; <xref ref-type="bibr" rid="B4">Diaz et al., 2018</xref>) have successfully used deep learning and visual data for accurate terrain identification, but these techniques require expansive data sets for their training. Combining multiple data modalities, such as depth sensors (<xref ref-type="bibr" rid="B24">Zhang et al., 2011</xref>; <xref ref-type="bibr" rid="B25">Zhang et al., 2019</xref>) and IMUs (<xref ref-type="bibr" rid="B21">Shah et al., 2022</xref>) show promise, though more comprehensive studies (e.g., <xref ref-type="bibr" rid="B24">Zhang et al., 2011</xref>), and publicly available datasets remain scarce. While databases hosting visual images or depth data are available, few if any exist in this domain that provide both modalities of data captured of a single scene (<xref ref-type="bibr" rid="B13">Laschowski et al., 2020</xref>). Further, none could be found that combine these modalities of data with simultaneous measurement of the device orientation sensors.</p>
<p>The proliferation of interpreted (stereo camera) or directly measured (LiDAR) depth data into a modern smartphone&#x2019;s sensor capabilities presents a new opportunity for terrain classification. Gyroscopes, magnetometer (compass), and accelerometer sensors are now included in even modest smartphones and allow capture of device orientation and inertia at the same instant as image and depth data. This synergy of visual, depth, and orientation data can potentially enhance environmental recognition accuracy.</p>
<p>This paper presents a method to capture a high-resolution, multi-modal dataset in real-time without expensive, professional grade equipment, and an novel dataset that can be used in the aforementioned domains. We expect that providing access to a dataset that hosts multiple modalities of simultaneously captured data will prove invaluable to many groups, from engineers developing high level control systems for exoskeletons, to researchers studying gait. Whether using visual, depth, or other modalities of data, terrain classification systems based on deep learning techniques require voluminous training data to produce models that are accurate and generalize well.</p>
<p>We therefore present <italic>L-AVATeD: the Lidar And Visual wAlking Terrain Dataset</italic>, a novel, open-source database of visual and LiDAR image pairs of human walking terrain with simultaneously captured device orientations.</p>
<p>In this data report, we provide a detailed description of the dataset, our hardware and methods for collecting the data, and post-processing steps taken to improve the utility and accessibility of the data. We conclude with suggestions for how other researchers may use this dataset. Analyses of the database for walking terrain classification will be presented in future work.</p>
</sec>
<sec sec-type="methods" id="s2">
<title>2 Methods</title>
<sec id="s2-1">
<title>2.1 Selection and definition of terrain types</title>
<p>To accurately reflect surfaces common in the built environment inside and surrounding typical North-American academic and healthcare institutions, nine terrain classes were chosen for this investigation: banked-left, banked-right, irregular, flat-even, grass, sloped-up, sloped-down, stairs-up and stairs-down.</p>
<p>While the class names were designed to be as descriptive as possible, some explanation is warranted.</p>
<p>The banked-[left, right] labels were applied in cases where the terrain declined significantly, perpendicular to the direction of motion. In cases where another class might also apply (e.g., grass, irregular), these labels were given precedence.</p>
<p>The flat-even class was a base-case; indicating any terrain that was generally smooth and solid, and neither sloped, banked nor grassy. Samples may include any material (e.g., concrete, tile) or color (there are many examples exhibiting bright colours and patterns).</p>
<p>Irregular surfaces were defined as those that had no slope up/down or left/right and had enough irregularity that they might be expected to materially affect gait. Examples include cobblestone, gravel, and rough mud.</p>
<p>Surfaces were labelled as <italic>grass</italic> if they were generally flat, similar to irregular in that they should be neither sloped nor banked, and consisted primarily of short grasses found on typical found in North American lawns.</p>
<p>sloped-[up,down] were defined as any surface (including grassy ones) which had a significant incline or decline in the immediate direction of motion.</p>
<p>stairs-[up,down] were the easiest to label, and consisted of stairs of any material, indoors or out.</p>
<p>Examples of visual and LiDAR image from each class can be seen in <xref ref-type="fig" rid="F1">Figure 1</xref>.</p>
<fig id="F1" position="float">
<label>FIGURE 1</label>
<caption>
<p>
<bold>(A)</bold> Data collection harness with iPhone sensor rig. <bold>(B)</bold> Approximate data capture field. <bold>(C)</bold> Example dataset image pairs by class. RGB (top), LiDAR (bottom) from left: banked-left, banked-right, flat-even, irregular, grass, sloped-down, sloped-up, stairs-down, stairs-up.</p>
</caption>
<graphic xlink:href="frobt-11-1384575-g001.tif"/>
</fig>
</sec>
<sec id="s2-2">
<title>2.2 Data-collection hardware</title>
<p>Apple&#x2019;s (Apple Inc., Cupertino, USA) iPhone (iOS) was chosen as the platform for collecting our walking terrain dataset due to their built-in LiDAR sensor as well as extensive, well-documented APIs for capturing and manipulating depth data. Data were captured on three physical devices (1 iPhone 12 Pro and two iPhone 14 Pros). While it is possible to extract depth data from mobile devices which interpret depth data using stereoscopic methods <xref ref-type="bibr" rid="B23">Wang et al. (2019)</xref>, we chose devices which contain a built-in LiDAR scanner, for accuracy and consistency.</p>
<p>Visual (RGB) images were captured in landscape orientation using the front-facing cameras of each of the iPhone 12 and 14. Data were captured at the native resolution of the sensor: 4,032<inline-formula id="inf1">
<mml:math id="m1">
<mml:mrow>
<mml:mo>&#xd7;</mml:mo>
</mml:mrow>
</mml:math>
</inline-formula> 3,024 pixels, RGB, 8-bit per channel and were stored using the Apple High Efficiency Image File Format (HEIF).</p>
<p>The specifications for the LiDAR scanner are not available directly from the manufacturer, but <xref ref-type="bibr" rid="B17">Luetzenburg et al. (2021)</xref> performed and in-depth analysis of the iPhone 12/14 LiDAR hardware (the devices share the same sensor) for applications in geosciences. Their conclusions coincide with the output LiDAR depth map dimensions we recorded, at 768 &#xd7; 576 pixels. In this dataset, each LiDAR &#x201c;pixel&#x201d; uses a full 32-bits to store depth data.</p>
<p>It should be noted that while differing device models were used to collect samples, both the visual and LiDAR data were identical in spatial and depth resolution.</p>
</sec>
<sec id="s2-3">
<title>2.3 Data collection</title>
<p>Data were collected over 6 months during the spring and summer of 2023, by three young healthy adult members of our research lab. Participants were fitted with a chest-mounted mobile-phone harness. The harness allowed for hands free data capture and provided some consistency in data capture across participants. iPhones were clipped to the harness horizontally (i.e., in &#x201c;landscape orientation&#x201d;) using the built-in mount, with an initial angle between 30&#x2013;50&#xb0; downward from the horizon. This orientation provided a wide view of the local walking surface within about a meter in front of the participant and up to roughly 5 m away, depending on the local terrain. The variation in mounting angle was similar in magnitude to the small up and down variation of camera field of view introduced by simple act of walking. This small amount of noise acts, in effect, as a natural regularizer for the dataset, and will help, e.g., a Convolutional Neural Network trained on it to better generalize on unseen data.</p>
<p>A custom iOS application was written to simplify simultaneous capture and labelling of visual, LiDAR, and device orientation data at 1 Hz. This capture frequency provided a balance between volume of data recorded while preventing too many captures of the same visual scene (i.e., walking terrain). In an individual capture session, participants would select the terrain type (based on their visual interpretation upcoming terrain and the definitions above) in the capture application, tap begin data capture, and terminate capture before the terrain class changed. The data would then be automatically labelled and stored on the device. Data were imported from each device into a central repository and individually reviewed (by D.W.) for labelling errors. All members who participated in data collection were briefed on use of the system prior to data collection.</p>
<p>Data were labelled in sequence, with a numeric prefix indicating the order of capture [000-999], and a unique suffix in the form of a universally unique identifier (UUID). Every image pair has the same file name, apart from an additional suffix &#x201c;_depth&#x201d; on the LiDAR disparity map and differing file extension. The gravity vector data was captured into individual comma-separated-values files, each with the matching UUID suffix.</p>
<p>7,968 RGB/LiDAR image pairs were captured along with the device orientation gravity vector, but due to availability of suitable terrain, class data is imbalanced (<xref ref-type="table" rid="T1">Table 1</xref>). Class imbalance, while not ideal, can be easily handled using any number of techniques when actually making use of the data. While training a Convolutional Neural Network, for example, oversampling of the minority class, adding class weights to the loss function, or using ensemble methods such as bagging and boosting are some common solutions.</p>
<table-wrap id="T1" position="float">
<label>TABLE 1</label>
<caption>
<p>Per-class characteristics of the Lidar And Visual wAlking Terrain Dataset.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="left">Class</th>
<th align="right">Images</th>
<th align="right">Proportion</th>
<th align="center">Min. Pitch</th>
<th align="center">Max. Pitch</th>
<th align="center">Min. Roll</th>
<th align="center">Max. Roll</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="left">Banked-left</td>
<td align="right">466</td>
<td align="right">5.8%</td>
<td align="center">43.2&#xb0;</td>
<td align="center">70.2&#xb0;</td>
<td align="center">&#x2212;11.4&#xb0;</td>
<td align="center">26.6&#xb0;</td>
</tr>
<tr>
<td align="left">Banked-right</td>
<td align="right">459</td>
<td align="right">5.8%</td>
<td align="center">43.1&#xb0;</td>
<td align="center">70.7&#xb0;</td>
<td align="center">&#x2212;23.3&#xb0;</td>
<td align="center">11.5&#xb0;</td>
</tr>
<tr>
<td align="left">Flat-level</td>
<td align="right">2,339</td>
<td align="right">29.4%</td>
<td align="center">16.4&#xb0;</td>
<td align="center">78.9&#xb0;</td>
<td align="center">&#x2212;18.9&#xb0;</td>
<td align="center">18.2&#xb0;</td>
</tr>
<tr>
<td align="left">Irregular</td>
<td align="right">920</td>
<td align="right">11.5%</td>
<td align="center">20.5&#xb0;</td>
<td align="center">73.1&#xb0;</td>
<td align="center">&#x2212;15.2&#xb0;</td>
<td align="center">16.1&#xb0;</td>
</tr>
<tr>
<td align="left">Grass</td>
<td align="right">1,217</td>
<td align="right">15.3%</td>
<td align="center">17.6&#xb0;</td>
<td align="center">72.0&#xb0;</td>
<td align="center">&#x2212;15.6&#xb0;</td>
<td align="center">17.0&#xb0;</td>
</tr>
<tr>
<td align="left">Sloped-down</td>
<td align="right">510</td>
<td align="right">6.4%</td>
<td align="center">25.6&#xb0;</td>
<td align="center">73.0&#xb0;</td>
<td align="center">&#x2212;19.4&#xb0;</td>
<td align="center">19.7&#xb0;</td>
</tr>
<tr>
<td align="left">Sloped-up</td>
<td align="right">428</td>
<td align="right">5.4%</td>
<td align="center">27.7&#xb0;</td>
<td align="center">68.4&#xb0;</td>
<td align="center">&#x2212;14.8&#xb0;</td>
<td align="center">12.4&#xb0;</td>
</tr>
<tr>
<td align="left">Stairs-down</td>
<td align="right">792</td>
<td align="right">9.9%</td>
<td align="center">22.3&#xb0;</td>
<td align="center">74.9&#xb0;</td>
<td align="center">&#x2212;28.8&#xb0;</td>
<td align="center">21.6&#xb0;</td>
</tr>
<tr>
<td align="left">Stairs-up</td>
<td align="right">837</td>
<td align="right">10.5%</td>
<td align="center">17.6&#xb0;</td>
<td align="center">86.7&#xb0;</td>
<td align="center">&#x2212;34.3&#xb0;</td>
<td align="center">19.6&#xb0;</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec id="s2-4">
<title>2.4 Post-processing</title>
<p>With the goal of making the dataset more manageable (the raw dataset is almost 45GBs in size), RGB images were pre-processed by reducing their size from 4032 &#xd7; 3024 pixels to 512 &#xd7; 384 pixels. Resizing used Apple&#x2019;s Scriptable Image Processing System (SIPS) and simultaneously converted images to a more standard JPEG format from their native HEIF format.</p>
<p>Processing the LiDAR data required special care, as it is captured in a format not readily consumed by typical image libraries. LiDAR data were captured and stored natively as 32-bit disparity maps (disparity &#x3d; 1/distance) and saved to TIFF files. These files were down sampled and normalized using OpenCV <xref ref-type="bibr" rid="B11">Itseez (2015)</xref>, converted to 8-bit grayscale, and exported to JPEG. At 768 &#xd7; 576 pixels, the spatial resolution of the built-in LiDAR scanner is much lower than the visual camera, and so these images were not pre-scaled.</p>
<p>A device orientation vector was also recorded at the instant both the visual and LiDAR data were captured. This orientation vector is recorded relative to gravity and the iPhone device axes. It is normalized in each direction in units of the accepted acceleration due to gravity at Earth&#x2019;s surface, i.e., <inline-formula id="inf2">
<mml:math id="m2">
<mml:mrow>
<mml:mn>9.8</mml:mn>
<mml:mtext>&#x2009;</mml:mtext>
<mml:mi>m</mml:mi>
<mml:mo>/</mml:mo>
<mml:msup>
<mml:mrow>
<mml:mi>s</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:msup>
</mml:mrow>
</mml:math>
</inline-formula>. These vectors are difficult to use in their raw <inline-formula id="inf3">
<mml:math id="m3">
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mi>x</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>y</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>z</mml:mi>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:math>
</inline-formula> vector format, and were translated to more understandable <italic>pitch</italic> and <italic>roll</italic> (<xref ref-type="disp-formula" rid="e1">Equations 1</xref>, <xref ref-type="disp-formula" rid="e2">2</xref> respectively) values using the <italic>two-element arc-tangent</italic> function:<disp-formula id="e1">
<mml:math id="m4">
<mml:mrow>
<mml:mi>p</mml:mi>
<mml:mi>i</mml:mi>
<mml:mi>t</mml:mi>
<mml:mi>c</mml:mi>
<mml:mi>h</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mi>a</mml:mi>
<mml:mi>t</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>n</mml:mi>
<mml:mn>2</mml:mn>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:mi>z</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>x</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:math>
<label>(1)</label>
</disp-formula>and,<disp-formula id="e2">
<mml:math id="m5">
<mml:mrow>
<mml:mi>r</mml:mi>
<mml:mi>o</mml:mi>
<mml:mi>l</mml:mi>
<mml:mi>l</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mi>a</mml:mi>
<mml:mi>t</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>n</mml:mi>
<mml:mn>2</mml:mn>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:mi>z</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>y</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:math>
<label>(2)</label>
</disp-formula>where pitch is defined as the angle down from the horizon, and roll as deviation either clockwise or anti-clockwise from the horizon line itself.</p>
</sec>
</sec>
<sec sec-type="discussion" id="s3">
<title>3 Discussion</title>
<p>A comprehensive understanding of the local walking context is crucial both for human gait analysis and the advanced control systems of robotic prostheses. Relying solely on laboratory analysis limits researchers&#x2019; capacity to accurately evaluate clinical gait in patients and hinders the optimal functioning of robotic prostheses in varied ecological contexts. Consequently, conducting studies across a spectrum of walking terrains is essential to address these limitations. A ground-truth understanding of walking terrain is traditionally identified using simple visual data and Deep Neural Networks, in particular Convolutional Neural Networks (CNNs). While CNNs are ideal for this classification task, they require extensive training data which may not be readily available.</p>
<p>Multiple studies ((<xref ref-type="bibr" rid="B19">Ophoff et al., 2019</xref>; <xref ref-type="bibr" rid="B18">Melotti et al., 2018</xref>; <xref ref-type="bibr" rid="B2">AlDahoul et al., 2021</xref>)) have revealed that classifiers trained using multi-modal data (and in particular a combination of depth and visual data) can outperform simple visual classification in object-detection and image classification (<xref ref-type="bibr" rid="B20">Schwarz et al., 2015</xref>). Previous datasets in the context of walking terrain have provided only a single modality of data (visual), or have been limited in their terrain classes. A novel dataset providing walking terrain data, simultaneously captured in both visual and depth modalities could therefore provide huge advantages to researchers training visual classifiers (in particular CNNs) for use in these domains. The L-AVATeD dataset fills this gap, improving on existing datasets by providing researchers and engineers a baseline set of multi-class, multi-sensor, walking terrain data in multiple modalities.</p>
<p>L-AVATeD is notable for being the only open-source database of its kind to provide not just two, but three modalities of data captured simultaneously for a given scene of walking terrain. While not as large as the largest available walking terrain datasets reviewed by <xref ref-type="bibr" rid="B13">Laschowski et al. (2020)</xref>, at almost 8,000 samples L-AVATeD matches the median dataset size. Further, the spatial resolution of the published RGB images in L-AVATeD (at 4032 &#xd7; 3024 pixels) is more than 9 times higher than the highest resolution presented. Images at resolutions of this magnitude may prove unwieldy for most neural networks, especially the low-resource-optimized architectures available for use in mobile and edge-computation hardware. It remains important however to preserve as much signal as possible to not limit future research, and so full-resolution images are provided.</p>
<p>Depth data, captured via LiDAR were in fact saved as <italic>disparity</italic> maps (i.e., <inline-formula id="inf4">
<mml:math id="m6">
<mml:mrow>
<mml:mn>1</mml:mn>
<mml:mo>/</mml:mo>
<mml:mi>d</mml:mi>
<mml:mi>i</mml:mi>
<mml:mi>s</mml:mi>
<mml:mi>t</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>n</mml:mi>
<mml:mi>c</mml:mi>
<mml:mi>e</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>). Notably, signal decreases as distance increases. At 768 &#xd7; 576 &#x201c;pixels&#x201d; of native spatial resolution, these too are the most information-dense depth measurements in any of the environment recognition systems reviewed in <xref ref-type="bibr" rid="B13">Laschowski et al. (2020)</xref>. To ensure we did not introduce any algorithmic bias to the depth data, the capture program was instructed to replace missing signal (depth &#x201c;pixels&#x201d;) with a placeholder <inline-formula id="inf5">
<mml:math id="m7">
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>N</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> value rather than automatically interpolating such points. This allows researchers to handle the missing data however they see fit, for example,: replacing with zeros, average depth, or an interpolation scheme. The raw depth maps are available as TIFF files. To make the depth data easily accessible, a copy is provided created using a very simple algorithm of replacing missing data with zeros, normalizing, and then compressing the depth dimension into 8 bits before saving the file as a JPEG. This format provides a representation suitable for use in most deep neural network software packages. It should be noted however that this simple algorithm is effectively a lossy compression of the data, and while useful, should not be considered as the most information-rich representation. Examples of these lossy depth-data representations can be seen in <xref ref-type="fig" rid="F1">Figure 1</xref> (bottom row).</p>
<p>Device orientation data were captured as a gravity vector. The raw data are available as comma-separated values, indexed with the sample UUID. This raw format is not easily interpreted by humans, and so each was converted to a more easily understood <italic>pitch</italic> (camera angle up and down relative to the horizon) and <italic>roll</italic> (camera rotation clockwise and anti-clockwise relative to the horizon). Per-class pitch and roll statistics are provided in <xref ref-type="table" rid="T1">Table 1</xref>. Preliminary <italic>post hoc</italic> analyses of these statistics do not immediately reveal significant signal, but used in combination with the RGB and LiDAR data counterparts, in a deep learning context in particular could prove fruitful. For example, researchers might use this data to &#x201c;de-rotate&#x201d; an image using the inverse device rotation angle before passing it into the classifier, potentially improving classifications where horizon angle of the scene may be important.</p>
<p>The number of samples in this dataset count almost two orders of magnitude smaller than the largest datasets available <xref ref-type="bibr" rid="B13">Laschowski et al. (2020)</xref>. Its usefulness however lies not in its number of samples, but in the diversity of information in those samples. The fusion of multi-modal sensor data has been shown to enhance task accuracy in many domains <xref ref-type="bibr" rid="B7">El Madawi et al. (2019)</xref>, <xref ref-type="bibr" rid="B2">AlDahoul et al. (2021)</xref>, <xref ref-type="bibr" rid="B8">Gao et al. (2018)</xref>. This dataset in particular will be useful in training accurate control systems for robotic prostheses <xref ref-type="bibr" rid="B4">Diaz et al. (2018)</xref>, locomotion modes for wearable robotics <xref ref-type="bibr" rid="B16">Li et al. (2022)</xref>, and mobile gait analysis systems. More specifically, a deep neural network combining multiple CNNs (for visual and depth data) modulated by device pitch and roll values could be trained to accurately classify terrain in real time <xref ref-type="bibr" rid="B14">Laschowski et al. (2022)</xref>.</p>
<p>These data have been made available through IEEE DataPort, and users of the L-AVATeD are requested to reference this report.</p>
</sec>
</body>
<back>
<sec sec-type="data-availability" id="s4">
<title>Data availability statement</title>
<p>The datasets presented in this study can be found in online repositories. The names of the repository/repositories and accession number(s) can be found below: <ext-link ext-link-type="uri" xlink:href="https://ieee-dataport.org/documents/l-avated-lidar-and-visual-walking-terrain-dataset">https://ieee-dataport.org/documents/l-avated-lidar-and-visual-walking-terrain-dataset</ext-link>.</p>
</sec>
<sec sec-type="author-contributions" id="s5">
<title>Author contributions</title>
<p>DW: Conceptualization, Data curation, Formal Analysis, Investigation, Methodology, Project administration, Software, Writing&#x2013;original draft, Writing&#x2013;review and editing, Validation, Visualization. PI: Investigation, Project administration, Supervision, Writing&#x2013;review and editing, Methodology. PD: Conceptualization, Funding acquisition, Project administration, Resources, Supervision, Writing&#x2013;review and editing.</p>
</sec>
<sec sec-type="funding-information" id="s6">
<title>Funding</title>
<p>The author(s) declare that financial support was received for the research, authorship, and/or publication of this article. PD acknowledges support from the Natural Sciences and Engineering Research Council of Canada (NSERC) Discovery grant (RRGPIN-2022-04217 and the fonds de recherche Qu&#xe9;bec Sant&#xe9; (FRQS) research scholar award (Junior 1).</p>
</sec>
<sec sec-type="COI-statement" id="s7">
<title>Conflict of interest</title>
<p>Author DW was employed by Simulation Curriculum Corp during part of the study.</p>
<p>The remaining authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec sec-type="disclaimer" id="s8">
<title>Publisher&#x2019;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<ref-list>
<title>References</title>
<ref id="B1">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Adamczyk</surname>
<given-names>P. G.</given-names>
</name>
<name>
<surname>Harper</surname>
<given-names>S. E.</given-names>
</name>
<name>
<surname>Reiter</surname>
<given-names>A. J.</given-names>
</name>
<name>
<surname>Roembke</surname>
<given-names>R. A.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Nichols</surname>
<given-names>K. M.</given-names>
</name>
<etal/>
</person-group> (<year>2023</year>). <article-title>Wearable sensing for understanding and influencing human movement in ecological contexts</article-title>. <source>Curr. Opin. Biomed. Eng.</source> <volume>28</volume>, <fpage>100492</fpage>. <pub-id pub-id-type="doi">10.1016/j.cobme.2023.100492</pub-id>
</citation>
</ref>
<ref id="B2">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>AlDahoul</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Karim</surname>
<given-names>H. A.</given-names>
</name>
<name>
<surname>Momo</surname>
<given-names>M. A.</given-names>
</name>
</person-group> (<year>2021</year>). &#x201c;<article-title>Rgb-d based multimodal convolutional neural networks for spacecraft recognition</article-title>,&#x201d; in <source>2021 IEEE international conference on image processing challenges (ICIPC)</source>, <fpage>1</fpage>&#x2013;<lpage>5</lpage>. <pub-id pub-id-type="doi">10.1109/ICIPC53495.2021.9620192</pub-id>
</citation>
</ref>
<ref id="B3">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>[Dataset] Hu</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Rouse</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Hargrove</surname>
<given-names>L.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>Benchmark datasets for bilateral lower-limb neuromechanical signals from wearable sensors during unassisted locomotion in able-bodied individuals</article-title>. <source>Front. Robot. AI</source> <volume>5</volume>, <fpage>14</fpage>. <pub-id pub-id-type="doi">10.3389/frobt.2018.00014</pub-id>
</citation>
</ref>
<ref id="B4">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Diaz</surname>
<given-names>J. P.</given-names>
</name>
<name>
<surname>da Silva</surname>
<given-names>R. L.</given-names>
</name>
<name>
<surname>Zhong</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Huang</surname>
<given-names>H. H.</given-names>
</name>
<name>
<surname>Lobaton</surname>
<given-names>E.</given-names>
</name>
</person-group> (<year>2018</year>). &#x201c;<article-title>Visual terrain identification and surface inclination estimation for improving human locomotion with a lower-limb prosthetic</article-title>,&#x201d; in <source>2018 40th annual international conference of the IEEE engineering in medicine and biology society (EMBC)</source>, <fpage>1817</fpage>&#x2013;<lpage>1820</lpage>. <pub-id pub-id-type="doi">10.1109/EMBC.2018.8512614</pub-id>
</citation>
</ref>
<ref id="B5">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Dixon</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Sch&#xfc;tte</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Vanwanseele</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Jacobs</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Dennerlein</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Schiffman</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>Gait adaptations of older adults on an uneven brick surface can be predicted by age-related physiological changes in strength</article-title>. <source>Gait and Posture</source> <volume>61</volume>, <fpage>257</fpage>&#x2013;<lpage>262</lpage>. <pub-id pub-id-type="doi">10.1016/j.gaitpost.2018.01.027</pub-id>
</citation>
</ref>
<ref id="B6">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Dixon</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Shah</surname>
<given-names>V.</given-names>
</name>
<name>
<surname>Willmott</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Effects of outdoor walking surface and slope on hip and knee joint angles in the sagittal plane</article-title>. <source>Gait and Posture</source> <volume>90</volume>, <fpage>48</fpage>&#x2013;<lpage>49</lpage>. <pub-id pub-id-type="doi">10.1016/j.gaitpost.2021.09.025</pub-id>
</citation>
</ref>
<ref id="B7">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>El Madawi</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Rashed</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>El Sallab</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Nasr</surname>
<given-names>O.</given-names>
</name>
<name>
<surname>Kamel</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Yogamani</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2019</year>). &#x201c;<article-title>Rg and lidar fusion based 3d semantic segmentation for autonomous driving</article-title>,&#x201d; in <source>2019 IEEE intelligent transportation systems conference (ITSC)</source>, <fpage>7</fpage>&#x2013;<lpage>12</lpage>. <pub-id pub-id-type="doi">10.1109/ITSC.2019.8917447</pub-id>
</citation>
</ref>
<ref id="B8">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Gao</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Cheng</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Zhao</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>D.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>Object classification using conn-based fusion of vision and lidar in autonomous vehicle environment</article-title>. <source>IEEE Trans. Industrial Inf.</source> <volume>14</volume>, <fpage>4224</fpage>&#x2013;<lpage>4231</lpage>. <pub-id pub-id-type="doi">10.1109/TII.2018.2822828</pub-id>
</citation>
</ref>
<ref id="B9">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ippersiel</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Robbins</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Dixon</surname>
<given-names>P.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Lower-limb coordination and variability during gait: the effects of age and walking surface</article-title>. <source>Gait and Posture</source> <volume>85</volume>, <fpage>251</fpage>&#x2013;<lpage>257</lpage>. <pub-id pub-id-type="doi">10.1016/j.gaitpost.2021.02.009</pub-id>
</citation>
</ref>
<ref id="B10">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ippersiel</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Shah</surname>
<given-names>V.</given-names>
</name>
<name>
<surname>Dixon</surname>
<given-names>P.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>The impact of outdoor walking surfaces on lower-limb coordination and variability during gait in healthy adults</article-title>. <source>Gait and Posture</source> <volume>91</volume>, <fpage>7</fpage>&#x2013;<lpage>13</lpage>. <pub-id pub-id-type="doi">10.1016/j.gaitpost.2021.09.176</pub-id>
</citation>
</ref>
<ref id="B11">
<citation citation-type="web">
<collab>Itseez</collab> (<year>2015</year>). <article-title>Open source computer vision library</article-title>. <comment>Available at: <ext-link ext-link-type="uri" xlink:href="https://github.com/opencv/opencv">https://github.com/opencv/opencv</ext-link>.</comment>
</citation>
</ref>
<ref id="B12">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Krizhevsky</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Sutskever</surname>
<given-names>I.</given-names>
</name>
<name>
<surname>Hinton</surname>
<given-names>G.</given-names>
</name>
</person-group> (<year>2012</year>). <article-title>ImageNet classification with deep convolutional neural networks</article-title>. <source>Neural Inf. Process. Syst.</source> <volume>25</volume>. <pub-id pub-id-type="doi">10.1145/3065386</pub-id>
</citation>
</ref>
<ref id="B13">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Laschowski</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>McNally</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Wong</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>McPhee</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Exonet database: wearable camera images of human locomotion environments</article-title>. <source>Front. Robot. AI</source> <volume>7</volume>, <fpage>562061</fpage>. <pub-id pub-id-type="doi">10.3389/frobt.2020.562061</pub-id>
</citation>
</ref>
<ref id="B14">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Laschowski</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>McNally</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Wong</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>McPhee</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Environment classification for robotic leg prostheses and exoskeletons using deep convolutional neural networks</article-title>. <source>Front. Neurorobotics</source> <volume>15</volume>, <fpage>730965</fpage>. <pub-id pub-id-type="doi">10.3389/fnbot.2021.730965</pub-id>
</citation>
</ref>
<ref id="B15">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lay</surname>
<given-names>A. N.</given-names>
</name>
<name>
<surname>Hass</surname>
<given-names>C. J.</given-names>
</name>
<name>
<surname>Gregor</surname>
<given-names>R. J.</given-names>
</name>
</person-group> (<year>2006</year>). <article-title>The effects of sloped surfaces on locomotion: a kinematic and kinetic analysis</article-title>. <source>J. biomechanics</source> <volume>39</volume>, <fpage>1621</fpage>&#x2013;<lpage>1628</lpage>. <pub-id pub-id-type="doi">10.1016/j.jbiomech.2005.05.005</pub-id>
</citation>
</ref>
<ref id="B16">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Zhong</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Lobaton</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Huang</surname>
<given-names>H.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Fusion of human gaze and machine vision for predicting intended locomotion mode</article-title>. <source>IEEE Trans. Neural Syst. Rehabilitation Eng.</source> <volume>30</volume>, <fpage>1103</fpage>&#x2013;<lpage>1112</lpage>. <pub-id pub-id-type="doi">10.1109/TNSRE.2022.3168796</pub-id>
</citation>
</ref>
<ref id="B17">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Luetzenburg</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Kroon</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Bj&#xf8;rk</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Evaluation of the apple iphone 12 pro lidar for an application in geosciences</article-title>. <source>Sci. Rep.</source> <volume>11</volume>, <fpage>22221</fpage>. <pub-id pub-id-type="doi">10.1038/s41598-021-01763-9</pub-id>
</citation>
</ref>
<ref id="B18">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Melotti</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Premebida</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Gon&#xe7;alves</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Nunes</surname>
<given-names>U.</given-names>
</name>
<name>
<surname>Faria</surname>
<given-names>D.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>Multimodal conn pedestrian classification: a study on combining lidar and camera data</article-title>. <fpage>3138</fpage>&#x2013;<lpage>3143</lpage>. <pub-id pub-id-type="doi">10.1109/ITSC.2018.8569666</pub-id>
</citation>
</ref>
<ref id="B19">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ophoff</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Beeck</surname>
<given-names>K. V.</given-names>
</name>
<name>
<surname>Goedem&#xe9;</surname>
<given-names>T.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Exploring rgb&#x2b;depth fusion for real-time object detection</article-title>. <source>Sensors Basel, Switz.</source> <volume>19</volume>, <fpage>866</fpage>. <pub-id pub-id-type="doi">10.3390/s19040866</pub-id>
</citation>
</ref>
<ref id="B20">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Schwarz</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Schulz</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Behnke</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2015</year>). &#x201c;<article-title>Rg-d object recognition and pose estimation based on pre-trained convolutional neural network features</article-title>,&#x201d; in <source>2015 IEEE international conference on robotics and automation (ICRA)</source>, <fpage>1329</fpage>&#x2013;<lpage>1335</lpage>. <pub-id pub-id-type="doi">10.1109/ICRA.2015.7139363</pub-id>
</citation>
</ref>
<ref id="B21">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Shah</surname>
<given-names>V.</given-names>
</name>
<name>
<surname>Flood</surname>
<given-names>M. W.</given-names>
</name>
<name>
<surname>Grimm</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Dixon</surname>
<given-names>P. C.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Generalizability of deep learning models for predicting outdoor irregular walking surfaces</article-title>. <source>J. Biomechanics</source> <volume>139</volume>, <fpage>111159</fpage>. <pub-id pub-id-type="doi">10.1016/j.jbiomech.2022.111159</pub-id>
</citation>
</ref>
<ref id="B22">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Voloshina</surname>
<given-names>A. S.</given-names>
</name>
<name>
<surname>Ferris</surname>
<given-names>D. P.</given-names>
</name>
</person-group> (<year>2015</year>). <article-title>Biomechanics and energetics of running on uneven terrain</article-title>. <source>J. Exp. Biol.</source> <volume>218</volume>, <fpage>711</fpage>&#x2013;<lpage>719</lpage>. <pub-id pub-id-type="doi">10.1242/jeb.106518</pub-id>
</citation>
</ref>
<ref id="B23">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Wang</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Lai</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Huang</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>B. H.</given-names>
</name>
<name>
<surname>van der Maaten</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Campbell</surname>
<given-names>M.</given-names>
</name>
<etal/>
</person-group> (<year>2019</year>). &#x201c;<article-title>Anytime stereo image depth estimation on mobile devices</article-title>,&#x201d; in <source>2019 international conference on robotics and automation (ICRA)</source>, <fpage>5893</fpage>&#x2013;<lpage>5900</lpage>. <pub-id pub-id-type="doi">10.1109/ICRA.2019.8794003</pub-id>
</citation>
</ref>
<ref id="B24">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Zhang</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Fang</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Huang</surname>
<given-names>H.</given-names>
</name>
</person-group> (<year>2011</year>). &#x201c;<article-title>Preliminary design of a terrain recognition system</article-title>,&#x201d; in <source>2011 annual international conference of the</source> (<publisher-name>IEEE Engineering in Medicine and Biology Society</publisher-name>), <fpage>5452</fpage>&#x2013;<lpage>5455</lpage>. <pub-id pub-id-type="doi">10.1109/IEMBS.2011.6091391</pub-id>
</citation>
</ref>
<ref id="B25">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhang</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Fu</surname>
<given-names>C.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Directional pointnet: 3d environmental classification for wearable robotics</article-title>
</citation>
</ref>
</ref-list>
</back>
</article>