<?xml version="1.0" encoding="UTF-8" standalone="no"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance" article-type="research-article" dtd-version="2.3" xml:lang="EN">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Mar. Sci.</journal-id>
<journal-title>Frontiers in Marine Science</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Mar. Sci.</abbrev-journal-title>
<issn pub-type="epub">2296-7745</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="doi">10.3389/fmars.2022.870005</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Marine Science</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Content-Aware Segmentation of Objects Spanning a Large Size Range: Application to Plankton Images</article-title>
</title-group>
<contrib-group>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Pana&#xef;otis</surname><given-names>Thelma</given-names>
</name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="author-notes" rid="fn001"><sup>*</sup></xref>
<uri xlink:href="https://loop.frontiersin.org/people/1498153"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Caray&#x2013;Counil</surname><given-names>Louis</given-names>
</name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Woodward</surname><given-names>Ben</given-names>
</name>
<xref ref-type="aff" rid="aff2"><sup>2</sup></xref>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Schmid</surname><given-names>Moritz S.</given-names>
</name>
<xref ref-type="aff" rid="aff3"><sup>3</sup></xref>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Daprano</surname><given-names>Dominic</given-names>
</name>
<xref ref-type="aff" rid="aff4"><sup>4</sup></xref>
<uri xlink:href="https://loop.frontiersin.org/people/1666855"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Tsai</surname><given-names>Sheng Tse</given-names>
</name>
<xref ref-type="aff" rid="aff4"><sup>4</sup></xref>
<uri xlink:href="https://loop.frontiersin.org/people/1726340"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Sullivan</surname><given-names>Christopher M.</given-names>
</name>
<xref ref-type="aff" rid="aff4"><sup>4</sup></xref>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Cowen</surname><given-names>Robert K.</given-names>
</name>
<xref ref-type="aff" rid="aff3"><sup>3</sup></xref>
<uri xlink:href="https://loop.frontiersin.org/people/686608"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Irisson</surname><given-names>Jean-Olivier</given-names>
</name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<uri xlink:href="https://loop.frontiersin.org/people/581061"/>
</contrib>
</contrib-group>
<aff id="aff1"><sup>1</sup><institution>Laboratoire d&#x2019;Oc&#xe9;anographie de Villefranche, Sorbonne Universit&#xe9;</institution>, <addr-line>Villefranche-sur-Mer</addr-line>, <country>France</country></aff>
<aff id="aff2"><sup>2</sup><institution>CVision AI</institution>, <addr-line>Medford, MA</addr-line>, <country>United States</country></aff>
<aff id="aff3"><sup>3</sup><institution>Hatfield Marine Science Center, Oregon State University</institution>, <addr-line>Newport, OR</addr-line>, <country>United States</country></aff>
<aff id="aff4"><sup>4</sup><institution>Center for Quantitative and Life Science, Oregon State University</institution>, <addr-line>Corvallis, OR</addr-line>, <country>United States</country></aff>
<author-notes>
<fn fn-type="edited-by">
<p>Edited by: Mark C. Benfield, Louisiana State University, United States</p>
</fn>
<fn fn-type="edited-by">
<p>Reviewed by: Damianos Chatzievangelou, Self-employed, Greece; Francesco Pomati, Swiss Federal Institute of Aquatic Science and Technology, Switzerland</p>
</fn>
<fn fn-type="corresp" id="fn001">
<p>*Correspondence: Thelma Pana&#xef;otis, <email xlink:href="mailto:thelma.panaiotis@imev-mer.fr">thelma.panaiotis@imev-mer.fr</email></p>
</fn>
<fn fn-type="other" id="fn002">
<p>This article was submitted to Ocean Observation, a section of the journal Frontiers in Marine Science</p>
</fn>
</author-notes>
<pub-date pub-type="epub">
<day>06</day>
<month>06</month>
<year>2022</year>
</pub-date>
<pub-date pub-type="collection">
<year>2022</year>
</pub-date>
<volume>9</volume>
<elocation-id>870005</elocation-id>
<history>
<date date-type="received">
<day>05</day>
<month>02</month>
<year>2022</year>
</date>
<date date-type="accepted">
<day>04</day>
<month>05</month>
<year>2022</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#xa9; 2022 Pana&#xef;otis, Caray&#x2013;Counil, Woodward, Schmid, Daprano, Tsai, Sullivan, Cowen and Irisson</copyright-statement>
<copyright-year>2022</copyright-year>
<copyright-holder>Pana&#xef;otis, Caray&#x2013;Counil, Woodward, Schmid, Daprano, Tsai, Sullivan, Cowen and Irisson</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<abstract>
<p>As the basis of oceanic food webs and a key component of the biological carbon pump, planktonic organisms play major roles in the oceans. Their study benefited from the development of <italic>in situ</italic> imaging instruments, which provide higher spatio-temporal resolution than previous tools. But these instruments collect huge quantities of images, the vast majority of which are of marine snow particles or imaging artifacts. Among them, the <italic>In Situ</italic> Ichthyoplankton Imaging System (ISIIS) samples the largest water volumes (&gt; 100 L s<sup>-1</sup>) and thus produces particularly large datasets. To extract manageable amounts of ecological information from <italic>in situ</italic> images, we propose to focus on planktonic organisms early in the data processing pipeline: at the segmentation stage. We compared three segmentation methods, particularly for smaller targets, in which plankton represents less than 1% of the objects: (i) a traditional thresholding over the background, (ii) an object detector based on maximally stable extremal regions (MSER), and (iii) a content-aware object detector, based on a Convolutional Neural Network (CNN). These methods were assessed on a subset of ISIIS data collected in the Mediterranean Sea, from which a ground truth dataset of &gt; 3,000 manually delineated organisms is extracted. The naive thresholding method captured 97.3% of those but produced ~340,000 segments, 99.1% of which were therefore not plankton (i.e. recall = 97.3%, precision = 0.9%). Combining thresholding with a CNN missed a few more planktonic organisms (recall = 91.8%) but the number of segments decreased 18-fold (precision increased to 16.3%). The MSER detector produced four times fewer segments than thresholding (precision = 3.5%), missed more organisms (recall = 85.4%), but was considerably faster. Because naive thresholding produces ~525,000 objects from 1 minute of ISIIS deployment, the more advanced segmentation methods significantly improve ISIIS data handling and ease the subsequent taxonomic classification of segmented objects. The&#xa0;cost in terms of recall is limited, particularly for the CNN object detector. These approaches are now standard in computer vision and could be applicable to other plankton imaging devices, the majority of which pose a data management problem.</p>
</abstract>
<kwd-group>
<kwd>plankton images</kwd>
<kwd>ISIIS</kwd>
<kwd>image processing</kwd>
<kwd>image segmentation</kwd>
<kwd>object detection</kwd>
<kwd>convolutional neural network</kwd>
<kwd>computer vision</kwd>
</kwd-group>
<counts>
<fig-count count="6"/>
<table-count count="3"/>
<equation-count count="1"/>
<ref-count count="77"/>
<page-count count="16"/>
<word-count count="8773"/>
</counts>
</article-meta>
</front>
<body>
<sec id="s1" sec-type="intro">
<title>1 Introduction</title>
<sec id="s1_1">
<title>1.1. Plankton Imaging Enables Fine Scale Studies</title>
<p>Planktonic organisms play crucial roles in the ocean: photosynthetic phytoplankton is responsible for about half of the primary production of the biosphere (<xref ref-type="bibr" rid="B22">Field et&#xa0;al., 1998</xref>) and is the basis of oceanic food webs (<xref ref-type="bibr" rid="B21">Falkowski, 2012</xref>); zooplankton acts as a trophic link between phytoplankton and higher trophic levels (<xref ref-type="bibr" rid="B76">Ware and Thomson, 2005</xref>; <xref ref-type="bibr" rid="B24">Frederiksen et&#xa0;al., 2006</xref>) and is a key component of the biological carbon pump, sequestering organic carbon at depth (<xref ref-type="bibr" rid="B43">Longhurst and Glen Harrison, 1989</xref>). Plankton comprises organisms from very diverse taxonomic groups (<xref ref-type="bibr" rid="B17">de Vargas et&#xa0;al., 2015</xref>) that span from micrometer scale picoplankton to meter-long Cnidarians (<xref ref-type="bibr" rid="B42">Lombard et&#xa0;al., 2019</xref>). Given this very wide size range, plankton sampling instruments cannot tackle all organisms at once and typically target a reduced size range instead (<xref ref-type="bibr" rid="B42">Lombard et&#xa0;al., 2019</xref>).</p>
<p>The power law underlying plankton or marine snow particle size spectra means that concentration drastically increases when size decreases: the relationship is linear in log-log form (<xref ref-type="bibr" rid="B65">Sheldon and Parsons, 1967</xref>; <xref ref-type="bibr" rid="B66">Sheldon et&#xa0;al., 1972</xref>; <xref ref-type="bibr" rid="B69">Stemmann and Boss, 2012</xref>; <xref ref-type="bibr" rid="B42">Lombard et&#xa0;al., 2019</xref>). The larger organisms, which each contribute significantly to biomass, are rare but easy to detect. Yet, it is critical to also focus on the smaller objects, to avoid artificially cutting the effective size range of any instrument, thus potentially discarding the most numerous objects in the sample (<xref ref-type="bibr" rid="B42">Lombard et&#xa0;al., 2019</xref>). Moreover, as marine snow particles cannot grow past a few centimeters because of disaggregation (<xref ref-type="bibr" rid="B2">Alldredge and Silver, 1988</xref>; <xref ref-type="bibr" rid="B1">Alldredge et&#xa0;al., 1990</xref>), the ratio of particles to plankton also decreases with increasing size. Therefore, while targeting small planktonic organisms is desirable, it comes with the difficulty of separating them from the largely dominant particles within the same size range.</p>
<p>While large scale plankton distribution patterns are resolved to a certain extent (<xref ref-type="bibr" rid="B60">Rutherford et&#xa0;al., 1999</xref>; <xref ref-type="bibr" rid="B59">Rombouts et&#xa0;al., 2009</xref>; <xref ref-type="bibr" rid="B73">Tittensor et&#xa0;al., 2010</xref>; <xref ref-type="bibr" rid="B36">Ibarbalz et&#xa0;al., 2019</xref>; <xref ref-type="bibr" rid="B10">Brand&#xe3;o et&#xa0;al., 2021</xref>), much remains to be discovered regarding fine scale distribution, in particular for zooplankton. For phytoplankton, submesoscale dynamics are known to influence their distribution and concentration: vertical currents may affect nutrient and cell distribution relative to the euphotic zone, thus affecting growth rate, horizontal currents can stir patches into filaments. These changes are expected to propagate to higher trophic levels (zooplankton, fish, etc.) (<xref ref-type="bibr" rid="B40">L&#xe9;vy et&#xa0;al., 2018</xref>). Indeed, the trophic and reproductive interactions of zooplankton occur at the scale of organisms (&#xb5;m to cm). Therefore, a local concentration of phytoplankton, in a thin layer for example, has more immediate consequences on the survival and development of zooplanktonic grazers than the average chlorophyll <italic>a</italic> concentration in the region. Thus, studying zooplankton distribution at fine scales, in relation with submesoscale dynamics, becomes relevant to understand the processes driving its distribution at regional scale.</p>
<p>Our lack of knowledge regarding the fine scale distribution of plankton partly stems from the difficulty to adequately sample it at such a small scale. Traditional plankton collection methods such as pumps, nets, and bottles typically integrate organisms over some vertical and/or horizontal distance and make it difficult to associate organism concentrations with their immediate environmental context (<xref ref-type="bibr" rid="B57">Remsen et al, 2004</xref>; <xref ref-type="bibr" rid="B6">Benfield et&#xa0;al., 2007</xref>; <xref ref-type="bibr" rid="B42">Lombard et&#xa0;al., 2019</xref>). Moreover, most damage fragile organisms and fail to sample some of them properly (<xref ref-type="bibr" rid="B57">Remsen et&#xa0;al., 2004</xref>).</p>
<p>As an alternative, <italic>in situ</italic> pelagic imaging instruments such as the Imaging FlowCytoBot (IFCB) (<xref ref-type="bibr" rid="B50">Olson and Sosik, 2007</xref>), the <italic>In Situ</italic> Ichthyoplankton Imaging System (ISIIS) (<xref ref-type="bibr" rid="B14">Cowen and Guigand, 2008</xref>), the Underwater Vision Profiler (UVP) (<xref ref-type="bibr" rid="B56">Picheral et&#xa0;al., 2010</xref>), and the Scripps Plankton Camera (SPC) (<xref ref-type="bibr" rid="B52">Orenstein et&#xa0;al., 2020</xref>) (see <xref ref-type="bibr" rid="B42">Lombard et&#xa0;al. (2019)</xref> for a detailed list) allow studying plankton distribution at all scales: from the fine ones they resolve in each sample to long time scales and global spatial coverage through the accumulation of individual samples (<xref ref-type="bibr" rid="B70">Stemmann et&#xa0;al., 2008</xref>; <xref ref-type="bibr" rid="B23">Forest et&#xa0;al., 2012</xref>; <xref ref-type="bibr" rid="B58">Robinson et&#xa0;al., 2021</xref>; <xref ref-type="bibr" rid="B37">Irisson et&#xa0;al., 2022</xref>). As a non-destructive sampling approach, these instruments allow investigating fragile planktonic objects, such as Rhizaria (<xref ref-type="bibr" rid="B16">Dennett et&#xa0;al., 2002</xref>; <xref ref-type="bibr" rid="B8">Biard et&#xa0;al., 2016</xref>; <xref ref-type="bibr" rid="B7">Biard and Ohman, 2020</xref>), Cnidaria and Ctenophora (<xref ref-type="bibr" rid="B44">Luo et&#xa0;al., 2014</xref>), or marine snow aggregates (<xref ref-type="bibr" rid="B33">Guidi et&#xa0;al., 2008</xref>; <xref ref-type="bibr" rid="B34">Guidi et&#xa0;al., 2015</xref>). Still, <italic>in situ</italic> imaging systems typically sample smaller volumes than plankton nets (<xref ref-type="bibr" rid="B42">Lombard et&#xa0;al., 2019</xref>), limiting their quantitative application to abundant taxa. To quantify rarer planktonic groups, sampling effort has to be increased to improve the chances of detection. For example, the ISIIS was initially developed with a very high sampling volume to study the very sparsely distributed fish larvae. Because of this, all <italic>in situ</italic> imaging instruments collect vast amounts of data, although the acquisition rate varies from one instrument to the next. ISIIS, for instance, collects up to 11 million objects per hour of sampling, while IFCB collects images at a rate of ~10,000 per hour (<xref ref-type="bibr" rid="B68">Sosik and Olson, 2007</xref>). Thus all these systems need efficient and automated data processing approaches, albeit with different stringency.</p>
<p>In addition, high resolution sampling is required to tackle questions that used to be out of reach, such as fine-scale plankton distribution in relation with environmental conditions (<xref ref-type="bibr" rid="B47">McClatchie et&#xa0;al., 2012</xref>; <xref ref-type="bibr" rid="B29">Greer et&#xa0;al., 2015</xref>; <xref ref-type="bibr" rid="B11">Brise&#xf1;o-Avena et&#xa0;al., 2020</xref>), plankton patch structure (<xref ref-type="bibr" rid="B58">Robinson et&#xa0;al., 2021</xref>), interactions between zooplankton and phytoplankton fine layers (<xref ref-type="bibr" rid="B31">Greer et&#xa0;al., 2013</xref>; <xref ref-type="bibr" rid="B26">Greer et&#xa0;al., 2020a</xref>; <xref ref-type="bibr" rid="B63">Schmid and Fortiers, 2019</xref>) or co-occurrences revealing biological interactions such as predation (<xref ref-type="bibr" rid="B30">Greer et&#xa0;al., 2014</xref>; <xref ref-type="bibr" rid="B61">Schmid et&#xa0;al., 2020</xref>; <xref ref-type="bibr" rid="B72">Swieca et&#xa0;al., 2020</xref>; <xref ref-type="bibr" rid="B28">Greer et&#xa0;al., 2021</xref>).</p>
</sec>
<sec id="s1_2">
<title>1.2. Objects Need to be Extracted Automatically From Pelagic Images</title>
<p>The first data processing step is separating relevant organisms and particles from the background in raw images, i.e. image segmentation. Various segmentation methods have been applied for images collected by commonly used <italic>in situ</italic> imaging devices: the UVP relies on a fixed gray level threshold (<xref ref-type="bibr" rid="B56">Picheral et&#xa0;al., 2010</xref>), the IFCB uses an algorithm based on edge detection (<xref ref-type="bibr" rid="B50">Olson and Sosik, 2007</xref>), the SPC (<xref ref-type="bibr" rid="B52">Orenstein et&#xa0;al., 2020</xref>) runs a canny edge detector to initialize the segmentation of its dark-field microscopy images. To segment images generated by the Zooglider, a glider equipped with a shadowgraph, <xref ref-type="bibr" rid="B49">Ohman et&#xa0;al. (2019)</xref> also applied a canny edge detector. Finally, to segment shadowgrams from the ISIIS, <xref ref-type="bibr" rid="B74">Tsechpenakis et&#xa0;al. (2007)</xref> and <xref ref-type="bibr" rid="B38">Iyer, (2012)</xref> used statistical modeling of the background of the image and identified anomalies over this background as objects of interest.</p>
<p>The ISIIS is deployed in an undulating manner, between the surface and a given depth (<xref ref-type="bibr" rid="B14">Cowen and Guigand, 2008</xref>). It targets organisms in the range 250 &#xb5;m - 10&#xa0;cm. Together with grayscale images, it continually records environmental variables (temperature, salinity, fluorescence, dissolved oxygen and irradiance). The use of shadowgraphy combined with a specific lens and lighting system provide a large depth of field and allow a high sampling rate (28 kHz line scan camera). Therefore, the ISIIS is capable of sampling volumes of waters larger than all other <italic>in situ</italic> imaging instruments [&gt; 100 L s<sup>-1</sup>; <xref ref-type="bibr" rid="B42">Lombard et&#xa0;al. (2019)</xref>]. This optical design also ensures that the organism&#x2019;s size is not affected by its position within the depth of field. Shadowgraphs are also able to detect heterogeneities in the medium that is traversed by the light, which makes them excellent to image transparent organisms such as plankton, gelatinous organisms in particular. But it also makes them sensitive to other sources of heterogeneity, such as suspended particles or water density changes. ISIIS may thus generate noisy images when deployed in turbid waters (<xref ref-type="bibr" rid="B45">Luo et&#xa0;al., 2018</xref>; <xref ref-type="bibr" rid="B27">Greer et&#xa0;al., 2018</xref>) or across strong density gradients (<xref ref-type="fig" rid="f1"><bold>Figures&#xa0;1D&#x2013;F</bold></xref>) (<xref ref-type="bibr" rid="B20">Faillettaz et&#xa0;al., 2016</xref>). Furthermore, the use of a line scan camera means that marks or dust on the lens cause continuous streaks in the generated images (the line continuously scans the same speckle; <xref ref-type="fig" rid="f1"><bold>Figures&#xa0;1A, D</bold></xref>). Those can be partially removed by applying a flat-fielding procedure, whereby the average gray value computed per row over a few thousand scanned lines is subtracted from the incoming new values (<xref ref-type="fig" rid="f1"><bold>Figures&#xa0;1B, E</bold></xref>) (<xref ref-type="bibr" rid="B20">Faillettaz et&#xa0;al., 2016</xref>; <xref ref-type="bibr" rid="B45">Luo et&#xa0;al., 2018</xref>; <xref ref-type="bibr" rid="B27">Greer et&#xa0;al., 2018</xref>).</p>
<fig id="f1" position="float">
<label>Figure&#xa0;1</label>
<caption>
<p>ISIIS frames in clean waters <bold>(A&#x2013;C)</bold> and across a density change <bold>(D&#x2013;F)</bold>. The signature of this density change is similar to what a shadowgraph would image in air, above a burning candle. The panels are: <bold>(A, D)</bold> raw output; <bold>(B, E)</bold> after flat-fielding; <bold>(C, F)</bold> after contrasting. The camera scans vertically and the image is acquired from the right edge, as ISIIS moves through the water. In panel <bold>(A)</bold>, the scale bar represents 1&#xa0;cm and is applicable to other panels.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fmars-09-870005-g001.tif"/>
</fig>
<p>The very characteristics that give the ISIIS its qualities as a plankton imager (large sampling volume, high speed, ability to detect transparent objects) also mean that it creates a huge amount of images, the background of which is often non-uniform. This makes segmentation of planktonic objects from raw images far from trivial. To perform this segmentation, the processing pipeline was initially based on anomalies from a gaussian mixture model of the background gray levels without flat-fielding (<xref ref-type="bibr" rid="B74">Tsechpenakis et&#xa0;al., 2007</xref>) and later on k-harmonic means clustering on flat-fielded images (<xref ref-type="bibr" rid="B38">Iyer, 2012</xref>). This latter method was used in several studies (<xref ref-type="bibr" rid="B45">Luo et&#xa0;al., 2018</xref>; <xref ref-type="bibr" rid="B27">Greer et al., 2018</xref>; <xref ref-type="bibr" rid="B61">Schmid et&#xa0;al., 2020</xref>) and the full pipeline was open sourced in order to make plankton imaging more accessible and lower entry barriers (<xref ref-type="bibr" rid="B62">Schmid et&#xa0;al., 2021</xref>). Other studies relied on flat-fielding followed by segmentation above a fixed gray level (<xref ref-type="bibr" rid="B20">Faillettaz et&#xa0;al., 2016</xref>; <xref ref-type="bibr" rid="B26">Greer et&#xa0;al., 2020a</xref>; <xref ref-type="bibr" rid="B32">Greer et&#xa0;al., 2020b</xref>). However, most of these studies focused on the larger end of size range targeted by the ISIIS, by considering only objects above a given size threshold (<xref ref-type="table" rid="T1"><bold>Table&#xa0;1</bold></xref>), often because those were desirable targets, not noise. Similarly, for their canny edge detector applied to ZooGlider images, <xref ref-type="bibr" rid="B49">Ohman et&#xa0;al. (2019)</xref> considered objects larger than 100 pixels (Equivalent Spherical Diameter, or ESD of 0.45&#xa0;mm). However, the algorithm failed when too many particles were present and had to fall back to a less sensitive (i.e. higher) gray threshold. As shown above, both planktonic organisms and particles are much more abundant towards the smaller end of the spectrum, meaning that such methods had to ignore a non-negligible part of planktonic organisms and marine snow in order to discard the background noise.</p>
<table-wrap id="T1" position="float">
<label>Table&#xa0;1</label>
<caption>
<p>Threshold in object area (number of pixels considered as part of the object) in studies exploiting ISIIS data.</p>
</caption>
<table frame="hsides">
<thead>
<tr>
<th valign="top" align="left">Reference </th>
<th valign="top" align="center">Area threshold (px) </th>
<th valign="top" align="center">ESD (mm) </th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="left">
<xref ref-type="bibr" rid="B61">Schmid et&#xa0;al. (2020)</xref>
</td>
<td valign="top" align="center">7</td>
<td valign="top" align="center">0.2</td>
</tr>
<tr>
<td valign="top" align="left">
<xref ref-type="bibr" rid="B45">Luo et&#xa0;al. (2018)</xref>
</td>
<td valign="top" align="center">50</td>
<td valign="top" align="center">0.53</td>
</tr>
<tr>
<td valign="top" align="left">
<xref ref-type="bibr" rid="B20">Faillettaz et&#xa0;al. (2016)</xref>
</td>
<td valign="top" align="center">250</td>
<td valign="top" align="center">0.92</td>
</tr>
<tr>
<td valign="top" align="left">
<xref ref-type="bibr" rid="B26">Greet er&#xa0;al. (2020b)</xref>
</td>
<td valign="top" align="center">400</td>
<td valign="top" align="center">0.95</td>
</tr>
<tr>
<td valign="top" align="left">
<xref ref-type="bibr" rid="B26">Greer et&#xa0;al. (2020a)</xref>
</td>
<td valign="top" align="center">900</td>
<td valign="top" align="center">1.4</td>
</tr>
<tr>
<td valign="top" align="left">
<xref ref-type="bibr" rid="B28">Greer et&#xa0;al. (2021)</xref>
</td>
<td valign="top" align="center">2000</td>
<td valign="top" align="center">3.0</td>
</tr>
<tr>
<td valign="top" align="left">
<xref ref-type="bibr" rid="B27">Greer et&#xa0;al. (2018)</xref>
</td>
<td valign="top" align="center">5000</td>
<td valign="top" align="center">5.4</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<fn id="fnT1_1">
<label>a</label>
<p>The conversion factor from area (px) to Equivalent Spherical Diameter (ESD, mm) depends on the ISIIS configuration.</p>
</fn>
</table-wrap-foot>
</table-wrap>
</sec>
<sec id="s1_3">
<title>1.3. Marine Snow and Imaging Artifacts Dominate <italic>In Situ</italic> Images and Complicate Plankton Detection</title>
<p>Marine snow particles are much more abundant than plankton in the ocean (<xref ref-type="bibr" rid="B42">Lombard et&#xa0;al., 2019</xref>), which means that the vast majority (often &gt; 85%) of images captured by <italic>in situ</italic> plankton imaging instruments are actually of various marine snow items (fecal pellets, large aggregates, small organism pieces, etc.; (<xref ref-type="bibr" rid="B71">Stemmann et&#xa0;al., 2000</xref>; <xref ref-type="bibr" rid="B56">Picheral et&#xa0;al., 2010</xref>; <xref ref-type="bibr" rid="B69">Stemmann and Boss, 2012</xref>)). Therefore, for plankton ecology studies, the bottleneck has often become the processing and filtering of collected images (<xref ref-type="bibr" rid="B37">Irisson et&#xa0;al., 2022</xref>). To reduce the proportion of detrital particles and focus on photosynthetic plankton, the IFCB and the FlowCam can use fluorescence image triggering, hence imaging only items that contain chlorophyll (<xref ref-type="bibr" rid="B67">Sieracki et&#xa0;al., 1998</xref>; <xref ref-type="bibr" rid="B68">Sosik and Olson, 2007</xref>). This is not possible over the large volumes and for the non-photosynthetic organisms that ISIIS or other zooplankton imagers target. Furthermore, density anomalies lead to the characteristically noisy shadowgrams presented above (<xref ref-type="fig" rid="f1"><bold>Figures&#xa0;1D&#x2013;F</bold></xref>), from which numerous artifactual &#x201c;particles&#x201d; are detected by the usual image processing pipelines. Those artifacts or noise, together with marine snow, can constitute 99% of the objects detected. Such an extreme class imbalance makes the automatic classification of these objects through machine learning a very arduous task (<xref ref-type="bibr" rid="B39">Lee et&#xa0;al., 2016</xref>).</p>
<p>Even for a trained human operator, the differentiation of some planktonic classes from the proteiform marine snow aggregates and noise, as well as distinction between marine snow and noise themselves, can be very challenging. Towards the smaller end of the size spectrum it becomes virtually impossible. Indeed, once these small objects are segmented out, the low pixel count combined with the lack of information regarding their context in the image makes their identification very difficult, for humans and computers alike (<xref ref-type="bibr" rid="B54">Parikh et&#xa0;al., 2012</xref>). Hence, one solution could be to focus solely on planktonic organisms from the segmentation step already and try to avoid segmenting non-planktonic objects, thanks to their broader context in the image, still accessible at this step. This should result in a much more manageable amount of data to classify and a lesser class imbalance. This approach requires the development of specific and &#x201c;intelligent&#x201d; segmentation methods that target specific objects only. The purpose of this work was (i) to develop such &#x201c;intelligent&#x201d; segmentation approaches and (ii) to compare them with classic methods to test whether they significantly improve the data processing pipeline. With this in mind, we benchmarked three segmentation methods against a ground-truth human segmentation using a dataset collected by the ISIIS in the North-Western Mediterranean Sea.</p>
</sec>
</sec>
<sec id="s2" sec-type="materials|methods">
<title>2 Materials and Methods</title>
<sec id="s2_1">
<title>2.1. Image Segmentation Methods</title>
<sec id="s2_1_1">
<title>2.1.1. Threshold-Based Segmentation</title>
<p>The simplest segmentation method is to threshold pixels below a given gray level: adjoining pixels darker than the threshold are considered as segments. This threshold can be a value fixed <italic>a priori</italic> or dynamically computed from the properties of each image. For example, the classic method of <xref ref-type="bibr" rid="B53">Otsu (1979)</xref> is to examine the histogram of intensity levels and define the threshold so that it separates pixels into two relatively homogeneous intensity classes. Here either a fixed threshold was set or the threshold was defined based on a quantile of the histogram of gray levels. This quantile-based approach resulted in a darker segmentation threshold on noisy images, such as those captured around the strong density gradient induced by the thermocline (<xref ref-type="fig" rid="f1"><bold>Figures&#xa0;1D&#x2013;F</bold></xref>), which were richer in dark pixels. It was well adapted to limit the number of artifact segments generated from these images. Moreover, the first quartile is barely affected by the presence of relatively large dark objects such as jellyfish tentacles, making the segmentation threshold robust to these natural occurrences. After thresholding, segments defined by connected components were dilated by 3 pixels and eroded by 2 pixels to fill potential holes in transparent organisms and reconnect thin appendages to the organisms bodies. Finally, only segments larger than 50 pixels (400 &#xb5;m in ESD) were retained, because it was the minimum size at which taxonomists could recognise organisms.</p>
</sec>
<sec id="s2_1_2">
<title>2.1.2. Threshold-MSER (T-MSER) Segmentation</title>
<p>This approach uses a signal-to-noise ratio (SNR) cutoff, calculated on images after flat-fielding, to determine whether the frame should be segmented using a Maximally Stable Extremal Region approach (MSER, <xref ref-type="bibr" rid="B46">Matas et&#xa0;al. (2004)</xref>), or if areas of high noise should first be filtered out using a naive thresholding approach before applying MSER. MSER was successfully applied to the segmentation of ZOOVIS imagery (<xref ref-type="bibr" rid="B9">Bi et&#xa0;al., 2015</xref>; <xref ref-type="bibr" rid="B13">Cheng et&#xa0;al., 2019</xref>). SNR can be used to determine the relative noise level in an image and was computed as<disp-formula>
<label>(1)</label>
<mml:math display="block" id="M1">
<mml:mrow>
<mml:mtext>SNR</mml:mtext>
<mml:mo>=</mml:mo>
<mml:mn>20</mml:mn>
<mml:mo>&#xd7;</mml:mo>
<mml:mtext>log&#xa0;</mml:mtext>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:mfrac>
<mml:mi>S</mml:mi>
<mml:mi>N</mml:mi>
</mml:mfrac>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</disp-formula>
where S is the signal, defined as the mean of the input data, and N is the noise, computed as the standard deviation around that mean. Here, flat-fielded frames with low SNR (i.e. high noise) were binarized using a fixed thresholding in order to extract continuous regions of interest with darker pixel values. The regions identified in this way were then extracted using a mask and subsequently re-segmented using the MSER approach. MSER detects stable connected regions in images, which are areas that stay nearly unchanged over a wide range of grayscale thresholds. MSER can be tuned to allow for varying degrees of stable region area and the range of pixel gray values tested in the dynamic thresholding. High SNR frames are directly segmented using the MSER approach (<xref ref-type="fig" rid="f2"><bold>Figure&#xa0;2</bold></xref> skip from step B to step D). Going from a pure MSER approach to the threshold+MSER (T-MSER) on low SNR (&lt; 50) frames increased the recall on the test data from 65% to 85%, while also substantially increasing precision. This SNR and MSER method is written in C++17. The OpenCV and OpenMP Python packages were used for general computer vision and parallel processing for high processing efficiency, respectively.</p>
<fig id="f2" position="float">
<label>Figure&#xa0;2</label>
<caption>
<p>Example MSER segmentation of a noisy raw frame (with low SNR). <bold>(A)</bold> Raw output; <bold>(B)</bold> after flat-fielding; (<bold>C)</bold> regions of interest created through naive thresholding; <bold>(D)</bold> regions of interest and their bounding boxes created by applying MSER to <bold>(C)</bold>. In a low SNR frame such as the one above the processing steps are <bold>(A&#x2013;D)</bold>, while in a high SNR frame the processing steps are <bold>(A, B, D)</bold>. In panel <bold>(A)</bold>, the scale bar represents 1&#xa0;cm and is applicable to other panels.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fmars-09-870005-g002.tif"/>
</fig>
</sec>
<sec id="s2_1_3">
<title>2.1.3. Threshold-CNN (T-CNN) Segmentation</title>
<p>Another solution is to use Convolutional Neural Networks to either detect (i.e. define bounding boxes around) or segment (i.e. define a pixel mask of) objects of interest. Such approaches open the possibility to focus the detection on some types of objects (here, plankton) and ignore others (here, marine snow and artifacts); this is also called content-aware object detection or segmentation. However, CNNs tend to underperform at detecting objects across a large size range, especially for objects starting from a few dozen pixels (<xref ref-type="bibr" rid="B12">Cai et&#xa0;al., 2016</xref>). They work best when the target objects are of the same size as the receptive field of the model (<xref ref-type="bibr" rid="B19">Eggert et&#xa0;al., 2016</xref>). Thus, the development of detectors implementing receptive fields of various sizes constituted a major improvement, as they allowed detecting objects across a larger size range (<xref ref-type="bibr" rid="B12">Cai et&#xa0;al., 2016</xref>). In particular, we chose the Detectron2 library (<xref ref-type="bibr" rid="B77">Wu et&#xa0;al., 2019</xref>) developed by Facebook AI Research, which provides state-of-the-art object detection and segmentation algorithms, as well as pre-trained models for such tasks. Detectron2 includes a feature pyramid network (<xref ref-type="bibr" rid="B41">Lin et&#xa0;al., 2017</xref>) backbone that extracts feature maps across multiple scales to enable the detection of objects of various sizes, which was critical in our application to plankton images. Yet, this was not enough to cover the very large size range of organisms imaged by the ISIIS (from 50 to hundreds of thousands of pixels in area).</p>
<p>As explained above, marine snow particles and density-induced imaging artifacts are especially dominant compared to plankton in the smaller size classes. Therefore, our CNN pipeline was set up to segment the smaller objects, from 50 to 400 pixels in area, where the ability to specifically segment plankton makes the most difference. Above 400 pixels, the quantile-based threshold approach, with dilation and erosion, was used because it was simple and did not generate too many non-plankton segments.</p>
<p>In Detectron2, we used Mask R-CNN (<xref ref-type="bibr" rid="B35">He et&#xa0;al., 2017</xref>), which allows simultaneous bounding box detection and instance segmentation. The model was initialized with weights trained on the COCO reference dataset<xref ref-type="fn" rid="fn1"><sup>1</sup></xref> but, for it to detect planktonic organisms on ISIIS images, it has to be fine-tuned on a dataset of ground truth bounding boxes and masks of such organisms. This dataset was generated by manually delineating all recognizable planktonic organisms in a set of ISIIS images, using a digital pen on a tablet computer. This produced 23,197 ground truth masks, from which bounding boxes were computed. Among those, 10,878 object were in the 50-400 pixels area range and usable. A 524&#xd7;524 pixels crop was generated around every ground truth object (pushing the crop back inside the image when it crossed the edges). The choice of this particular size is a tradeoff between the maximum size of planktonic organisms that can be detected and the memory available on the graphics card. Moreover, it is in the line with common input sizes for segmentation models and was convenient to generate a tiling on ISIIS images. Several objects could be present in a crop. The crops were then split into 70% for training, 15% for validation, and 15% for testing. This split was stratified by the average gray level of the crop to ensure that both noisy (darker) and clean (lighter) images were present in each split, so that the model was presented with all kinds of images during training. Indeed, a model trained on clean images only would have performed poorly on noisy ones.</p>
<p>Detectron2 can perform multiclass object detection or segmentation, meaning that objects are both detected/segmented and classified in a single step. However, it requires sufficient examples in each class for training. This condition could not be satisfied here, given how time-consuming it was to obtain pixel-level masks for every object and because plankton samples are usually dominated by a few abundant taxa while most others are very rare (<xref ref-type="bibr" rid="B64">Ser-Giacomi et&#xa0;al., 2018</xref>). Since the focus of this study is on segmentation, we decided to perform one-class object detection/segmentation, thus training the model to recognize planktonic organisms of any taxon. This implies that classification needs to be done after segmentation. Once an object is detected, this sequential, rather than concurrent, approach does not affect the result of the classification, since the same information is available to the subsequent classifier as to the concurrent one. Furthermore, focusing on segmentation only is also more comparable with the two other methods described above.</p>
<p>The model was trained for 30,000 iterations, and evaluation was run on the validation set every 1,000 iterations to ensure that the validation loss reached a plateau. The learning rate was set to 0.0005 initially and decreased 10 fold after 10,000 and 20,000 iterations. To increase the generality of the detector, data augmentation was used in the form of random resizing of the 524 pixels crops (to 640, 672, 704, 736, 768 or 800 pixels) and random horizontal flipping. The test set was used to assess theoretical performance after training and guide the choice of model settings; the actual performance was assessed on a separate, real-world dataset (presented below).</p>
<p>To apply the trained model to new images, a tiling of 524&#xd7;524 pixels crops (the size used during model training) was generated over each input image, resulting in an overlap of 143 pixels vertically and 135 pixels horizontally. The overlap ensured that detectable objects spread over two crops were not missed. Crops were upscaled to 900&#xd7;900 pixels to improve detection of small objects (<xref ref-type="bibr" rid="B19">Eggert et&#xa0;al., 2016</xref>). For each crop, the model predicted the bounding boxes of objects and their masks. We only considered the boxes, resolved overlaps in detections caused by overlapping crops, and submitted each box to exactly the same quantile-based thresholding as what was used above 400 pixels. This was preferred over using Detectron&#x2019;s mask proposals because their outline was not as detailed or replicable as the threshold-based ones. Furthermore, it also ensured that morphometric measurements performed on the masks (area in particular) were exactly comparable between the objects that went through the CNN and those above 400 pixels that were defined by simple thresholding. For each bounding box proposal, the model computes a confidence score. We retained all boxes with a score over 0.1, which is a quite low confidence threshold designed to increase the chance of detecting all objects of interest (i.e. favor recall) at the cost of some false positive detections (i.e. lower precision). Those false positives (i.e. segmented objects that are not plankton) will have the opportunity to be eliminated later, when segments are classified taxonomically.</p>
<p>The CNN was coded in Python with PyTorch, the original implementation library for Detectron2. Training was conducted on an Nvidia Quadro RTX 8000 GPU and the code is available at <uri xlink:href="https://github.com/ThelmaPana/Detectron2_plankton_training">https://github.com/ThelmaPana/Detectron2_plankton_training</uri>. The combined CNN and threshold segmentation pipeline is implemented in <uri xlink:href="https://github.com/jiho/apeep">https://github.com/jiho/apeep</uri> and this was run in several Linux-based environments, using various Nvidia GPUs.</p>
</sec>
</sec>
<sec id="s2_2">
<title>2.2. Application to ISIIS Data from VISUFRONT Campaign</title>
<p>We evaluated these segmentation methods on ISIIS data from the VISUFRONT campaign, which sampled the Ligurian current front (North Western Mediterranean Sea), in the 0-100&#xa0;m depth range, during summer 2013. Towed at a speed of 2&#xa0;m s<sup>-1</sup> (4kts) and set for a 28 kHz scanning rate, the ISIIS sampled 108 L per second. The 2048 pixels high continuous image strip created by the line scan camera moving in the water was cut in 2048&#xd7;2048 pixels frames for storage. The ISIIS captured marked volutes caused by water density variations (<xref ref-type="fig" rid="f1"><bold>Figures&#xa0;1D&#x2013;F</bold></xref>), mostly driven by temperature changes around the thermocline, previously described by <xref ref-type="bibr" rid="B20">Faillettaz et&#xa0;al. (2016)</xref>.</p>
<p>The continuous image strip was reassembled from the stored 2048&#xd7;2048 pixels frames. Each line of pixels was flat-fielded by subtracting the row-wise average over a 8000 pixels moving window, hence removing streaks (<xref ref-type="fig" rid="f1"><bold>Figures&#xa0;1A, B, D, E</bold></xref>). The cleaned image was cut into 10,240 pixels long images (5 frames, instead of 1) to reduce the probability of cutting objects across images while keeping the memory footprint of each image manageable. Finally, the image was contrasted by stretching the intensity range between percentiles 0 and 40 (<xref ref-type="fig" rid="f1"><bold>Figures&#xa0;1B, C, E, F</bold></xref>). These values were chosen by iteration, through discussions with the taxonomist in charge of delineating planktonic organisms from raw images, as to achieve the highest distinguishability for those.</p>
<p>A ground truth dataset was generated by manually delineating all planktonic organisms (using a digital pen and tablet) in 106 10,240&#xd7;2048 pixels images, regularly spread across a full transect, hence representative of different environments. This resulted in 3,356 objects that were later taxonomically sorted into 24 taxa (<xref ref-type="fig" rid="f3"><bold>Figure&#xa0;3</bold></xref>), in the Ecotaxa web application (<xref ref-type="bibr" rid="B55">Picheral et&#xa0;al., 2017</xref>). This dataset was completely independent from the one that was used to train, validate and test the Detectron2 model. Some images were checked by two independent operators to check their consistency; when this was done, no differences were found.</p>
<fig id="f3" position="float">
<label>Figure&#xa0;3</label>
<caption>
<p>Examples of planktonic organisms imaged by the ISIIS. <bold>(A)</bold> Acantharea; <bold>(B)</bold> Actinopterygii; <bold>(C)</bold> Annelida; <bold>(D)</bold> Appendicularia; <bold>(E)</bold> Appendicularia (house only); <bold>(F)</bold> Appendicularia (body only); <bold>(G)</bold> Aulacanthidae; <bold>(H)</bold> Bacillariophyceae; <bold>(I)</bold> Chaetognatha; <bold>(J)</bold> solitary Collodaria; <bold>(K)</bold> Hydrozoa; <bold>(L)</bold> Cnidaria (other than Hydrozoa); <bold>(M)</bold> Crustacea (other than Harpacticoida, Copepoda and Eumalacostraca); <bold>(N)</bold> Harpacticoida; <bold>(O)</bold> Copepoda (other than Harpacticoida); <bold>(P)</bold> Eumalacostraca; <bold>(Q)</bold> Echinodermata (pluteus larva); <bold>(R)</bold> colonial Collodaria; <bold>(S)</bold> Ctenophora; <bold>(T)</bold> Doliolida; <bold>(U)</bold> Mollusca; <bold>(V)</bold> Pyrocystis; <bold>(W)</bold> Rhizaria (other than Acantharea; Aulacanthidae and Collodaria); <bold>(X)</bold> Siphonophorae.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fmars-09-870005-g003.tif"/>
</fig>
<p>Segments from each of the three automated methods were matched with ground truth segments of the same image. A bounding box intersection over union (IoU) score higher than 10% was considered as a match between segments. This threshold was set after manually inspecting a set of potential matches with various IoU values and was found to be the best value to discriminate between true and false matches. In case a ground truth segment matched multiple automatic segments, only one match was retained, to avoid inflating artificially the number of matches from the automated pipelines. In case an automatic segment matched multiple ground truth segments, the match was not counted either because it corresponded to a large segment that encompassed several organisms likely belonging to different taxa, which would make it unexploitable ecologically. Both choices made the match metrics conservative.</p>
<p>From these matches, global precision and recall were computed to summarize performance. Precision was computed as the proportion of automatic segments that matched ground truth segments. A 100% precision means that the algorithm only extracted ground truth segments. Recall was computed as the proportion of ground truth segments detected by the automated segmentation algorithm. A 100% recall means the algorithm did segment every manually delineated organism. Precision and recall scores were also computed per size class, where size was defined as the length of the diagonal of the bounding box; size classes were defined as intervals of 10 pixels, from 10 to 100 pixels, plus a class &gt; 100 pixels. These size classes do not aim at reflecting any ecological groups but were designed to split segments into roughly balanced classes. Recall was also computed for each taxonomic group defined in the ground truth segments. Precision does not make sense for taxonomic groups since it would only reflect the performance of the classification, not of the segmentation. The particle matching and metric computation code is available at <uri xlink:href="https://github.com/ThelmaPana/segmentation_benchmark">https://github.com/ThelmaPana/segmentation_benchmark</uri>.</p>
</sec>
</sec>
<sec id="s3">
<title>3 Results</title>
<sec id="s3_1">
<title>3.1. Number and Size Distribution of Segments</title>
<p>On the 106 images of the segmentation benchmark dataset, 3,356 organisms were manually segmented, whereas the automated pipelines generated many more segments, especially the threshold-based one (<xref ref-type="table" rid="T2"><bold>Table&#xa0;2</bold></xref>).</p>
<table-wrap id="T2" position="float">
<label>Table&#xa0;2</label>
<caption>
<p>Number of segments generated by each pipeline on the 106 benchmark images and estimation of the amount of segments they would produce on one minute of ISIIS data.</p>
</caption>
<table frame="hsides">
<thead>
<tr>
<th valign="top" align="left">Segmentation pipeline </th>
<th valign="top" align="center">Number of segments on benchmark images </th>
<th valign="top" align="center">Average number of segments per minute of ISIIS deployment </th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="left">Ground truth</td>
<td valign="top" align="center">3,356</td>
<td valign="top" align="center">~5,000</td>
</tr>
<tr>
<td valign="top" align="left">Threshold</td>
<td valign="top" align="center">339,907</td>
<td valign="top" align="center">~525,000</td>
</tr>
<tr>
<td valign="top" align="left">Threshold-MSER</td>
<td valign="top" align="center">82,731</td>
<td valign="top" align="center">~130,000</td>
</tr>
<tr>
<td valign="top" align="left">Threshold-CNN</td>
<td valign="top" align="center">19,048</td>
<td valign="top" align="center">~30,000</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>The normalized abundance size spectra (NASS) (<xref ref-type="fig" rid="f4"><bold>Figure&#xa0;4</bold></xref>) display the expected linear decrease of abundance with size in log-log scale. For the ground truth segments, the curve dips below this linear relationship for objects of 25 pixels in diagonal and smaller (dotted vertical line on <xref ref-type="fig" rid="f4"><bold>Figure&#xa0;4</bold></xref>). Since this dataset specifically targeted recognisable planktonic organisms, this dip highlights that not all organisms below this size could be detected by a human taxonomist upon detailed examination of the images (<xref ref-type="bibr" rid="B42">Lombard et&#xa0;al., 2019</xref>). The discontinuity is towards smaller diagonal sizes in the automated pipelines, but likely because many of the small segments are of non-plankton objects.</p>
<fig id="f4" position="float">
<label>Figure&#xa0;4</label>
<caption>
<p>Normalized abundance size spectra (NASS) of all segments generated by the benchmarked pipelines and ground truth segmentation. To compute the NASS, segments were grouped into size classes on a log2 scale, each class size being two times wider than the previous one. Normalized abundance was computed by dividing the number of segments in each class by the size class width, resulting in an adimensional quantity (number of segments) divided by a length (mm here). The double x-axis is the length of the diagonal bounding box displayed both in pixels and after conversion in mm. The dotted vertical line highlights the slope discontinuity in the size spectrum of ground truth segments. Note that both axes use log10 scaling. T, threshold-based; T-MSER, threshold-MSER; T-CNN, threshold-CNN.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fmars-09-870005-g004.tif"/>
</fig>
<p>All automated pipelines have NASS curves above the ground truth, which highlights the fact that they segmented non-plankton objects. This was true over the entire size range but was particularly pronounced for the smaller size classes. Above 10 mm/200 pixels in diagonal, the T-MSER pipeline produced a number of segments comparable to the ground truth, which is satisfying, although it does not guarantee that those are of the same objects (it might have missed some plankton and segmented marine snow/artifacts in the same size range; see precision and recall performances for the largest size class in <xref ref-type="fig" rid="f5"><bold>Figure&#xa0;5</bold></xref> below). From the maximal size down to ~70 pixels in diagonal, the T and T-CNN pipelines produced the same segments. This coincides with the critical size of 400 pixels in area at which the segmentation method switched from threshold-based to content-aware. Indeed, the conversion from area to bounding box diagonal is not linear because it depends on the shape of the objects. For an object of 400 pixels in area, the bounding box diagonal is between 30 and 70 pixels. This shows that the T-CNN pipeline was effective in reducing the number of segments compared to naive thresholding, because the NASS diverges below that size.</p>
<fig id="f5" position="float">
<label>Figure&#xa0;5</label>
<caption>
<p>Precision <bold>(A)</bold> and recall <bold>(B)</bold> scores per size class. In <bold>(B)</bold>, n indicates the number of segments per size class for the ground truth dataset. T, threshold-based; T-MSER, threshold-MSER; T-CNN, threshold-CNN.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fmars-09-870005-g005.tif"/>
</fig>
<p>A linear regression performed on the linear portion of the NASS (diagonal values between 30 and 500 pixels) followed by an analysis of covariance demonstrated significant difference in slopes between the segmentation methods: F(3,105) = 133.07; p &lt; 0.001 (<xref ref-type="supplementary-material" rid="SM1"><bold>Table S1</bold></xref>). <italic>Post hoc</italic> analysis showed a significant difference between all segmentation methods (p &lt; 0.001 for all pairs) (<xref ref-type="supplementary-material" rid="SM1"><bold>Table S2</bold></xref>).</p>
</sec>
<sec id="s3_2">
<title>3.2. Global Performance Statistics</title>
<p>Overall, the three pipelines demonstrate good recall: when looking at the total number of segments, they all captured over 85% of the ground truth organisms. The T-CNN pipeline largely outperformed both the threshold-based and T-MSER pipelines in terms of precision (<xref ref-type="table" rid="T3"><bold>Table&#xa0;3</bold></xref>). In other words, although it segmented almost all planktonic objects, the threshold-based pipeline generated mostly non-plankton segments (~99%), composed of both marine snow and density volutes artifacts. The T-CNN pipeline also produced non-planktonic segments but they &#x201c;only&#x201d; represented 84% of segments, while still segmenting a good proportion of planktonic objects. The T-MSER performed somewhere in between those two extremes.</p>
<table-wrap id="T3" position="float">
<label>Table&#xa0;3</label>
<caption>
<p>Precision and recall values of the automated pipelines evaluated against the 3,356 ground truth organisms.</p>
</caption>
<table frame="hsides">
<thead>
<tr>
<th valign="top" align="left">Pipeline </th>
<th valign="top" align="center">Precision </th>
<th valign="top" align="center">Recall </th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="left">Threshold</td>
<td valign="top" align="right">0.9%</td>
<td valign="top" align="center">97.3%</td>
</tr>
<tr>
<td valign="top" align="left">Threshold-MSER</td>
<td valign="top" align="right">3.5%</td>
<td valign="top" align="center">85.4%</td>
</tr>
<tr>
<td valign="top" align="left">Threshold-CNN</td>
<td valign="top" align="right">16.3%</td>
<td valign="top" align="center">91.9%</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec id="s3_3">
<title>3.3. Performances Per Size Class</title>
<p>Because the behavior of the pipelines seems to vary with size (<xref ref-type="fig" rid="f4"><bold>Figure&#xa0;4</bold></xref>), it seems relevant to break down the matching statistics per size class. With the threshold-based pipeline, precision decreased with size: smaller segments included a lower proportion of planktonic organisms than larger ones (<xref ref-type="fig" rid="f5"><bold>Figure&#xa0;5A</bold></xref>). The T-CNN pipeline had better precision than the others for small segments while T-MSER had a better precision for larger segments. In terms of recall, the threshold-based pipeline always performed better than the others, regardless of size class (<xref ref-type="fig" rid="f5"><bold>Figure&#xa0;5B</bold></xref>). The T-MSER pipeline performed as well as the T-CNN pipeline on middle size classes, but achieved a lower recall for both very small and very large segments.</p>
</sec>
<sec id="s3_4">
<title>3.4. Performances Per Taxonomic Group</title>
<p>In the ground truth dataset, half of the 24 detected taxa were represented by fewer than 18 individuals (median is 18.5), hence inducing little resolution and large variance in the performance statistics of segmentation pipelines. Among the other half of the taxa, the recall of the T-CNN pipeline was lower than that of the threshold pipeline by more than 10% for only two taxa (Bacillaryophycea and Doliolida) and for only four in the case of the T-MSER pipeline (Bacillariophyceae, Ctenophora, Acantharea, and other Rhizaria; <xref ref-type="fig" rid="f6"><bold>Figure&#xa0;6</bold></xref>). The lowest recall values were reached for Bacillariophyceae and Ctenophora, for all pipelines. In concordance with the consistent recall performance across size classes, taxa-wise recall performance of the T-CNN pipeline do not seem linked to organism size: small organisms (e.g. Acantharea, Pyrocystis) were accurately detected.</p>
<fig id="f6" position="float">
<label>Figure&#xa0;6</label>
<caption>
<p>Recall scores per taxon. n is the number of individuals from each taxon in the 106 benchmark images and taxa are sorted in decreasing order of abundance. T, threshold-based; T-MSER, threshold-MSER; T-CNN, threshold-CNN.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fmars-09-870005-g006.tif"/>
</fig>
</sec>
</sec>
<sec id="s4">
<title>4 Discussion</title>
<sec id="s4_1">
<title>4.1. Summary of Results</title>
<p>The threshold-based pipeline performed an exhaustive segmentation: planktonic organisms were almost all properly detected, yet they were drowned in the overwhelming majority of non-planktonic objects (<xref ref-type="table" rid="T2"><bold>Table&#xa0;2</bold></xref>). The T-CNN pipeline reduced this problem, significantly increasing precision (<xref ref-type="table" rid="T3"><bold>Table&#xa0;3</bold></xref> and <xref ref-type="fig" rid="f5"><bold>Figure&#xa0;5A</bold></xref>) while still achieving a very good detection of plankton across the entire size range targeted by ISIIS. The T-MSER pipeline also reduced the segmentation of non-planktonic objects, especially at the top-end of the size range, but detected fewer planktonic organisms than the other pipelines (<xref ref-type="fig" rid="f5"><bold>Figure&#xa0;5B</bold></xref>). Despite the large decrease in number of segmented objects, for most taxa, the MSER or CNN pipelines reduced recall by less than 10% (<xref ref-type="fig" rid="f6"><bold>Figure&#xa0;6</bold></xref>). One explanation for these differences is that naive thresholding captured a lot of noise (i.e. density volutes) and, additionally, broke it into many small segments. The use of either MSER or a CNN allowed ignoring these noise segments and/or not breaking them apart, hence producing much fewer non-planktonic segments. The decrease in abundance below the expected slope at the smaller end of the size spectrum of ground truth segments (<xref ref-type="fig" rid="f4"><bold>Figure&#xa0;4</bold></xref>) suggests that identification of planktonic organisms becomes non-exhaustive below 25 pixels in bounding box diagonal. Below this size, which amounts to 600 &#xb5;m in ESD on average, some organisms can still be detected. This means that relative concentrations between locations/times can likely be exploited within a taxon but that further filtering and corrections are needed to reach absolute concentrations.</p>
<p>The statistical difference between NASS slopes (<xref ref-type="fig" rid="f4"><bold>Figure&#xa0;4</bold></xref>) indicates that they segment different kinds and amounts of non-planktonic objects, compared to the all-plankton ground truth. This implies that the output of different segmentation approaches should not be directly compared in terms of size distribution. Segmentation methods were already shown to have an impact on the definition of particle size and shape, which propagates to subsequent analyses such as particle flux estimates (<xref ref-type="bibr" rid="B25">Giering et&#xa0;al., 2020</xref>). This slope discrepancy as well as the vastly larger intercept of the NASS of automated pipelines compared to the ground truth means that the computation of an appropriate plankton size spectrum requires a classification step that would exclude non-planktonic objects.</p>
</sec>
<sec id="s4_2">
<title>4.2. Targeted Organisms</title>
<p>Some taxa were systematically less often detected than others. Some of the not detected Bacillariophyceae were large, blurry, and too translucent (<xref ref-type="fig" rid="f3"><bold>Figure&#xa0;3H</bold></xref>) to be caught by the threshold-based branch of the T-CNN pipeline or by the T-MSER method. The other, smaller, ones that were missed by the content-aware branch of T-CNN were not detected because they were quite different from the ones used during training (blurrier). Integrating more representative examples of Bacillariophyceae for CNN training could have improved performance on this taxon. Similarly, doliolids (<xref ref-type="fig" rid="f3"><bold>Figure&#xa0;3T)</bold></xref>, that were often large, should have been segmented by the threshold-based branch of T-CNN as well as by T-MSER. The ones missed, mostly by T-CNN, were also blurry and too translucent for intensity-based thresholding with a single threshold. Ctenophores (likely of the Mertensiidae family, <xref ref-type="fig" rid="f3"><bold>Figure&#xa0;3S</bold></xref>) displayed thin, translucent tentacles that were often missed by threshold-based methods. Therefore, only the body was segmented, which resulted in a bounding box IoU value &lt; 0.1, too low to be considered a match with the ground truth segment that included the tentacles. Still, a later CNN classifier should be able to correctly identify even such portions of organisms, as CNNs were shown to mostly rely on local shape and texture features instead of on the global shape (<xref ref-type="bibr" rid="B3">Baker et&#xa0;al., 2018</xref>; <xref ref-type="bibr" rid="B4">Baker et&#xa0;al., 2020</xref>). Finally, the T-MSER pipeline resulted in a lower recall for Acantharea and other Rhizaria (<xref ref-type="fig" rid="f3"><bold>Figures&#xa0;3A, W</bold></xref>). This seems to stem from a too aggressive thresholding step in low SNR high noise frames, the pre-processing step before MSER is applied. Further fine-tuning would likely allow it to retain more or all Acantharea and other Rhizaria images.</p>
<p>In the present study, we aimed at performing an exhaustive detection of every planktonic organism across the size range targeted by the ISIIS. However, in general, the segmentation algorithm should be chosen according to the target organisms. For example, to focus on organisms towards the larger end of the ISIIS size range (e.g. &gt; 10&#xa0;mm), where particles &#x2014; mostly marine snow aggregates &#x2014; are much less abundant, a simple gray-level threshold seems sufficient.</p>
</sec>
<sec id="s4_3">
<title>4.3. Processing Time and Cost</title>
<p>The quantile-based thresholding pipeline ran on a single CPU core at a rate of 30 minutes of processing for 1 minute of ISIIS data (0.03x), on an Intel Xeon E5-2643 v3 (3.40 GHz). Its memory requirements were limited so it was easy to run simultaneous processing of multiple batches of data on a multi-core/multi-processor machine, but the treatment of ISIIS data as a continuous stream for flat-fielding prevented automatic multithreading. The T-CNN pipeline required a GPU with sufficient memory (48 GB, on aNvidia Quadro RTX 8000 in our case) to efficiently train the CNN portion and to fit ISIIS images in at evaluation time. It processed data at the same rate as the threshold-based pipeline (30&#xa0;min processing for 1&#xa0;min of data, or 0.03x). The T-MSER pipeline was optimized for speed and utilized the 8 cores of an AMD Ryzen 3700, processing one minute of ISIIS data in 50 seconds (1.2x), or 6&#xa0;min 40 s of processing for 1&#xa0;min of ISIIS data (0.15x) when considering running on one core.</p>
<p>The MSER implementation followed <xref ref-type="bibr" rid="B46">Matas et&#xa0;al. (2004)</xref> closely. The optimization of the T-MSER approach stems from adding the SNR switch, which leads to the pre-processing of high-noise images with naive thresholding, while going straight to the MSER-based detection in low noise images. Adding these changes increased segmentation recall from 65% to 85%. Further optimization included making the code multi-thread ready for deployment on High Performance Computing infrastructures. Using the specialized CPUs of these infrastructures, such as the AMD EPYC 7742 (64 cores, 128 threads) performance could improve well above 1.2x. At current data collection rates of 75-100&#xa0;h of ISIIS data per scientific cruise, a real time or faster than real time segmentation approach constitutes a substantial benefit.</p>
<p>At first glance, the T-CNN pipeline seems expensive in terms of set up and architecture: it requires a GPU with sufficient memory to operate, implies the use of relatively new deep learning coding frameworks and the preparation of a training set with manual delineation of thousands of planktonic organisms. But these costs are offset by the time gained not processing a multitude of particles in each image, resulting in a processing rate comparable to that of the pure threshold-based pipeline, as stated above. Furthermore, the fact that T-CNN produced 20 times fewer segments will also considerably reduce the classification time (often CNN based too). Finally, since recall barely decreased, the objects ignored were mostly the dominant non-plankton objects, as per design; this will diminish the imbalance among classes that classifiers are sensitive too, further improving the classification step. Moreover, both the Detectron2 library and the baseline model on which the T-CNN pipeline relies are easily downloadable and well documented<sup>2</sup>. With GPU resources becoming increasingly available for scientific research and the associated frameworks becoming easier to use, such tools are poised to become more powerful and accessible.</p>
</sec>
<sec id="s4_4">
<title>4.4. Detection of Small Objects by CNN Models</title>
<p>The detection of objects measuring just a few pixels is still a research problem in its own right in computer sciences (<xref ref-type="bibr" rid="B18">Eggert et&#xa0;al., 2017</xref>), coined very low resolution recognition problems (<xref ref-type="bibr" rid="B75">Wang et&#xa0;al., 2016</xref>). They are characterized by targets smaller than 16&#xd7;16 pixels, which can be challenging even for the perceptual abilities of human experts. They target applications for company logo detection (<xref ref-type="bibr" rid="B19">Eggert et&#xa0;al., 2016</xref>; <xref ref-type="bibr" rid="B18">Eggert et&#xa0;al., 2017</xref>), face recognition from video surveillance, or text recognition (<xref ref-type="bibr" rid="B75">Wang et&#xa0;al., 2016</xref>). The receptive fields of common object detection architectures match the target object size and range from 50&#xd7;50 to 450&#xd7;450 pixels which is much larger than the small objects targeted in low resolution studies (<xref ref-type="bibr" rid="B18">Eggert et&#xa0;al., 2017</xref>). Here, the smallest organisms targeted had an area of 50 pixels, which corresponded to a bounding box diagonal of 12 pixels, or an 8x8 pixels square. Thus the exhaustive detection of plankton organisms in ISIIS images, including the smaller ones, clearly falls in the domain of very low resolution recognition. A common solution is image upscaling, as highlighted by <xref ref-type="bibr" rid="B19">Eggert et&#xa0;al. (2016)</xref>, which we implemented in the present work. The 524&#xd7;524 pixels crops were upscaled to 900&#xd7;900 pixels before evaluation in the Detectron2 model. The 900 pixels size is a compromise between detection accuracy, usage of the GPU memory, and processing time. Other approaches for multi-scale object detection are described by (<xref ref-type="bibr" rid="B12">Cai et&#xa0;al., 2016</xref>) and include magnification of regions susceptible to contain small objects (<xref ref-type="bibr" rid="B19">Eggert et&#xa0;al., 2016</xref>) or the integration of contextual information outside of regions of interest (<xref ref-type="bibr" rid="B5">Bell et al., 2016</xref>).</p>
<p>No automated segmentation method is perfect; depending on their settings, they either avoid objects other than their targets but miss some objects of interest (high precision, low recall) or detect most objects of interest but also many others (high recall, low precision). If the segmentation or object detection task is followed by a classification step, which is always the case for plankton imaging, we advocate in favor of recall over precision during segmentation, provided that the amount of data remains manageable. Hence, a maximum number of planktonic objects have the opportunity to be classified. The precision can be improved after classification, by filtering out low confidence, usually error prone, predictions based on the score given by the classifier (<xref ref-type="bibr" rid="B20">Faillettaz et&#xa0;al., 2016</xref>; <xref ref-type="bibr" rid="B45">Luo et&#xa0;al., 2018</xref>).</p>
<p>To extract planktonic organisms of various taxa from ISIIS images, full instance segmentation would have been the most elegant approach, outputting classified mask instances in a single step (<xref ref-type="bibr" rid="B15">Dai et&#xa0;al., 2016</xref>). Several obstacles still lay ahead for this approach to be applicable. First, training an instance segmentation model to recognize each taxonomic group would require hundreds to thousands of ground truth (i.e. human-produced) masks of all taxa. Given the long tailed distribution of taxa concentrations in the planktonic world, with many rare taxa, in particular the largest ones, this would require a considerable amount of searching and labeling effort. Indeed, assembling enough examples to train classifications models is already challenging (<xref ref-type="bibr" rid="B37">Irisson et&#xa0;al., 2022</xref>) and manual delineation of each organism is much more time consuming than manual classification. A second obstacle is the size range of organisms imaged by ISIIS. Although Detectron2 does produce multi-scale feature maps through a Feature Pyramid Network in order to apply receptive fields of multiple size, the ratio between the largest and the smallest feature maps is only 16. Here, the ratio between the smallest and largest bounding box diagonals of manually segmented organisms is 65 and can reach &gt; 180 in more exhaustive ISIIS datasets. To tackle this span, one could theoretically set up an ensemble of detectors, fed with crops of different sizes, each one targeting a restricted size range. Yet, this would be a particularly computationally demanding and complex set up, for a gain yet to be determined since, for larger sizes, the proportion of non-plankton objects, and therefore the advantage of a CNN-based segmentation, diminishes. Finally, masks generated by instance segmentation models currently lack both precision (their outline is smoothed, not matching the fine appendages of plankton) and reproducibility (because of the randomness included during training to avoid overfitting, two models trained on the same data will output different masks). These drawbacks are particularly critical for plankton application, where the size of the organisms, computed from their masks, is often of interest.</p>
</sec>
</sec>
<sec id="s5">
<title>5 Conclusion and Perspectives</title>
<p>We developed combined segmentation pipelines able to detect planktonic organisms spanning a broad size range. The fact that all methods comprised a deterministic, threshold-based segmentation ensured that particle shapes and measurement were consistent over the whole size range. Still, the segmentation method affected the shape of the size spectrum and additional processing steps (including classification) are needed to extract the correct size structure of living organisms. The MSER method limited over-segmentation of background noise objects and extracted more consistent segments, at a very high processing rate. This speed opens the possibility for near-real time processing, which is particularly relevant for adaptive sampling during a cruise or an early warning system in a time series context. Although at the lower limit of the detection capabilities of CNNs, our content-aware approach was able to detect planktonic organisms among an overwhelming number of marine snow and noise images, exhibiting the best recall of the three methods. Therefore, the ideal segmentation approach depends on the study objectives and operational constraints.</p>
<p>These approaches seem relevant for imaging studies focused on living planktonic organisms, since they reduce the number of objects from non-plankton classes that are extracted. In turn, this dampens the imbalance towards these classes, laying the foundations for easier, faster, and more accurate subsequent object classification by (i) reducing the amount of work needed to generate a training set with similar class distribution, which is essential to avoid the caveat of dataset shift (<xref ref-type="bibr" rid="B48">Moreno-Torres et&#xa0;al., 2012</xref>); (ii) decreasing the computation time because there are fewer objects; and (iii) limiting the contamination of the rare planktonic classes by the dominant, non-plankton, ones.</p>
<p>Although CNN-based object detection may seem overwhelming at first, both in terms of set up and processing time, it actually is fast enough and within the reach of marine ecologists, particularly now that artificial intelligence frameworks and GPU computing are being made more accessible. This work constitutes a step towards the &#x201c;intelligent&#x201d; segmentation of ecological images, even at low resolution, which could find even wider applications such as the automated separation of objects overlapping onto each other on an image for more accurate species counts, the detection and classification in a single step for more automated surveys, or the extraction of individual-level traits to track e.g., reproductive organs development, for a richer exploitation of ecological images (<xref ref-type="bibr" rid="B51">Orenstein et&#xa0;al., 2021</xref>). Such tasks are in no way limited to plankton images and are common in data collected by trawl cameras, benthic observations or surveying cameras, vessel monitoring cameras, etc.</p>
<p>In this era of data-driven oceanography, the volume of data collected is increasing sharply, thanks to technological advances such as high frequency imagery, autonomous instruments (e.g. floats, gliders), satellite-based methods as well as environmental -omics approaches permitted by high throughput sequencing. In this context of abundant data, the development of automated and efficient data processing techniques becomes a key element in drawing a holistic understanding of oceanic ecosystems; it is needed to provide an extensive description of biodiversity, including species distributions as well as estimates of biomass and abundance.</p>
</sec>
<sec id="s6" sec-type="data-availability">
<title>Data Availability Statement</title>
<p>The raw data supporting the conclusions of this article will be made available by the authors, without undue reservation.</p>
</sec>
<sec id="s7" sec-type="author-contributions">
<title>Author Contributions</title>
<p>J-OI and TP conceptualized the study. LC-C and TP generated and taxonomically sorted ground truth plankton segments. BW and TP developed the CNN segmentation method. J-OI and TP developed the threshold-based and the T-CNN processing pipelines. MS, DD, ST, CS, and RC set up and ran the T-MSER method. TP prepared the original draft. All co-authors proof-read the manuscript prior to submission. All authors read and approved the final manuscript.</p>
</sec>
<sec id="s8" sec-type="funding-information">
<title>Funding</title>
<p>This study is part of project &#x201c;World Wide Web of Plankton Image Curation&#x201d;, funded by the Belmont Forum through the Agence Nationale de la Recherche ANR-18-BELM-0003-01 and the National Science Foundation (NSF) #ICER1927710. Funding also came from NSF #OCE2125407. The Extreme Science and Engineering Discovery Environment (XSEDE) provided computing resources to the US team through grant #TG-OCE170012. Data acquisition during the VISUFRONT cruise was funded by the Partner University Fund and supported by the French Oceanographic Fleet through ship time. TP&#x2019;s doctoral fellowship was granted by the French Ministry of Higher Education, Research and Innovation (#3500/2019).</p>
</sec>
<sec id="s9" sec-type="COI-statement">
<title>Conflict of Interest</title>
<p>Author Ben Woodward was employed by company CVision AI.</p>
<p>The remaining authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec id="s10" sec-type="disclaimer">
<title>Publisher&#x2019;s Note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
</body>
<back>
<ack>
<title>Acknowledgments</title>
<p>The authors would like to thank the officers and crew of the R/V Tethys 2 who made the VISUFRONT campaign a success, as well as the additional scientists who took part in the cruise: R Faillettaz, C M Guigand, F Lombard, M Lilley, and J Luo.</p>
</ack>
<sec id="s11" sec-type="supplementary-material">
<title>Supplementary Material</title>
<p>The Supplementary Material for this article can be found online at: <ext-link ext-link-type="uri" xlink:href="https://www.frontiersin.org/articles/10.3389/fmars.2022.870005/full#supplementary-material">https://www.frontiersin.org/articles/10.3389/fmars.2022.870005/full#supplementary-material</ext-link>
</p>
<supplementary-material xlink:href="Table_1.docx" id="SM1" mimetype="application/vnd.openxmlformats-officedocument.wordprocessingml.document"/>
</sec>
<fn-group>
<fn id="fn1">
<label>1</label>
<p>
<uri xlink:href="https://github.com/facebookresearch/detectron2/blob/main/configs/COCO-InstanceSegmentation/mask_rcnn_R_50_FPN_3x.yaml/">https://github.com/facebookresearch/detectron2/blob/main/configs/COCO-InstanceSegmentation/mask_rcnn_R_50_FPN_3x.yaml/</uri>
</p>
</fn>
</fn-group>
<ref-list>
<title>References</title>
<ref id="B1">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Alldredge</surname> <given-names>A. L.</given-names>
</name>
<name>
<surname>Granata</surname> <given-names>T. C.</given-names>
</name>
<name>
<surname>Gotschalk</surname> <given-names>C. C.</given-names>
</name>
<name>
<surname>Dickey</surname> <given-names>T. D.</given-names>
</name>
</person-group> (<year>1990</year>). <article-title>The Physical Strength of Marine Snow and Its Implications for Particle Disaggregation in the Ocean</article-title>. <source>Limnol. Oceanog.</source> <volume>35</volume>, <fpage>1415</fpage>&#x2013;<lpage>1428</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.4319/lo.1990.35.7.1415</pub-id>
</citation>
</ref>
<ref id="B2">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Alldredge</surname> <given-names>A. L.</given-names>
</name>
<name>
<surname>Silver</surname> <given-names>M. W.</given-names>
</name>
</person-group> (<year>1988</year>). <article-title>Characteristics, Dynamics and Significance of Marine Snow</article-title>. <source>Prog. Oceanog.</source> <volume>20</volume>, <fpage>41</fpage>&#x2013;<lpage>82</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/0079-6611(88)90053-5</pub-id>
</citation>
</ref>
<ref id="B3">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Baker</surname> <given-names>N.</given-names>
</name>
<name>
<surname>Lu</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Erlikhman</surname> <given-names>G.</given-names>
</name>
<name>
<surname>Kellman</surname> <given-names>P. J.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>Deep Convolutional Networks do Not Classify Based on Global Object Shape</article-title>. <source>PloS Comput. Biol.</source> <volume>14</volume>, <elocation-id>e1006613</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1371/journal.pcbi.1006613</pub-id>
</citation>
</ref>
<ref id="B4">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Baker</surname> <given-names>N.</given-names>
</name>
<name>
<surname>Lu</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Erlikhman</surname> <given-names>G.</given-names>
</name>
<name>
<surname>Kellman</surname> <given-names>P. J.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Local Features and Global Shape Information in Object Classification by Deep Convolutional Neural Networks</article-title>. <source>Vision Res.</source> <volume>172</volume>, <fpage>46</fpage>&#x2013;<lpage>61</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.visres.2020.04.003</pub-id>
</citation>
</ref>
<ref id="B5">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Bell</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Zitnick</surname> <given-names>C. L.</given-names>
</name>
<name>
<surname>Bala</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Girshick</surname> <given-names>R.</given-names>
</name>
</person-group> (<year>2016</year>). &#x201c;<article-title>Inside-Outside Net: Detecting Objects in Context With Skip Pooling and Recurrent Neural Networks</article-title>,&#x201d; in <conf-name>Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition (CVPR)</conf-name>, <fpage>2874</fpage>&#x2013;<lpage>2883</lpage>.</citation>
</ref>
<ref id="B6">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Benfield</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Grosjean</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Culverhouse</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Irigolen</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Sieracki</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Lopez-Urrutia</surname> <given-names>A.</given-names>
</name>
<etal/>
</person-group>. (<year>2007</year>). <article-title>RAPID: Research on Automated Plankton Identification</article-title>. <source>Oceanography</source> <volume>20</volume>, <fpage>172</fpage>&#x2013;<lpage>187</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.5670/oceanog.2007.63</pub-id>
</citation>
</ref>
<ref id="B7">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Biard</surname> <given-names>T.</given-names>
</name>
<name>
<surname>Ohman</surname> <given-names>M. D.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Vertical Niche Definition of Test-Bearing Protists (Rhizaria) Into the Twilight Zone Revealed by <italic>in Situ</italic> Imaging</article-title>. <source>Limnol. Oceanog.</source> <volume>65</volume>, <fpage>2583</fpage>&#x2013;<lpage>2602</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1002/lno.11472</pub-id>
</citation>
</ref>
<ref id="B8">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Biard</surname> <given-names>T.</given-names>
</name>
<name>
<surname>Stemmann</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Picheral</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Mayot</surname> <given-names>N.</given-names>
</name>
<name>
<surname>Vandromme</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Hauss</surname> <given-names>H.</given-names>
</name>
<etal/>
</person-group>. (<year>2016</year>). <article-title><italic>In Situ</italic> Imaging Reveals the Biomass of Giant Protists in the Global Ocean</article-title>. <source>Nature</source> <volume>532</volume>, <fpage>504</fpage>&#x2013;<lpage>507</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1038/nature17652</pub-id>
</citation>
</ref>
<ref id="B9">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Bi</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Guo</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Benfield</surname> <given-names>M. C.</given-names>
</name>
<name>
<surname>Fan</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Ford</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Shahrestani</surname> <given-names>S.</given-names>
</name>
<etal/>
</person-group>. (<year>2015</year>). <article-title>). A Semi-Automated Image Analysis Procedure for <italic>In Situ</italic> Plankton Imaging Systems</article-title>. <source>PloS One</source> <volume>10</volume>, <elocation-id>e0127121</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1371/journal.pone.0127121</pub-id>
</citation>
</ref>
<ref id="B10">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Brand&#xe3;o</surname> <given-names>M. C.</given-names>
</name>
<name>
<surname>Benedetti</surname> <given-names>F.</given-names>
</name>
<name>
<surname>Martini</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Soviadan</surname> <given-names>Y. D.</given-names>
</name>
<name>
<surname>Irisson</surname> <given-names>J.-O.</given-names>
</name>
<name>
<surname>Romagnan</surname> <given-names>J.-B.</given-names>
</name>
<etal/>
</person-group>. (<year>2021</year>). <article-title>Macroscale Patterns of Oceanic Zooplankton Composition and Size Structure</article-title>. <source>Sci. Rep.</source> <volume>11</volume>, <fpage>15714</fpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1038/s41598-021-94615-5</pub-id>
</citation>
</ref>
<ref id="B11">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Brise&#xf1;o-Avena</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Schmid</surname> <given-names>M. S.</given-names>
</name>
<name>
<surname>Swieca</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Sponaugle</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Brodeur</surname> <given-names>R. D.</given-names>
</name>
<name>
<surname>Cowen</surname> <given-names>R. K.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Three-Dimensional Cross-Shelf Zooplankton Distributions Off the Central Oregon Coast During Anomalous Oceanographic Conditions</article-title>. <source>Prog. Oceanog.</source> <volume>188</volume>, <elocation-id>102436</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.pocean.2020.102436</pub-id>
</citation>
</ref>
<ref id="B12">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Cai</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Fan</surname> <given-names>Q.</given-names>
</name>
<name>
<surname>Feris</surname> <given-names>R. S.</given-names>
</name>
<name>
<surname>Vasconcelos</surname> <given-names>N.</given-names>
</name>
</person-group> (<year>2016</year>). &#x201c;<article-title>A Unified Multi-Scale Deep Convolutional Neural Network for Fast Object Detectionv</article-title>,&#x201d; in <conf-name>European Conference on Computer Vision (ECCV)</conf-name>, <fpage>1607.07155</fpage>.</citation>
</ref>
<ref id="B13">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Cheng</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Cheng</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Bi</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Benfield</surname> <given-names>M. C.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Enhanced Convolutional Neural Network for Plankton Identification and Enumeration</article-title>. <source>PloS One</source> <volume>14</volume>, <elocation-id>e0219570</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1371/journal.pone.0219570</pub-id>
</citation>
</ref>
<ref id="B14">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Cowen</surname> <given-names>R. K.</given-names>
</name>
<name>
<surname>Guigand</surname> <given-names>C. M.</given-names>
</name>
</person-group> (<year>2008</year>). <article-title><italic>In Situ</italic> Ichthyoplankton Imaging System (ISIIS): System Design and Preliminary Results</article-title>. <source>Limnol. Oceanog.: Methods</source> <volume>6</volume>, <fpage>126</fpage>&#x2013;<lpage>132</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.4319/lom.2008.6.126</pub-id>
</citation>
</ref>
<ref id="B15">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Dai</surname> <given-names>J.</given-names>
</name>
<name>
<surname>He</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Sun</surname> <given-names>J.</given-names>
</name>
</person-group> (<year>2016</year>). &#x201c;<article-title>Instance-Aware Semantic Segmentation <italic>via</italic> Multi-Task Network Cascades</article-title>,&#x201d; in <conf-name>Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition (CVPR)</conf-name>, <fpage>3150</fpage>&#x2013;<lpage>3158</lpage>.</citation>
</ref>
<ref id="B16">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Dennett</surname> <given-names>M. R.</given-names>
</name>
<name>
<surname>Caron</surname> <given-names>D. A.</given-names>
</name>
<name>
<surname>Michaels</surname> <given-names>A. F.</given-names>
</name>
<name>
<surname>Gallager</surname> <given-names>S. M.</given-names>
</name>
<name>
<surname>Davis</surname> <given-names>C. S.</given-names>
</name>
</person-group> (<year>2002</year>). <article-title>Video Plankton Recorder Reveals High Abundances of Colonial Radiolaria in Surface Waters of the Central North Pacific</article-title>. <source>J. Plank. Res.</source> <volume>24</volume>, <fpage>797</fpage>&#x2013;<lpage>805</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1093/plankt/24.8.797</pub-id>
</citation>
</ref>
<ref id="B17">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>de Vargas</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Audic</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Henry</surname> <given-names>N.</given-names>
</name>
<name>
<surname>Decelle</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Mah&#xe9;</surname> <given-names>F.</given-names>
</name>
<name>
<surname>Logares</surname> <given-names>R.</given-names>
</name>
<etal/>
</person-group>. (<year>2015</year>). <article-title>Eukaryotic Plankton Diversity in the Sunlit Ocean</article-title>. <source>Science</source> <volume>348</volume>, <elocation-id>1261605</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1126/SCIENCE.1261605</pub-id>
</citation>
</ref>
<ref id="B18">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Eggert</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Brehm</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Winschel</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Zecha</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Lienhart</surname> <given-names>R.</given-names>
</name>
</person-group> (<year>2017</year>). &#x201c;<article-title>A Closer Look: Small Object Detection in Faster R-CNN</article-title>,&#x201d; in <conf-name>Proceedings of the IEEE International Conference on Multimedia and Expo (ICME)</conf-name>, <fpage>421</fpage>&#x2013;<lpage>426</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1109/ICME.2017.8019550</pub-id>
</citation>
</ref>
<ref id="B19">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Eggert</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Winschel</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Zecha</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Lienhart</surname> <given-names>R.</given-names>
</name>
</person-group> (<year>2016</year>). &#x201c;<article-title>Saliency-Guided Selective Magnification for Company Logo Detection,</article-title>&#x201d; in <conf-name>23rd International Conference on Pattern Recognition (ICPR)</conf-name>, <fpage>651</fpage>&#x2013;<lpage>656</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1109/ICPR.2016.7899708</pub-id>
</citation>
</ref>
<ref id="B20">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Faillettaz</surname> <given-names>R.</given-names>
</name>
<name>
<surname>Picheral</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Luo</surname> <given-names>J. Y.</given-names>
</name>
<name>
<surname>Guigand</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Cowen</surname> <given-names>R. K.</given-names>
</name>
<name>
<surname>Irisson</surname> <given-names>J.-O.</given-names>
</name>
</person-group> (<year>2016</year>). <article-title>Imperfect Automatic Image Classification Successfully Describes Plankton Distribution Patterns</article-title>. <source>Methods Oceanog.</source> <volume>15&#x2013;16</volume>, <fpage>60</fpage>&#x2013;<lpage>77</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/J.MIO.2016.04.003</pub-id>
</citation>
</ref>
<ref id="B21">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Falkowski</surname> <given-names>P.</given-names>
</name>
</person-group> (<year>2012</year>). <article-title>Ocean Science: The Power of Plankton</article-title>. <source>Nature</source> <volume>483</volume>, <fpage>S17</fpage>&#x2013;<lpage>S20</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1038/483S17a</pub-id>
</citation>
</ref>
<ref id="B22">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Field</surname> <given-names>C. B.</given-names>
</name>
<name>
<surname>Behrenfeld</surname> <given-names>M. J.</given-names>
</name>
<name>
<surname>Randerson</surname> <given-names>J. T.</given-names>
</name>
<name>
<surname>Falkowski</surname> <given-names>P.</given-names>
</name>
</person-group> (<year>1998</year>). <article-title>Primary Production of the Biosphere: Integrating Terrestrial and Oceanic Components</article-title>. <source>Science</source> <volume>281</volume>, <fpage>237</fpage>&#x2013;<lpage>240</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1126/science.281.5374.237</pub-id>
</citation>
</ref>
<ref id="B23">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Forest</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Stemmann</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Picheral</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Burdorf</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Robert</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Fortier</surname> <given-names>L.</given-names>
</name>
<etal/>
</person-group>. (<year>2012</year>). <article-title>Size Distribution of Particles and Zooplankton Across the Shelf-Basin System in Southeast Beaufort Sea: Combined Results From an Underwater Vision Profiler and Vertical Net Tows</article-title>. <source>Biogeosciences</source> <volume>9</volume>, <fpage>1301</fpage>&#x2013;<lpage>1320</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.5194/bg-9-1301-2012</pub-id>
</citation>
</ref>
<ref id="B24">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Frederiksen</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Edwards</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Richardson</surname> <given-names>A. J.</given-names>
</name>
<name>
<surname>Halliday</surname> <given-names>N. C.</given-names>
</name>
<name>
<surname>Wanless</surname> <given-names>S.</given-names>
</name>
</person-group> (<year>2006</year>). <article-title>From Plankton to Top Predators: Bottom-Up Control of a Marine Food Web Across Four Trophic Levels</article-title>. <source>J. Anim. Ecol.</source> <volume>75</volume>, <fpage>1259</fpage>&#x2013;<lpage>1268</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1111/j.1365-2656.2006.01148.x</pub-id>
</citation>
</ref>
<ref id="B25">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Giering</surname> <given-names>S. L. C.</given-names>
</name>
<name>
<surname>Hosking</surname> <given-names>B.</given-names>
</name>
<name>
<surname>Briggs</surname> <given-names>N.</given-names>
</name>
<name>
<surname>Iversen</surname> <given-names>M. H.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>The Interpretation of Particle Size, Shape, and Carbon Flux of Marine Particle Images Is Strongly Affected by the Choice of Particle Detection Algorithm</article-title>. <source>Front. Mar. Sci.</source> <volume>7</volume>. doi:&#xa0;<pub-id pub-id-type="doi">10.3389/fmars.2020.00564</pub-id>
</citation>
</ref>
<ref id="B26">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Greer</surname> <given-names>A. T.</given-names>
</name>
<name>
<surname>Boyette</surname> <given-names>A. D.</given-names>
</name>
<name>
<surname>Cruz</surname> <given-names>V. J.</given-names>
</name>
<name>
<surname>Cambazoglu</surname> <given-names>M. K.</given-names>
</name>
<name>
<surname>Dzwonkowski</surname> <given-names>B.</given-names>
</name>
<name>
<surname>Chiaverano</surname> <given-names>L. M.</given-names>
</name>
<etal/>
</person-group>. (<year>2020</year>a). <article-title>Contrasting Fine-Scale Distributional Patterns of Zooplankton Driven by the Formation of a Diatom-Dominated Thin Layer</article-title>. <source>Limnol. Oceanog.</source> <volume>65</volume>, <fpage>2236</fpage>&#x2013;<lpage>2258</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1002/lno.11450</pub-id>
</citation>
</ref>
<ref id="B27">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Greer</surname> <given-names>A. T.</given-names>
</name>
<name>
<surname>Chiaverano</surname> <given-names>L. M.</given-names>
</name>
<name>
<surname>Luo</surname> <given-names>J. Y.</given-names>
</name>
<name>
<surname>Cowen</surname> <given-names>R. K.</given-names>
</name>
<name>
<surname>Graham</surname> <given-names>W. M.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>Ecology and Behaviour of Holoplanktonic Scyphomedusae and Their Interactions With Larval and Juvenile Fishes in the Northern Gulf of Mexico</article-title>. <source>ICES. J. Mar. Sci.</source> <volume>75</volume>, <fpage>751</fpage>&#x2013;<lpage>763</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1093/icesjms/fsx168</pub-id>
</citation>
</ref>
<ref id="B28">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Greer</surname> <given-names>A. T.</given-names>
</name>
<name>
<surname>Chiaverano</surname> <given-names>L. M.</given-names>
</name>
<name>
<surname>Treible</surname> <given-names>L. M.</given-names>
</name>
<name>
<surname>Brise&#xf1;o-Avena</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Hernandez</surname> <given-names>F. J.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>From Spatial Pattern to Ecological Process Through Imaging Zooplankton Interactions</article-title>. <source>ICES. J. Mar. Sci</source> <volume>78</volume>(<issue>8</issue>):<page-range>2664&#x2013;74</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.1093/icesjms/fsab149</pub-id>
</citation>
</ref>
<ref id="B29">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Greer</surname> <given-names>A. T.</given-names>
</name>
<name>
<surname>Cowen</surname> <given-names>R. K.</given-names>
</name>
<name>
<surname>Guigand</surname> <given-names>C. M.</given-names>
</name>
<name>
<surname>Hare</surname> <given-names>J. A.</given-names>
</name>
</person-group> (<year>2015</year>). <article-title>Fine-Scale Planktonic Habitat Partitioning at a Shelf-Slope Front Revealed by a High-Resolution Imaging System</article-title>. <source>J. Mar. Syst.</source> <volume>142</volume>, <fpage>111</fpage>&#x2013;<lpage>125</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.jmarsys.2014.10.008</pub-id>
</citation>
</ref>
<ref id="B30">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Greer</surname> <given-names>A. T.</given-names>
</name>
<name>
<surname>Cowen</surname> <given-names>R. K.</given-names>
</name>
<name>
<surname>Guigand</surname> <given-names>C. M.</given-names>
</name>
<name>
<surname>Hare</surname> <given-names>J. A.</given-names>
</name>
<name>
<surname>Tang</surname> <given-names>D.</given-names>
</name>
</person-group> (<year>2014</year>). <article-title>The Role of Internal Waves in Larval Fish Interactions With Potential Predators and Prey</article-title>. <source>Prog. Oceanog.</source> <volume>127</volume>, <fpage>47</fpage>&#x2013;<lpage>61</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.pocean.2014.05.010</pub-id>
</citation>
</ref>
<ref id="B31">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Greer</surname> <given-names>A. T.</given-names>
</name>
<name>
<surname>Cowen</surname> <given-names>R. K.</given-names>
</name>
<name>
<surname>Guigand</surname> <given-names>C. M.</given-names>
</name>
<name>
<surname>McManus</surname> <given-names>M. A.</given-names>
</name>
<name>
<surname>Sevadjian</surname> <given-names>J. C.</given-names>
</name>
<name>
<surname>Timmerman</surname> <given-names>A. H. V.</given-names>
</name>
</person-group> (<year>2013</year>). <article-title>Relationships Between Phytoplankton Thin Layers and the Fine-Scale Vertical Distributions of Two Trophic Levels of Zooplankton</article-title>. <source>J. Plank. Res.</source> <volume>35</volume>, <fpage>939</fpage>&#x2013;<lpage>956</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1093/plankt/fbt056</pub-id>
</citation>
</ref>
<ref id="B32">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Greer</surname> <given-names>A. T.</given-names>
</name>
<name>
<surname>Lehrter</surname> <given-names>J. C.</given-names>
</name>
<name>
<surname>Binder</surname> <given-names>B. M.</given-names>
</name>
<name>
<surname>Nayak</surname> <given-names>A. R.</given-names>
</name>
<name>
<surname>Barua</surname> <given-names>R.</given-names>
</name>
<name>
<surname>Rice</surname> <given-names>A. E.</given-names>
</name>
<etal/>
</person-group>. (<year>2020</year>b). <article-title>High-Resolution Sampling of a Broad Marine Life Size Spectrum Reveals Differing Size- and Composition-Based Associations With Physical Oceanographic Structure</article-title>. <source>Front. Mar. Sci.</source> <volume>7</volume> (<volume>8</volume>), <elocation-id>2664</elocation-id>&#x2013;<lpage>2674</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.3389/fmars.2020.542701</pub-id>
</citation>
</ref>
<ref id="B33">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Guidi</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Jackson</surname> <given-names>G. A.</given-names>
</name>
<name>
<surname>Stemmann</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Miquel</surname> <given-names>J. C.</given-names>
</name>
<name>
<surname>Picheral</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Gorsky</surname> <given-names>G.</given-names>
</name>
</person-group> (<year>2008</year>). <article-title>Relationship Between Particle Size Distribution and Flux in the Mesopelagic Zone</article-title>. <source>Deep. Sea. Res. Part I.: Oceanog. Res. Pap.</source> <volume>55</volume>, <fpage>1364</fpage>&#x2013;<lpage>1374</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/J.DSR.2008.05.014</pub-id>
</citation>
</ref>
<ref id="B34">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Guidi</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Legendre</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Reygondeau</surname> <given-names>G.</given-names>
</name>
<name>
<surname>Uitz</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Stemmann</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Henson</surname> <given-names>S. A.</given-names>
</name>
</person-group> (<year>2015</year>). <article-title>A New Look at Ocean Carbon Remineralization for Estimating Deepwater Sequestration</article-title>. <source>Global Biogeochem. Cycle.</source> <volume>29</volume>, <fpage>1044</fpage>&#x2013;<lpage>1059</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1002/2014GB005063</pub-id>
</citation>
</ref>
<ref id="B35">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>He</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Gkioxari</surname> <given-names>G.</given-names>
</name>
<name>
<surname>Dollar</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Girshick</surname> <given-names>R.</given-names>
</name>
<collab>Mask R-CNN</collab>
</person-group> (<year>2017</year>). <conf-name>Proceedings of the IEEE International Conference on Computer Vision (ICCV)</conf-name>. <fpage>2961</fpage>&#x2013;<lpage>2969</lpage>.</citation>
</ref>
<ref id="B36">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ibarbalz</surname> <given-names>F. M.</given-names>
</name>
<name>
<surname>Henry</surname> <given-names>N.</given-names>
</name>
<name>
<surname>Brand&#xe3;o</surname> <given-names>M. C.</given-names>
</name>
<name>
<surname>Martini</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Busseni</surname> <given-names>G.</given-names>
</name>
<name>
<surname>Byrne</surname> <given-names>H.</given-names>
</name>
<etal/>
</person-group>. (<year>2019</year>). <article-title>Global Trends in Marine Plankton Diversity Across Kingdoms of Life</article-title>. <source>Cell</source> <volume>179</volume>, <fpage>1084</fpage>&#x2013;<lpage>1097.e21</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.cell.2019.10.008</pub-id>
</citation>
</ref>
<ref id="B37">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Irisson</surname> <given-names>J.-O.</given-names>
</name>
<name>
<surname>Ayata</surname> <given-names>S.-D.</given-names>
</name>
<name>
<surname>Lindsay</surname> <given-names>D. J.</given-names>
</name>
<name>
<surname>Karp-Boss</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Stemmann</surname> <given-names>L.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Machine Learning for the Study of Plankton and Marine Snow From Images</article-title>. <source>Ann. Rev. Mar. Sci.</source> <volume>14</volume>, <fpage>277</fpage>&#x2013;<lpage>301</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1146/annurev-marine-041921-013023</pub-id>
</citation>
</ref>
<ref id="B38">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Iyer</surname> <given-names>N.</given-names>
</name>
</person-group> (<year>2012</year>). <source>Machine Vision Assisted <italic>in Situ</italic> Ichthyoplankton Imaging System</source> (<publisher-name>Purdue University</publisher-name>).</citation>
</ref>
<ref id="B39">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Lee</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Park</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Kim</surname> <given-names>J.</given-names>
</name>
</person-group> (<year>2016</year>). &#x201c;<article-title>Plankton Classification on Imbalanced Large Scale Database <italic>via</italic> Convolutional Neural Networks With Transfer Learning,</article-title>&#x201d; in <conf-name>IEEE International Conference on Image Processing (ICIP)</conf-name>, <fpage>3713</fpage>&#x2013;<lpage>3717</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1109/ICIP.2016.7533053</pub-id>
</citation>
</ref>
<ref id="B40">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>L&#xe9;vy</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Franks</surname> <given-names>P. J. S.</given-names>
</name>
<name>
<surname>Smith</surname> <given-names>K. S.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>The Role of Submesoscale Currents in Structuring Marine Ecosystems</article-title>. <source>Nat. Commun.</source> <volume>9</volume>, <fpage>4758</fpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1038/s41467-018-07059-3</pub-id>
</citation>
</ref>
<ref id="B41">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lin</surname> <given-names>T.-Y.</given-names>
</name>
<name>
<surname>Doll&#xe1;r</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Girshick</surname> <given-names>R.</given-names>
</name>
<name>
<surname>He</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Hariharan</surname> <given-names>B.</given-names>
</name>
<name>
<surname>Belongie</surname> <given-names>S.</given-names>
</name>
</person-group> (<year>2017</year>). &#x201c;<article-title>Feature Pyramid Networks for Object Detection,</article-title>&#x201d; in <conf-name>Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition (CVPR)</conf-name>, <fpage>1612.03144</fpage>.</citation>
</ref>
<ref id="B42">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lombard</surname> <given-names>F.</given-names>
</name>
<name>
<surname>Boss</surname> <given-names>E.</given-names>
</name>
<name>
<surname>Waite</surname> <given-names>A. M.</given-names>
</name>
<name>
<surname>Vogt</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Uitz</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Stemmann</surname> <given-names>L.</given-names>
</name>
<etal/>
</person-group>. (<year>2019</year>). <article-title>Globally Consistent Quantitative Observations of Planktonic Ecosystems</article-title>. <source>Front. Mar. Sci.</source> <volume>6</volume>, <elocation-id>196</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.3389/fmars.2019.00196</pub-id>
</citation>
</ref>
<ref id="B43">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Longhurst</surname> <given-names>A. R.</given-names>
</name>
<name>
<surname>Glen Harrison</surname> <given-names>W.</given-names>
</name>
</person-group> (<year>1989</year>). <article-title>The Biological Pump: Profiles of Plankton Production and Consumption in the Upper Ocean</article-title>. <source>Prog. Oceanog.</source> <volume>22</volume>, <fpage>47</fpage>&#x2013;<lpage>123</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/0079-6611(89)90010-4</pub-id>
</citation>
</ref>
<ref id="B44">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Luo</surname> <given-names>J. Y.</given-names>
</name>
<name>
<surname>Grassian</surname> <given-names>B.</given-names>
</name>
<name>
<surname>Tang</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Irisson</surname> <given-names>J.-O.</given-names>
</name>
<name>
<surname>Greer</surname> <given-names>A. T.</given-names>
</name>
<name>
<surname>Guigand</surname> <given-names>C. M.</given-names>
</name>
<etal/>
</person-group>. (<year>2014</year>). <article-title>Environmental Drivers of the Fine-Scale Distribution of a Gelatinous Zooplankton Community Across a Mesoscale Front</article-title>. <source>Mar. Ecol. Prog. Ser.</source> <volume>510</volume>, <fpage>129</fpage>&#x2013;<lpage>149</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.3354/meps10908</pub-id>
</citation>
</ref>
<ref id="B45">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Luo</surname> <given-names>J. Y.</given-names>
</name>
<name>
<surname>Irisson</surname> <given-names>J.-O.</given-names>
</name>
<name>
<surname>Graham</surname> <given-names>B.</given-names>
</name>
<name>
<surname>Guigand</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Sarafraz</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Mader</surname> <given-names>C.</given-names>
</name>
<etal/>
</person-group>. (<year>2018</year>). <article-title>Automated Plankton Image Analysis Using Convolutional Neural Networks</article-title>. <source>Limnol. Oceanog.: Methods</source> <volume>16</volume>, <fpage>814</fpage>&#x2013;<lpage>827</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1002/lom3.10285</pub-id>
</citation>
</ref>
<ref id="B46">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Matas</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Chum</surname> <given-names>O.</given-names>
</name>
<name>
<surname>Urban</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Pajdla</surname> <given-names>T.</given-names>
</name>
</person-group> (<year>2004</year>). <article-title>Robust Wide-Baseline Stereo From Maximally Stable Extremal Regions</article-title>. <source>Imag. Vision Comput.</source> <volume>22</volume>, <fpage>761</fpage>&#x2013;<lpage>767</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.imavis.2004.02.006</pub-id>
</citation>
</ref>
<ref id="B47">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>McClatchie</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Cowen</surname> <given-names>R.</given-names>
</name>
<name>
<surname>Nieto</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Greer</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Luo</surname> <given-names>J. Y.</given-names>
</name>
<name>
<surname>Guigand</surname> <given-names>C.</given-names>
</name>
<etal/>
</person-group>. (<year>2012</year>). <article-title>Resolution of Fine Biological Structure Including Small Narcomedusae Across a Front in the Southern California Bight</article-title>. <source>J. Geophys. Res.: Ocean.</source> <volume>117</volume>, <fpage>C04020</fpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1029/2011JC007565</pub-id>
</citation>
</ref>
<ref id="B48">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Moreno-Torres</surname> <given-names>J. G.</given-names>
</name>
<name>
<surname>Raeder</surname> <given-names>T.</given-names>
</name>
<name>
<surname>Alaiz-Rodr&#xed;guez</surname> <given-names>R.</given-names>
</name>
<name>
<surname>Chawla</surname> <given-names>N. V.</given-names>
</name>
<name>
<surname>Herrera</surname> <given-names>F.</given-names>
</name>
</person-group> (<year>2012</year>). <article-title>A Unifying View on Dataset Shift in Classification</article-title>. <source>Pattern Recognit.</source> <volume>45</volume>, <fpage>521</fpage>&#x2013;<lpage>530</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.patcog.2011.06.019</pub-id>
</citation>
</ref>
<ref id="B49">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ohman</surname> <given-names>M. D.</given-names>
</name>
<name>
<surname>Davis</surname> <given-names>R. E.</given-names>
</name>
<name>
<surname>Sherman</surname> <given-names>J. T.</given-names>
</name>
<name>
<surname>Grindley</surname> <given-names>K. R.</given-names>
</name>
<name>
<surname>Whitmore</surname> <given-names>B. M.</given-names>
</name>
<name>
<surname>Nickels</surname> <given-names>C. F.</given-names>
</name>
<etal/>
</person-group>. (<year>2019</year>). <article-title>Zooglider: An Autonomous Vehicle for Optical and Acoustic Sensing of Zooplankton</article-title>. <source>Limnol. Oceanog.: Methods</source> <volume>17</volume>, <fpage>69</fpage>&#x2013;<lpage>86</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1002/lom3.10301</pub-id>
</citation>
</ref>
<ref id="B50">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Olson</surname> <given-names>R. J.</given-names>
</name>
<name>
<surname>Sosik</surname> <given-names>H. M.</given-names>
</name>
</person-group> (<year>2007</year>). <article-title>A Submersible Imaging-in-Flow Instrument to Analyze Nano-and Microplankton: Imaging FlowCytobot</article-title>. <source>Limnol. Oceanog.: Methods</source> <volume>5</volume>, <fpage>195</fpage>&#x2013;<lpage>203</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.4319/lom.2007.5.195</pub-id>
</citation>
</ref>
<ref id="B51">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Orenstein</surname> <given-names>E. C.</given-names>
</name>
<name>
<surname>Ayata</surname> <given-names>S.-D.</given-names>
</name>
<name>
<surname>Maps</surname> <given-names>F.</given-names>
</name>
<name>
<surname>Biard</surname> <given-names>T.</given-names>
</name>
<name>
<surname>Becker</surname> <given-names>&#xc9;.C.</given-names>
</name>
<name>
<surname>Benedetti</surname> <given-names>F.</given-names>
</name>
<etal/>
</person-group>. (<year>2021</year>). <article-title>Machine Learning Techniques to Characterize Functional Traits of Plankton From Image Data</article-title>. <fpage>hal&#x2013;03482282</fpage>
</citation>
</ref>
<ref id="B52">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Orenstein</surname> <given-names>E. C.</given-names>
</name>
<name>
<surname>Ratelle</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Brise&#xf1;o-Avena</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Carter</surname> <given-names>M. L.</given-names>
</name>
<name>
<surname>Franks</surname> <given-names>P. J. S.</given-names>
</name>
<name>
<surname>Jaffe</surname> <given-names>J. S.</given-names>
</name>
<etal/>
</person-group>. (<year>2020</year>). <article-title>The Scripps Plankton Camera System: A Framework and Platform for <italic>in Situ</italic> Microscopy</article-title>. <source>Limnol. Oceanog.: Methods</source> <volume>18</volume>, <fpage>681</fpage>&#x2013;<lpage>695</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1002/lom3.10394</pub-id>
</citation>
</ref>
<ref id="B53">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Otsu</surname> <given-names>N.</given-names>
</name>
</person-group> (<year>1979</year>). <article-title>A Threshold Selection Method From Gray-Level Histograms</article-title>. <source>IEEE Trans. Sys. Man. Cybernet.</source> <volume>9</volume>, <fpage>62</fpage>&#x2013;<lpage>66</lpage>.</citation>
</ref>
<ref id="B54">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Parikh</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Zitnick</surname> <given-names>C. L.</given-names>
</name>
<name>
<surname>Chen</surname> <given-names>T.</given-names>
</name>
</person-group> (<year>2012</year>). <article-title>Exploring Tiny Images: The Roles of Appearance and Contextual Information for Machine and Human Object Recognition</article-title>. <source>IEEE Trans. Pattern Anal. Mach. Intell.</source> <volume>34</volume>, <fpage>1978</fpage>&#x2013;<lpage>1991</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1109/TPAMI.2011.276</pub-id>
</citation>
</ref>
<ref id="B55">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Picheral</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Colin</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Irisson</surname> <given-names>J.-O.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>EcoTaxa, a Tool for the Taxonomic Classification of Images</article-title>.</citation>
</ref>
<ref id="B56">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Picheral</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Guidi</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Stemmann</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Karl</surname> <given-names>D. M.</given-names>
</name>
<name>
<surname>Iddaoud</surname> <given-names>G.</given-names>
</name>
<name>
<surname>Gorsky</surname> <given-names>G.</given-names>
</name>
</person-group> (<year>2010</year>). <article-title>The Underwater Vision Profiler 5: An Advanced Instrument for High Spatial Resolution Studies of Particle Size Spectra and Zooplankton</article-title>. <source>Limnol. Oceanog.: Methods</source> <volume>8</volume>, <fpage>462</fpage>&#x2013;<lpage>473</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.4319/lom.2010.8.462</pub-id>
</citation>
</ref>
<ref id="B57">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Remsen</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Hopkins</surname> <given-names>T. L.</given-names>
</name>
<name>
<surname>Samson</surname> <given-names>S.</given-names>
</name>
</person-group> (<year>2004</year>). <article-title>What You See is Not What You Catch: A Comparison of Concurrently Collected Net, Optical Plankton Counter, and Shadowed Image Particle Profiling Evaluation Recorder Data From the Northeast Gulf of Mexico</article-title>. <source>Deep. Sea. Res. Part I.: Oceanog. Res. Pap.</source> <volume>51</volume>, <fpage>129</fpage>&#x2013;<lpage>151</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/J.DSR.2003.09.008</pub-id>
</citation>
</ref>
<ref id="B58">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Robinson</surname> <given-names>K. L.</given-names>
</name>
<name>
<surname>Sponaugle</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Luo</surname> <given-names>J. Y.</given-names>
</name>
<name>
<surname>Gleiber</surname> <given-names>M. R.</given-names>
</name>
<name>
<surname>Cowen</surname> <given-names>R. K.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Big or Small, Patchy All: Resolution of Marine Plankton Patch Structure at Micro- to Submesoscales for 36 Taxa</article-title>. <source>Sci. Adv.</source> <volume>7</volume>, <fpage>eabk2904</fpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1126/sciadv.abk2904</pub-id>
</citation>
</ref>
<ref id="B59">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Rombouts</surname> <given-names>I.</given-names>
</name>
<name>
<surname>Beaugrand</surname> <given-names>G.</given-names>
</name>
<name>
<surname>Iba&#x148;ez</surname> <given-names>F.</given-names>
</name>
<name>
<surname>Gasparini</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Chiba</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Legendre</surname> <given-names>L.</given-names>
</name>
</person-group> (<year>2009</year>). <article-title>Global Latitudinal Variations in Marine Copepod Diversity and Environmental Factors</article-title>. <source>Proc. R. Soc. B.: Biol. Sci.</source> <volume>276</volume>, <fpage>3053</fpage>&#x2013;<lpage>3062</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1098/rspb.2009.0742</pub-id>
</citation>
</ref>
<ref id="B60">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Rutherford</surname> <given-names>S.</given-names>
</name>
<name>
<surname>D&#x2019;Hondt</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Prell</surname> <given-names>W.</given-names>
</name>
</person-group> (<year>1999</year>). <article-title>Environmental Controls on the Geographic Distribution of Zooplankton Diversity</article-title>. <source>Nature</source> <volume>400</volume>, <fpage>749</fpage>&#x2013;<lpage>753</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1038/23449</pub-id>
</citation>
</ref>
<ref id="B61">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Schmid</surname> <given-names>M. S.</given-names>
</name>
<name>
<surname>Cowen</surname> <given-names>R. K.</given-names>
</name>
<name>
<surname>Robinson</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Luo</surname> <given-names>J. Y.</given-names>
</name>
<name>
<surname>Brise&#xf1;o-Avena</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Sponaugle</surname> <given-names>S.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Prey and Predator Overlap at the Edge of a Mesoscale Eddy: Fine-Scale, <italic>in-Situ</italic> Distributions to Inform Our Understanding of Oceanographic Processes</article-title>. <source>Sci. Rep.</source> <volume>10</volume>, <fpage>1</fpage>&#x2013;<lpage>16</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1038/s41598-020-57879-x</pub-id> schmid2020Prey
</citation>
</ref>
<ref id="B62">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Schmid</surname> <given-names>M. S.</given-names>
</name>
<name>
<surname>Daprano</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Jacobson</surname> <given-names>K. M.</given-names>
</name>
<name>
<surname>Sullivan</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Brise&#xf1;o-Avena</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Luo</surname> <given-names>J. Y.</given-names>
</name>
<etal/>
</person-group>. (<year>2021</year>). <article-title>A Convolutional Neural Network Based High-Throughput Image Classification Pipeline - Code and Documentation to Process Plankton Underwater Imagery Using Local HPC Infrastructure and NSF&#x2019;s XSEDE</article-title>. <source>Zenodo</source>. doi:&#xa0;<pub-id pub-id-type="doi">10.5281/zenodo.4641158</pub-id>
</citation>
</ref>
<ref id="B63">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Schmid</surname> <given-names>M. S.</given-names>
</name>
<name>
<surname>Fortier</surname> <given-names>L.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>The Intriguing Co-Distribution of the Copepods Calanus Hyperboreus and Calanus Glacialis in the Subsurface Chlorophyll Maximum of Arctic Seas</article-title>. <source>Element.: Sci. Anthropocene.</source> <volume>7</volume>, <fpage>50</fpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1525/elementa.388</pub-id>
</citation>
</ref>
<ref id="B64">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ser-Giacomi</surname> <given-names>E.</given-names>
</name>
<name>
<surname>Zinger</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Malviya</surname> <given-names>S.</given-names>
</name>
<name>
<surname>De Vargas</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Karsenti</surname> <given-names>E.</given-names>
</name>
<name>
<surname>Bowler</surname> <given-names>C.</given-names>
</name>
<etal/>
</person-group>. (<year>2018</year>). <article-title>Ubiquitous Abundance Distribution of non-Dominant Plankton Across the Global Ocean</article-title>. <source>Nat. Ecol. Evol.</source> <volume>2</volume>, <fpage>1243</fpage>&#x2013;<lpage>1249</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1038/s41559-018-0587-2</pub-id>
</citation>
</ref>
<ref id="B65">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sheldon</surname> <given-names>R. W.</given-names>
</name>
<name>
<surname>Parsons</surname> <given-names>T. R.</given-names>
</name>
</person-group> (<year>1967</year>). <article-title>A Continuous Size Spectrum for Particulate Matter in the Sea</article-title>. <source>J. Fish. Res. Board. Canada.</source> <volume>24</volume>, <fpage>909</fpage>&#x2013;<lpage>915</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1139/f67-081</pub-id>
</citation>
</ref>
<ref id="B66">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sheldon</surname> <given-names>R. W.</given-names>
</name>
<name>
<surname>Prakash</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Sutcliffe</surname> <given-names>W. H.</given-names>
</name>
</person-group> (<year>1972</year>). <article-title>The Size Distribution of Particles in the Ocean</article-title>. <source>Limnol. Oceanog.</source> <volume>17</volume>, <fpage>327</fpage>&#x2013;<lpage>340</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.4319/lo.1972.17.3.0327</pub-id>
</citation>
</ref>
<ref id="B67">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sieracki</surname> <given-names>C. K.</given-names>
</name>
<name>
<surname>Sieracki</surname> <given-names>M. E.</given-names>
</name>
<name>
<surname>Yentsch</surname> <given-names>C. S.</given-names>
</name>
</person-group> (<year>1998</year>). <article-title>An Imaging-in-Flow System for Automated Analysis of Marine Microplankton</article-title>. <source>Mar. Ecol. Prog. Ser.</source> <volume>168</volume>, <fpage>285</fpage>&#x2013;<lpage>296</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.3354/meps168285</pub-id>
</citation>
</ref>
<ref id="B68">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sosik</surname> <given-names>H. M.</given-names>
</name>
<name>
<surname>Olson</surname> <given-names>R. J.</given-names>
</name>
</person-group> (<year>2007</year>). <article-title>Automated Taxonomic Classification of Phytoplankton Sampled With Imaging-in-Flow Cytometry</article-title>. <source>Limnol. Oceanog.: Methods</source> <volume>5</volume>, <fpage>204</fpage>&#x2013;<lpage>216</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.4319/lom.2007.5.204</pub-id>
</citation>
</ref>
<ref id="B69">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Stemmann</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Boss</surname> <given-names>E.</given-names>
</name>
</person-group> (<year>2012</year>). <article-title>Plankton and Particle Size and Packaging: From Determining Optical Properties to Driving the Biological Pump</article-title>. <source>Annu. Rev. Mar. Sci.</source> <volume>4</volume>, <fpage>263</fpage>&#x2013;<lpage>290</lpage>.</citation>
</ref>
<ref id="B70">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Stemmann</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Hosia</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Youngbluth</surname> <given-names>M. J.</given-names>
</name>
<name>
<surname>S&#xf8;iland</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Picheral</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Gorsky</surname> <given-names>G.</given-names>
</name>
</person-group> (<year>2008</year>). <article-title>Vertical Distribution (0&#x2013;1000 M) of Macrozooplankton, Estimated Using the Underwater Video Profiler, in Different Hydrographic Regimes Along the Northern Portion of the Mid-Atlantic Ridge</article-title>. <source>Deep. Sea. Res. Part II.: Top. Stud. Oceanog.</source> <volume>55</volume>, <fpage>94</fpage>&#x2013;<lpage>105</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/J.DSR2.2007.09.019</pub-id>
</citation>
</ref>
<ref id="B71">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Stemmann</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Picheral</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Gorsky</surname> <given-names>G.</given-names>
</name>
</person-group> (<year>2000</year>). <article-title>Diel Variation in the Vertical Distribution of Particulate Matter (&gt;0.15mm) in the NW Mediterranean Sea Investigated With the Underwater Video Profiler</article-title>. <source>Deep. Sea. Res. Part I.: Oceanog. Res. Pap.</source> <volume>47</volume>, <fpage>505</fpage>&#x2013;<lpage>531</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/S0967-0637(99)00100-4</pub-id>
</citation>
</ref>
<ref id="B72">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Swieca</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Sponaugle</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Brise&#xf1;o-Avena</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Schmid</surname> <given-names>M. S.</given-names>
</name>
<name>
<surname>Brodeur</surname> <given-names>R. D.</given-names>
</name>
<name>
<surname>Cowen</surname> <given-names>R. K.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Changing With the Tides: Fine-Scale Larval Fish Prey Availability and Predation Pressure Near a Tidally Modulated River Plume</article-title>. <source>Mar. Ecol. Prog. Ser.</source> <volume>650</volume>, <fpage>217</fpage>&#x2013;<lpage>238</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.3354/meps13367</pub-id>
</citation>
</ref>
<ref id="B73">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tittensor</surname> <given-names>D. P.</given-names>
</name>
<name>
<surname>Mora</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Jetz</surname> <given-names>W.</given-names>
</name>
<name>
<surname>Lotze</surname> <given-names>H. K.</given-names>
</name>
<name>
<surname>Ricard</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Berghe</surname> <given-names>E. V.</given-names>
</name>
<etal/>
</person-group>. (<year>2010</year>). <article-title>Global Patterns and Predictors of Marine Biodiversity Across Taxa</article-title>. <source>Nature</source> <volume>466</volume>, <fpage>1098</fpage>&#x2013;<lpage>1101</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1038/nature09329</pub-id>
</citation>
</ref>
<ref id="B74">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tsechpenakis</surname> <given-names>G.</given-names>
</name>
<name>
<surname>Guigand</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Cowen</surname> <given-names>R. K.</given-names>
</name>
</person-group> (<year>2007</year>). <article-title>Image Analysis Techniques to Accompany a New <italic>In Situ</italic> Ichthyoplankton Imaging System</article-title>. <conf-name>OCEANS 2007 - Europe</conf-name>, <fpage>1</fpage>&#x2013;<lpage>6</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1109/OCEANSE.2007.4302271</pub-id>
</citation>
</ref>
<ref id="B75">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Wang</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Chang</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Yang</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Liu</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Huang</surname> <given-names>T. S.</given-names>
</name>
</person-group> (<year>2016</year>). &#x201c;<article-title>Studying Very Low Resolution Recognition Using Deep Networks</article-title>,&#x201d; in <conf-name>Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition (CVPR)</conf-name>, <fpage>4792</fpage>&#x2013;<lpage>4800</lpage>.</citation>
</ref>
<ref id="B76">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ware</surname> <given-names>D. M.</given-names>
</name>
<name>
<surname>Thomson</surname> <given-names>R. E.</given-names>
</name>
</person-group> (<year>2005</year>). <article-title>Bottom-Up Ecosystem Trophic Dynamics Determine Fish Production in the Northeast Pacific</article-title>. <source>Science</source> <volume>308</volume>, <fpage>1280</fpage>&#x2013;<lpage>1284</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1126/SCIENCE.1109049</pub-id>
</citation>
</ref>
<ref id="B77">
<citation citation-type="other">
<person-group person-group-type="author">
<name>
<surname>Wu</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Kirillov</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Massa</surname> <given-names>F.</given-names>
</name>
<name>
<surname>Lo</surname> <given-names>W.-Y.</given-names>
</name>
<name>
<surname>Girshick</surname> <given-names>R.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Detectron2</article-title>.</citation>
</ref>
</ref-list>
</back>
</article>