<?xml version="1.0" encoding="UTF-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article article-type="research-article" dtd-version="2.3" xml:lang="EN" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Genet.</journal-id>
<journal-title>Frontiers in Genetics</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Genet.</abbrev-journal-title>
<issn pub-type="epub">1664-8021</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="publisher-id">864092</article-id>
<article-id pub-id-type="doi">10.3389/fgene.2022.864092</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Genetics</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Neural Networks for Classification and Image Generation of Aging in Genetic Syndromes</article-title>
<alt-title alt-title-type="left-running-head">Duong et al.</alt-title>
<alt-title alt-title-type="right-running-head">Neural Networks in Syndromic Aging</alt-title>
</title-group>
<contrib-group>
<contrib contrib-type="author">
<name>
<surname>Duong</surname>
<given-names>Dat</given-names>
</name>
<uri xlink:href="https://loop.frontiersin.org/people/1689338/overview"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Hu</surname>
<given-names>Ping</given-names>
</name>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Tekendo-Ngongang</surname>
<given-names>Cedrik</given-names>
</name>
<uri xlink:href="https://loop.frontiersin.org/people/642262/overview"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Hanchard</surname>
<given-names>Suzanna E. Ledgister</given-names>
</name>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Liu</surname>
<given-names>Simon</given-names>
</name>
<uri xlink:href="https://loop.frontiersin.org/people/1697227/overview"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Solomon</surname>
<given-names>Benjamin D.</given-names>
</name>
<xref ref-type="corresp" rid="c001">&#x2a;</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1670757/overview"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Waikel</surname>
<given-names>Rebekah L.</given-names>
</name>
<xref ref-type="corresp" rid="c001">&#x2a;</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1653255/overview"/>
</contrib>
</contrib-group>
<aff>
<institution>Medical Genomics Unit</institution>, <institution>National Human Genome Research Institute</institution>, <addr-line>Bethesda</addr-line>, <addr-line>MD</addr-line>, <country>United States</country>
</aff>
<author-notes>
<fn fn-type="edited-by">
<p>
<bold>Edited by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/578052/overview">Gavin R. Oliver</ext-link>, Mayo Clinic, United States</p>
</fn>
<fn fn-type="edited-by">
<p>
<bold>Reviewed by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/478805/overview">Manhua Liu</ext-link>, Shanghai Jiao Tong University, China</p>
<p>
<ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/307438/overview">Jan Egger</ext-link>, University Hospital Essen, Germany</p>
</fn>
<corresp id="c001">&#x2a;Correspondence: Benjamin D. Solomon, <email>solomonb@mail.nih.gov</email>; Rebekah L. Waikel, <email>rebekah.waikel@nih.gov</email>
</corresp>
<fn fn-type="other">
<p>This article was submitted to Human and Medical Genomics, a section of the journal Frontiers in Genetics</p>
</fn>
</author-notes>
<pub-date pub-type="epub">
<day>11</day>
<month>04</month>
<year>2022</year>
</pub-date>
<pub-date pub-type="collection">
<year>2022</year>
</pub-date>
<volume>13</volume>
<elocation-id>864092</elocation-id>
<history>
<date date-type="received">
<day>28</day>
<month>01</month>
<year>2022</year>
</date>
<date date-type="accepted">
<day>28</day>
<month>02</month>
<year>2022</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#xa9; 2022 Duong, Hu, Tekendo-Ngongang, Hanchard, Liu, Solomon and Waikel.</copyright-statement>
<copyright-year>2022</copyright-year>
<copyright-holder>Duong, Hu, Tekendo-Ngongang, Hanchard, Liu, Solomon and Waikel</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<abstract>
<p>
<bold>Background:</bold> In medical genetics, one application of neural networks is the diagnosis of genetic diseases based on images of patient faces. While these applications have been validated in the literature with primarily pediatric subjects, it is not known whether these applications can accurately diagnose patients across a lifespan. We aimed to extend previous works to determine whether age plays a factor in facial diagnosis as well as to explore other factors that may contribute to the overall diagnostic accuracy.</p>
<p>
<bold>Methods:</bold> To investigate this, we chose two relatively common conditions, Williams syndrome and 22q11.2 deletion syndrome. We built a neural network classifier trained on images of affected and unaffected individuals of different ages and compared classifier accuracy to clinical geneticists. We analyzed the results of saliency maps and the use of generative adversarial networks to boost accuracy.</p>
<p>
<bold>Results:</bold> Our classifier outperformed clinical geneticists at recognizing face images of these two conditions within each of the age groups (the performance varied between the age groups): 1) under 2&#xa0;years old, 2) 2&#x2013;9&#xa0;years old, 3) 10&#x2013;19&#xa0;years old, 4) 20&#x2013;34&#xa0;years old, and 5) &#x2265;35&#xa0;years old. The overall accuracy improvement by our classifier over the clinical geneticists was 15.5 and 22.7% for Williams syndrome and 22q11.2 deletion syndrome, respectively. Additionally, comparison of saliency maps revealed that key facial features learned by the neural network differed with respect to age. Finally, joint training real images with multiple different types of fake images created by a <ext-link ext-link-type="uri" xlink:href="https://en.wikipedia.org/wiki/Generative_adversarial_network">generative adversarial network showed</ext-link> up to 3.25% accuracy gain in classification accuracy.</p>
<p>
<bold>Conclusion:</bold> The ability of clinical geneticists to diagnose these conditions is influenced by the age of the patient. Deep learning technologies such as our classifier can more accurately identify patients across the lifespan based on facial features. Saliency maps of computer vision reveal that the syndromic facial feature attributes change with the age of the patient. Modest improvements in the classifier accuracy were observed when joint training was carried out with both real and fake images. Our findings highlight the need for a greater focus on age as a confounder in facial diagnosis.</p>
</abstract>
<kwd-group>
<kwd>deep learning</kwd>
<kwd>generative adversarial networks</kwd>
<kwd>22q11.2 deletion syndrome</kwd>
<kwd>aging</kwd>
<kwd>Williams syndrome</kwd>
<kwd>facial recognition</kwd>
<kwd>facial diagnosis</kwd>
</kwd-group>
</article-meta>
</front>
<body>
<sec id="s1">
<title>Background</title>
<p>Neural networks are emerging as powerful tools in many areas of biomedical research and are starting to impact clinical care. In the field of genomics, these methods are applied in multiple ways, including generating differential diagnoses for patients with a possible genetic syndrome based on images, (<xref ref-type="bibr" rid="B15">Gurovich et al., 2019</xref>; <xref ref-type="bibr" rid="B17">Hsieh et al., 2021</xref>; <xref ref-type="bibr" rid="B30">Porras et al., 2021</xref>), analysis of DNA sequencing data (<xref ref-type="bibr" rid="B23">Luo et al., 2019</xref>) including phenotype-based annotation (<xref ref-type="bibr" rid="B7">Clark et al., 2019</xref>) and variant classification (<xref ref-type="bibr" rid="B13">Frazer et al., 2021</xref>), and prediction of the protein structure (<xref ref-type="bibr" rid="B2">Baek et al., 2021</xref>; <xref ref-type="bibr" rid="B19">Jumper et al., 2021</xref>).</p>
<p>In the field of clinical genetics, clinicians typically encounter many different conditions that are individually rare and which can be difficult to differentiate (<xref ref-type="bibr" rid="B34">Solomon et al., 2013</xref>; <xref ref-type="bibr" rid="B11">Ferreira, 2019</xref>). This complexity, coupled with a lack of trained experts (<xref ref-type="bibr" rid="B24">Maiese et al., 2019</xref>; <xref ref-type="bibr" rid="B18">Jenkins et al., 2021</xref>), can lead to delayed diagnosis and suboptimal management for affected people (<xref ref-type="bibr" rid="B14">Gonzaludo et al., 2019</xref>). Such challenges can disproportionately impact older patients, as many clinical geneticists are initially trained in pediatric medicine and tend to focus on pediatric diagnosis (<xref ref-type="bibr" rid="B18">Jenkins et al., 2021</xref>). Despite these issues, previous large-scale clinical genetic applications of neural networks studied populations affected by many different genetic conditions and yielded impressive results (<xref ref-type="bibr" rid="B15">Gurovich et al., 2019</xref>; <xref ref-type="bibr" rid="B30">Porras et al., 2021</xref>). In the current study, we endeavored to build upon these existing works by collating our own age-annotated datasets. These datasets were designed to allow further study of the impact of patient age on facial diagnosis as well as to perform additional neural network analyses, which can also be extended to larger datasets or applied to different conditions.</p>
<p>We chose two distinct genetic conditions for further study: Williams syndrome (WS) (MIM 194050), which affects approximately one in 7,500 live births, and 22q11.2 deletion syndrome (22q), sometimes imperfectly referred to as &#x201c;DiGeorge syndrome&#x201d; (MIM 188400), which affects approximately one in 4,000&#x2013;7,000 live births (<xref ref-type="bibr" rid="B35">Stromme et al., 2002</xref>; <xref ref-type="bibr" rid="B4">Botto et al., 2003</xref>; <xref ref-type="bibr" rid="B29">Oskarsdottir et al., 2004</xref>). We selected these conditions as they may be recognizable from facial features (in addition to other manifestations) (<xref ref-type="bibr" rid="B35">Stromme et al., 2002</xref>; <xref ref-type="bibr" rid="B4">Botto et al., 2003</xref>; <xref ref-type="bibr" rid="B6">Campbell et al., 2018</xref>; <xref ref-type="bibr" rid="B26">Morris et al., 2020</xref>) and based on relative data availability, which is still very limited compared to more common health conditions. Additionally, these two conditions represent varying ease of diagnosis based on facial appearances: people with WS may have more consistently recognizable facial features, whereas people with 22q may have a more subtle facial presentation, which likely contributes to underdiagnosis for this as well as many other conditions without obvious or overtly pathognomonic signs.</p>
<p>To examine the influence of age on facial recognition, we evaluated how well clinical geneticists and our classifier recognize these conditions based on facial images of varying ages. We further explored additional neural networks applications, including saliency maps and <ext-link ext-link-type="uri" xlink:href="https://en.wikipedia.org/wiki/Generative_adversarial_network">generative adversarial networks (GANs</ext-link>), to study both facial recognition as a whole and as a function of age.</p>
</sec>
<sec sec-type="materials|methods" id="s2">
<title>Materials and Methods</title>
<sec id="s2-1">
<title>Data Collection</title>
<p>We searched Google and PubMed using the disease names of interest to select publicly available images depicting individuals with WS, 22q, or other genetic conditions that may resemble WS or 22q (see <xref ref-type="sec" rid="s11">Supplementary Table S1</xref> for more details about these conditions). After that, when the context is clear, we refer to these &#x201c;other genetic conditions&#x201d; as the control group. From the available source information for each image, we categorized the images into five age brackets: 1) infant (under 2&#xa0;years old), 2) child (2&#x2013;9&#xa0;years old), 3) adolescent (10&#x2013;19&#xa0;years old), 4) young adult (20&#x2013;34&#xa0;years old), and 5) older adult (&#x2265;35&#xa0;years old). We attempted to collect images of individuals from diverse ancestral backgrounds, though standardized and complete information regarding race and ethnicity was often unavailable (see <xref ref-type="sec" rid="s11">Supplementary Table S2</xref>). In total, we collected 1,894 images and partitioned them into 1,713 and 181 train and test images, respectively (see <xref ref-type="sec" rid="s11">Supplementary Table S3</xref>). The image sets included both color and black and white images with varying image resolution. The test images were selected from color images subjectively judged to have adequate resolution for human viewing and included representations of both sexes and of apparently ancestrally diverse individuals, though we recognized many challenges in these and related areas (<xref ref-type="bibr" rid="B5">Byeon et al., 2021</xref>). The control group test images included individuals with other genetic and congenital conditions, including those with overlapping facial features with WS or 22q (i.e., conditions that are sometimes considered in the differential diagnosis of WS or 22q). We applied the StyleGAN face detector and image preprocessing to rotate and center our images and manually aligned images that failed this preprocessing step (<xref ref-type="fig" rid="F1">Figure 1</xref>).</p>
<fig id="F1" position="float">
<label>FIGURE 1</label>
<caption>
<p>Centered and aligned images of real individuals (different individuals are shown at different ages) affected with 22q (top row) and WS (bottom row). These images have been previously published and are granted to be freely distributed for noncommercial research purposes.</p>
</caption>
<graphic xlink:href="fgene-13-864092-g001.tif"/>
</fig>
</sec>
<sec id="s2-2">
<title>Classifier</title>
<p>We selected the EfficientNet-B4 classifier, which obtained high performance on the ImageNet data with a relatively low number of parameters (<xref ref-type="bibr" rid="B36">Tan and Le, 2019</xref>). We loaded the weights pretrained on ImageNet and continued training EfficientNet-B4 end-to-end. Combining and then jointly training a small dataset of interest with a larger auxiliary dataset often increases the prediction accuracy (<xref ref-type="bibr" rid="B1">Ahmad et al., 2018</xref>; <xref ref-type="bibr" rid="B25">Meftah et al., 2020</xref>). Our auxiliary dataset is the FairFace dataset, which contains 108,000 &#x201c;in-the-wild&#x201d; faces (i.e., faces oriented in various angles and/or partially covered with hands, hats, or sunglasses) of equal ratio from (using definitions in FairFace) white, black, Latino, East/Southeast Asian, Indian, and Middle Eastern populations (<xref ref-type="bibr" rid="B20">K&#xe4;rkk&#xe4;inen and Joo, 2019</xref>).</p>
<p>The StyleGAN face detector and image preprocessing, such as rotating and centering faces, were applied to the in-the-wild FairFace faces, resulting in 62,088 usable images. We partitioned these 62,088 images into the age groups as described previously via the FairFace age classifier. One-sixth of the images (N &#x3d; 10,348) in each age category were randomly chosen as test images, which we evaluated with our own test images. The remaining 51,740 images were used with our images to train EfficientNet-B4.</p>
<p>Because of our small dataset and the assumption that at least some features persisted across age groups, we trained EfficientNet-B4 on our images (regardless of age) and FairFace images, to recognize the four labels: WS, 22q, other genetic conditions (control), and unaffected. We included unaffected individuals as an important consideration in clinical practice, which is the ability to differentiate a potentially affected from an unaffected person, especially as some genetic conditions can have subtle findings often missed by general clinicians as well as subspecialists. The classifier was trained with cross entropy loss function in which one-hot encodings represent true image labels. We rescaled all the images into resolution 448 &#xd7; 448 pixels when training EfficientNet-B4. Image resolution was chosen to maximize GPU usage (two Nvidia P100, training batch size 64).</p>
<p>We trained five classifiers via 5-fold cross-validation (one for each fold) and then created an ensemble predictor by averaging the predicted label probabilities of an image from these five classifiers. When averaging, we considered only the classifiers that produced a maximum predicted probability (over all the labels) of at least 0.5.</p>
</sec>
<sec id="s2-3">
<title>Comparison to Clinicians</title>
<p>We compared our classifier to board-certified or board-eligible clinical geneticist physicians via surveys sent by Qualtrics (Provo, Utah, United States). As WS and 22q syndromes are relatively distinct, we felt that it was more meaningful to evaluate WS test images against their own controls and likewise for 22q test images. We emphasize that for nontrivial comparisons, the control test images were of conditions resembling WS and 22q. For WS surveys, there were 50 WS (10 images per age group) and 50 corresponding control test images. To keep the survey length reasonable, each participant went through a random subset of 25 WS (5 images per age category) and 25 age category-matched control images. The ordering of the selected images in a survey was randomized, and the answer choice for a question was either &#x201c;Williams Syndrome&#x201d; or &#x201c;Other Condition.&#x201d; The same setup was also employed for 22q surveys. In addition to asking clinical geneticists to classify images, we also asked questions about the impact of patients&#x2019; age on diagnosis to determine attitudes and opinions on the age in the diagnosis process. Example surveys can be found at <ext-link ext-link-type="uri" xlink:href="https://github.com/datduong/Classify-WS-22q-Img">https://github.com/datduong/Classify-WS-22q-Img</ext-link>.</p>
<p>Following previous methods (<xref ref-type="bibr" rid="B38">Tschandl et al., 2018</xref>; <xref ref-type="bibr" rid="B37">Tschandl et al., 2019</xref>; <xref ref-type="bibr" rid="B9">Duong et al., 2021b</xref>), we estimated that 30 participants would provide a statistical power of 95% to detect a 10% difference. The participants were recruited via email. To identify survey respondents, we obtained email addresses through professional networks, departmental websites, journal publications, and other web-available lists. A total of 225 clinical geneticists were contacted, of which 36 completed the 22q survey and 34 completed the WS survey. If multiple respondents completed the same survey, only the first survey was used for analysis (see <xref ref-type="sec" rid="s11">Supplementary Table S4</xref> for the description of survey respondents).</p>
</sec>
<sec id="s2-4">
<title>Generative Adversarial Network (GAN)</title>
<p>We trained a GAN for each data partition from the 5-fold cross-validation in section &#x201c;<italic>Classifier</italic>.<italic>&#x201d;</italic> We describe the GAN training and image generation for a data partition <italic>p</italic>, which also will apply to the other partitions (see <xref ref-type="sec" rid="s11">Supplementary Figure S1</xref> for the flowchart of our GAN image production and its application with the real images to train the disease classifier). The partition <italic>p</italic> contains our images of affected individuals and FairFace unaffected individuals. Ideally, we would want to train the GAN model on all FairFace images and our dataset as blending the features of different racial/ethnic groups in FairFace with our dataset would generate diverse images of affected individuals. However, the partition <italic>p</italic> has approximately 41,392 images of unaffected FairFace individuals, and our preliminary GAN experiments required a large amount of computational power. The larger FairFace dataset also often skewed GAN output in which the generated images of affected individuals looked more like the unaffected subset. Therefore, in the partition <italic>p</italic>, we trained GAN on our images and a fixed subset of FairFace. This subset was randomly chosen with 500 individuals in each age bracket. In each training batch, we selected an equal number of affected and unaffected individuals.</p>
<p>Our GAN is based on the conditional StyleGAN2-ADA and generates images using both disease statuses and age categories (<xref ref-type="bibr" rid="B21">Karras et al., 2020</xref>). We made the following key modification to StyleGAN2-ADA. The default label embedding is L x 512, where L is the number of labels and produces a vector of length 512 for each label. Training this embedding requires many people with a specific disease in a certain age category. However, some disease and age label combinations have small sample sizes; for example, our dataset has 39 WS and 35 22q individuals older than 35 years of age. We replaced the default label embedding with two smaller matrices, namely 4 &#xd7; 256 and 5 &#xd7; 256 to represent the four diseases (WS, 22q, other conditions, and unaffected) and five age categories. Then, training the 4 &#xd7; 256 disease embedding uses all affected individuals in every age category. Likewise, training the 5 &#xd7; 256 age embedding uses all the images in our dataset and in FairFace. The outputs of these two components are concatenated to a vector of size 512 to match the rest of the StyleGAN2-ADA architecture. Hence, except for the label embeddings, we initialized all the StyleGAN2-ADA weights with the pretrained values on FFHQ dataset at resolution 256 &#xd7; 256 pixels (<xref ref-type="bibr" rid="B21">Karras et al., 2020</xref>); all fake images were also generated in the resolution 256 &#xd7; 256 pixels. Image resolution was chosen to maximize GPU usage (two Nvidia P100).</p>
<p>After training GAN on the data partition <italic>p</italic>, we generated four types of fake images: 1) unrelated faces for a specific disease and age group; 2) similar faces at different ages for a specific disease; 3) the same face at different ages for a specific disease; and 4) faces containing characteristics of two different conditions, which we hypothesized could aid classifier accuracy.</p>
<p>Given the large number of unaffected people in FairFace, we generated just images of affected individuals from the disease label <inline-formula id="inf1">
<mml:math id="m1">
<mml:mrow>
<mml:mi>d</mml:mi>
<mml:mo>&#x2208;</mml:mo>
</mml:mrow>
</mml:math>
</inline-formula> {WS, 22q, other condition} and age bracket <inline-formula id="inf2">
<mml:math id="m2">
<mml:mrow>
<mml:mi>a</mml:mi>
<mml:mo>&#x2208;</mml:mo>
</mml:mrow>
</mml:math>
</inline-formula> {infant, child, adolescent, young adult, older adult}. The number of generated images for each pair (<italic>d,a</italic>) is equal to the average count of all age groups with a specific disease <italic>d</italic> in the data partition. Thus, for a specific condition, respectively, we made more and fewer images for the uncommon (less represented) and common (more represented) age groups. In total, each image type has the same count as the size of the affected individuals in the data partition in which the GAN was trained.</p>
<p>For type 1, we generated a fake image <italic>i</italic> by concatenating the random vector <italic>r</italic>
<sub>
<italic>iad</italic>
</sub> with the label embedding <italic>e</italic>
<sub>
<italic>d</italic>
</sub> and <italic>e</italic>
<sub>
<italic>a</italic>
</sub>, denoted as [<italic>r</italic>
<sub>
<italic>iad,</italic>
</sub> <italic>e</italic>
<sub>
<italic>a,</italic>
</sub> <italic>e</italic>
<sub>
<italic>d</italic>
</sub>] and then passing this new vector to our GAN image generator. Each disease <italic>d</italic> and age group <italic>a</italic> combination has images generated from their own unique random vector <italic>r</italic>
<sub>
<italic>iad</italic>
</sub>, so that all the fake images are theoretically unique (<xref ref-type="fig" rid="F2">Figure 2A</xref>).</p>
<fig id="F2" position="float">
<label>FIGURE 2</label>
<caption>
<p>Examples of GAN fake images for 22q and WS. Type 1 fake images of <bold>(A)</bold> 22q (top row) and WS (bottom row) were generated with GAN and are all theoretically unique. For type 2 fake images <bold>(B)</bold> of 22q (top row) and WS (bottom row), general features, such as skin tone and hair color, are roughly preserved. For type 3 <bold>(C)</bold>, the generated images of 22q (top row) and WS (bottom row) look consistent at depicting the same &#x201c;person&#x201d; progressing through different age groups. Type 4 fake images <bold>(D)</bold> were created with blended facial characteristics of two disease labels. The main disease condition (22q or WS) represents 55% of the facial phenotype, and the added disease condition represents 45% of the facial phenotype. For example, WS:unaffected is a blend of 55% WS facial features and 45% unaffected facial features. Only the blended images were used for training, and the left most images are shown here as references.</p>
</caption>
<graphic xlink:href="fgene-13-864092-g002.tif"/>
</fig>
<p>Type 2 images are generated by varying the age embedding <italic>e</italic>
<sub>
<italic>a</italic>
</sub>, where <inline-formula id="inf3">
<mml:math id="m3">
<mml:mrow>
<mml:mi>a</mml:mi>
<mml:mo>&#x2208;</mml:mo>
</mml:mrow>
</mml:math>
</inline-formula> {infant, child, adolescent, young adult, older adult}, while fixing the random vector <italic>r</italic>
<sub>
<italic>id</italic>
</sub> and the disease embedding <italic>e</italic>
<sub>
<italic>d</italic>
</sub> constant. The vector <italic>r</italic>
<sub>
<italic>id</italic>
</sub> is unique to the <italic>i</italic>th image having disease <italic>d</italic>. Following this, every disease has own unique images, but within the same disease the images at each age category have similar facial features such as skin tone and hair color (<xref ref-type="fig" rid="F2">Figure 2B</xref>).</p>
<p>For type 3 images, we interpolated three equally spaced vectors between the age embedding e<sub>
<italic>infant</italic>
</sub> and e<sub>
<italic>older adult</italic>
</sub>
<italic>.</italic> For a disease <italic>d</italic>, we generated five fake images from a random vector <italic>r</italic>
<sub>
<italic>id</italic>
</sub> by passing these five inputs, namely [<italic>r</italic>
<sub>
<italic>id</italic>
</sub>
<italic>, e</italic>
<sub>
<italic>infant</italic>
</sub>
<italic>, e</italic>
<sub>
<italic>d</italic>
</sub>], [<italic>r</italic>
<sub>
<italic>id</italic>
</sub>
<italic>, 0.75e</italic>
<sub>
<italic>infant</italic>
</sub>
<italic>&#x2b;0.25e</italic>
<sub>
<italic>older adult</italic>
</sub>
<italic>, e</italic>
<sub>
<italic>d</italic>
</sub>], [<italic>r</italic>
<sub>
<italic>id</italic>
</sub>
<italic>, 0.5e</italic>
<sub>
<italic>infant</italic>
</sub>
<italic>&#x2b;0.5e</italic>
<sub>
<italic>older adult</italic>
</sub>
<italic>, e</italic>
<sub>
<italic>d</italic>
</sub>], [<italic>r</italic>
<sub>
<italic>id</italic>
</sub>
<italic>, 0.25e</italic>
<sub>
<italic>infant</italic>
</sub> <italic>&#x2b;0.75e</italic>
<sub>
<italic>older adult</italic>
</sub>
<italic>, e</italic>
<sub>
<italic>d</italic>
</sub>], and [<italic>r</italic>
<sub>
<italic>id</italic>
</sub>
<italic>, e</italic>
<sub>
<italic>older adult</italic>
</sub>
<italic>, e</italic>
<sub>
<italic>d</italic>
</sub>] into the GAN generator. These images closely represent the same person affected with disease <italic>d</italic> at five different age groups (<xref ref-type="fig" rid="F2">Figure 2C</xref>). There are additional potential approaches for depicting age progression, which we may explore in future studies (<xref ref-type="bibr" rid="B28">Or-El et al., 2020</xref>). Of note, the previous work (<xref ref-type="bibr" rid="B15">Gurovich et al., 2019</xref>; <xref ref-type="bibr" rid="B30">Porras et al., 2021</xref>) used different and/or additional age brackets, some of which may not involve sufficient numbers of images for robust analyses, at least in our datasets.</p>
<p>Type 4 images were generated like type 3 images; however, we reversed the roles of disease <italic>d</italic> and age label <italic>a</italic>. With a random vector <italic>r</italic>
<sub>
<italic>ia</italic>
</sub>, we generated three fake images from the inputs: [<italic>r</italic>
<sub>
<italic>ia</italic>
</sub>
<italic>, e</italic>
<sub>
<italic>a</italic>
</sub>
<italic>, ce</italic>
<sub>
<italic>WS</italic>
</sub> <italic>&#x2b;</italic> (<italic>1-c</italic>)<italic>e</italic>
<sub>
<italic>22q</italic>
</sub>], [<italic>r</italic>
<sub>
<italic>ia</italic>
</sub>
<italic>, e</italic>
<sub>
<italic>a</italic>
</sub>
<italic>, ce</italic>
<sub>
<italic>WS</italic>
</sub> <italic>&#x2b;</italic> (<italic>1-c</italic>)<italic>e</italic>
<sub>
<italic>control</italic>
</sub>], and [<italic>r</italic>
<sub>
<italic>ia</italic>
</sub>
<italic>, e</italic>
<sub>
<italic>a</italic>
</sub>
<italic>, ce</italic>
<sub>
<italic>WS</italic>
</sub> <italic>&#x2b;</italic> (<italic>1-c</italic>)<italic>e</italic>
<sub>
<italic>unaffected</italic>
</sub>], where <italic>c</italic> is a predefined fraction between 0 and 1. These images represent a person at age <italic>a</italic> having facial characteristics of two different diseases (<xref ref-type="fig" rid="F2">Figure 2D</xref>). The true labels for type 4 images are soft labels; for example, the image created from the vector [<italic>r</italic>
<sub>
<italic>ia</italic>
</sub>
<italic>, e</italic>
<sub>
<italic>a</italic>
</sub>
<italic>, ce</italic>
<sub>
<italic>WS</italic>
</sub> <italic>&#x2b;</italic> (<italic>1-c</italic>)<italic>e</italic>
<sub>
<italic>22q</italic>
</sub>] would have the label encoding [<italic>c, 1-c, 0, 0</italic>] instead of the traditional one-hot encoding. Here, the training loss function is still cross-entry but for soft label.</p>
<p>Next, we created four new larger datasets by combining partition <italic>p</italic> with each of the four fake image types. We then trained EfficientNet-B4 on each of these new larger datasets. For each type of new dataset, we created the ensemble predictor over all the data partitions following the approach mentioned in section <italic>&#x201c;Classifier.&#x201d;</italic>
</p>
</sec>
<sec id="s2-5">
<title>Attribution Analysis for Features in Different Age Groups</title>
<p>To visualize which facial features of an image the classifier considered to be important, we produced saliency maps using a window size 20 &#xd7; 20 pixels and stride 10 &#xd7; 10 pixels using the occlusion attribution method (<xref ref-type="bibr" rid="B39">Zeiler and Fergus, 2014</xref>). For a test image, we averaged the saliency maps of the classifiers in the ensemble predictor. We used the permutation test to measure how much the facial features identified by the classifiers differ with respect to age. Our Qualtrics surveys had 10 test images for each disease and age label combination. Conditioned on a disease and two age groups, we permuted the 20 images into two sets and repeated this permutation 100 times. Each time, we averaged the saliency map over the 10 images in each set and then retrieved the embedding of this average attribution via the EfficientNet-B4 trained on ImageNet. In each permutation, we computed the Euclidean distance between the embeddings of these two sets. If the observed Euclidean distance is smaller than 5% (or some other threshold) of the permutation values, then the two age groups of a specific disease were defined as not statistically different. That is, the key facial features identified by the classifier do not differ with respect to age.</p>
</sec>
</sec>
<sec sec-type="results" id="s3">
<title>Results</title>
<sec id="s3-1">
<title>Classifier Accuracy</title>
<p>Our classifier, which was trained on images of individuals with WS, 22q, other genetic conditions, and individuals who are presumably unaffected, correctly classified unique test images 68&#x2013;100% of the time, with the lowest accuracy for 22q and the highest accuracy for unaffected individuals (<xref ref-type="fig" rid="F3">Figure 3</xref>). While unaffected individuals are not misclassified as affected, the opposite is not typically true, presumably as some affected individuals may show only subtle features of the condition. A similar situation happens in clinical situations, where it can sometime be difficult to tell whether a person may be affected by a genetic condition and what that condition might be based on physical examination features (or other information). Classification accuracy of WS (86%) was the highest among the affected individuals we examined; our results suggest that these individuals often clearly display key findings (e.g., the dysmorphology or distinctive features affecting the eyes and mouth) compared to individuals with 22q, who are frequently misclassified (30%) as the control group (genetic or congenital conditions other than WS or 22q).</p>
<fig id="F3" position="float">
<label>FIGURE 3</label>
<caption>
<p>Confusion matrix of accuracy of the classifier trained on real images. Rows represent the correct label, while columns represent the label chosen by the classifier. The diagonal numbers represent the percent accuracy for each category (the percentage of images when the correct label was identified), while the off-diagonal numbers represent misclassification percentage ascribed to an incorrect category. Accuracy is based on 50 test images of WS, 50 of 22q, and 81 of other conditions.</p>
</caption>
<graphic xlink:href="fgene-13-864092-g003.tif"/>
</fig>
</sec>
<sec id="s3-2">
<title>Comparison to Clinical Geneticists</title>
<p>Clinical geneticists completed the survey(s) classifying images of individuals with WS versus other conditions and 22q versus other conditions. We intentionally used two separate surveys, based on preliminary testing, as we felt that WS and 22q would typically be considered a distinct condition by clinical geneticists and that it would be more meaningful to evaluate these conditions separately. The statistical differences between clinical geneticists and our model were measured via the paired <italic>t</italic>-test. As our model was trained with four labels, we took the highest prediction probability between WS and control for a test image in surveys containing WS and control images. The same idea applied to surveys containing 22q and their corresponding controls. Our model outperformed clinical geneticists by 15.5% (77.5 vs. 93%, <italic>p</italic> &#x3d; 6.828e<sup>&#x2212;11</sup>) for WS and by 22.7% (59.3 vs. 82%, <italic>p</italic> &#x3d; 3.203e<sup>&#x2212;13</sup>) for 22q.</p>
<p>To determine whether patient age affects accuracy, we determined the average accuracy for each age group (<xref ref-type="table" rid="T1">Table 1</xref>). On an average, the clinical geneticists had the most difficulty identifying infants affected with WS (67.3%) and the greatest accuracy with adolescents (80.7%). The clinical geneticists had the most difficulty classifying older adults with 22q (50.7%) and the greatest accuracy classifying adolescents (67.3%). Our classifier outperformed the clinical geneticists in all age groups (see <xref ref-type="sec" rid="s11">Supplementary Table S5</xref>). However, we emphasize that because each age group has 10 images, performance differences may represent only a few images. Despite the small test size per condition and age bracket, our results suggest clinical diagnosis may be more difficult in some age groups.</p>
<table-wrap id="T1" position="float">
<label>TABLE 1</label>
<caption>
<p>Average accuracy over all 10 test images in each age group. Each test image was viewed and classified by 15 clinical geneticists. Our classifier, either trained on real images alone or on both real and GAN age progression (age prog) images, obtains higher accuracy for each age group, except for the oldest 22q cohort. <italic>P</italic>-values in parentheses comparing the human against model were computed via the permutation test.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th rowspan="3" colspan="2" align="center">Age</th>
<th colspan="3" align="center">22q</th>
<th colspan="3" align="center">WS</th>
</tr>
<tr>
<th/>
<th colspan="2" align="center">Model</th>
<th/>
<th colspan="2" align="center">Model</th>
</tr>
<tr>
<th align="center">Human</th>
<th align="center">Real images</th>
<th align="center">&#x2b; Age prog</th>
<th align="center">Human</th>
<th align="center">Real images</th>
<th align="center">&#x2b; Age prog</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td rowspan="5" align="left">Disease</td>
<td align="left">Infant</td>
<td align="char" char=".">0.54</td>
<td align="char" char="(">0.7 (0.18)</td>
<td align="char" char="(">0.8 (0.05)</td>
<td align="char" char=".">0.673</td>
<td align="char" char="(">1 (0)</td>
<td align="center">1 (0</td>
</tr>
<tr>
<td align="left">Child</td>
<td align="char" char=".">0.527</td>
<td align="char" char="(">0.8 (0.04)</td>
<td align="char" char="(">0.9 (0)</td>
<td align="char" char=".">0.707</td>
<td align="char" char="(">1 (0)</td>
<td align="center">1 (0)</td>
</tr>
<tr>
<td align="left">Adolescent</td>
<td align="char" char=".">0.673</td>
<td align="char" char="(">0.7 (0.44)</td>
<td align="char" char="(">0.8 (0.21)</td>
<td align="char" char=".">0.0807</td>
<td align="char" char="(">0.(0.24)</td>
<td align="center">1 (0)</td>
</tr>
<tr>
<td align="left">Young adult</td>
<td align="char" char=".">0.613</td>
<td align="char" char="(">0.8 (0.10)</td>
<td align="char" char="(">0.9 (0.01)</td>
<td align="char" char=".">0.74</td>
<td align="char" char="(">1 (0)</td>
<td align="center">1 (0)</td>
</tr>
<tr>
<td align="left">Older adult</td>
<td align="char" char=".">0.507</td>
<td align="char" char="(">0.5 (0.50)</td>
<td align="char" char="(">0.5 (0.50)</td>
<td align="char" char=".">0.713</td>
<td align="char" char="(">0.9 (0.09)</td>
<td align="center">1 (0)</td>
</tr>
<tr>
<td rowspan="5" align="left">Other conditions</td>
<td align="left">Infant</td>
<td align="char" char=".">0.753</td>
<td align="char" char="(">1 (0)</td>
<td align="char" char="(">1 (0)</td>
<td align="char" char=".">0.84</td>
<td align="char" char="(">0.9 (0.35)</td>
<td align="center">0.9 (0.35</td>
</tr>
<tr>
<td align="left">Child</td>
<td align="char" char=".">0.6</td>
<td align="char" char="(">0.9 (0)</td>
<td align="char" char="(">0.9 (0)</td>
<td align="char" char=".">0.767</td>
<td align="char" char="(">0.8 (0.40)</td>
<td align="center">0.8 (0.40)</td>
</tr>
<tr>
<td align="left">Adolescent</td>
<td align="char" char=".">0.52</td>
<td align="char" char="(">1 (0)</td>
<td align="char" char="(">1 (0)</td>
<td align="char" char=".">0.86</td>
<td align="char" char="(">0.9 (0.39)</td>
<td align="center">0.9 (0.39</td>
</tr>
<tr>
<td align="left">Young adult</td>
<td align="char" char=".">0.6</td>
<td align="char" char="(">0.9 (0.01)</td>
<td align="char" char="(">1 (0)</td>
<td align="char" char=".">0.787</td>
<td align="char" char="(">0.9 (0.19)</td>
<td align="center">0.9 (0.19</td>
</tr>
<tr>
<td align="left">Older adult</td>
<td align="char" char=".">0.593</td>
<td align="char" char="(">0.9 (0.01)</td>
<td align="char" char="(">0.9 (0.01)</td>
<td align="char" char=".">0.86</td>
<td align="char" char="(">1 (0)</td>
<td align="center">1 (0)</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec id="s3-3">
<title>Estimating Facial Differences Across Age Groups</title>
<p>Conditioned on a disease, we compared the saliency maps of the test images in different age groups based on the method in &#x201c;<italic>Attribution Analysis for Features in Different Age Groups</italic>.<italic>&#x201d;</italic> The composite images of saliency maps averaged over all the test images were generated for each age group. <xref ref-type="fig" rid="F4">Figure 4</xref> provides qualitative descriptions for the differences among the key facial features identified by the classifier for each age bracket. <xref ref-type="fig" rid="F5">Figure 5</xref> quantitatively compares these differences by measuring the Euclidean distances among the embeddings of these composite saliency maps. Assuming the standard 5% statistical significant threshold, there were significant differences among the five age brackets. For example, the observed distance between the embeddings of the 22q infant and child composite saliency maps ranks higher than 43 of the 100 permutation values (<xref ref-type="fig" rid="F5">Figure 5</xref>). The greatest differences are seen for the infant and older adults in both WS and 22q. Compared to 22q, WS facial features identified by our model differ more with respect to age, which may show how facial features of a condition can be age-specific. Again, we emphasize that there can be confounding factors. For example, a person&#x2019;s facial expression (e.g., whether a person is smiling) may explain why features of adolescent WS test images are different from those of the other age groups (<xref ref-type="fig" rid="F5">Figure 5</xref>). Clinical geneticists may also rely on nonspecific facial clues (as well as other clinical manifestations) to classify syndromes and conditions. For example, high sociability and friendliness are common features of people with WS (<xref ref-type="bibr" rid="B27">Morris et al., 1993</xref>; <xref ref-type="bibr" rid="B26">Morris et al., 2020</xref>). While we did not intentionally select images based on facial expression (e.g., whether they were smiling or were not smiling), we found that more WS test images (60%) had a partial or full smile than other conditions (44%). Clinical geneticists were more likely to misclassify a Williams syndrome image if the image showed a person who was not smiling (58.3%) vs. smiling (82.4%) (see <xref ref-type="sec" rid="s11">Supplementary Table S6</xref>). The presence or absence of a smile did not appear to impact classification of other conditions.</p>
<fig id="F4" position="float">
<label>FIGURE 4</label>
<caption>
<p>Occlusion analysis of 22q and WS facial images across the lifespan. Composite saliency maps were made by averaging over all the test images in each age group: infant, child, adolescent, young adult, and older adult (reading left to right) for both 22q <bold>(A)</bold> and WS <bold>(B)</bold>. Green indicates positive contribution, and red indicates negative contribution to the correct label. For 22q, specific regions of interest (e.g., periorbital regions, glabella, nasal bridge, and the mandible) subjectively appear to be consistently important at all ages analyzed. However, there appear to be some areas that are more specific for people in certain age groups, such as the areas superior to the lateral mandibular region in the youngest age group. Additionally, the periorbital region and nasal root appear to be more important in older age groups. For WS, facial features of interest across all age categories include the eyes (possibly due to the stellate iris or other important ocular features; we note that this pattern was not seen in people with 22q) and the mouth. Our composite WS images suggest that as aging progresses, the positive attribution present during infancy in the nasal root and periorbital region, as well as the eyes to some degree, decreases through older adulthood.</p>
</caption>
<graphic xlink:href="fgene-13-864092-g004.tif"/>
</fig>
<fig id="F5" position="float">
<label>FIGURE 5</label>
<caption>
<p>Quantitative comparison of key facial features during aging. Rank of the observed Euclidean distance (in fraction out of 100 permutations) between embeddings of the averages of occlusion analysis for two age groups. A small number indicates that key features identified by the neural network for two age groups are more statistically similar, whereas a larger number indicates that key features are more statistically different. Possible key facial feature differences across age categories are qualitatively explained in <xref ref-type="fig" rid="F4">Figure 4</xref>.</p>
</caption>
<graphic xlink:href="fgene-13-864092-g005.tif"/>
</fig>
</sec>
<sec id="s3-4">
<title>Classifiers Trained on Real and GAN Images</title>
<p>
<xref ref-type="table" rid="T2">Table 2</xref> shows the accuracies of the classifiers trained on real images and different types of GAN images [see <xref ref-type="sec" rid="s11">Supplementary Figure S2</xref> for the corresponding confusion matrices]. Although the improvements are minimal, all four types of GAN images obtain slightly higher average accuracy than the base classifier, with type 3 GAN images (showing age progression for the same person) performing the best.</p>
<table-wrap id="T2" position="float">
<label>TABLE 2</label>
<caption>
<p>Accuracy of the classifier trained on real images compared to classifiers trained on both real images and each of the four types of fake images. Column names unrelated, age related, age progression (age prog), and blended correspond to fake GAN images of type 1, 2, 3, and 4, respectively. The greatest improvement in accuracy was observed with the addition of age progression fake images.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="left"/>
<th align="center">Real images</th>
<th align="center">&#x2b; Unrelated</th>
<th align="center">&#x2b; Age related</th>
<th align="center">&#x2b; Age prog</th>
<th align="center">&#x2b; Blended 55&#x2013;45%</th>
<th align="center">&#x2b; Blended 75&#x2013;25%</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="left">22q</td>
<td align="char" char=".">68</td>
<td align="char" char=".">76</td>
<td align="char" char=".">74</td>
<td align="char" char=".">76</td>
<td align="char" char=".">68</td>
<td align="char" char=".">76</td>
</tr>
<tr>
<td align="left">Controls</td>
<td align="char" char=".">86.4</td>
<td align="char" char=".">81.5</td>
<td align="char" char=".">82.5</td>
<td align="char" char=".">82.7</td>
<td align="char" char=".">85.2</td>
<td align="char" char=".">85.2</td>
</tr>
<tr>
<td align="left">Unaffected</td>
<td align="char" char=".">100</td>
<td align="char" char=".">100</td>
<td align="char" char=".">100</td>
<td align="char" char=".">100</td>
<td align="char" char=".">100</td>
<td align="char" char=".">100</td>
</tr>
<tr>
<td align="left">WS</td>
<td align="char" char=".">94</td>
<td align="char" char=".">96</td>
<td align="char" char=".">96</td>
<td align="char" char=".">100</td>
<td align="char" char=".">96</td>
<td align="char" char=".">96</td>
</tr>
<tr>
<td align="left">Average</td>
<td align="char" char=".">87.1</td>
<td align="char" char=".">88.375</td>
<td align="char" char=".">88.175</td>
<td align="char" char=".">90.35</td>
<td align="char" char=".">87.3</td>
<td align="char" char=".">89.3</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>We also compared the classifier trained with type 3 GAN images against the clinical geneticists in each age bracket (<xref ref-type="table" rid="T1">Table 1</xref>). The cumulative improvements over humans are 95 vs. 77.5% (<italic>p</italic> &#x3d; 1.846e<sup>&#x2212;11</sup>) for WS and 87 vs. 59.3% (<italic>p</italic> &#x3d; 1.703e<sup>&#x2212;15</sup>) for 22q.</p>
<p>We suspect that type 3 GAN images improved the base classifier more because key facial features varied with respect to age (<xref ref-type="fig" rid="F4">Figure 4</xref>, <xref ref-type="fig" rid="F5">5</xref>). By conditioning on the same person, we may capture more specific details of how a condition progresses with time.</p>
</sec>
</sec>
<sec sec-type="discussion" id="s4">
<title>Discussion</title>
<p>The practice of medical genetics has shifted considerably in the last several decades. One major reason is the growing availability of high-throughput genetic/genomic testing, such as exome and genome sequencing. These testing methods allow more precise diagnosis and have changed the approach to phenotyping (<xref ref-type="bibr" rid="B16">Hennekam and Biesecker, 2012</xref>). However, access to these testing technologies is uneven, and it remains important to be able to quickly recognize patients who may be affected by certain conditions, especially those with near-term management implications (<xref ref-type="bibr" rid="B34">Solomon et al., 2013</xref>; <xref ref-type="bibr" rid="B3">Bick et al., 2021</xref>). For example, people with WS are prone to infantile electrolyte abnormalities and immunologic dysfunction, and people with 22q may be affected by endocrine, immunologic, cardiovascular, and other sequelae that require immediate attention (<xref ref-type="bibr" rid="B6">Campbell et al., 2018</xref>; <xref ref-type="bibr" rid="B26">Morris et al., 2020</xref>). Recognizing the likelihood of these conditions quickly&#x2014;before the results of even the fastest genetic/genomic tests may be available&#x2014;can be important for these and other conditions.</p>
<p>To provide examples of ways to bolster the standard diagnostic process as well as to build on the impressive findings of previous, related studies, (<xref ref-type="bibr" rid="B15">Gurovich et al., 2019</xref>; <xref ref-type="bibr" rid="B9">Duong et al., 2021b</xref>; <xref ref-type="bibr" rid="B17">Hsieh et al., 2021</xref>; <xref ref-type="bibr" rid="B30">Porras et al., 2021</xref>), we analyzed and provided a larger dataset of WS and 22q individuals (although these other studies contained a much larger total number of individuals having multiple other diseases). We also compared results for different ages of individuals. Our classifier outperformed clinical geneticists at identifying WS and 22q individuals by large margins (15.5 and 22.7%, respectively). This was consistently true for each age group.</p>
<p>We hypothesized that because geneticists overall often have more clinical experience with children, and as textbooks and the overall medical literature tend to focus more on pediatric presentations of congenital disorders, respondents would feel the most confident about diagnosis in younger age groups and would also perform best with images of younger patients. However, for WS, our results show that respondents&#x2019; accuracy did not correlate with their confidence level in diagnosing the conditions at various ages. For example, 46.7% (14/30) and 50% (15/30) of clinical geneticists surveyed reported that infants and older adults with WS are difficult to classify based on facial features, respectively, but the geneticists were able to classify these patients with similar accuracy to those of other ages (see <xref ref-type="sec" rid="s11">Supplementary Table S7</xref>). This may imply other explanations. For example, clinicians may feel the most confident considering patients at ages they most often see in clinical practice, but this confidence may not be reflected in their performance. The features of WS may also be more pronounced with age such that clinicians can more readily recognize the condition in older patients, even when they have less real-life experience with patients at older ages. On the other hand, 60% (18/30) and 40% (12/30) of clinical geneticists, respectively, reported that infants and older adults affected with 22q are difficult to classify based on facial features, which aligns better with their performance in the survey we administered. There are again multiple explanations, but one possibility is that 22q may simply be a more subtle condition based on facial features, or that facial features in people with 22q do not become more obvious with age. Our saliency maps suggest that age-specific changes in key facial features exist in both 22q and WS. While saliency maps provide insights into the behavior of a neural network, these approaches have not been fully standardized or validated yet for the interpretation of medical data (<xref ref-type="bibr" rid="B32">Saporta et al., 2021</xref>). To explore these and other questions further, we plan to extend our analyses in the future to additional images and conditions, including by determining which particular features are objectively assessed by humans. This may help reveal underlying reasons for diagnostic patterns.</p>
<p>Intuitively, due to sample size difference, a classifier trained on fake and real images should outperform the one trained on just real images. Interestingly, this approach does not always improve the prediction outcome in previous works from other disciplines (<xref ref-type="bibr" rid="B12">Finlayson et al., 2018</xref>; <xref ref-type="bibr" rid="B31">Qin et al., 2020</xref>). Our results also showed that there was a small improvement with the incorporation of images (up to 3.25% accuracy gain). In the future, we plan to evaluate whether GAN images may be useful in other applications, for example, the generated images could help as educational tools. The GAN images could also be used to generate realistic images to obviate data sharing and privacy concerns. Along these lines, our results suggest areas of weakness that could be targeted for the generation of GANs, such as images of infants and older individuals, which could be used for medical training purposes. It would also be informative to conduct a meta-analysis on the existing literature across different disciplines to estimate the improvement of training GAN images and real images.</p>
<p>Our study has multiple limitations. First, our dataset is small compared to other datasets used for image recognition and may involve biases. Since collecting publicly free images of confirmed cases is challenging, we did not have balanced numbers of images for each condition and age bracket combination (<xref ref-type="sec" rid="s11">Supplementary Table S3</xref>), and the types of images may have differed in certain categories. For example, we included some gray-scale images; having different numbers of these in some subsets could affect the color consistency for GAN-based transformation (see <xref ref-type="sec" rid="s11">Supplementary Figure S3</xref>). However, the average age for each age grouping was consistent (see <xref ref-type="sec" rid="s11">Supplementary Table S8</xref>). For example, the average age for the child age grouping was 5.54, 5.01, and 5.33 years for 22q, WS, and controls, respectively. Second, but also related to our sample size, we may have suboptimal grouping. For example, grouping all individuals older than a certain age into the oldest age group may have obscured differences within that group.</p>
</sec>
<sec sec-type="conclusion" id="s5">
<title>Conclusion</title>
<p>Our contributions and findings can be summarized in the following four points.</p>
<p>First, we collected a dataset of publicly available WS and 22q images, which may be larger than others previously studied (<xref ref-type="bibr" rid="B15">Gurovich et al., 2019</xref>; <xref ref-type="bibr" rid="B22">Liu et al., 2021</xref>; <xref ref-type="bibr" rid="B30">Porras et al., 2021</xref>). Second, beyond the dataset, our approaches (and available code) may be used as subcomponents of other algorithms (<xref ref-type="bibr" rid="B9">Duong et al., 2021b</xref>). We trained a neural network classifier on our dataset (N &#x3d; 1,894), which is still small compared to many other deep learning datasets, thus pushing the capability of the neural network model. Our classifier consistently outperformed clinical geneticists at recognizing individuals in the test set with these two syndromes for individuals in all five age brackets. Third, we show that key facial features (analyzed via saliency maps) identified by the classifier differ with respect to age. This type of approach is important for DL in biomedical contexts (<xref ref-type="bibr" rid="B8">DeGrave et al., 2021</xref>), including related to disease progression and other temporal factors. Fourth, there is a modest prediction accuracy increment by jointly training real images with different types of fake images created via GAN, in which including the fake images illustrating age progression for the same person yielded the best improvement.</p>
<p>Despite the rarity (and therefore lack of data availability) of many genetic conditions, neural networks have high potential in this area, due to both the ability to accurately categorize patients based on underlying molecular causes and the lack of trained experts throughout the world such that these tools could be highly valuable (<xref ref-type="bibr" rid="B33">Solomon, 2021</xref>). This area provides a ripe opportunity for patients, clinicians, researchers, and others to collaborate for the good of the impacted community. Privacy and data handling issues must be taken seriously; we hope that obstacles around data and code sharing can be addressed so as not to impose undue barriers for helping affected individuals and families.</p>
</sec>
</body>
<back>
<sec id="s6">
<title>Data Availability Statement</title>
<p>The datasets presented in this study can be found in online repositories. We make the versions of the publicly available images included in our analyses available (via CC0 license) for the purpose of reproducibility and research, with the assumption that these would only be used for purposes that would be considered fair use. These data were compiled to produce a new, derivative work, which we offer as a whole. The names of the repository/repositories and accession number(s) can be found via the link below: <ext-link ext-link-type="uri" xlink:href="https://github.com/datduong/Classify-WS-22q-Img">https://github.com/datduong/Classify-WS-22q-Img</ext-link>.</p>
</sec>
<sec id="s7">
<title>Ethics Statement</title>
<p>The study was discussed with National Human Genome Research Institute (NHGRI) bioethicists and formally reviewed by the National Institutes of Health (NIH) Institutional Review Board (IRB). The main analyses were considered not human subjects research; a waiver of consent was granted by the NIH IRB (NIH protocol: 000537) for the work involving the surveys of medical professionals. The study contains human images, which have been previously published and/or made publicly available and are granted to be freely distributed (as a new, derivative whole) via CC0 license for noncommercial research purposes.</p>
</sec>
<sec id="s8">
<title>Author Contributions</title>
<p>DD provided oversight, helped to analyze data, and helped to write the manuscript. PH collected data and helped to critically revise the manuscript. CT-N collected and analyzed data and helped to critically revise the manuscript. SH collected data and helped to critically revise the manuscript. SL collected data and helped to critically revise the manuscript. BS collected and analyzed data and helped to draft the manuscript. RW provided oversight, collected, and analyzed data and helped to critically revise the manuscript.</p>
</sec>
<sec id="s9">
<title>Funding</title>
<p>This research was supported by the Intramural Research Program of the National Human Genome Research Institute, National Institutes of Health.</p>
</sec>
<sec sec-type="COI-statement" id="s10">
<title>Conflict of Interest</title>
<p>BS is the editor-in-chief of the American Journal of Medical Genetics.</p>
<p>The remaining authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec sec-type="disclaimer" id="s11">
<title>Publisher&#x2019;s Note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors, and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<ack>
<p>This work utilized the computational resources of the NIH HPC Biowulf cluster (<ext-link ext-link-type="uri" xlink:href="http://hpc.nih.gov/">http://hpc.nih.gov</ext-link>). An earlier version of this work has also been deposited in the preprint server, MedRxiv (see citation) (<xref ref-type="bibr" rid="B10">Duong et al., 2021a</xref>).</p>
</ack>
<sec id="s12">
<title>Supplementary Material</title>
<p>The Supplementary Material for this article can be found online at: <ext-link ext-link-type="uri" xlink:href="https://www.frontiersin.org/articles/10.3389/fgene.2022.864092/full#supplementary-material">https://www.frontiersin.org/articles/10.3389/fgene.2022.864092/full&#x23;supplementary-material</ext-link>
</p>
<supplementary-material xlink:href="Table1.docx" id="SM1" mimetype="application/docx" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="Presentation1.PPTX" id="SM2" mimetype="application/PPTX" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="Table8.DOCX" id="SM3" mimetype="application/DOCX" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="Table9.XLSX" id="SM4" mimetype="application/XLSX" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="Table7.DOCX" id="SM5" mimetype="application/DOCX" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="Table2.docx" id="SM6" mimetype="application/docx" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="Table5.DOCX" id="SM7" mimetype="application/DOCX" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="Table3.DOCX" id="SM8" mimetype="application/DOCX" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="Table4.DOCX" id="SM9" mimetype="application/DOCX" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="Presentation2.PPTX" id="SM10" mimetype="application/PPTX" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="Image1.PNG" id="SM11" mimetype="application/PNG" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="Table6.DOCX" id="SM12" mimetype="application/DOCX" xmlns:xlink="http://www.w3.org/1999/xlink"/>
</sec>
<ref-list>
<title>References</title>
<ref id="B1">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Ahmad</surname>
<given-names>W. U.</given-names>
</name>
<name>
<surname>Bai</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Huang</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Peng</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Chang</surname>
<given-names>K.-W.</given-names>
</name>
</person-group> (<year>2018</year>). <source>Multi-task Learning for Universal Sentence Embeddings: A Thorough Evaluation Using Transfer and Auxiliary Tasks</source>. <comment>arXiv preprint arXiv:1804.07911</comment>. </citation>
</ref>
<ref id="B2">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Baek</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Dimaio</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Anishchenko</surname>
<given-names>I.</given-names>
</name>
<name>
<surname>Dauparas</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Ovchinnikov</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Lee</surname>
<given-names>G. R.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>Accurate Prediction of Protein Structures and Interactions Using a Three-Track Neural Network</article-title>. <source>Science</source> <volume>373</volume>, <fpage>871</fpage>&#x2013;<lpage>876</lpage>. <pub-id pub-id-type="doi">10.1126/science.abj8754</pub-id> </citation>
</ref>
<ref id="B3">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Bick</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Bick</surname>
<given-names>S. L.</given-names>
</name>
<name>
<surname>Dimmock</surname>
<given-names>D. P.</given-names>
</name>
<name>
<surname>Fowler</surname>
<given-names>T. A.</given-names>
</name>
<name>
<surname>Caulfield</surname>
<given-names>M. J.</given-names>
</name>
<name>
<surname>Scott</surname>
<given-names>R. H.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>An Online Compendium of Treatable Genetic Disorders</article-title>. <source>Am. J. Med. Genet.</source> <volume>187</volume>, <fpage>48</fpage>&#x2013;<lpage>54</lpage>. <pub-id pub-id-type="doi">10.1002/ajmg.c.31874</pub-id> </citation>
</ref>
<ref id="B4">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Botto</surname>
<given-names>L. D.</given-names>
</name>
<name>
<surname>May</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Fernhoff</surname>
<given-names>P. M.</given-names>
</name>
<name>
<surname>Correa</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Coleman</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Rasmussen</surname>
<given-names>S. A.</given-names>
</name>
<etal/>
</person-group> (<year>2003</year>). <article-title>A Population-Based Study of the 22q11.2 Deletion: Phenotype, Incidence, and Contribution to Major Birth Defects in the Population</article-title>. <source>Pediatrics</source> <volume>112</volume>, <fpage>101</fpage>&#x2013;<lpage>107</lpage>. <pub-id pub-id-type="doi">10.1542/peds.112.1.101</pub-id> </citation>
</ref>
<ref id="B5">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Byeon</surname>
<given-names>Y. J. J.</given-names>
</name>
<name>
<surname>Islamaj</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Yeganova</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Wilbur</surname>
<given-names>W. J.</given-names>
</name>
<name>
<surname>Lu</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Brody</surname>
<given-names>L. C.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>Evolving Use of Ancestry, Ethnicity, and Race in Genetics Research-A Survey Spanning Seven Decades</article-title>. <source>Am. J. Hum. Genet.</source> <volume>108</volume>, <fpage>2215</fpage>&#x2013;<lpage>2223</lpage>. <pub-id pub-id-type="doi">10.1016/j.ajhg.2021.10.008</pub-id> </citation>
</ref>
<ref id="B6">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Campbell</surname>
<given-names>I. M.</given-names>
</name>
<name>
<surname>Sheppard</surname>
<given-names>S. E.</given-names>
</name>
<name>
<surname>Crowley</surname>
<given-names>T. B.</given-names>
</name>
<name>
<surname>Mcginn</surname>
<given-names>D. E.</given-names>
</name>
<name>
<surname>Bailey</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Mcginn</surname>
<given-names>M. J.</given-names>
</name>
<etal/>
</person-group> (<year>2018</year>). <article-title>What Is New with 22q? an Update from the 22q and You Center at the Children&#x27;s Hospital of Philadelphia</article-title>. <source>Am. J. Med. Genet.</source> <volume>176</volume>, <fpage>2058</fpage>&#x2013;<lpage>2069</lpage>. <pub-id pub-id-type="doi">10.1002/ajmg.a.40637</pub-id> </citation>
</ref>
<ref id="B7">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Clark</surname>
<given-names>M. M.</given-names>
</name>
<name>
<surname>Hildreth</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Batalov</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Ding</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Chowdhury</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Watkins</surname>
<given-names>K.</given-names>
</name>
<etal/>
</person-group> (<year>2019</year>). <article-title>Diagnosis of Genetic Diseases in Seriously Ill Children by Rapid Whole-Genome Sequencing and Automated Phenotyping and Interpretation</article-title>. <source>Sci. Transl Med.</source> <volume>11</volume>. <pub-id pub-id-type="doi">10.1126/scitranslmed.aat6177</pub-id> </citation>
</ref>
<ref id="B8">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>DeGrave</surname>
<given-names>A. J.</given-names>
</name>
<name>
<surname>Janizek</surname>
<given-names>J. D.</given-names>
</name>
<name>
<surname>Lee</surname>
<given-names>S. I.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>AI for Radiographic COVID-19 Detection Selects Shortcuts over Signal</article-title>. <source>Nat. Mach. Intell.</source> <volume>3</volume> (<issue>7</issue>), <fpage>610</fpage>&#x2013;<lpage>619</lpage>. </citation>
</ref>
<ref id="B9">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Duong</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Waikel</surname>
<given-names>R. L.</given-names>
</name>
<name>
<surname>Hu</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Tekendo-Ngongang</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Solomon</surname>
<given-names>B. D.</given-names>
</name>
</person-group> (<year>2022b</year>). <article-title>Neural Network Classifiers for Images of Genetic Conditions with Cutaneous Manifestations</article-title>. <source>HGG Adv.</source> <volume>3</volume>, <fpage>100053</fpage>. <pub-id pub-id-type="doi">10.1016/j.xhgg.2021.100053</pub-id> </citation>
</ref>
<ref id="B10">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Duong</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Hu</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Tekendo-Ngongang</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Hanchard</surname>
<given-names>S. L.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Solomon</surname>
<given-names>B. D.</given-names>
</name>
<etal/>
</person-group> (<year>2021a</year>). <article-title>Neural Networks for Classification and Image Generation of Aging in Genetic Syndromes</article-title>. <source>medRxiv</source> <volume>2012</volume>, <fpage>21267472</fpage>. </citation>
</ref>
<ref id="B11">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ferreira</surname>
<given-names>C. R.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>The burden of Rare Diseases</article-title>. <source>Am. J. Med. Genet.</source> <volume>179</volume>, <fpage>885</fpage>&#x2013;<lpage>892</lpage>. <pub-id pub-id-type="doi">10.1002/ajmg.a.61124</pub-id> </citation>
</ref>
<ref id="B12">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Finlayson</surname>
<given-names>S. G.</given-names>
</name>
<name>
<surname>Lee</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Kohane</surname>
<given-names>I. S.</given-names>
</name>
<name>
<surname>Oakden-Rayner</surname>
<given-names>L.</given-names>
</name>
</person-group> (<year>2018</year>). <source>Towards Generative Adversarial Networks as a New Paradigm for Radiology Education</source>. <comment>arXiv preprint arXiv:1812.01547</comment>. </citation>
</ref>
<ref id="B13">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Frazer</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Notin</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Dias</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Gomez</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Min</surname>
<given-names>J. K.</given-names>
</name>
<name>
<surname>Brock</surname>
<given-names>K.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>Disease Variant Prediction with Deep Generative Models of Evolutionary Data</article-title>. <source>Nature</source>. <pub-id pub-id-type="doi">10.1038/s41586-021-04043-8</pub-id> </citation>
</ref>
<ref id="B14">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Gonzaludo</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Belmont</surname>
<given-names>J. W.</given-names>
</name>
<name>
<surname>Gainullin</surname>
<given-names>V. G.</given-names>
</name>
<name>
<surname>Taft</surname>
<given-names>R. J.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Estimating the burden and Economic Impact of Pediatric Genetic Disease</article-title>. <source>Genet. Med.</source> <volume>21</volume>, <fpage>1781</fpage>&#x2013;<lpage>1789</lpage>. <pub-id pub-id-type="doi">10.1038/s41436-018-0398-5</pub-id> </citation>
</ref>
<ref id="B15">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Gurovich</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Hanani</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Bar</surname>
<given-names>O.</given-names>
</name>
<name>
<surname>Nadav</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Fleischer</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Gelbman</surname>
<given-names>D.</given-names>
</name>
<etal/>
</person-group> (<year>2019</year>). <article-title>Identifying Facial Phenotypes of Genetic Disorders Using Deep Learning</article-title>. <source>Nat. Med.</source> <volume>25</volume>, <fpage>60</fpage>&#x2013;<lpage>64</lpage>. <pub-id pub-id-type="doi">10.1038/s41591-018-0279-0</pub-id> </citation>
</ref>
<ref id="B16">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hennekam</surname>
<given-names>R. C. M.</given-names>
</name>
<name>
<surname>Biesecker</surname>
<given-names>L. G.</given-names>
</name>
</person-group> (<year>2012</year>). <article-title>Next-generation Sequencing Demands Next-Generation Phenotyping</article-title>. <source>Hum. Mutat.</source> <volume>33</volume>, <fpage>884</fpage>&#x2013;<lpage>886</lpage>. <pub-id pub-id-type="doi">10.1002/humu.22048</pub-id> </citation>
</ref>
<ref id="B17">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Hsieh</surname>
<given-names>T.-C.</given-names>
</name>
<name>
<surname>Bar-Haim</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Moosa</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Ehmke</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Gripp</surname>
<given-names>K. W.</given-names>
</name>
<name>
<surname>Pantel</surname>
<given-names>J. T.</given-names>
</name>
<etal/>
</person-group> (<year>20212020</year>). <source>GestaltMatcher: Overcoming the Limits of Rare Disease Matching Using Facial Phenotypic Descriptors</source>. <publisher-loc>Cold Spring Harbor, NY</publisher-loc>: <publisher-name>medRxiv</publisher-name>, <fpage>2028.20248193</fpage>. </citation>
</ref>
<ref id="B18">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Jenkins</surname>
<given-names>B. D.</given-names>
</name>
<name>
<surname>Fischer</surname>
<given-names>C. G.</given-names>
</name>
<name>
<surname>Polito</surname>
<given-names>C. A.</given-names>
</name>
<name>
<surname>Maiese</surname>
<given-names>D. R.</given-names>
</name>
<name>
<surname>Keehn</surname>
<given-names>A. S.</given-names>
</name>
<name>
<surname>Lyon</surname>
<given-names>M.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>The 2019 US Medical Genetics Workforce: a Focus on Clinical Genetics</article-title>. <source>Genet. Med.</source> <pub-id pub-id-type="doi">10.1038/s41436-021-01162-5</pub-id> </citation>
</ref>
<ref id="B19">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Jumper</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Evans</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Pritzel</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Green</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Figurnov</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Ronneberger</surname>
<given-names>O.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>Highly Accurate Protein Structure Prediction with AlphaFold</article-title>. <source>Nature</source> <volume>596</volume>, <fpage>583</fpage>&#x2013;<lpage>589</lpage>. <pub-id pub-id-type="doi">10.1038/s41586-021-03819-2</pub-id> </citation>
</ref>
<ref id="B20">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>K&#xe4;rkk&#xe4;inen</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Joo</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2019</year>). <source>Fairface: Face Attribute Dataset for Balanced Race, Gender, and Age</source>. <comment>arXiv preprint arXiv:1908.04913</comment>. </citation>
</ref>
<ref id="B21">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Karras</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Aittala</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Hellsten</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Laine</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Lehtinen</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Aila</surname>
<given-names>T.</given-names>
</name>
</person-group> (<year>2020</year>). <source>Training Generative Adversarial Networks with Limited Data</source>. <comment>arXiv preprint arXiv:2006.06676</comment>. </citation>
</ref>
<ref id="B22">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Liu</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Mo</surname>
<given-names>Z.-H.</given-names>
</name>
<name>
<surname>Yang</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>Z.-F.</given-names>
</name>
<name>
<surname>Hong</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Wen</surname>
<given-names>L.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>Automatic Facial Recognition of Williams-Beuren Syndrome Based on Deep Convolutional Neural Networks</article-title>. <source>Front. Pediatr.</source> <volume>9</volume>, <fpage>648255</fpage>. <pub-id pub-id-type="doi">10.3389/fped.2021.648255</pub-id> </citation>
</ref>
<ref id="B23">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Luo</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Sedlazeck</surname>
<given-names>F. J.</given-names>
</name>
<name>
<surname>Lam</surname>
<given-names>T.-W.</given-names>
</name>
<name>
<surname>Schatz</surname>
<given-names>M. C.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>A Multi-Task Convolutional Deep Neural Network for Variant Calling in Single Molecule Sequencing</article-title>. <source>Nat. Commun.</source> <volume>10</volume>, <fpage>998</fpage>. <pub-id pub-id-type="doi">10.1038/s41467-019-09025-z</pub-id> </citation>
</ref>
<ref id="B24">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Maiese</surname>
<given-names>D. R.</given-names>
</name>
<name>
<surname>Keehn</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Lyon</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Flannery</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Watson</surname>
<given-names>M.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Working Groups of the National Coordinating Center for Seven Regional Genetics Service Current Conditions in Medical Genetics Practice</article-title>. <source>Genet. Med.</source> <volume>21</volume>, <fpage>1874</fpage>&#x2013;<lpage>1877</lpage>. <pub-id pub-id-type="doi">10.1038/s41436-018-0417-6</pub-id> </citation>
</ref>
<ref id="B25">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Meftah</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Semmar</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Tahiri</surname>
<given-names>M.-A.</given-names>
</name>
<name>
<surname>Tamaazousti</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Essafi</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Sadat</surname>
<given-names>F.</given-names>
</name>
</person-group> (<year>2020</year>). &#x201c;<article-title>Multi-Task Supervised Pretraining for Neural Domain Adaptation</article-title>,&#x201d; in <source>Proceedings of the Eighth International Workshop on Natural Language Processing for Social Media</source>, <fpage>61</fpage>&#x2013;<lpage>71</lpage>. <pub-id pub-id-type="doi">10.18653/v1/2020.socialnlp-1.8</pub-id> </citation>
</ref>
<ref id="B26">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Morris</surname>
<given-names>C. A.</given-names>
</name>
<name>
<surname>Braddock</surname>
<given-names>S. R.</given-names>
</name>
<name>
<surname>Council On</surname>
<given-names>G.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Health Care Supervision for Children with Williams Syndrome</article-title>. <source>Pediatrics</source> <volume>145</volume>. <pub-id pub-id-type="doi">10.1542/peds.2019-3761</pub-id> </citation>
</ref>
<ref id="B27">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Morris</surname>
<given-names>C. A.</given-names>
</name>
<name>
<surname>Adam</surname>
<given-names>M. P.</given-names>
</name>
<name>
<surname>Ardinger</surname>
<given-names>H. H.</given-names>
</name>
<name>
<surname>Pagon</surname>
<given-names>R. A.</given-names>
</name>
<name>
<surname>Wallace</surname>
<given-names>S. E.</given-names>
</name>
<name>
<surname>Bean</surname>
<given-names>L. J. H.</given-names>
</name>
<etal/>
</person-group> (<year>1993</year>). &#x201c;<article-title>Williams Syndrome</article-title>,&#x201d; in <source>GeneReviews</source> (<publisher-loc>Seattle</publisher-loc>). </citation>
</ref>
<ref id="B28">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Or-El</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Sengupta</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Fried</surname>
<given-names>O.</given-names>
</name>
<name>
<surname>Shechtman</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Kemelmacher-Shlizerman</surname>
<given-names>I.</given-names>
</name>
</person-group> (<year>2020</year>). &#x201c;<article-title>Lifespan Age Transformation Synthesis</article-title>,&#x201d; in <source>European Conference on Computer Vision</source> (<publisher-name>Springer</publisher-name>), <fpage>739</fpage>&#x2013;<lpage>755</lpage>. </citation>
</ref>
<ref id="B29">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Oskarsdottir</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Vujic</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Fasth</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2004</year>). <article-title>Incidence and Prevalence of the 22q11 Deletion Syndrome: a Population-Based Study in Western Sweden</article-title>. <source>Arch. Dis. Child.</source> <volume>89</volume>, <fpage>148</fpage>&#x2013;<lpage>151</lpage>. <pub-id pub-id-type="doi">10.1136/adc.2003.026880</pub-id> </citation>
</ref>
<ref id="B30">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Porras</surname>
<given-names>A. R.</given-names>
</name>
<name>
<surname>Rosenbaum</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Tor-Diez</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Summar</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Linguraru</surname>
<given-names>M. G.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Development and Evaluation of a Machine Learning-Based point-of-care Screening Tool for Genetic Syndromes in Children: a Multinational Retrospective Study</article-title>. <source>Lancet Digit Health</source>. <pub-id pub-id-type="doi">10.1016/s2589-7500(21)00137-0</pub-id> </citation>
</ref>
<ref id="B31">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Qin</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Zhu</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Xue</surname>
<given-names>Y.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>A GAN-based Image Synthesis Method for Skin Lesion Classification</article-title>. <source>Comp. Methods Programs Biomed.</source> <volume>195</volume>, <fpage>105568</fpage>. <pub-id pub-id-type="doi">10.1016/j.cmpb.2020.105568</pub-id> </citation>
</ref>
<ref id="B32">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Saporta</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Gui</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Agrawal</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Pareek</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Nguyen</surname>
<given-names>D. T.</given-names>
</name>
<name>
<surname>Ngo</surname>
<given-names>V. D.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>Benchmarking Saliency Methods for Chest X-ray Interpretation</article-title>. <source>medRxiv</source>, <fpage>2021</fpage>&#x2013;<lpage>2022</lpage>. </citation>
</ref>
<ref id="B33">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Solomon</surname>
<given-names>B. D.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Can Artificial Intelligence Save Medical Genetics?</article-title> <source>Am. J. Med. Genet. A.</source> </citation>
</ref>
<ref id="B34">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Solomon</surname>
<given-names>B. D.</given-names>
</name>
<name>
<surname>Nguyen</surname>
<given-names>A.-D.</given-names>
</name>
<name>
<surname>Bear</surname>
<given-names>K. A.</given-names>
</name>
<name>
<surname>Wolfsberg</surname>
<given-names>T. G.</given-names>
</name>
</person-group> (<year>2013</year>). <article-title>Clinical Genomic Database</article-title>. <source>Proc. Natl. Acad. Sci.</source> <volume>110</volume>, <fpage>9851</fpage>&#x2013;<lpage>9855</lpage>. <pub-id pub-id-type="doi">10.1073/pnas.1302575110</pub-id> </citation>
</ref>
<ref id="B35">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Str&#xf8;mme</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Bj&#xf8;rnstad</surname>
<given-names>P. G.</given-names>
</name>
<name>
<surname>Ramstad</surname>
<given-names>K.</given-names>
</name>
</person-group> (<year>2002</year>). <article-title>Prevalence Estimation of Williams Syndrome</article-title>. <source>J. Child. Neurol.</source> <volume>17</volume>, <fpage>269</fpage>&#x2013;<lpage>271</lpage>. <pub-id pub-id-type="doi">10.1177/088307380201700406</pub-id> </citation>
</ref>
<ref id="B36">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Tan</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Le</surname>
<given-names>Q.</given-names>
</name>
</person-group> (<year>2019</year>). &#x201c;<article-title>Efficientnet: Rethinking Model Scaling for Convolutional Neural Networks</article-title>,&#x201d; in <source>International Conference on Machine Learning</source> (<publisher-loc>Brookline, MA</publisher-loc>: <publisher-name>PMLR</publisher-name>), <fpage>6105</fpage>&#x2013;<lpage>6114</lpage>. </citation>
</ref>
<ref id="B37">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tschandl</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Codella</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Akay</surname>
<given-names>B. N.</given-names>
</name>
<name>
<surname>Argenziano</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Braun</surname>
<given-names>R. P.</given-names>
</name>
<name>
<surname>Cabo</surname>
<given-names>H.</given-names>
</name>
<etal/>
</person-group> (<year>2019</year>). <article-title>Comparison of the Accuracy of Human Readers versus Machine-Learning Algorithms for Pigmented Skin Lesion Classification: an Open, Web-Based, International, Diagnostic Study</article-title>. <source>Lancet Oncol.</source> <volume>20</volume>, <fpage>938</fpage>&#x2013;<lpage>947</lpage>. <pub-id pub-id-type="doi">10.1016/s1470-2045(19)30333-x</pub-id> </citation>
</ref>
<ref id="B38">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tschandl</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Rosendahl</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Kittler</surname>
<given-names>H.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>The HAM10000 Dataset, a Large Collection of Multi-Source Dermatoscopic Images of Common Pigmented Skin Lesions</article-title>. <source>Sci. Data</source> <volume>5</volume>, <fpage>180161</fpage>. <pub-id pub-id-type="doi">10.1038/sdata.2018.161</pub-id> </citation>
</ref>
<ref id="B39">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Zeiler</surname>
<given-names>M. D.</given-names>
</name>
<name>
<surname>Fergus</surname>
<given-names>R.</given-names>
</name>
</person-group> (<year>2014</year>). &#x201c;<article-title>Visualizing and Understanding Convolutional Networks</article-title>,&#x201d; in <source>European Conference on Computer Vision</source> (<publisher-name>Springer</publisher-name>), <fpage>818</fpage>&#x2013;<lpage>833</lpage>. <pub-id pub-id-type="doi">10.1007/978-3-319-10590-1_53</pub-id> </citation>
</ref>
</ref-list>
</back>
</article>