<?xml version="1.0" encoding="utf-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance" article-type="research-article" dtd-version="2.3" xml:lang="EN">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Psychol.</journal-id>
<journal-title>Frontiers in Psychology</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Psychol.</abbrev-journal-title>
<issn pub-type="epub">1664-1078</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="doi">10.3389/fpsyg.2024.1520131</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Psychology</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Color and tone color: audiovisual crossmodal correspondences with musical instrument timbre</article-title>
</title-group>
<contrib-group>
<contrib contrib-type="author" corresp="yes"><name><surname>Reymore</surname> <given-names>Lindsey</given-names></name><xref ref-type="aff" rid="aff1"><sup>1</sup></xref><xref ref-type="aff" rid="aff2"><sup>2</sup></xref><xref ref-type="corresp" rid="c001"><sup>&#x002A;</sup></xref>
<uri xlink:href="https://loop.frontiersin.org/people/848801/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/funding-acquisition/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author"><name><surname>Lindsey</surname> <given-names>Delwin T.</given-names></name><xref ref-type="aff" rid="aff3"><sup>3</sup></xref><xref ref-type="aff" rid="aff4"><sup>4</sup></xref>
<uri xlink:href="https://loop.frontiersin.org/people/2902807/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
</contrib-group>
<aff id="aff1"><sup>1</sup><institution>School of Music, Dance and Theatre, Herberger Institute for Design and the Arts, Arizona State University</institution>, <addr-line>Tempe, AZ</addr-line>, <country>United States</country></aff>
<aff id="aff2"><sup>2</sup><institution>School of Music, The Ohio State University</institution>, <addr-line>Columbus, OH</addr-line>, <country>United States</country></aff>
<aff id="aff3"><sup>3</sup><institution>Department of Psychology, The Ohio State University</institution>, <addr-line>Columbus, OH</addr-line>, <country>United States</country></aff>
<aff id="aff4"><sup>4</sup><institution>College of Optometry, The Ohio State University</institution>, <addr-line>Columbus, OH</addr-line>, <country>United States</country></aff>
<author-notes>
<fn fn-type="edited-by" id="fn0001">
<p>Edited by: Asterios Zacharakis, Aristotle University of Thessaloniki, Greece</p>
</fn>
<fn fn-type="edited-by" id="fn0002">
<p>Reviewed by: Nicola Di Stefano, National Research Council (CNR), Italy</p>
<p>Emilios Cambouropoulos, Aristotle University of Thessaloniki, Greece</p>
</fn>
<corresp id="c001">&#x002A;Correspondence: Lindsey Reymore, <email>lreymore@asu.edu</email></corresp>
</author-notes>
<pub-date pub-type="epub">
<day>07</day>
<month>01</month>
<year>2025</year>
</pub-date>
<pub-date pub-type="collection">
<year>2024</year>
</pub-date>
<volume>15</volume>
<elocation-id>1520131</elocation-id>
<history>
<date date-type="received">
<day>30</day>
<month>10</month>
<year>2024</year>
</date>
<date date-type="accepted">
<day>17</day>
<month>12</month>
<year>2024</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#x00A9; 2025 Reymore and Lindsey.</copyright-statement>
<copyright-year>2025</copyright-year>
<copyright-holder>Reymore and Lindsey</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<abstract>
<p>Crossmodal correspondences, or widely shared tendencies for mapping experiences across sensory domains, are revealed in common descriptors of musical timbre such as <italic>bright</italic>, <italic>dark</italic>, and <italic>warm</italic>. Two experiments are reported in which participants listened to recordings of musical instruments playing major scales, selected colors to match the timbres, and rated the timbres on crossmodal semantic scales. Experiment A used three different keyboard instruments, each played in three pitch registers. Stimuli in Experiment B, representing six different orchestral instruments, were similar to those in Experiment A but were controlled for pitch register. Overall, results were consistent with hypothesized concordances between ratings on crossmodal timbre descriptors and participants&#x2019; color associations. Semantic ratings predicted the lightness and saturation of colors matched to instrument timbres; effects were larger when both pitch register and instrument type varied (Experiment A) but were still evident when pitch register was held constant (Experiment B). We also observed a weak relationship between participant ratings of musical stimuli on the terms <italic>warm</italic> and <italic>cool</italic> and the warmth-coolness of selected colors in Experiment B only. Results were generally consistent with the hypothesis that instrument type and pitch register are related to color choice, though we speculate that these associations may only be relevant for certain instruments. Overall, the results have implications for our understanding the relationship between music and color, suggesting that while timbre/color matching behavior is in many ways diverse, observable trends in strategy can in part be linked to crossmodal timbre semantics.</p>
</abstract>
<kwd-group>
<kwd>timbre</kwd>
<kwd>color</kwd>
<kwd>crossmodal correspondences</kwd>
<kwd>timbre semantics</kwd>
<kwd>music and color</kwd>
</kwd-group>
<counts>
<fig-count count="10"/>
<table-count count="3"/>
<equation-count count="0"/>
<ref-count count="69"/>
<page-count count="19"/>
<word-count count="14540"/>
</counts>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-at-acceptance</meta-name>
<meta-value>Auditory Cognitive Neuroscience</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec sec-type="intro" id="sec1">
<label>1</label>
<title>Introduction</title>
<p>Musical sounds are often described using terms common in other sensory domains; for example, sounds may be bright (visual), sweet, (gustatory), or rough (tactile). Crossmodal correspondences, as distinct from synesthesia, refer to tendencies in the general population for certain stimuli or features of stimuli in one sensory domain to be associated with stimuli or features in another sensory domain (<xref ref-type="bibr" rid="ref56">Spence and Sathian, 2020</xref>; see also <xref ref-type="bibr" rid="ref7">Deroy and Spence, 2013</xref>; <xref ref-type="bibr" rid="ref37">Motoki et al., 2023</xref>; <xref ref-type="bibr" rid="ref40">Parise and Spence, 2013</xref>; <xref ref-type="bibr" rid="ref52">Spence, 2011</xref>). These tendencies manifest in both literary and vernacular language, and previous research provides evidence consistent with the theory that crossmodal correspondences are relevant to perception as well as language (<xref ref-type="bibr" rid="ref29">Marks, 1996</xref>; <xref ref-type="bibr" rid="ref30">Marks, 2013</xref>). Complementary to previous literature on audio-visual crossmodal correspondences (see <xref ref-type="bibr" rid="ref54">Spence, 2020b</xref> for a review), which has primarily focused on the basic feature of pitch (e.g., <xref ref-type="bibr" rid="ref33">Martino and Marks, 1999</xref>; <xref ref-type="bibr" rid="ref59">Walker et al., 2012</xref>) or on complex musical compositions as color-evocative (see <xref ref-type="bibr" rid="ref53">Spence, 2020a</xref>), the experiments reported in this paper address correspondences between color and musical instrument timbre, engaging the question of how these correspondences may be reflected in crossmodal semantic timbre descriptors.</p>
<p>The documented widespread use of crossmodal terms in timbre lexicons (e.g., <xref ref-type="bibr" rid="ref49">Saitis and Weinzierl, 2019</xref>; <xref ref-type="bibr" rid="ref60">Wallmark, 2019a</xref>; <xref ref-type="bibr" rid="ref45">Reymore and Huron, 2020</xref>) suggests that timbre-color correspondences are a promising area of investigation. <xref ref-type="bibr" rid="ref50">Saitis et al. (2020)</xref> argue that studying crossmodal correspondences with timbre offers a new way of approaching questions about the mechanisms of auditory semantics, while at the same time, providing more general insight into timbre perception and human semantic processing. Relatively few previous studies have considered timbre as a contributing feature to crossmodal correspondences between sound/music and color. <xref ref-type="bibr" rid="ref1">Adeli et al. (2014)</xref> asked participants to match timbres to colored shapes and found that the &#x201C;softer&#x201D; timbres of marimba and piano were associated with blue and green shapes, while the &#x201C;harsher&#x201D; timbres of crash cymbals, gong, and triangle were matched with red or yellow shapes. An experiment by <xref ref-type="bibr" rid="ref43">Reuter et al. (2018)</xref> tested associations between single note samples at varying pitch heights from a variety of Western musical instruments and found that selected colors appeared to be related to both pitch height and spectral features, while also identifying a few associations between hue and instrument. Participants in <xref ref-type="bibr" rid="ref41">Qi et al. (2020)</xref> selected sounds produced by various Chinese instruments to match given colors; results suggested significant pitch-color associations as well as some select instrument timbre-color associations. Though these studies address color and timbre, they do not investigate the relevance of crossmodal linguistics. A recent study by <xref ref-type="bibr" rid="ref15">Gurman et al. (2021)</xref>, which explores correspondences between timbre and shape, demonstrates the potential for semantic tasks to supplement assessments of crossmodal correspondences. The authors argue that parallel semantic tasks can augment our understanding of crossmodal correspondences, offering insight into participants&#x2019; interpretation of the meaning of stimuli.</p>
<p>The present experiments were designed to examine the relationships between timbre-color associations and crossmodal semantic ratings. Participants first listened to a set of recordings of musical instruments presented in random order. Following each recording, they selected from a color palette all those colors that they felt were most consistent with the sound quality or characteristic of the musical instrument. Participants then listened to the recordings again in a random order, this time rating the degree to which each musical stimulus was associated with various crossmodal timbre descriptors. The principal aim of this study is to map relationships between people&#x2019;s use of crossmodal timbre descriptors and their individual color choices, assessing how closely participants&#x2019; conceptions of timbral characteristics align with their corresponding color choices. This approach can be distinguished from prior research, which has analyzed timbre-color correspondences using computationally derived audio descriptors (<xref ref-type="bibr" rid="ref22">Lindborg and Friberg, 2015</xref>; <xref ref-type="bibr" rid="ref43">Reuter et al., 2018</xref>) or researcher-generated descriptions (<xref ref-type="bibr" rid="ref1">Adeli et al., 2014</xref>).</p>
<p>Close interactions between timbre and pitch height merit careful consideration in timbre research. Although timbre contributes to sound source identification, a single sound source can produce many timbres (<xref ref-type="bibr" rid="ref51">Siedenburg and McAdams, 2017</xref>), and timbre varies notably across the pitch range of musical instruments (<xref ref-type="bibr" rid="ref44">Reymore, 2021</xref>; <xref ref-type="bibr" rid="ref46">Reymore et al., 2023</xref>). Furthermore, changes in fundamental frequency (F0) have been shown to affect perceived timbral brightness (<xref ref-type="bibr" rid="ref32">Marozeau and de Cheveign&#x00E9;, 2007</xref>). Many previous crossmodal experiments have used carefully controlled stimuli to attempt to isolate fundamental frequency and its perceptual correlate, pitch height, as parameters. However, pitch height and timbre are usually inextricable in everyday musical contexts: as pitch height changes, so does timbre. Our experimental design frames instrument type and pitch height as two sources of timbral variation, recognizing that pitch height and timbre cannot be fully disentangled in ecologically valid contexts. In Experiment A, we systematically vary both pitch height and instrument type, whereas Experiment B holds pitch height constant and offers comparison across a more timbrally diverse set of instruments.</p>
<p>To formulate hypotheses about lightness and saturation of color choices, we began by assembling an initial set of crossmodal terms that (1) have been previously identified in scholarship as common descriptors of timbre and (2) have been implicated in prior studies on audio-visual crossmodal research. We then expanded the initial set of descriptors by extrapolating based on prothetic (magnitude-related) relationships and polar alignments among the terms (see <xref ref-type="bibr" rid="ref52">Spence, 2011</xref>). The complete set of 12 timbre descriptors rated in the experiments includes the terms <italic>high</italic>, <italic>low</italic>, <italic>bright</italic>, <italic>dark</italic>, <italic>small</italic>, <italic>big</italic>, <italic>light in weight</italic>, <italic>heavy</italic>, <italic>happy</italic>, <italic>sad</italic>, <italic>warm</italic>, and <italic>cool</italic>.</p>
<p>We hypothesized the first 10 terms in the above list to predict the lightness and saturation of matched colors, whereas we hypothesized <italic>warm</italic> and <italic>cool</italic> to predict the warmth-coolness of matched colors. In the following sections, we review literature motivating our selection of semantic terms.</p>
<sec id="sec2">
<label>1.1</label>
<title>Brightness, lightness, and pitch</title>
<p>The terms <italic>bright</italic> and <italic>dark</italic>, used commonly to describe colors, are also widely acknowledged as some of the most frequently used timbre descriptors, and the relationship between visual and timbral brightness has been investigated through numerous perceptual studies. For example, <xref ref-type="bibr" rid="ref63">Wallmark et al. (2021)</xref> observed via a speeded classification task that incongruity between visual and timbral brightness increased error rate, though this did not affect response time; <xref ref-type="bibr" rid="ref48">Saitis and Wallmark (2024)</xref> observed that timbral brightness modulated the perception of pitch and possibly visual brightness. More generally, past research has provided robust evidence for both pitch-brightness and timbre-brightness correspondences (e.g., <xref ref-type="bibr" rid="ref26">Marks, 1982</xref>; <xref ref-type="bibr" rid="ref31">Marks et al., 1987</xref>; <xref ref-type="bibr" rid="ref61">Wallmark, 2019b</xref>; <xref ref-type="bibr" rid="ref62">Wallmark and Allen, 2020</xref>). While the mechanisms behind such congruences are not fully understood, the audio-visual connection is readily apparent.</p>
<p>A relationship between pitch height and lightness is also well-documented: higher pitches are associated with lighter colors, while lower pitches are associated with darker colors (<xref ref-type="bibr" rid="ref25">Marks, 1974</xref>; <xref ref-type="bibr" rid="ref27">Marks, 1987</xref>; <xref ref-type="bibr" rid="ref28">Marks, 1989</xref>; <xref ref-type="bibr" rid="ref35">Melara and Marks, 1990</xref>; <xref ref-type="bibr" rid="ref33">Martino and Marks, 1999</xref>; <xref ref-type="bibr" rid="ref36">Mondloch and Maurer, 2004</xref>; <xref ref-type="bibr" rid="ref65">Ward et al., 2006</xref>). Given that timbral variation can affect perception of pitch height (e.g., <xref ref-type="bibr" rid="ref66">Warrier and Zatorre, 2002</xref>; <xref ref-type="bibr" rid="ref20">Kuang et al., 2016</xref>), we included the terms <italic>high</italic> and <italic>low</italic> among our timbre semantic descriptors. We anticipated potential differences in semantic judgments on these terms even as pitch register is held constant (as in Experiment B)&#x2014;and that these differences would influence lightness of color choice. Taken together, previous research on relationships among pitch, timbre, visual lightness, and visual brightness indicates that timbres perceived as <italic>brighter</italic> and <italic>higher</italic> will tend to be matched with lighter colors, whereas <italic>darker-and lower</italic>-sounding timbres will tend to be matched with darker colors.</p>
</sec>
<sec id="sec3">
<label>1.2</label>
<title>Sound-size symbolism</title>
<p>The terms <italic>big</italic>, <italic>small</italic>, <italic>heavy</italic>, and <italic>light (in weight)</italic> are colloquially used in describing instrument timbre (<xref ref-type="bibr" rid="ref45">Reymore and Huron, 2020</xref>). Previous research has not directly explored correspondences between timbral characteristics and size or weight; however, we can extrapolate predictions from past experiments that have demonstrated robust pitch-size associations, where lower sounds are linked to larger size (<xref ref-type="bibr" rid="ref58">Walker and Smith, 1984</xref>; <xref ref-type="bibr" rid="ref36">Mondloch and Maurer, 2004</xref>; <xref ref-type="bibr" rid="ref13">Gallace and Spence, 2006</xref>; <xref ref-type="bibr" rid="ref12">Evans and Treisman, 2010</xref>; <xref ref-type="bibr" rid="ref5">Bien et al., 2012</xref>; <xref ref-type="bibr" rid="ref59">Walker et al., 2012</xref>).</p>
<p>Given (1) the close connections described in the previous section among visual lightness/brightness, timbral brightness, and pitch height, and (2) the observation that contrasts of <italic>bright-dark, high-low, light&#x2013;dark,</italic> and <italic>small-big</italic> are considered to be examples of polar dimensions that align with one another (<xref ref-type="bibr" rid="ref40">Parise and Spence, 2013</xref>), we included the terms <italic>small</italic>, <italic>big</italic>, <italic>heavy</italic>, and <italic>light in weight</italic> among our semantic descriptors. Triangulating timbre-brightness, pitch-brightness, and pitch-size correspondences led us to hypothesize that timbres judged to be <italic>smaller</italic> and <italic>lighter in weight</italic> will be matched to lighter colors, whereas <italic>bigger</italic> and <italic>heavier</italic> timbres will be matched to darker colors.</p>
<p>Across work on audiovisual crossmodal correspondences, the role of saturation has been studied less frequently than lightness. Here, research on heaviness, saturation, and pitch guided our predictions. <xref ref-type="bibr" rid="ref2">Alexander and Shansky (1976)</xref> observed that participants assigned heavier weights to darker and more saturated colors, while <xref ref-type="bibr" rid="ref57">Walker et al. (2017)</xref> found that objects perceived as heavier were judged to be darker and to make lower-pitched sounds. From here, we used prothetic relationships and polar alignments to extrapolate to other terms in our set. For example, <xref ref-type="bibr" rid="ref58">Walker and Smith (1984)</xref> found <italic>happy</italic> to align with <italic>high</italic>, <italic>little</italic>, and <italic>light (in weight)</italic> in an interference task; similarly, <xref ref-type="bibr" rid="ref11">Eitan and Timmers (2010)</xref> observed that <italic>light in weight-heavy</italic> were consistently matched with <italic>high-low</italic> for both English-and Hebrew-speaking participants. Based on these findings, we hypothesized that timbres rated as <italic>heavier</italic> would not only be matched to darker colors, but to more saturated colors, guided principally by the findings of <xref ref-type="bibr" rid="ref2">Alexander and Shansky (1976)</xref>. In parallel, we anticipated <italic>big</italic> would also correspond to more saturated colors, whereas <italic>small</italic> and <italic>light (in weight)</italic> would map to less saturated colors.</p>
<p>Notably, an argument could also be made in favor of a relationship with saturation in the opposite direction, based on other previous studies. From a series of experiments using the implicit associations test, <xref ref-type="bibr" rid="ref3">Anikin and Johansson (2019)</xref> found saturation to be positively associated with frequency and spectral centroid. <xref ref-type="bibr" rid="ref16">Hamilton-Fletcher et al. (2017)</xref> asked participants to adjust hue and chroma to heard sounds while luminance was held constant. They similarly found that frequency, as well as a timbral manipulation caused by adjusting the spectral center of gravity, was positively correlated with saturation.</p>
</sec>
<sec id="sec4">
<label>1.3</label>
<title>Emotion</title>
<p>Open-ended interviews suggest that musicians readily use emotion-related words in timbral discourse (<xref ref-type="bibr" rid="ref45">Reymore and Huron, 2020</xref>). <xref ref-type="bibr" rid="ref8">Di Stefano (2023)</xref> presents an account of how timbre can be associated with emotional qualities via the concept of atmosphere, which has an essential affective component. He notes that the relationship between timbre and perceived emotion does not depend on high-level cognitive mediation but rather is motivated by acoustic features. Indeed, behavioral evidence implicates instrument timbre in ratings of valence, tension arousal, and energy arousal, with increases on these scales corresponding to changes in acoustical descriptors such as fundamental frequency and spectral centroid (e.g., <xref ref-type="bibr" rid="ref34">McAdams et al., 2017</xref>; <xref ref-type="bibr" rid="ref10">Eerola et al., 2012</xref>; <xref ref-type="bibr" rid="ref18">Korsmit et al., 2023</xref>; see also <xref ref-type="bibr" rid="ref19">Korsmit et al., 2024</xref>).</p>
<p><xref ref-type="bibr" rid="ref55">Spence and Di Stefano (2022)</xref> argue for an emotional mediation account of color-sound correspondences, positing that ultimately, emotion explains these correspondences, &#x201C;no matter whether the stimuli are simple or complex&#x201D; (p. 30). <xref ref-type="bibr" rid="ref39">Palmer et al. (2013)</xref> found strong evidence consistent with an emotion mediation account in the context of matching excerpts of classical orchestral music to color. The emotion mediation hypothesis has been supported by findings from subsequent studies using single-line piano melodies (<xref ref-type="bibr" rid="ref38">Palmer et al., 2016</xref>), excerpts from J.S. Bach&#x2019;s <italic>Well-Tempered Clavier</italic> (<xref ref-type="bibr" rid="ref17">Isbilen and Krumhansl, 2016</xref>), and musical excerpts from a variety of genres (<xref ref-type="bibr" rid="ref68">Whiteford et al., 2018</xref>). Should emotion mapping be the primary mechanism at work in timbre-color matching when musical content is controlled, we expect to observe ratings on the emotion words <italic>happy</italic> and <italic>sad</italic>, as semantic descriptions of timbre, to explain as much as or more variance in color choice as the other words in our set.</p>
</sec>
<sec id="sec5">
<label>1.4</label>
<title>Warm and cool</title>
<p>From previous research in timbre semantics, two other descriptors stand out as ostensibly related to color associations: <italic>warm</italic> and <italic>cool</italic>. With respect to timbre description, <italic>warm</italic> is one of the most common descriptors in English and has been included in dimension labels across many timbre studies (see Saitis and Weinzierl for a review). <italic>Cool</italic>, or <italic>cold</italic>, as timbre descriptors, are far less common, but these terms do also emerge in timbre description (<xref ref-type="bibr" rid="ref45">Reymore and Huron, 2020</xref>). In everyday parlance, these words are most commonly used to describe temperature, but <italic>warm</italic> and <italic>cool</italic> are also applied readily to colors, with <italic>warm</italic> suggesting colors like red, orange, and yellow, and <italic>cool</italic> suggesting colors like green and blue (likely related to temperature-associated imagery like fire and ice). This shared vocabulary across visual and auditory modalities suggests a potential association between warm and cool sounds and warm and cool colors; however, no previous studies have investigated this correspondence.</p>
</sec>
<sec id="sec6">
<label>1.5</label>
<title>Summary</title>
<p>In sum, to generate hypotheses about relationships between timbre description and timbre-color matching, we drew on common crossmodal timbre terms, choosing specific terms related to previously established correspondences between pitch-brightness/lightness, timbre-brightness, pitch-size/weight, and music-emotion. Motivated by the literature reviewed above, we assembled a set of ten terms that were used to predict both lightness and saturation of colors matched to timbres. Our experiments also address whether the terms <italic>warm</italic> and <italic>cool</italic>, often used to describe both timbre and color, are related in a crossmodal context.</p>
</sec>
</sec>
<sec id="sec7">
<label>2</label>
<title>Hypotheses</title>
<p>We report the results of two experiments in which participants rated instrument timbres on 12 crossmodal semantic descriptors and then matched the same heard excerpts with colors. Specifically, these experiments tested hypothesized relationships between individual semantic terms and the lightness, saturation, and warmth-coolness of the matched color choices. Experiment A systematically varied pitch register across three different keyboard instruments; Experiment B tested a more diverse set of six different musical instruments while holding pitch register constant.</p>
<p>The principal aim of this research is to examine relationships between crossmodal semantic descriptions of timbre and timbre-color associations. The experimental design also presented the opportunities to assess the relevance of sound source identity (musical instrument) and fundamental frequency to timbre-color associations and to explore trends in timbre-color matching across a range of musical instruments.</p>
<disp-quote>
<p><italic>H1</italic>: When asked to match timbres with colors, participants&#x2019; semantic assessments of timbres will be systematically associated with color choices.</p>
</disp-quote>
<disp-quote>
<p><italic>H1a</italic>: Semantic ratings on the terms <italic>bright</italic>, <italic>dark</italic>, <italic>high, low</italic>, <italic>light in weight</italic>, <italic>heavy</italic>, <italic>small</italic>, <italic>big</italic>, <italic>happy</italic>, and <italic>sad</italic> will predict the lightness of matched colors. Specifically, ratings on the terms <italic>bright, high, light in weight, small,</italic> and <italic>happy</italic> will be positively associated with lightness; ratings on <italic>dark</italic>, <italic>low</italic>, <italic>heavy</italic>, <italic>big</italic>, and <italic>sad</italic> will negatively correlate with lightness.</p>
</disp-quote>
<disp-quote>
<p><italic>H1b</italic>: Semantic ratings on the terms <italic>bright</italic>, <italic>dark</italic>, <italic>high, low</italic>, <italic>light in weight</italic>, <italic>heavy</italic>, <italic>small</italic>, <italic>big</italic>, <italic>happy</italic>, and <italic>sad</italic> are associated with saturation of matched colors. Specifically, ratings on the terms <italic>bright, high, light in weight, small,</italic> and <italic>happy</italic> will be negatively associated with saturation; ratings on <italic>dark</italic>, <italic>low</italic>, <italic>heavy</italic>, <italic>big</italic>, and <italic>sad</italic> will positively correlate with saturation.</p>
</disp-quote>
<disp-quote>
<p><italic>H1c</italic>: Semantic ratings on the terms <italic>warm</italic> and <italic>cool</italic> are associated with perceived warmth-coolness of matched colors.</p>
</disp-quote>
<disp-quote>
<p><italic>H2</italic>: When asked to match timbres with colors, lightness, saturation, and warmth-coolness of participants&#x2019; color choices will vary systematically across different musical instrument timbres.</p>
</disp-quote>
<disp-quote>
<p><italic>H3</italic>: When asked to match timbres with colors, lightness, saturation, and warmth-coolness of participants&#x2019; color choices will vary systematically with pitch register. Higher pitch register will result in increased lightness and decreased saturation.</p>
</disp-quote>
</sec>
<sec id="sec8">
<label>3</label>
<title>Experiment A</title>
<sec id="sec9">
<label>3.1</label>
<title>Methods</title>
<sec id="sec10">
<label>3.1.1</label>
<title>Participants</title>
<p>Participants (<italic>n</italic>&#x202F;=&#x202F;96; 41&#x202F;M, 52 F, 1 agender, 2 not reporting) were recruited from the Center for Science and Industry (COSI) in Columbus, Ohio, USA. Participant age ranged from 18 to 70 (<italic>M</italic>&#x202F;=&#x202F;30.2; <italic>SD</italic>&#x202F;=&#x202F;11.0). Three participants reported synesthesia, including one color-sound, one color-key, and one not specified; nine participants indicated they were not sure whether they experienced synesthesia. We did not exclude data from these participants.</p>
</sec>
<sec id="sec11">
<label>3.1.2</label>
<title>Stimuli</title>
<sec id="sec12">
<label>3.1.2.1</label>
<title>Recordings</title>
<p>Recordings were made on three different keyboard instruments using a Zoom H4N recorder: a Bl&#x00FC;thner grand piano, built in 1906, a Flemish harpsichord after the Colmar Ruckers 1624 original, built by Keith Hill, Op. 486, 2016, and a Lautenwerk or lute harpsichord, built by Keith Hill, Op. 510, 2018. The lute harpsichord uses gut strings, while the Flemish harpsichord uses metal; this difference results in the former instrument producing a softer, more mellow sound and the latter producing a brighter, more nasal sound. Recordings were made on-site at Hill&#x2019;s workshop. The recording level was kept constant for all recordings, and the microphone was placed at roughly equal distances away from the instruments. All instruments were tuned within the half hour prior to their recordings. Three ascending major scales were recorded on each of the three instruments, beginning on F2, F3, and F4. Performers referred to a silent, blinking metronome, set at 88 beats per minute, to maintain a consistent tempo, holding each pitch for two beats. By using complete major scales as stimuli in the experiments reported here (as opposed to single notes or composed music) we consider stimuli that while for the most part simple, cover part of the middle ground between simple and complex, conferring an additional dimension of musicality while maintaining control of the musical content, which in composed music may introduce numerous confounding variables.</p>
</sec>
<sec id="sec13">
<label>3.1.2.2</label>
<title>Color palette</title>
<p>An approximation of the color palette used for our control study and all subsequent experiments described below is shown in <xref ref-type="fig" rid="fig1">Figure 1</xref>. The colors, which averaged 258&#x202F;cd/m<sup>2</sup> in maximum luminance, were spatially arranged in a 20&#x202F;&#x00D7;&#x202F;8 palette of color samples organized according to hue (palette columns spanning the color circle) and saturation from most saturated to almost white, displayed on an 85&#x202F;cd/m<sup>2</sup> 5000K gray background. Also included in each display were black, white, and gray samples located in a row immediately above the color palette.</p>
<fig position="float" id="fig1">
<label>Figure 1</label>
<caption>
<p>Color palette used in Experiments A and B. Colors are arranged in 20&#x202F;&#x00D7;&#x202F;8 matrix of chromatic colors plus black, white and gray. Here, hue varies horizontally left-to-right from reds through greens and blues to violets; hue in any column is constant but saturation increases from top to bottom. In order to minimize bias in color selections due to unintended configural cues in the palette, hues were rotated randomly from trial-to-trial and rows were &#x201C;flipped&#x201D; so that on some trials, saturation increased bottom to top.</p>
</caption>
<graphic xlink:href="fpsyg-15-1520131-g001.tif"/>
</fig>
<p>Colors for elements in the bottom row of the color palette as oriented in <xref ref-type="fig" rid="fig1">Figure 1</xref> were chosen to correspond to high-chroma hues spanning the color circle and to include examples of the lexical categories red, green, blue, yellow, orange and purple. Greens and blue-greens, and blues were somewhat less saturated than other colors, while blues were also somewhat less luminous than the other colors. Variations in saturation and luminance were dictated by constraints on color reproduction imposed by the iPad display&#x2019;s color gamut (P3). Finally, yellows were somewhat more luminous than average, as less luminous yellows were informally judged by us to be poor examples of this color category.</p>
<p>Once the bottom row of color palette colors had been established, the colorimetric purities of colors in subsequent rows were chosen to produce approximately equal steps in saturation for each hue from its maximum (bottom row) to a minimal (top row) difference from white (see <xref ref-type="bibr" rid="ref69">Wyszecki and Stiles, 1982</xref> for formal definitions). The colorimetric purities of the minimum-saturation colored samples were chosen to create predominantly white samples, slightly tinged with the corresponding hue for that column.</p>
<p>Participants had access to a slider simulated on the iPad touch screen that they could use to manually adjust the mean luminance of the color palette in a quasi-continuous fashion before making their color sample selections. As the slider was moved, the resulting color samples maintained their respective chromaticities; only the luminances of the samples relative to their respective maxima changed.</p>
<p>To address the concern that absolute spatial correspondences in the arrangement of palette colors could influence color choice (for example, less saturated samples always appearing toward the top of the palette), we introduced a random trial-to-trial variation in the absolute spatial arrangement of the palette samples displayed on the iPad. For example, on some trials, sample saturation increased from top-to-bottom, as shown in <xref ref-type="fig" rid="fig1">Figure 1</xref>, while on other trials, saturation increased from bottom-to-top. Additionally, mapping of the hue circle onto the matrix shown in <xref ref-type="fig" rid="fig1">Figure 1</xref> was randomly rotated from one trial to the next, while preserving hue order.</p>
</sec>
</sec>
<sec id="sec14">
<label>3.1.3</label>
<title>Procedure</title>
<p>Following a brief demographic survey, participants used a custom-programmed and calibrated iPad and headphones to listen to the musical stimuli and view the visual stimuli. In the first part of the experiment, participants were asked to listen to each excerpt and then to choose the color or colors that they felt best represented the sound quality and character of the musical instrument. Participants were provided with an array of colors (<xref ref-type="fig" rid="fig1">Figure 1</xref>). They were instructed to set the slider to the desired lightness and then use the touch screen of the iPad to select a single color or multiple colors, with no limit on the number of colors that could be selected.</p>
<p>In the second part of the experiment, participants listened to the same excerpts in a different random order and were asked to rate the appropriateness of various adjectives for describing the quality and character of the sound, including <italic>high, low, bright, dark, light in weight</italic>, <italic>heavy</italic>, <italic>small, big, happy, sad, warm</italic> and <italic>cool,</italic> using a continuous slider that ranged from &#x201C;not appropriate&#x201D; to &#x201C;very appropriate.&#x201D; The presentation order of the adjectives and of stimuli were randomized.</p>
</sec>
</sec>
<sec id="sec15">
<label>3.2</label>
<title>Results</title>
<p>Consensus plots (<xref ref-type="fig" rid="fig2">Figure 2</xref>) demonstrate relative agreement across subjects in their color selections for each keyboard instrument and pitch register tested in Experiment A; these plots map onto the palette used in Experiment A, as shown in the figure. Each palette shows contour plots of consensus relative to the maximum consensus for that instrument. The maximum consensus, indicated in the upper-right corner of each panel, is expressed as the proportion of participants selecting the most frequently-chosen color square. Note that participants were able to select multiple color squares. In Experiment A, the median number of color samples chosen for a given stimulus was 4 (IQR&#x202F;=&#x202F;6), with a minimum of 1 and a maximum of 44.</p>
<fig position="float" id="fig2">
<label>Figure 2</label>
<caption>
<p>Contour plots of relative consensus in color selections (irrespective of luminance settings) for keyboard instruments played in each of three pitch registers (F2, F3, F4) in Experiment A. Maximum consensus for each instrument/pitch condition is indicated in the upper-right corner of each panel. Contours have been adjusted in contrast to highlight regions of color plots with highest consensus values.</p>
</caption>
<graphic xlink:href="fpsyg-15-1520131-g002.tif"/>
</fig>
<p>Stimuli in the lowest octave (F2) were concentrated toward greater saturation levels with selections moving toward less saturation as pitch register increases, at least from F2 to F3. Overall, there was a bias toward more saturated colors. All three keyboard instruments in their lowest octaves tended to map onto blues and purples, while the Lautenwerk and the Flemish harpsichord also include a second concentration of responses in highly saturated reds, oranges, and browns. In the upper two octaves, selections for the piano show a diverse range of hues, while selections for the Lautenwerk and Flemish harpsichord show a predilection for reds, oranges, yellows, and some greens. The highest octaves in both harpsichords show higher consensus, particularly on yellows. The piano does not appear to have a unique hue profile that remains consistent as register changes; both harpsichords demonstrate concentrations of warmer responses in all three octaves, though at F2, they also demonstrate a second concentration in blues and purples that is consistent across the three instruments.</p>
</sec>
<sec id="sec16">
<label>3.3</label>
<title>Analysis</title>
<p>To assess relationships between semantic descriptors and color choices (H1), we constructed three sets of regression models to predict color-related dependent variables (lightness, saturation, warmth-coolness) from semantic ratings of individual terms. The effects of musical instrument (H2) and pitch register (H3) on color choices were tested via three two-way ANOVAs with lightness, saturation, and warmth-coolness as dependent variables.</p>
<sec id="sec17">
<label>3.3.1</label>
<title>Semantic ratings (H1)</title>
<p>To test hypothesized relationships between specific semantic terms and the lightness, saturation, and warmth-coolness of color choices, we evaluated significance for each semantic term with separate models. Modeling individual terms allowed us to (1) identify which terms have significant relationships to particular dimensions of matched colors and (2) quantify relative effect sizes.</p>
<p>In sum, three sets of models were built to test H1. The purpose of the first set was to assess the relationship of semantic terms with lightness (H1a); this set included ten models, one for each hypothesized term. Here, lightness was predicted from semantic ratings on the given word, as a fixed effect, with participant ID included as a random intercept. A similar process yielded the second set of models, now with saturation as the dependent variable (H1b). To test relationships between ratings on the terms <italic>warm</italic> and <italic>cool</italic> and the warmth-coolness of selected colors (H1c), we first calculated a measure of warmth-coolness (the warm-cool index), based on a control study (see <xref ref-type="supplementary-material" rid="SM1">Supplementary material S1</xref>). The third set of models included two models, one for <italic>warm</italic> and the other for <italic>cool</italic>, which predicted the warm-cool index from ratings on each term.</p>
<p>All models were built with R Statistical Software (v4.3.0, <xref ref-type="bibr" rid="ref42">R Core Team, 2023</xref>). We noted deviations from the assumptions of normality for linear regression, the extent of which varied across the models reported in this paper. Some of these deviations appear to be related to a ceiling effect for the lightness slider. In order to be able to compare results within and across experiments, it was necessary to use the same modeling techniques throughout. As linear mixed models are easier to interpret and thus preferable to report, we opted to first model all hypothesized relationships using both parametric (linear mixed models) and nonparametric methods (cumulative link mixed models) with the intention of comparing results: should the two methods yield converging results, we would report the more easily interpretable linear mixed models. Because results from both parametric and nonparametric versions of the models were indeed highly similar, we report results of the linear mixed models in the manuscript of the paper, noting two minor discrepancies accordingly. For all models, we include plots of the distributions of the residuals, qq plots, and full parametric and nonparametric model summary details in the <xref ref-type="supplementary-material" rid="SM1">Supplementary materials S2&#x2013;S4</xref>.</p>
<p>Linear mixed models were built using the lme4 package (<xref ref-type="bibr" rid="ref4">Bates et al., 2015</xref>); cumulative link mixed models were built with the <italic>clmm</italic> function in the ordinal package (<xref ref-type="bibr" rid="ref6">Christensen, 2023</xref>). Pseudo-<italic>R<sup>2</sup></italic> values were calculated using the <italic>r2_nakagawa</italic> function from the performance package (<xref ref-type="bibr" rid="ref23">L&#x00FC;decke et al., 2021</xref>). All variables were scaled and centered prior to modeling.</p>
<p>Finally, note that we report unadjusted <italic>p-</italic>values throughout the paper. Given the number of models reported in this paper, we anticipate that some results are spurious; effects with <italic>p-</italic>values close to 0.05 should be interpreted with caution. We make note of this where relevant in the Discussion.</p>
<sec id="sec18">
<label>3.3.1.1</label>
<title>Lightness (H1a)</title>
<p>Participants were able to choose as many or as few samples as they liked; for modeling, we computed the average lightness of a participant&#x2019;s chosen samples for each stimulus. A subset of participants (<italic>n</italic>&#x202F;=&#x202F;13) did not use the slider when choosing colors, defaulting to the lightest setting, while others (<italic>n</italic>&#x202F;=&#x202F;11) adjusted the lightness slider for only one stimulus. Because there was little or no variance on luminance for these participants, models of lightness are reported for the subset of participants (n&#x202F;=&#x202F;72) who made use of the slider for two or more trials.</p>
<p>H1 predicts significant relationships between lightness values of matched colors and timbre semantic ratings on 10 terms. Correlations between participants&#x2019; ratings and corresponding lightness values were in the predicted directions: that is, higher ratings on <italic>bright</italic>, <italic>high</italic>, <italic>light in weight</italic>, <italic>small</italic>, and <italic>happy</italic> corresponded with lighter colors, whereas higher ratings on <italic>dark</italic>, <italic>low</italic>, <italic>heavy</italic>, <italic>big</italic>, and <italic>sad</italic> corresponded with darker colors.</p>
<p>All 10 linear mixed models, which included random intercepts for participant ID, yielded significant results for the fixed effect of semantic rating. To provide a metric of effect size that can be used to compare results within and across experiments in this paper, we calculated pseudo-<italic>R<sup>2</sup></italic> values. <xref ref-type="table" rid="tab1">Table 1</xref> reports <italic>p</italic> values for the fixed effect in each model (i.e., ratings of the semantic term), along with marginal and conditional <italic>R<sup>2</sup></italic> values for each model. Here, marginal <italic>R<sup>2</sup></italic> represents an approximation of variance explained by fixed effects (in this case, semantic ratings for a model&#x2019;s given term), whereas conditional <italic>R<sup>2</sup></italic> represents an approximation of the total variance explained by both the fixed and random effects. These pseudo-<italic>R<sup>2</sup></italic> values are not equivalent to <italic>R<sup>2</sup></italic> in a standard linear model, but rather provide useful context for comparing across models within the paper and interpreting our overall results.</p>
<table-wrap position="float" id="tab1">
<label>Table 1</label>
<caption>
<p>Results of linear mixed models predicting the lightness of chosen colors from semantic ratings of 10 terms in both Experiment A and Experiment B.</p>
</caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="top" colspan="7">Lightness ~ Semantic Rating&#x202F;+&#x202F;(1 | Participant)</th>
</tr>
<tr>
<th/>
<th align="center" valign="top" colspan="3">Experiment A (Keyboards)</th>
<th align="center" valign="top" colspan="3">Experiment B (Orchestral Instruments)</th>
</tr>
<tr>
<th/>
<th align="center" valign="middle">
<italic>p</italic>
</th>
<th align="center" valign="middle">Marginal <italic>R<sup>2</sup></italic></th>
<th align="center" valign="middle">Conditional <italic>R<sup>2</sup></italic></th>
<th align="center" valign="middle">
<italic>p</italic>
</th>
<th align="center" valign="middle">Marginal <italic>R<sup>2</sup></italic></th>
<th align="center" valign="middle">Conditional <italic>R<sup>2</sup></italic></th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="middle">Bright</td>
<td align="center" valign="middle">&#x003C;0.001</td>
<td align="center" valign="middle">0.14</td>
<td align="center" valign="middle">0.21</td>
<td align="center" valign="middle">0.003</td>
<td align="center" valign="middle">0.02</td>
<td align="center" valign="middle">0.16</td>
</tr>
<tr>
<td align="left" valign="middle">Dark</td>
<td align="center" valign="middle">&#x003C;0.001</td>
<td align="center" valign="middle">0.14</td>
<td align="center" valign="middle">0.19</td>
<td align="center" valign="middle">0.01</td>
<td align="center" valign="middle">0.02</td>
<td align="center" valign="middle">0.15</td>
</tr>
<tr>
<td align="left" valign="middle">High</td>
<td align="center" valign="middle">&#x003C;0.001</td>
<td align="center" valign="middle">0.15</td>
<td align="center" valign="middle">0.19</td>
<td align="center" valign="middle">0.01</td>
<td align="center" valign="middle">0.02</td>
<td align="center" valign="middle">0.16</td>
</tr>
<tr>
<td align="left" valign="middle">Low</td>
<td align="center" valign="middle">&#x003C;0.001</td>
<td align="center" valign="middle">0.16</td>
<td align="center" valign="middle">0.22</td>
<td align="center" valign="middle">0.001</td>
<td align="center" valign="middle">0.03</td>
<td align="center" valign="middle">0.16</td>
</tr>
<tr>
<td align="left" valign="middle">Light in weight</td>
<td align="center" valign="middle">&#x003C;0.001</td>
<td align="center" valign="middle">0.15</td>
<td align="center" valign="middle">0.23</td>
<td align="center" valign="middle">0.04</td>
<td align="center" valign="middle">0.01</td>
<td align="center" valign="middle">0.16</td>
</tr>
<tr>
<td align="left" valign="middle">Heavy</td>
<td align="center" valign="middle">&#x003C;0.001</td>
<td align="center" valign="middle">0.15</td>
<td align="center" valign="middle">0.21</td>
<td align="center" valign="middle">&#x003C;0.001</td>
<td align="center" valign="middle">0.03</td>
<td align="center" valign="middle">0.17</td>
</tr>
<tr>
<td align="left" valign="middle">Small</td>
<td align="center" valign="middle">&#x003C;0.001</td>
<td align="center" valign="middle">0.10</td>
<td align="center" valign="middle">0.16</td>
<td align="center" valign="middle">0.62</td>
<td align="center" valign="middle">0.00</td>
<td align="center" valign="middle">0.15</td>
</tr>
<tr>
<td align="left" valign="middle">Big</td>
<td align="center" valign="middle">&#x003C;0.001</td>
<td align="center" valign="middle">0.12</td>
<td align="center" valign="middle">0.17</td>
<td align="center" valign="middle">0.73</td>
<td align="center" valign="middle">0.00</td>
<td align="center" valign="middle">0.15</td>
</tr>
<tr>
<td align="left" valign="middle">Happy</td>
<td align="center" valign="middle">&#x003C;0.001</td>
<td align="center" valign="middle">0.09</td>
<td align="center" valign="middle">0.14</td>
<td align="center" valign="middle">0.02</td>
<td align="center" valign="middle">0.01</td>
<td align="center" valign="middle">0.16</td>
</tr>
<tr>
<td align="left" valign="middle">Sad</td>
<td align="center" valign="middle">&#x003C;0.001</td>
<td align="center" valign="middle">0.08</td>
<td align="center" valign="middle">0.12</td>
<td align="center" valign="middle">0.001</td>
<td align="center" valign="middle">0.03</td>
<td align="center" valign="middle">0.18</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<p>For each model, the table includes <italic>p-</italic>values of the fixed effect along with corresponding marginal and conditional <italic>R<sup>2</sup></italic> values. For individual model summaries and additional detail, see <xref ref-type="supplementary-material" rid="SM1">Supplementary materials S4.3, S4.4</xref>. Summary of models predicting lightness from semantic ratings.</p>
</table-wrap-foot>
</table-wrap>
</sec>
<sec id="sec19">
<label>3.3.1.2</label>
<title>Saturation (H1b)</title>
<p>In cases where participants made multiple color selections, saturation values were averaged for modeling, in parallel to our treatment of the lightness data. In deciding to use the mean for analysis, we first evaluated the data and found that color selections were almost always made in single clusters of contiguous color samples, for which the mean was a reasonable representation.</p>
<p>All correlations between semantic ratings and saturation were in the predicted directions, where higher ratings on <italic>bright</italic>, <italic>high</italic>, <italic>light in weight</italic>, <italic>small</italic>, and <italic>happy</italic> corresponded with less saturated colors, whereas higher ratings on <italic>dark</italic>, <italic>low</italic>, <italic>heavy</italic>, <italic>big</italic>, and <italic>sad</italic> corresponded with more saturated colors.</p>
<p>Structures for the linear mixed models predicting saturation were identical to those built for lightness, except that the average saturation of chosen colors replaced lightness as the dependent variable. We included all participant observations in these models, as use of the slider affected variance in lightness only.</p>
<p>Most models yielded a significant effect of semantic rating, with generally higher <italic>p</italic> values and lower pseudo-<italic>R<sup>2</sup></italic> values in comparison to the lightness models. Ratings on the term <italic>sad</italic> were not significant; models for two other terms&#x2014;<italic>dark</italic> and <italic>small</italic>&#x2014;achieved significance when modeled using nonparametric methods, but not in linear mixed models (see <xref ref-type="supplementary-material" rid="SM1">Supplementary material S4.5</xref>). Generally, while most terms are significant in predicting saturation, effect sizes are small. <xref ref-type="table" rid="tab2">Table 2</xref> below reports <italic>p</italic> values for fixed effects, along with the marginal and conditional <italic>R<sup>2</sup></italic> values for each model.</p>
<table-wrap position="float" id="tab2">
<label>Table 2</label>
<caption>
<p>Results of linear mixed models predicting the saturation of chosen colors from semantic ratings of 10 terms, including the <italic>p</italic> values of the fixed effect along with marginal and conditional <italic>R<sup>2</sup></italic> values.</p>
</caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="top" colspan="7">Saturation&#x202F;~&#x202F;Semantic Rating&#x202F;+&#x202F;(1 | Participant)</th>
</tr>
<tr>
<th/>
<th align="center" valign="top" colspan="3">Experiment A (Keyboards)</th>
<th align="center" valign="top" colspan="3">Experiment B (Orchestral Instruments)</th>
</tr>
<tr>
<th/>
<th align="center" valign="middle">
<italic>p</italic>
</th>
<th align="center" valign="middle">Marginal <italic>R<sup>2</sup></italic></th>
<th align="center" valign="middle">Conditional <italic>R<sup>2</sup></italic></th>
<th align="center" valign="middle">
<italic>p</italic>
</th>
<th align="center" valign="middle">Marginal <italic>R<sup>2</sup></italic></th>
<th align="center" valign="middle">Conditional <italic>R<sup>2</sup></italic></th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="middle">Bright</td>
<td align="center" valign="middle">0.01</td>
<td align="center" valign="middle">0.01</td>
<td align="center" valign="middle">0.14</td>
<td align="center" valign="middle">0.17</td>
<td align="center" valign="middle">0.00</td>
<td align="center" valign="middle">0.32</td>
</tr>
<tr>
<td align="left" valign="middle">Dark&#x002A;</td>
<td align="center" valign="middle">0.11</td>
<td align="center" valign="middle">0.00</td>
<td align="center" valign="middle">0.14</td>
<td align="center" valign="middle">0.94</td>
<td align="center" valign="middle">0.00</td>
<td align="center" valign="middle">0.32</td>
</tr>
<tr>
<td align="left" valign="middle">High</td>
<td align="center" valign="middle">0.01</td>
<td align="center" valign="middle">0.01</td>
<td align="center" valign="middle">0.14</td>
<td align="center" valign="middle">0.15</td>
<td align="center" valign="middle">0.00</td>
<td align="center" valign="middle">0.32</td>
</tr>
<tr>
<td align="left" valign="middle">Low</td>
<td align="center" valign="middle">0.006</td>
<td align="center" valign="middle">0.01</td>
<td align="center" valign="middle">0.14</td>
<td align="center" valign="middle">0.21</td>
<td align="center" valign="middle">0.00</td>
<td align="center" valign="middle">0.32</td>
</tr>
<tr>
<td align="left" valign="middle">Light in weight</td>
<td align="center" valign="middle">0.003</td>
<td align="center" valign="middle">0.01</td>
<td align="center" valign="middle">0.14</td>
<td align="center" valign="middle">0.02</td>
<td align="center" valign="middle">0.01</td>
<td align="center" valign="middle">0.33</td>
</tr>
<tr>
<td align="left" valign="middle">Heavy</td>
<td align="center" valign="middle">0.01</td>
<td align="center" valign="middle">0.01</td>
<td align="center" valign="middle">0.14</td>
<td align="center" valign="middle">0.08</td>
<td align="center" valign="middle">0.01</td>
<td align="center" valign="middle">0.32</td>
</tr>
<tr>
<td align="left" valign="middle">Small&#x002A;</td>
<td align="center" valign="middle">0.05</td>
<td align="center" valign="middle">0.00</td>
<td align="center" valign="middle">0.14</td>
<td align="center" valign="middle">0.12</td>
<td align="center" valign="middle">0.00</td>
<td align="center" valign="middle">0.33</td>
</tr>
<tr>
<td align="left" valign="middle">Big</td>
<td align="center" valign="middle">0.01</td>
<td align="center" valign="middle">0.01</td>
<td align="center" valign="middle">0.14</td>
<td align="center" valign="middle">&#x003C;0.001</td>
<td align="center" valign="middle">0.02</td>
<td align="center" valign="middle">0.23</td>
</tr>
<tr>
<td align="left" valign="middle">Happy</td>
<td align="center" valign="middle">0.01</td>
<td align="center" valign="middle">0.01</td>
<td align="center" valign="middle">0.14</td>
<td align="center" valign="middle">0.06</td>
<td align="center" valign="middle">0.01</td>
<td align="center" valign="middle">0.32</td>
</tr>
<tr>
<td align="left" valign="middle">Sad</td>
<td align="center" valign="middle">0.25</td>
<td align="center" valign="middle">0.00</td>
<td align="center" valign="middle">0.14</td>
<td align="center" valign="middle">0.27</td>
<td align="center" valign="middle">0.00</td>
<td align="center" valign="middle">0.33</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<p>For individual model summaries and additional detail, see <xref ref-type="supplementary-material" rid="SM1">Supplementary materials S4.5, S4.6</xref>. Summary of models predicting saturation from semantic ratings. &#x002A;Among the saturation models, the terms <italic>dark</italic> and <italic>small</italic> did not reach significance in the linear mixed models but are significant when modeled using cumulative link mixed models. See <xref ref-type="supplementary-material" rid="SM1">Supplementary material S4.5</xref>.</p>
</table-wrap-foot>
</table-wrap>
</sec>
<sec id="sec20">
<label>3.3.1.3</label>
<title>Warm-cool (H1c)</title>
<p>The panels in <xref ref-type="fig" rid="fig3">Figure 3</xref> show the centroids (circular mean hue and mean saturation) for participants&#x2019; color selections for each instrument/pitch register condition in Experiment A, irrespective of luminance settings. Centroids are plotted on top of a background depicting the regions of warm (reddish) and cool (bluish) palette colors based on a template, derived from our control study (see <xref ref-type="supplementary-material" rid="SM1">Supplementary material S1</xref>). Circular means (see <xref ref-type="bibr" rid="ref24">Mardia and Jupp, 2009</xref>) were computed as the mean angular location of the color selections on a color circle, assuming all 20 hues represented in our color palette are equally spaced on a circle. These angles were then mapped onto the horizontal dimension of each panel shown in <xref ref-type="fig" rid="fig3">Figure 3</xref>. The saturation component (vertical dimension) of each centroid was calculated as the arithmetic mean of saturations of the colors selected by a participant.</p>
<fig position="float" id="fig3">
<label>Figure 3</label>
<caption>
<p>Centroids of participants&#x2019; color selections for keyboard instruments in Experiment A. Each white dot in each panel plots one centroid for one participant. Plotting convention is by hue vs. saturation, as depicted in reference panel at top of figure. The background shades of red and blue shown in the nine lower panels depict warm-cool consensus for the palette colors obtained in a separate control experiment (details in <xref ref-type="supplementary-material" rid="SM1">Supplementary material S1</xref>).</p>
</caption>
<graphic xlink:href="fpsyg-15-1520131-g003.tif"/>
</fig>
<p>Note that while participants&#x2019; centroids exhibit considerable dispersion across the individual panels, there is a tendency in some panels, such as with the Lautenwerk at F2, for centroids to aggregate in multiple regions.</p>
<p>To operationalize the degree of warmth vs. coolness of color choices, we calculated a &#x201C;warm-cool index,&#x201D; <italic>I<sup>m</sup></italic>, based on the results of the control study (see <xref ref-type="supplementary-material" rid="SM1">Supplementary material S1</xref>). This statistic quantifies, for each stimulus, the average bias in each participant&#x2019;s color selections toward either warm (positive values) or cool (negative values) color selections.</p>
<p>The warm-cool index, <italic>I<sup>m</sup></italic>, was used in statistical tests of H1c, which predicted that ratings of timbre on the terms <italic>warm</italic> and <italic>cool</italic> are associated with the warmth-coolness of matched colors. Two regression models were built predicting the warm-cool index from ratings of the terms <italic>warm</italic> and <italic>cool</italic>, with participant ID included as a random intercept. No significant effect of ratings on the terms <italic>warm</italic> or <italic>cool</italic> were observed on the warmth-coolness of the selected corresponding colors (see <xref ref-type="supplementary-material" rid="SM1">Supplementary material S4.7</xref> for model summaries). <xref ref-type="table" rid="tab3">Table 3</xref> includes results, with <italic>p-</italic>values as well as marginal and conditional R2 values for <italic>warm</italic> and <italic>cool</italic> models.</p>
<table-wrap position="float" id="tab3">
<label>Table 3</label>
<caption>
<p>Results of linear mixed models predicting the warm-cool index of chosen colors from semantic ratings of <italic>warm</italic> and <italic>cool</italic>, including the <italic>p-</italic>values of the fixed effect along with marginal and conditional <italic>R<sup>2</sup></italic> values.</p>
</caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="top" colspan="7">Warm-Cool Index ~ Semantic Rating&#x202F;+&#x202F;(1 | Participant)</th>
</tr>
<tr>
<th/>
<th align="center" valign="top" colspan="3">Experiment A (Keyboards)</th>
<th align="center" valign="top" colspan="3">Experiment B (Orchestral Instruments)</th>
</tr>
<tr>
<th/>
<th align="center" valign="middle">
<italic>p</italic>
</th>
<th align="center" valign="middle">Marginal <italic>R<sup>2</sup></italic></th>
<th align="center" valign="middle">Conditional <italic>R<sup>2</sup></italic></th>
<th align="center" valign="middle">
<italic>p</italic>
</th>
<th align="center" valign="middle">Marginal <italic>R<sup>2</sup></italic></th>
<th align="center" valign="middle">Conditional <italic>R<sup>2</sup></italic></th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="middle">Warm</td>
<td align="center" valign="middle">0.72</td>
<td align="center" valign="middle">0.00</td>
<td align="center" valign="middle">0.04</td>
<td align="center" valign="middle">0.01</td>
<td align="center" valign="middle">0.01</td>
<td align="center" valign="middle">0.01</td>
</tr>
<tr>
<td align="left" valign="middle">Cool</td>
<td align="center" valign="middle">0.82</td>
<td align="center" valign="middle">0.00</td>
<td align="center" valign="middle">0.04</td>
<td align="center" valign="middle">&#x003C;0.001</td>
<td align="center" valign="middle">0.02</td>
<td align="center" valign="middle">0.02</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<p>Table includes results for both Experiments A and B. For individual model summaries and additional detail, see <xref ref-type="supplementary-material" rid="SM1">Supplementary material S4.7</xref>. Summary of models predicting warm-cool index from semantic ratings.</p>
</table-wrap-foot>
</table-wrap>
</sec>
</sec>
<sec id="sec21">
<label>3.3.2</label>
<title>Musical instrument (H2) and pitch register (H3)</title>
<p>To assess whether musical instrument (H2) and pitch register (H3) are related to color choices, we ran three two-way ANOVAs with lightness, saturation, and warmth-coolness as dependent variables. Independent variables for each model included pitch register (whether the scale began on F2, F3, or F4) and instrument type (piano, Flemish harpsichord, and Lautenwerk) as fixed effects, with participant ID as a random intercept (see <xref ref-type="supplementary-material" rid="SM1">Supplementary material S4.8</xref> for further detail).</p>
<sec id="sec22">
<label>3.3.2.1</label>
<title>Lightness</title>
<p>A two-way ANOVA indicated differences in group lightness means for the three pitch registers [<italic>F</italic>(2, 568)&#x202F;=&#x202F;131.18, <italic>p</italic>&#x202F;&#x003C;&#x202F;0.001] and for the three instrument types [<italic>F</italic>(2, 568)&#x202F;=&#x202F;5.39, <italic>p</italic>&#x202F;=&#x202F;0.004]; their interaction was not significant [<italic>F</italic>(4, 568)&#x202F;=&#x202F;0.25, <italic>p</italic>&#x202F;=&#x202F;0.91]. As with the semantic models, only observations from participants who used the slider more than once are included. <xref ref-type="fig" rid="fig4">Figure 4</xref> shows lightness as a function of pitch register for each of the three keyboard instruments.</p>
<fig position="float" id="fig4">
<label>Figure 4</label>
<caption>
<p>Mean color selection lightness for keyboard instruments across pitch registers in Experiment A. F2, F3, and F4 indicate the starting pitches for major scales. Error bars are +/&#x2212; 1&#x202F;SE.</p>
</caption>
<graphic xlink:href="fpsyg-15-1520131-g004.tif"/>
</fig>
</sec>
<sec id="sec23">
<label>3.3.2.2</label>
<title>Saturation</title>
<p>A two-way ANOVA indicated differences in group means for the three pitch registers [<italic>F</italic>(2, 760)&#x202F;=&#x202F;12.43, <italic>p</italic>&#x202F;&#x003C;&#x202F;0.001] and instrument [<italic>F</italic>(2, 760)&#x202F;=&#x202F;8.03 <italic>p</italic>&#x202F;&#x003C;&#x202F;0.001], while their interaction was not significant [<italic>F</italic>(4, 760)&#x202F;=&#x202F;1.13, <italic>p</italic>&#x202F;=&#x202F;0.34]. <xref ref-type="fig" rid="fig5">Figure 5</xref> plots saturation as a function of pitch register and instrument.</p>
<fig position="float" id="fig5">
<label>Figure 5</label>
<caption>
<p>Saturation (in CIELAB units from white) for keyboard instruments across pitch registers. F2, F3, and F4 indicate the starting pitches for major scales. Error bars are +/&#x2212; 1&#x202F;SE.</p>
</caption>
<graphic xlink:href="fpsyg-15-1520131-g005.tif"/>
</fig>
</sec>
<sec id="sec24">
<label>3.3.2.3</label>
<title>Warm-cool</title>
<p>A two-way ANOVA identified significant effects of both instrument [<italic>F</italic>(2, 760)&#x202F;=&#x202F;19.71; <italic>p</italic>&#x202F;&#x003C;&#x202F;0.001] and pitch [<italic>F</italic>(2, 760)&#x202F;=&#x202F;11.74; <italic>p</italic>&#x202F;&#x003C;&#x202F;0.001] with no significant interaction [<italic>F</italic>(4, 760)&#x202F;=&#x202F;1.33; <italic>p</italic>&#x202F;=&#x202F;0.26]. <xref ref-type="fig" rid="fig6">Figure 6</xref> illustrates mean warm-cool indices across pitch registers and instruments.</p>
<fig position="float" id="fig6">
<label>Figure 6</label>
<caption>
<p>Mean warm-cool indices for each instrument/pitch register condition. Error bars are +/&#x2212; 1&#x202F;SE.</p>
</caption>
<graphic xlink:href="fpsyg-15-1520131-g006.tif"/>
</fig>
</sec>
</sec>
</sec>
</sec>
<sec id="sec25">
<label>4</label>
<title>Experiment B</title>
<sec id="sec26">
<label>4.1</label>
<title>Methods</title>
<p>Experiment A varied both instrument timbre and pitch register; Experiment B was designed to experimentally control for pitch register while investigating timbre-color correspondences across a more diverse range of instrument timbres.</p>
<sec id="sec27">
<label>4.1.1</label>
<title>Participants</title>
<p>Ninety-two participants (39&#x202F;M, 52 F, 1 genderqueer) were recruited from the Center for Science and Industry (COSI; <italic>n</italic>&#x202F;=&#x202F;70) and The Ohio State University music school subject pool (<italic>n</italic>&#x202F;=&#x202F;22), which is composed of second-year music students enrolled in Aural Skills. The participant group in Experiment B does not overlap with that of Experiment A. Two participants reported synesthesia, including one color-sound associator; three participants indicated they were unsure as to whether they experienced synesthesia. We did not exclude data from these participants.</p>
</sec>
<sec id="sec28">
<label>4.1.2</label>
<title>Stimuli</title>
<sec id="sec29">
<label>4.1.2.1</label>
<title>Recordings</title>
<p>Experiment B tested H1 and H2 with six orchestral instruments: flute, oboe, B&#x266D; clarinet, trumpet, violin, and viola. This set of instruments was selected based on a pilot study in which 12 musician participants were asked to fill out an open-response questionnaire that listed each semantic descriptor of interest (e.g., <italic>bright</italic>, <italic>dark</italic>) and asked for a single musical instrument whose sound was most representative of that descriptor. Because the purpose of Experiment B was to experimentally control for pitch across a larger group of instrument types than in Experiment A, we selected instruments that were able to play the same scale in a comfortable, middle range. One concern was avoiding extremes in range, as timbres in the extreme high or low register of an instrument are often noticeably different from the middle register; thus, we selected instruments whose overall ranges were as similar as possible.</p>
<p>Stimuli were recorded in a studio by professional musicians. Recordings used an AKG 414 microphone, set to cardioid pattern, through an API 3124 preamp, and Pro Tools, with an SSL Delta-Link interface, and were normalized to-18dBFS. The sampling rate was 44.1&#x202F;kHz/24. Stimuli for Experiment B included flute, oboe, B&#x266D; clarinet, trumpet, violin, and viola each playing one-octave major scales from F4 to F5, approximately 10&#x202F;s in length each.</p>
</sec>
<sec id="sec30">
<label>4.1.2.2</label>
<title>Color palette</title>
<p>The color palette was identical to that used in Experiment A.</p>
</sec>
</sec>
<sec id="sec31">
<label>4.1.3</label>
<title>Procedure</title>
<p>Aside from the use of different instruments in the musical stimuli, the procedure was identical to Experiment A.</p>
</sec>
</sec>
<sec id="sec32">
<label>4.2</label>
<title>Results</title>
<p>Consensus in participants&#x2019; color selections for the six orchestral instruments tested in Experiment B are shown in <xref ref-type="fig" rid="fig7">Figure 7</xref>. The format of each panel is the same as that for <xref ref-type="fig" rid="fig2">Figure 2</xref>, above, where contours indicate areas of relatively high consensus, and consensus is calculated irrespective of participants&#x2019; luminance settings. Maximum consensus values are displayed in the upper-right corner of each palette. In Experiment B, the median number of color samples chosen was 4 (IQR&#x202F;=&#x202F;4) with a minimum of 1 and a maximum of 49.</p>
<fig position="float" id="fig7">
<label>Figure 7</label>
<caption>
<p>Contour plots of relative consensus in color selections (irrespective of luminance settings) for orchestral instruments studied in Experiment B. Maximum consensus for each instrument is indicated in the upper-right corner of each panel. Contours have been adjusted in contrast to highlight regions of color plots with highest consensus values.</p>
</caption>
<graphic xlink:href="fpsyg-15-1520131-g007.tif"/>
</fig>
<p>While it is apparent that there is diversity in color selection, it is also clear that there is more agreement for some instruments than others. The flute and viola do not demonstrate obvious hue bias, although green and blue responses for the viola are less common compared to other hues. The oboe spans a number of hues, but color selections are more concentrated than for the flute or viola, avoiding teals and purples. Color selections for the trumpet show a hue bias toward reds, oranges, yellows. The violin demonstrates a tendency to be matched to reds and oranges. Clarinet selections are somewhat concentrated on blues.</p>
</sec>
<sec id="sec33">
<label>4.3</label>
<title>Analysis</title>
<p>Analysis for Experiment B replicates the procedures described for Experiment A.</p>
<sec id="sec34">
<label>4.3.1</label>
<title>Semantic ratings</title>
<sec id="sec35">
<label>4.3.1.1</label>
<title>Lightness</title>
<p>As in Experiment A, a subset of participants (<italic>n</italic>&#x202F;=&#x202F;22) did not use the slider when choosing colors, defaulting to the lightest setting, while another subset of participants (<italic>n</italic>&#x202F;=&#x202F;9) adjusted the lightness slider for only one of the stimuli, suggesting that they decided after the first trial not to adjust the lightness on subsequent trials. We report lightness models based on the subset of participants (<italic>n</italic>&#x202F;=&#x202F;61) who made use of the slider for more than one trial.</p>
<p>Most models yielded a significant effect of semantic ratings on lightness of color selections, except for the terms <italic>small</italic> and <italic>big</italic>. <xref ref-type="table" rid="tab1">Table 1</xref> includes the results across the 10 linear mixed models predicting lightness; summaries of all parametric and nonparametric models are included in <xref ref-type="supplementary-material" rid="SM1">Supplementary material S4.4</xref>.</p>
</sec>
<sec id="sec36">
<label>4.3.1.2</label>
<title>Saturation</title>
<p>In relating semantic ratings to the saturation of matched colors, only the terms <italic>light in weight</italic> and <italic>big</italic> yielded significant results, with small effect sizes. A summary of saturation models for Experiment B is included in <xref ref-type="table" rid="tab2">Table 2</xref>; summaries of all parametric and nonparametric models are included in <xref ref-type="supplementary-material" rid="SM1">Supplementary material S4.6</xref>.</p>
</sec>
<sec id="sec37">
<label>4.3.1.3</label>
<title>Warm-cool</title>
<p>Centroids for participants&#x2019; color selections for the six orchestral instruments tested in Experiment B are shown below in <xref ref-type="fig" rid="fig8">Figure 8</xref>. Background colors depict degrees of warmness (reddish areas) and coolness (blueish areas) associated with colors in the selection palette. Centroids are dispersed across both hue and saturation, particularly for the flute. As in Experiment A, some instruments (e.g., clarinet) show clusters in multiple regions.</p>
<fig position="float" id="fig8">
<label>Figure 8</label>
<caption>
<p>Centroids of participants&#x2019; color selections in Experiment B. Format of plots as in <xref ref-type="fig" rid="fig3">Figure 3</xref>.</p>
</caption>
<graphic xlink:href="fpsyg-15-1520131-g008.tif"/>
</fig>
<p>As in Experiment A, two regression models were built predicting the warm-cool index, with semantic terms <italic>warm</italic> and <italic>cool</italic> each as a fixed effect and participant ID as a random intercept in both models. Both yielded significant fixed effects with low effect sizes (<italic>warm</italic>: <italic>p</italic>&#x202F;=&#x202F;0.01, marginal <italic>R<sup>2</sup></italic>&#x202F;=&#x202F;0.01; <italic>cool</italic>: <italic>p</italic>&#x202F;&#x003C;&#x202F;0.001, marginal <italic>R<sup>2</sup></italic>&#x202F;=&#x202F;0.02; see <xref ref-type="table" rid="tab3">Table 3</xref>). Fit for the linear mixed models for both <italic>warm</italic> and <italic>cool</italic> was singular due to zero variance for the random intercept. Refitting the models without a random intercept resulted in highly similar results that were in agreement on the significance of terms.</p>
</sec>
</sec>
<sec id="sec38">
<label>4.3.2</label>
<title>Musical instrument (H2)</title>
<sec id="sec39">
<label>4.3.2.1</label>
<title>Lightness</title>
<p>Experiment B experimentally controlled for pitch register, allowing evaluation of H2 across a more diverse set of six different orchestral instrument types. To test whether instrument type has an effect on lightness of color selections, we fit a linear mixed model including instrument as a fixed effect and participant as a random intercept, including observations only from those participants who used the slider more than once during the experiment. A one-way ANOVA indicated no difference in group means within the six instruments [<italic>F</italic>(5, 300)&#x202F;=&#x202F;1.98, <italic>p</italic>&#x202F;=&#x202F;0.08].</p>
</sec>
<sec id="sec40">
<label>4.3.2.2</label>
<title>Saturation</title>
<p>To test whether instrument type has an effect on saturation of color selections, independent of pitch, we fit a linear mixed model including instrument as a fixed effect and participant as a random effect, including observations from all participants. A one-way ANOVA indicated differences in group means within the six instruments [<italic>F</italic>(5, 455)&#x202F;=&#x202F;2.89, <italic>p</italic>&#x202F;=&#x202F;0.01]. However, this difference appears to be solely driven by the trumpet in comparison to the other five instruments, as suggested by the plot in <xref ref-type="fig" rid="fig9">Figure 9</xref>.</p>
<fig position="float" id="fig9">
<label>Figure 9</label>
<caption>
<p>Average saturation (in CIELAB units from white) of color choices across participants for each instrument. Error bars are +/&#x2212; 1&#x202F;SE.</p>
</caption>
<graphic xlink:href="fpsyg-15-1520131-g009.tif"/>
</fig>
</sec>
<sec id="sec41">
<label>4.3.2.3</label>
<title>Warm-cool</title>
<p>A one-way ANOVA revealed an overall effect of instrument type on warmth-coolness [<italic>F</italic>(5, 455)&#x202F;=&#x202F;3.03, <italic>p</italic>&#x202F;=&#x202F;0.01]. The bar chart in <xref ref-type="fig" rid="fig10">Figure 10</xref> plots warm-cool indices <italic>I<sup>m</sup></italic> for the test instruments, arranged in ascending index order from average cool (clarinet) to average warm (trumpet) color selections (see 3.3.1.3 for an explanation of <italic>I<sup>m</sup></italic>). However, as was the case for the keyboard instruments tested in Experiment A, care must be taken in interpreting these results, given the substantial dispersion and, in some cases, bimodal aggregation, of the data.</p>
<fig position="float" id="fig10">
<label>Figure 10</label>
<caption>
<p>Average warm-cool indices by instrument, arranged in ascending warm-cool index order. Error bars are +/&#x2212; 1&#x202F;SE.</p>
</caption>
<graphic xlink:href="fpsyg-15-1520131-g010.tif"/>
</fig>
</sec>
</sec>
</sec>
</sec>
<sec id="sec42">
<label>5</label>
<title>General discussion</title>
<sec id="sec43">
<label>5.1</label>
<title>Semantic ratings (H1)</title>
<sec id="sec44">
<label>5.1.1</label>
<title>Lightness (H1a)</title>
<p>Overall, results are consistent with H1a, which posited that semantic characterizations of timbre predict lightness of timbre-color matching. All terms were statistically significant in Experiment A, which systematically varied pitch register; in Experiment B, where pitch register was experimentally controlled, all terms except <italic>small</italic> and <italic>big</italic> were significant. Effect sizes for significant terms, as quantified with pseudo-<italic>R<sup>2</sup></italic>, were larger in Experiment A, when pitch register varied (range&#x202F;=&#x202F;0.08&#x2013;0.16) than in Experiment B (range&#x202F;=&#x202F;0.01&#x2013;0.03), when pitch register was held constant.</p>
<p>Ratings on the semantic descriptors <italic>small</italic> and <italic>big</italic> were significant predictors of lightness when pitch register varied systematically, but not when it was held constant. Results of Experiment A are consistent with established understanding of sound-size symbolism, where higher pitch is perceived as corresponding to small objects and lower pitch is perceived as corresponding to bigger objects. That we did not observe the same relationship in Experiment B suggests that as descriptors of sound, these crossmodal terms are primarily or entirely driven by pitch register. Because we used ecologically relevant recordings of musical instruments as stimuli, we cannot fully disentangle change in fundamental frequency from timbral changes that correspond with increased pitch height. Because timbre varies with pitch height (<xref ref-type="bibr" rid="ref44">Reymore, 2021</xref>; <xref ref-type="bibr" rid="ref46">Reymore et al., 2023</xref>), crossmodal correspondences that vary with pitch register may be related to changes in both fundamental frequency and spectral components.</p>
<p>Yet, although <italic>small</italic> and <italic>big</italic> were not significant in Experiment B, the semantically related terms <italic>heavy</italic> and <italic>light in weight</italic> were. Among the eight significant terms in Experiment B, <italic>light in weight</italic> is ostensibly spurious, with a <italic>p</italic> value of 0.04. <italic>Heavy</italic> resulted in the lowest <italic>p</italic> value and highest marginal pseudo-<italic>R<sup>2</sup></italic> among the Experiment B lightness models, but notably did not reach significance for the corresponding saturation model. Perhaps timbral weight carries more relevance for lightness, independent of pitch register, than does size, but these relationships need further study for confirmation and clarity.</p>
</sec>
<sec id="sec45">
<label>5.1.2</label>
<title>Saturation (H1b)</title>
<p>Given established connections among saturation, pitch, and heaviness, we constructed H1b using the same set of terms derived for the lightness hypothesis. Results demonstrated that these terms were, for the most part, significantly related to saturation in Experiment A but not in Experiment B. Experiment A yielded clear significant relationships between saturation and ratings on <italic>bright</italic>, <italic>high</italic>, <italic>low</italic>, <italic>light in weight</italic>, <italic>heavy</italic>, <italic>big</italic>, and <italic>happy</italic>. The terms <italic>dark</italic> and <italic>small</italic> were significant in ordinal models but not the linear mixed models; <italic>sad</italic> did not reach significance in either type of model. In Experiment B, only two out of the 10 terms were significant: <italic>light in weight</italic> and <italic>big</italic>. Overall, marginal <italic>R<sup>2</sup></italic> values for the saturation models are much smaller than those describing the lightness models, suggesting that in timbre-color matching, crossmodal associations have greater influence on lightness than on saturation. In Experiment B, fewer terms reached significance, and effect sizes were smaller as compared to Experiment A, consistent with the idea that correspondences with saturation are for the most part dependent on pitch height. Perceptions of timbres as <italic>big</italic> or <italic>light in weight</italic> may relate to saturation of color choice across sound sources independently of pitch register, but the effects are small.</p>
<p>We based H1b, which predicted increased saturation with higher ratings on <italic>low</italic>, <italic>dark</italic>, <italic>heavy</italic>, etc., on previous work that suggested connections between saturation, heaviness, and pitch height (e.g., <xref ref-type="bibr" rid="ref2">Alexander and Shansky, 1976</xref>). In our experiments, we found results to be consistent with our stated hypothesis. As noted in the Introduction, however, other research (<xref ref-type="bibr" rid="ref16">Hamilton-Fletcher et al., 2017</xref>; <xref ref-type="bibr" rid="ref3">Anikin and Johansson, 2019</xref>) would suggest the opposite correlation&#x2014;that higher ratings on <italic>low</italic>, <italic>dark</italic>, <italic>heavy</italic>, etc. would align with decreased saturation. One possible reason for the discrepancies between studies could be that timbre-saturation correspondences interact with hue selection. For example, in <xref ref-type="bibr" rid="ref16">Hamilton-Fletcher et al. (2017)</xref>, higher pitch was associated with yellower hues: bright yellow is necessarily light and highly saturated, so it is possible that in this context, the specificities of this hue correspondence overshadowed a more general tendency to match weight and saturation. Then again, <xref ref-type="bibr" rid="ref3">Anikin and Johansson (2019)</xref> did not observe the association between frequency and yellow, instead capturing a weak correlation between frequency and blue, but they did observe higher pitch and higher spectral centroid to correlate with more saturated colors. Thus, it is also plausible that timbre-saturation correspondences are influenced by the experimental paradigm and/or the most salient features of the available auditory stimuli set.</p>
<p>Experiment B allowed for consideration of whether the correspondence between heaviness and saturation holds when timbral heaviness is considered independently of pitch height. Here, results were mixed, where ratings on the term <italic>heavy</italic> were not significant, but those on <italic>light in weight</italic> were&#x2014;although, with a <italic>p-</italic>value of 0.02, the <italic>light in weight</italic> result may be spurious. It is intriguing that the other term significantly predicting saturation in Experiment B was <italic>big</italic>. Size is often semantically related to weight; yet, <italic>small</italic> was not significant in Experiment A. It may be that stimuli with greater timbral variation could reveal clearer relationships among size, weight, and saturation, but it is also possible that any such relationships are relevant and meaningful only when pitch height is varied, and/or that our mixed results are due to noise in the data. For both lightness and saturation models, the size and weight descriptors yielded mixed results across both studies. To clarify potential nuances, more variance among timbres may be necessary, and a more constrained experimental design may be better suited to disentangle these particular relationships.</p>
</sec>
<sec id="sec46">
<label>5.1.3</label>
<title>Warmth-coolness (H1c)</title>
<p>Ratings on the terms <italic>warm</italic> and <italic>cool</italic> significantly predicted the warmth-coolness of color choices in Experiment B but not in Experiment A. This may be in part because Experiment B offered more types of instruments with a wider range of timbral characteristics than did Experiment A. Effect sizes in Experiment B were small; it is possible that including a wider and more diverse set of instrument timbres would yield a larger effect and could be explored in more depth in future studies. As discussed in more detail below in 5.4, there seems to be a meaningful difference between timbral warmth and warmth of color.</p>
</sec>
</sec>
<sec id="sec47">
<label>5.2</label>
<title>Effect of musical instrument (H2)</title>
<p>H2 proposed a relationship between musical instrument type and color choice, operationalizing color choice, as in H1, in three ways (lightness, saturation, warmth-coolness). Results linking instrument type and lightness were mixed: the relationship was significant in Experiment A but not in Experiment B. Both experiments demonstrated significant differences among instruments with respect to saturation and the warm-cool index.</p>
<p>In Experiment A, pairwise tests carried out via the <italic>emmeans</italic> function in the emmeans package (<xref ref-type="bibr" rid="ref21">Lenth, 2023</xref>) and adjusted using the Tukey method pointed to significant differences between the Flemish harpsichord and piano as well as between the Lautenwerk and piano, but not between the two harpsichords, for all three dependent variables. Among the orchestral instruments in Experiment B, only clarinet/trumpet and clarinet/violin contrasts were significantly different for saturation (see <xref ref-type="fig" rid="fig10">Figure 10</xref> for a visualization of all instruments&#x2019; mean warm-cool indices in Experiment B). With respect to the differences in warm-cool index, the trumpet appears to be solely responsible for the significant ANOVA result (see <xref ref-type="fig" rid="fig9">Figure 9</xref>). Taken together, it seems that there may be certain instruments that are characterized by relatively high or low values for one or more dimensions of color, such as the trumpet, but that these dimensions are not equally relevant for all instruments and/or may only be evident with particular types of timbral contrasts. The extent to which relationships between instrument type and individual dimensions of color are driven by, or interact with, the motivation to match an instrument to a certain hue (such as the trumpet, to red), calls for further clarification.</p>
</sec>
<sec id="sec48">
<label>5.3</label>
<title>Effect of pitch register (H3)</title>
<p>H3, which proposed a relationship between pitch register and color, was tested in Experiment A only. We observed significant differences in lightness, saturation, and the warm-cool index as functions of pitch register (see <xref ref-type="fig" rid="fig4">Figures 4</xref>&#x2013;<xref ref-type="fig" rid="fig6">6</xref>), and no interactions were found with instrument type. Replicating findings established in previous scholarship, increases in pitch register were associated with increasing lightness. Saturation decreased with increasing pitch register, consistent with our hypothesis but inconsistent with some previous findings in the literature (see 5.1.2 for a more in-depth discussion). The piano and Flemish harpsichord both showed a consistent increase of color warmth with pitch register, while the Lautenwerk&#x2019;s matched colors were warmest in the middle register. <italic>Post hoc</italic> Tukey tests of the warm-cool index model show significant increases in warm-cool index between F2 and F3, and F2 and F4, but not between F3 and F4, likely due to the Lautenwerk&#x2019;s divergent profile.</p>
</sec>
<sec id="sec49">
<label>5.4</label>
<title>Warmth-coolness: further discussion</title>
<p>Although our Experiment A findings suggest a positive relationship between pitch register and warmth of color choice, trends in the semantic ratings across pitch register reveal further complexity: ratings of <italic>warm</italic> are highest in the middle register (F3) for each of the keyboard instruments, while ratings of <italic>cool</italic> tend to increase with pitch register! The latter observation is consistent with the results of work by <xref ref-type="bibr" rid="ref64">Wang and Spence (2017)</xref>, where participants matched the experience of drinking cold water with higher pitch, as compared to room temperature and hot water. However, it is notable that as timbre semantic ratings, <italic>warm</italic> and <italic>cool</italic> are not simple opposites when it comes to their relationships with pitch height.</p>
<p>Research on timbre semantics provides some insight into these relationships. Through interviews and an online survey with expert participants, <xref ref-type="bibr" rid="ref47">Rosi et al. (2022)</xref> found that &#x201C;a <italic>bright</italic> sound has most of the spectral energy in the high frequencies. It is often a high-pitched sound, with clarity, definition, and similarities with a metallic sound&#x2026;A <italic>warm</italic> sound encloses substantial spectral energy in the low-mid frequencies. It is a rather low pitch sound&#x2026;A warm sound is pleasant, enveloping, and rich&#x201D; (480). That all three keyboards were rated as timbrally warmest in the middle register (rather than the lowest) suggests a kind of sweet spot for timbral warmth, rather than a simple linear relationship between pitch height and warmth. This middle register peak is not unexpected in light of findings from <xref ref-type="bibr" rid="ref46">Reymore et al. (2023)</xref> and <xref ref-type="bibr" rid="ref44">Reymore (2021)</xref> that other plausibly pleasant dimensions&#x2014;<italic>smooth/singing</italic> and <italic>watery/fluid</italic>&#x2014;show an inverted-U relationship with register. If <italic>warm</italic> timbres are associated with middle and lower registers, but color warmth is associated with higher pitch register, this could explain the null finding for H1c in Experiment A. Perhaps, participants&#x2019; gravitation toward warmer color choices in timbre-color matching is better explained by a combination of increased pitch and increased timbral brightness than by timbral warmth. In Experiment B, we noted that the instruments with higher warm-cool indices (oboe, violin, trumpet) seem to be brighter and more nasal than those with lower values (clarinet, flute, viola). Examination of average ratings on <italic>bright</italic> confirmed that this casual observation is consistent with participant ratings. Taken together, these <italic>post hoc</italic> observations suggest that timbral brightness may be more closely related to color-based warmth than is timbral warmth.</p>
<p>Thus, although we did find a significant relationship between ratings on <italic>warm</italic> and <italic>cool</italic> and warmth-coolness of color choice in Experiment B (but not Experiment A), results should be interpreted with caution. The significant effect may have been driven by other factors or may depend on the sample of instruments tested. Future work should assess judgments across a wider range of instruments to test generalizability and identify specificities that may interact with broader trends.</p>
</sec>
<sec id="sec50">
<label>5.5</label>
<title>Emotion mediation and semantic mediation</title>
<p><xref ref-type="bibr" rid="ref39">Palmer et al. (2013)</xref> posited that if correlations between color and musical excerpts were mediated by common emotional associations, they would find analogous results when asking participants to choose colors most/least consistent with music and with any other set of stimuli strongly associated with the same emotional dimensions. Results from both subsequent experiments designed to test this were consistent with the hypothesis. Previous research shows that participants can make judgments on perceived emotion of timbres based on short tones (e.g., <xref ref-type="bibr" rid="ref34">McAdams et al., 2017</xref>; <xref ref-type="bibr" rid="ref18">Korsmit et al., 2023</xref>), suggesting that it is possible for perceptions of emotion to drive timbre-color matching.</p>
<p>Should shared emotion, mood, or affect provide the best account for timbre-color matching, as posited in the emotion-mediation account, we could anticipate in our experiments that the terms <italic>happy</italic> and <italic>sad</italic> would provide equivalent or better explanatory power than other descriptors. In our experiments, ratings on the terms <italic>happy</italic> and <italic>sad</italic> were significantly related to lightness in both experiments, but <italic>sad</italic> was not significant for saturation in either experiment. For lightness, in Experiment A, <italic>happy</italic> and <italic>sad</italic> resulted in slightly lower marginal <italic>R<sup>2</sup></italic> values among the significant terms; in Experiment B, <italic>R<sup>2</sup></italic> values were equivalent to those for other terms. That is, our emotion terms seem to be less related to color choice than other terms in our set.</p>
<p>On one hand, it may be that <italic>happy</italic> and <italic>sad</italic> were not apt emotional descriptors for the available instrumental timbres, but that other emotional terms would have provided a better fit. However, this seems unlikely for the given musical context; <xref ref-type="bibr" rid="ref18">Korsmit et al. (2023)</xref> found that two dimensions are sufficient for capturing variance in emotional assessment of single-note and chromatic scale stimuli. On the other hand, it seems likely that emotional qualities of isolated timbres are not always the most relevant consideration for a given task. Other, lower-level perceptual features might be more immediately relevant for timbre-color mapping. Indeed, <xref ref-type="bibr" rid="ref53">Spence, (2020a)</xref> observes that emotion mediation tends to account for more variance for complex stimuli as compared to simpler and less emotionally valent stimuli. Similarly, from results of an auditory-conceptual association study, <xref ref-type="bibr" rid="ref9">Di Stefano et al. (2024)</xref> note that their findings support the view that complex stimuli are more likely to generate emotional meaning than are simpler stimuli (e.g., isolated sounds). Future research could test the emotion mediation hypothesis using <xref ref-type="bibr" rid="ref39">Palmer et al.&#x2019;s (2013)</xref> method with isolated timbres to directly address the claim that this explanation is equally applicable to complex and simple stimuli.</p>
<p><xref ref-type="bibr" rid="ref52">Spence (2011)</xref> proposed three categories motivating crossmodal matching, including semantic, physiological, and statistical, where semantic mediation relates to linguistic or lexical correspondence (e.g., we use <italic>high</italic> to describe both elevation and pitch). <xref ref-type="bibr" rid="ref37">Motoki et al. (2023)</xref> added the category of &#x201C;affective&#x201D; to this model to account for emotion mediation (see also <xref ref-type="bibr" rid="ref53">Spence, 2020a</xref>). These categories are not mutually exclusive, and a given correspondence may be motivated by multiple categories. As we chose our terms for this study based on crossmodal terms observed in previous timbre semantics research, it seems reasonable that semantic mediation provides an overall better explanation for our results, at least with respect to lightness and saturation. Notably, the terms <italic>happy</italic> and <italic>sad</italic> are used less often to describe timbre as compared to the other terms in our set&#x2014;from a semantic mediation perspective, this may explain why these terms explained relatively less variance.</p>
</sec>
<sec id="sec51">
<label>5.6</label>
<title>Hue</title>
<p>Participant color choices for each instrument were diverse. Some instruments demonstrated no specific trends in hue but did show heavier concentrations of responses at particular lightness levels, such as the flute. Some instruments demonstrated trends in hue, such as the trumpet, violin, and clarinet. While specific instrument-hue associations from our participants only seem apparent for a few instruments, the origins of these associations provide an intriguing subject for further research, particularly in the case of the trumpet, which has shown robust hue correlations in both historical and empirical studies (<xref ref-type="bibr" rid="ref43">Reuter et al., 2018</xref>).</p>
<p>In the case of the keyboard instruments, pitch register appears to have played an important role: the lowest octave showed a concentration on saturated blues and purples for all three instruments. The piano was associated with a dispersed range of hues in the top two octaves, but both harpsichords showed an increasing tendency toward yellows as pitch increased. For the harpsichords at F2 in Experiment A, we speculate that participants may have been responding to two different aspects of the sound&#x2014;some may have chosen colors primarily on the basis of pitch register (blues and purples), whereas another group may have been influenced by the brighter timbres of the harpsichord to select oranges and reds, explaining the bimodal concentration of responses in <xref ref-type="fig" rid="fig2">Figure 2</xref>. In general, color choices for the two types of harpsichords are far more similar to each other in each octave than they are to the piano and are generally associated with warmer colors as compared to the piano. Similarities in color choices between the two types of harpsichords may reflect their relative perceived similarity in timbre as compared to the piano.</p>
</sec>
<sec id="sec52">
<label>5.7</label>
<title>Limitations and considerations for future studies</title>
<p>One limitation of our approach to analysis was our use of the mean in quantifying the saturation and warmth-coolness of participant color choices. Participants often selected multiple colors for a stimulus; for such observations, the saturation and warm-cool values used in modeling were approximated using the mean. This becomes potentially problematic when participants selected colors in different areas of the arrays. For example, if a participant selected a low saturation color and a high saturation color, the average saturation of their response is in the middle, which may not be representative of their response tactic. However, before settling on the use of the mean for analysis, we manually reviewed the data and determined that color selections were usually made in single clusters of contiguous color samples, for which the mean was a reasonable representation.</p>
<p>As previously mentioned, participants were not required to use the slider during their selections, and the data reveal that a subset of participants did not use the slider when choosing colors, defaulting to the lightest setting. It is unclear whether participants were intentional about this and felt that the lightest palette best exemplified the colors they wanted, misunderstood the directions, or opted not to use the sliders in order to get through the experiment more quickly. In experiments using a similar interface, we recommend enforcing slider use by requiring participants to acknowledge the slider via touch, even when they prefer to leave it in its initial location. We also acknowledge the possibility that for some subjects and conditions, even our most luminous palettes might not have been sufficiently luminous to match the subject&#x2019;s subjective appraisal of the musical recording.</p>
<p>It should be noted while the use of the major scale for stimuli facilitated comparisons among instruments and allowed us to directly compare semantic ratings and color choices, using stimuli with different musical parameters (mode, articulation, rhythm, etc.) would likely result in somewhat different choices in both semantic descriptions and colors. While we hypothesize that the relationship between semantic descriptions and colors would hold in other simple musical contexts, future research might vary musical parameters to test the generalizability of our results.</p>
<p>Finally, though some portions of the variance in lightness and saturation were explained from the semantic ratings, the majority of variance in each case remains unexplained; a number of other factors likely influence timbre/color matching behavior. It should be noted that the purpose of the experiments reported here was not to thoroughly model the timbre-color matching process, but rather to test whether crossmodal language is related to timbre-color associations. Our results provide converging evidence in support of this theory but also suggest that the issue is multi-layered and complex, and that there are likely multiple principles guiding timbre-color matching behavior.</p>
<p>For example, interactions among the three color dimensions, particularly between hue and each of the other dimensions, may account for some of the variance in color choice. The trumpet provides a likely example of this. Color selections for the trumpet appear to be hue-focused: selections near red, orange, and yellow were most often selected. If a participant hears the trumpet and immediately thinks of a basic color category, such as that exemplified by a typical fire-engine red, the imagined hue necessitates a particular level of lightness&#x2014;the lightness of the choice is in some way an artifact of the hue. Similarly, if a participant associates a sound with a typical yellow, the selection will be necessarily on the higher end of the lightness scale because as yellow darkens, it becomes brown. However, <xref ref-type="bibr" rid="ref16">Hamilton-Fletcher et al. (2017)</xref> previously found crossmodal associations between non-musical sounds and color, including a relationship between yellow and high frequencies, even when controlling for the influence of lightness. Such observations raise the possibility that similar types of relationships may have been at play in the current study concerning instrumental timbre (note specifically the hue results among the keyboard instruments), though further research is needed to determine whether that is the case for the types of stimuli studied here and how such relationships may be connected to participants&#x2019; use of crossmodal language.</p>
<p>Although we did not observe that ratings on <italic>happy</italic> or <italic>sad</italic> were privileged among semantic terms in predicting color, timbre-color matching might be more generally a product of valence transfer (e.g., <xref ref-type="bibr" rid="ref67">Weinreich and Gollwitzer, 2016</xref>). Emotion might also influence color choice via participant mood. For example, perhaps participants in a happier mood might be generally more prone to choosing lighter or yellower colors. Another potential mediating variable could be personal preference. Participants may be more likely to choose colors that they like, but preference might also play a more complex role: for example, participants might match colors they like to timbres they like and colors they dislike to timbres they dislike, where preference interacts with valence or qualia transfer.</p>
<p>Additionally, some color matching choices for recognizable instruments may be primarily driven by semantic rather than perceptual associations derived from cultural tropes. This might be the case for the trumpet&#x2019;s association with red or the clarinet&#x2019;s association with blue. Even if this is sometimes the case, however, it may still be possible that some of these cultural associations have perceptual origins, and so the two cannot be completely separated without further investigation. Future work could assess participants&#x2019; familiarity with instruments, which may help identify such connections.</p>
<p>We did not observe any particular patterns or differences among participants with sound-related synesthesia, though there were only a few such participants. The question of how synesthesia might interact with timbre-color associations could be addressed in a subsequent study with a larger population of synesthetes. In general, further systematic investigation of effects of other individual or cultural differences could also help contextualize variability of responses.</p>
<p>Finally, although the number of stimuli in the current experiment limited the generalizability of audio feature analysis, future research might include a larger and more varied stimulus set to assess the impact of particular features, such as spectral centroid. A stimulus set with increased timbral diversity would help clarify such potential relationships and might reveal more pronounced differences in semantic ratings and/or color choices. Future work might expand beyond instruments in the Western orchestral tradition to include instruments from other musical traditions in order to widen the array of available timbres.</p>
</sec>
</sec>
<sec sec-type="conclusions" id="sec53">
<label>6</label>
<title>Conclusion</title>
<p>Implicit in our experimental design and those guiding other studies is the assumption that the three putative perceptual dimensions of color&#x2014;hue, saturation and lightness (brightness)&#x2014;are separable (see <xref ref-type="bibr" rid="ref14">Garner, 1974</xref>) and that we can probe for timbre/color correspondences related to each dimension independent of the other two. However, our results suggest that these correspondences may arise from more holistic mental representations of color and timbre. Though we find lightness is most closely associated with variations in timbre, there are clear interactions among the different dimensions of color that are evident in the data and complicate this picture. Moreover, for some experimental conditions, participants&#x2019; color selections tend to cluster in a few regions of our color palette. This suggests that for these conditions at least, individual differences in timbre-color correspondences are not entirely random but are guided by a few distinct mental representations of timbre-color correspondence.</p>
<p>Taken together, our findings demonstrate that crossmodal timbre semantic terms bear relation to timbre-color matching behavior, particularly in relation to the lightness of selected color samples. These results are consistent with semantic mediation (<xref ref-type="bibr" rid="ref52">Spence, 2011</xref>) for color-timbre correspondences. Effects are larger when both pitch register and sound source are varied, but they are observable even when pitch register is held constant. The 10 terms proposed in H1a and H1b are more robustly related to the lightness of selected colors than to saturation. With respect to saturation, only two terms reached significance when pitch height was controlled. We found a weak relationship between ratings on <italic>warm</italic> and <italic>cool</italic> with the warmth-coolness of matched colors among orchestral instruments in Experiment B, but not among the keyboard instruments in Experiment A. Overall, evidence is consistent with H2, which posited that color choice varies among musical instrument types, but that these relationships are complex and involve specificities. Lightness differences were found in Experiment A but not B, while saturation and warm-cool differences were found in both experiments. However, <italic>post hoc</italic> comparisons suggest that these significant results were typically motivated by contrasts between particular instruments. Finally, results supported H3, which predicted a relationship between pitch height and color. Significant differences across pitch registers were found with respect to lightness, saturation, and the warm-cool index.</p>
<p>Diverse music have diverse musical goals and consequently prioritize different musical parameters. The relative salience of various parameters is important in determining which aspects of the music provide listeners the strongest cues related to crossmodal associations. There remains work to be done on understanding the relationship between crossmodal associations with basic sensory features and the emergent emotional interpretation that seems to play such an important role in crossmodal associations with composed music. The major scales used in these experiments represent a middleground: they are more complex than single notes but less complex than composed music, introducing a dimension of musicality while maintaining control of musical content. Previous research on crossmodal associations with basic sensory features has focused heavily on pitch; in order to relate crossmodal associations with simple stimuli to associations with complex stimuli, further research is called for on parameters other than pitch, such as timbre.</p>
<p>An understanding of common trends in timbre-color associations is especially relevant for the field of music visualization, in which visual images, colors, and shapes are set or co-created with music, with possibilities for artists working in musical multimedia continuing to grow as technology advances. The current consideration of instrumental timbre as simpler stimuli, outside of the context of more complex musical stimuli, may be especially relevant for those contemporary composers who approach timbre as critical for or central to their compositional process. In general, composers often seek to affect their audiences through understanding preferences and manipulating expectations; thus, a theory of audiovisual art necessitates a thorough exploration and understanding of crossmodal preferences and expectations, which likely relate to both the experience and aesthetic appraisal of multimedia art.</p>
<p>This is the first study to establish the conceptual relationships between the crossmodal linguistics of timbre and crossmodal correspondences with all three dimensions of color, and it is the first to address the interaction and relative contributions of timbre and pitch register to color-matching behavior with a range of musical instrument timbres. Our use of multi-note stimuli controlled for musical content adds a dimension of ecological validity to our results as they relate to topics such as the analysis of multimedia art, and the color palette offers colorimetrically precise three-dimensional variation in color and an expansive array of color choices to participants that far exceeds choices presented to participants in many music-color studies, while allowing for multiple color selection.</p>
</sec>
</body>
<back>
<sec sec-type="data-availability" id="sec54">
<title>Data availability statement</title>
<p>The raw data supporting the conclusions of this article will be made available by the authors, without undue reservation.</p>
</sec>
<sec sec-type="ethics-statement" id="sec55">
<title>Ethics statement</title>
<p>The studies involving humans were approved by The Ohio State University Institutional Review Board. The studies were conducted in accordance with the local legislation and institutional requirements. The participants provided their written informed consent to participate in this study.</p>
</sec>
<sec sec-type="author-contributions" id="sec56">
<title>Author contributions</title>
<p>LR: Conceptualization, Data curation, Formal analysis, Funding acquisition, Investigation, Methodology, Visualization, Writing &#x2013; original draft, Writing &#x2013; review &#x0026; editing. DL: Conceptualization, Data curation, Formal analysis, Methodology, Software, Visualization, Writing &#x2013; original draft, Writing &#x2013; review &#x0026; editing.</p>
</sec>
<sec sec-type="funding-information" id="sec57">
<title>Funding</title>
<p>The author(s) declare that financial support was received for the research, authorship, and/or publication of this article. LR was supported by a graduate Summer Research Award from The Ohio State Center for Cognitive and Brain Sciences.</p>
</sec>
<ack>
<p>The authors would like to thank Mark Rubenstein, the audio engineer for the initial experiment and Experiment B, as well as each of the professional musicians who recorded the stimuli for all experiments. Special thanks to Keith Hill, who provided and tuned the instruments used in the recording for Experiment A. Thank you to Hannah Moore for assisting with data collection. We would also like to acknowledge the Center for Science and Industry (COSI) and the Buckeye Language Network (BLN) for providing the opportunity for data collection. Special thanks to Laura Wagner, whose collegiality and leadership greatly facilitated this opportunity. We extend our appreciation for the Summer Research Award from the Ohio State Center for Cognitive and Brain Sciences, which funded this work. Finally, we would like to thank Marcel Montrey for his consultation on the statistical analysis for this paper.</p>
</ack>
<sec sec-type="COI-statement" id="sec58">
<title>Conflict of interest</title>
<p>The authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec sec-type="ai-statement" id="sec59">
<title>Generative AI statement</title>
<p>The authors declare that no Generative AI was used in the creation of this manuscript.</p>
</sec>
<sec sec-type="disclaimer" id="sec60">
<title>Publisher&#x2019;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<sec sec-type="supplementary-material" id="sec61">
<title>Supplementary material</title>
<p>The Supplementary material for this article can be found online at: <ext-link xlink:href="https://www.frontiersin.org/articles/10.3389/fpsyg.2024.1520131/full#supplementary-material" ext-link-type="uri">https://www.frontiersin.org/articles/10.3389/fpsyg.2024.1520131/full#supplementary-material</ext-link></p>
<supplementary-material xlink:href="Data_Sheet_1.PDF" id="SM1" mimetype="application/pdf" xmlns:xlink="http://www.w3.org/1999/xlink"/>
</sec>
<ref-list>
<title>References</title>
<ref id="ref1"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Adeli</surname> <given-names>M.</given-names></name> <name><surname>Rouat</surname> <given-names>J.</given-names></name> <name><surname>Molotchnikoff</surname> <given-names>S.</given-names></name></person-group> (<year>2014</year>). <article-title>Audiovisual correspondence between musical timbre and visual shapes</article-title>. <source>Front. Hum. Neurosci.</source> <volume>8</volume>:<fpage>352</fpage>. doi: <pub-id pub-id-type="doi">10.3389/fnhum.2014.00352</pub-id>, PMID: <pub-id pub-id-type="pmid">24910604</pub-id></citation></ref>
<ref id="ref2"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Alexander</surname> <given-names>K. R.</given-names></name> <name><surname>Shansky</surname> <given-names>M. S.</given-names></name></person-group> (<year>1976</year>). <article-title>Influence of hue, value, and chroma on the perceived heaviness of colors</article-title>. <source>Percept. Psychophys.</source> <volume>19</volume>, <fpage>72</fpage>&#x2013;<lpage>74</lpage>. doi: <pub-id pub-id-type="doi">10.3758/BF03199388</pub-id></citation></ref>
<ref id="ref3"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Anikin</surname> <given-names>A.</given-names></name> <name><surname>Johansson</surname> <given-names>N.</given-names></name></person-group> (<year>2019</year>). <article-title>Implicit associations between individual properties of color and sound</article-title>. <source>Atten. Percept. Psychophys.</source> <volume>81</volume>, <fpage>764</fpage>&#x2013;<lpage>777</lpage>. doi: <pub-id pub-id-type="doi">10.3758/s13414-018-01639-7</pub-id>, PMID: <pub-id pub-id-type="pmid">30547381</pub-id></citation></ref>
<ref id="ref4"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bates</surname> <given-names>D.</given-names></name> <name><surname>Maechler</surname> <given-names>M.</given-names></name> <name><surname>Bolker</surname> <given-names>B.</given-names></name> <name><surname>Walker</surname> <given-names>S.</given-names></name></person-group> (<year>2015</year>). <article-title>Fitting linear mixed-effects models using lme4</article-title>. <source>J. Stat. Softw.</source> <volume>67</volume>, <fpage>1</fpage>&#x2013;<lpage>48</lpage>. doi: <pub-id pub-id-type="doi">10.18637/jss.v067.i01</pub-id></citation></ref>
<ref id="ref5"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bien</surname> <given-names>N.</given-names></name> <name><surname>ten Oever</surname> <given-names>S.</given-names></name> <name><surname>Goebel</surname> <given-names>R.</given-names></name> <name><surname>Sack</surname> <given-names>A. T.</given-names></name></person-group> (<year>2012</year>). <article-title>The sound of size</article-title>. <source>NeuroImage</source> <volume>59</volume>, <fpage>663</fpage>&#x2013;<lpage>672</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.neuroimage.2011.06.095</pub-id></citation></ref>
<ref id="ref6"><citation citation-type="other"><person-group person-group-type="author"><name><surname>Christensen</surname> <given-names>R.</given-names></name></person-group> (<year>2023</year>). Ordinal-regression models for ordinal data. R package version 2023.12&#x2013;4.1. Available at: <ext-link xlink:href="https://CRAN.R-project.org/package=ordinal" ext-link-type="uri">https://CRAN.R-project.org/package=ordinal</ext-link> (Accessed August 15, 2024).</citation></ref>
<ref id="ref7"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Deroy</surname> <given-names>O.</given-names></name> <name><surname>Spence</surname> <given-names>C.</given-names></name></person-group> (<year>2013</year>). <article-title>Why we are not all synesthetes (not even weakly so)</article-title>. <source>Psychon. Bull. Rev.</source> <volume>20</volume>, <fpage>643</fpage>&#x2013;<lpage>664</lpage>. doi: <pub-id pub-id-type="doi">10.3758/s13423-013-0387-2</pub-id>, PMID: <pub-id pub-id-type="pmid">23413012</pub-id></citation></ref>
<ref id="ref8"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Di Stefano</surname> <given-names>N.</given-names></name></person-group> (<year>2023</year>). <article-title>Musical emotions and timbre: from expressiveness to atmospheres</article-title>. <source>Philosophia</source> <volume>51</volume>, <fpage>2625</fpage>&#x2013;<lpage>2637</lpage>. doi: <pub-id pub-id-type="doi">10.1007/s11406-023-00700-6</pub-id></citation></ref>
<ref id="ref9"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Di Stefano</surname> <given-names>N.</given-names></name> <name><surname>Ansani</surname> <given-names>A.</given-names></name> <name><surname>Schiavio</surname> <given-names>A.</given-names></name> <name><surname>Spence</surname> <given-names>C.</given-names></name></person-group> (<year>2024</year>). <article-title>Prokofiev was (almost) right: a cross-cultural investigation of auditory-conceptual associations in Peter and the wolf</article-title>. <source>Psychon. Bull. Rev.</source> <volume>31</volume>, <fpage>1735</fpage>&#x2013;<lpage>1744</lpage>. doi: <pub-id pub-id-type="doi">10.3758/s13423-023-02435-7</pub-id>, PMID: <pub-id pub-id-type="pmid">38267741</pub-id></citation></ref>
<ref id="ref10"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Eerola</surname> <given-names>T.</given-names></name> <name><surname>Ferrer</surname> <given-names>R.</given-names></name> <name><surname>Alluri</surname> <given-names>V.</given-names></name></person-group> (<year>2012</year>). <article-title>Timbre and affect dimensions: evidence from affect and similarity ratings and acoustic correlates of isolated instrument sounds</article-title>. <source>Music Percept.</source> <volume>30</volume>, <fpage>49</fpage>&#x2013;<lpage>70</lpage>. doi: <pub-id pub-id-type="doi">10.1525/mp.2012.30.1.49</pub-id></citation></ref>
<ref id="ref11"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Eitan</surname> <given-names>Z.</given-names></name> <name><surname>Timmers</surname> <given-names>R.</given-names></name></person-group> (<year>2010</year>). <article-title>Beethoven&#x2019;s last piano sonata and those who follow crocodiles: cross-domain mappings of auditory pitch in a musical context</article-title>. <source>Cognition</source> <volume>114</volume>, <fpage>405</fpage>&#x2013;<lpage>422</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.cognition.2009.10.013</pub-id>, PMID: <pub-id pub-id-type="pmid">20036356</pub-id></citation></ref>
<ref id="ref12"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Evans</surname> <given-names>K. K.</given-names></name> <name><surname>Treisman</surname> <given-names>A.</given-names></name></person-group> (<year>2010</year>). <article-title>Natural cross-modal mappings between visual and auditory features</article-title>. <source>J. Vis.</source> <volume>10</volume>, <fpage>6.1</fpage>&#x2013;<lpage>6.12</lpage>. doi: <pub-id pub-id-type="doi">10.1167/10.1.6</pub-id>, PMID: <pub-id pub-id-type="pmid">20143899</pub-id></citation></ref>
<ref id="ref13"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Gallace</surname> <given-names>A.</given-names></name> <name><surname>Spence</surname> <given-names>C.</given-names></name></person-group> (<year>2006</year>). <article-title>Multisensory synesthetic interactions in the speeded classification of visual size</article-title>. <source>Percept. Psychophys.</source> <volume>68</volume>, <fpage>1191</fpage>&#x2013;<lpage>1203</lpage>. doi: <pub-id pub-id-type="doi">10.3758/BF03193720</pub-id>, PMID: <pub-id pub-id-type="pmid">17355042</pub-id></citation></ref>
<ref id="ref14"><citation citation-type="book"><person-group person-group-type="author"><name><surname>Garner</surname> <given-names>W. R.</given-names></name></person-group> (<year>1974</year>). <source>The processing of information and structure</source>. <publisher-loc>Potomac, MD</publisher-loc>: <publisher-name>Erlbaum</publisher-name>.</citation></ref>
<ref id="ref15"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Gurman</surname> <given-names>D.</given-names></name> <name><surname>McCormick</surname> <given-names>C. R.</given-names></name> <name><surname>Klein</surname> <given-names>R. M.</given-names></name></person-group> (<year>2021</year>). <article-title>Crossmodal correspondence between auditory timbre and visual shape</article-title>. <source>Multisens. Res.</source> <volume>35</volume>, <fpage>221</fpage>&#x2013;<lpage>241</lpage>. doi: <pub-id pub-id-type="doi">10.1163/22134808-bja10067</pub-id>, PMID: <pub-id pub-id-type="pmid">35065536</pub-id></citation></ref>
<ref id="ref16"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Hamilton-Fletcher</surname> <given-names>G.</given-names></name> <name><surname>Witzel</surname> <given-names>C.</given-names></name> <name><surname>Reby</surname> <given-names>D.</given-names></name> <name><surname>Ward</surname> <given-names>J.</given-names></name></person-group> (<year>2017</year>). <article-title>Sound properties associated with equiluminant colours</article-title>. <source>Multisens. Res.</source> <volume>30</volume>, <fpage>337</fpage>&#x2013;<lpage>362</lpage>. doi: <pub-id pub-id-type="doi">10.1163/22134808-00002567</pub-id>, PMID: <pub-id pub-id-type="pmid">31287083</pub-id></citation></ref>
<ref id="ref17"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Isbilen</surname> <given-names>E. S.</given-names></name> <name><surname>Krumhansl</surname> <given-names>C. L.</given-names></name></person-group> (<year>2016</year>). <article-title>The color of music: emotion-mediated associations to Bach&#x2019;s well-tempered clavier</article-title>. <source>Psychomusicology</source> <volume>26</volume>, <fpage>149</fpage>&#x2013;<lpage>161</lpage>. doi: <pub-id pub-id-type="doi">10.1037/pmu0000147</pub-id></citation></ref>
<ref id="ref18"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Korsmit</surname> <given-names>I. R.</given-names></name> <name><surname>Montrey</surname> <given-names>M.</given-names></name> <name><surname>Wong-Min</surname> <given-names>A. Y. T.</given-names></name> <name><surname>McAdams</surname> <given-names>S.</given-names></name></person-group> (<year>2023</year>). <article-title>A comparison of dimensional and discrete models for the representation of perceived and induced affect in response to short musical sounds</article-title>. <source>Front. Psychol.</source> <volume>14</volume>:<fpage>1287334</fpage>. doi: <pub-id pub-id-type="doi">10.3389/fpsyg.2023.1287334</pub-id>, PMID: <pub-id pub-id-type="pmid">38023037</pub-id></citation></ref>
<ref id="ref19"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Korsmit</surname> <given-names>I. R.</given-names></name> <name><surname>Montrey</surname> <given-names>M.</given-names></name> <name><surname>Wong-Min</surname> <given-names>A. Y. T.</given-names></name> <name><surname>McAdams</surname> <given-names>S.</given-names></name></person-group> (<year>2024</year>). <article-title>The acoustic properties of affective timbres: consistencies and discrepancies in a synthesis of multiple datasets</article-title>. <source>Music Sci.</source> <volume>7</volume>. doi: <pub-id pub-id-type="doi">10.1177/20592043241256012</pub-id></citation></ref>
<ref id="ref20"><citation citation-type="confproc"><person-group person-group-type="author"><name><surname>Kuang</surname> <given-names>J.</given-names></name> <name><surname>Guo</surname> <given-names>Y.</given-names></name> <name><surname>Liberman</surname> <given-names>M.</given-names></name></person-group> (<year>2016</year>). <article-title>Voice quality as a pitch-range indicator</article-title>. In <conf-name>Proceedings of Speech Prosody</conf-name>. <volume>8</volume>, <fpage>1061</fpage>&#x2013;<lpage>1065</lpage>. Available at: <ext-link xlink:href="https://www.isca-archive.org/speechprosody_2016/kuang16_speechprosody.pdf" ext-link-type="uri">https://www.isca-archive.org/speechprosody_2016/kuang16_speechprosody.pdf</ext-link></citation></ref>
<ref id="ref21"><citation citation-type="other"><person-group person-group-type="author"><name><surname>Lenth</surname> <given-names>R.</given-names></name></person-group> (<year>2023</year>). Emmeans: estimated marginal means, aka least-squares means. R package version 1.8.6. Available at: <ext-link xlink:href="https://CRAN.R-project.org/package=emmeans" ext-link-type="uri">https://CRAN.R-project.org/package=emmeans</ext-link> (Accessed August 15, 2024).</citation></ref>
<ref id="ref22"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Lindborg</surname> <given-names>P. M.</given-names></name> <name><surname>Friberg</surname> <given-names>A. K.</given-names></name></person-group> (<year>2015</year>). <article-title>Colour association with music is mediated by emotion: evidence from an experiment using a CIE lab interface and interviews</article-title>. <source>PLoS One</source> <volume>10</volume>:<fpage>e0144013</fpage>. doi: <pub-id pub-id-type="doi">10.1371/journal.pone.0144013</pub-id>, PMID: <pub-id pub-id-type="pmid">26642050</pub-id></citation></ref>
<ref id="ref23"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>L&#x00FC;decke</surname> <given-names>D.</given-names></name> <name><surname>Ben-Shachar</surname> <given-names>M. S.</given-names></name> <name><surname>Patil</surname> <given-names>I.</given-names></name> <name><surname>Waggoner</surname> <given-names>P.</given-names></name> <name><surname>Makowski</surname> <given-names>D.</given-names></name></person-group> (<year>2021</year>). <article-title>Performance: an R package for assessment, comparison and testing of statistical models</article-title>. <source>J. Open Source Soft.</source> <volume>6</volume>:<fpage>3139</fpage>. doi: <pub-id pub-id-type="doi">10.21105/joss.03139</pub-id></citation></ref>
<ref id="ref24"><citation citation-type="book"><person-group person-group-type="author"><name><surname>Mardia</surname> <given-names>K. V.</given-names></name> <name><surname>Jupp</surname> <given-names>P. E.</given-names></name></person-group> (<year>2009</year>). <source>Directional statistics</source>, vol. <volume>494</volume>. <publisher-loc>Chichester, England</publisher-loc>: <publisher-name>John Wiley &#x0026; Sons</publisher-name>.</citation></ref>
<ref id="ref25"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Marks</surname> <given-names>L. E.</given-names></name></person-group> (<year>1974</year>). <article-title>On associations of light and sound: the mediation of brightness, pitch, and loudness</article-title>. <source>Am. J. Psychol.</source> <volume>87</volume>, <fpage>173</fpage>&#x2013;<lpage>188</lpage>. doi: <pub-id pub-id-type="doi">10.2307/1422011</pub-id>, PMID: <pub-id pub-id-type="pmid">4451203</pub-id></citation></ref>
<ref id="ref26"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Marks</surname> <given-names>L. E.</given-names></name></person-group> (<year>1982</year>). <article-title>Bright sneezes and dark coughs, loud sunlight and soft moonlight</article-title>. <source>J. Exp. Psychol. Hum. Percept. Perform.</source> <volume>8</volume>, <fpage>177</fpage>&#x2013;<lpage>193</lpage>. doi: <pub-id pub-id-type="doi">10.1037//0096-1523.8.2.177</pub-id>, PMID: <pub-id pub-id-type="pmid">6461716</pub-id></citation></ref>
<ref id="ref27"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Marks</surname> <given-names>L. E.</given-names></name></person-group> (<year>1987</year>). <article-title>On cross-modal similarity: auditory&#x2013;visual interactions in speeded discrimination</article-title>. <source>J. Exp. Psychol. Hum. Percept. Perform.</source> <volume>13</volume>, <fpage>384</fpage>&#x2013;<lpage>394</lpage>. doi: <pub-id pub-id-type="doi">10.1037/0096-1523.13.3.384</pub-id>, PMID: <pub-id pub-id-type="pmid">2958587</pub-id></citation></ref>
<ref id="ref28"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Marks</surname> <given-names>L. E.</given-names></name></person-group> (<year>1989</year>). <article-title>On cross-modal similarity: the perceptual structure of pitch, loudness, and brightness</article-title>. <source>J. Exp. Psychol. Hum. Percept. Perform.</source> <volume>15</volume>, <fpage>586</fpage>&#x2013;<lpage>602</lpage>. doi: <pub-id pub-id-type="doi">10.1037/0096-1523.15.3.586</pub-id>, PMID: <pub-id pub-id-type="pmid">2527964</pub-id></citation></ref>
<ref id="ref29"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Marks</surname> <given-names>L. E.</given-names></name></person-group> (<year>1996</year>). <article-title>On perceptual metaphors</article-title>. <source>Metaphor Symbolic Activity</source> <volume>11</volume>, <fpage>39</fpage>&#x2013;<lpage>66</lpage>. doi: <pub-id pub-id-type="doi">10.1207/s15327868ms1101_3</pub-id></citation></ref>
<ref id="ref30"><citation citation-type="book"><person-group person-group-type="author"><name><surname>Marks</surname> <given-names>L. E.</given-names></name></person-group> (<year>2013</year>). &#x201C;<article-title>Audiovisual cross-modal correspondences in the general population</article-title>&#x201D; in <source>Oxford handbook of synesthesia</source>. eds. <person-group person-group-type="editor"><name><surname>Simner</surname> <given-names>J.</given-names></name> <name><surname>Hubbard</surname> <given-names>E.</given-names></name></person-group> (<publisher-loc>Oxford, United Kingdom</publisher-loc>: <publisher-name>Oxford University Press</publisher-name>), <fpage>761</fpage>&#x2013;<lpage>789</lpage>.</citation></ref>
<ref id="ref31"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Marks</surname> <given-names>L. E.</given-names></name> <name><surname>Hammeal</surname> <given-names>R. J.</given-names></name> <name><surname>Bornstein</surname> <given-names>M. H.</given-names></name> <name><surname>Smith</surname> <given-names>L. B.</given-names></name></person-group> (<year>1987</year>). <article-title>Perceiving similarity and comprehending metaphor</article-title>. <source>Monogr. Soc. Res. Child Dev.</source> <volume>52</volume>, <fpage>1</fpage>&#x2013;<lpage>102</lpage>. doi: <pub-id pub-id-type="doi">10.2307/1166084</pub-id>, PMID: <pub-id pub-id-type="pmid">3431563</pub-id></citation></ref>
<ref id="ref32"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Marozeau</surname> <given-names>J.</given-names></name> <name><surname>de Cheveign&#x00E9;</surname> <given-names>A.</given-names></name></person-group> (<year>2007</year>). <article-title>The effect of fundamental frequency on the brightness dimension of timbre</article-title>. <source>J. Acoust. Soc. Am.</source> <volume>121</volume>, <fpage>383</fpage>&#x2013;<lpage>387</lpage>. doi: <pub-id pub-id-type="doi">10.1121/1.2384910</pub-id>, PMID: <pub-id pub-id-type="pmid">17297793</pub-id></citation></ref>
<ref id="ref33"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Martino</surname> <given-names>G.</given-names></name> <name><surname>Marks</surname> <given-names>L. E.</given-names></name></person-group> (<year>1999</year>). <article-title>Perceptual and linguistic interactions in speeded classification: tests of the semantic coding hypothesis</article-title>. <source>Perception</source> <volume>28</volume>, <fpage>903</fpage>&#x2013;<lpage>923</lpage>. doi: <pub-id pub-id-type="doi">10.1068/p2866</pub-id>, PMID: <pub-id pub-id-type="pmid">10664781</pub-id></citation></ref>
<ref id="ref34"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>McAdams</surname> <given-names>S.</given-names></name> <name><surname>Douglas</surname> <given-names>C.</given-names></name> <name><surname>Vempala</surname> <given-names>N. N.</given-names></name></person-group> (<year>2017</year>). <article-title>Perception and modeling of affective qualities of musical instrument sounds across pitch registers</article-title>. <source>Front. Psychol.</source> <volume>8</volume>:<fpage>153</fpage>. doi: <pub-id pub-id-type="doi">10.3389/fpsyg.2017.00153</pub-id>, PMID: <pub-id pub-id-type="pmid">28228741</pub-id></citation></ref>
<ref id="ref35"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Melara</surname> <given-names>R. D.</given-names></name> <name><surname>Marks</surname> <given-names>L. E.</given-names></name></person-group> (<year>1990</year>). <article-title>Interaction among auditory dimensions: timbre, pitch, and loudness</article-title>. <source>Percept. Psychophys.</source> <volume>48</volume>, <fpage>169</fpage>&#x2013;<lpage>178</lpage>. doi: <pub-id pub-id-type="doi">10.3758/BF03207084</pub-id>, PMID: <pub-id pub-id-type="pmid">2385491</pub-id></citation></ref>
<ref id="ref36"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Mondloch</surname> <given-names>C. J.</given-names></name> <name><surname>Maurer</surname> <given-names>D.</given-names></name></person-group> (<year>2004</year>). <article-title>Do small white balls squeak? Pitch-object correspondences in young children</article-title>. <source>Cogn. Affect. Behav. Neurosci.</source> <volume>4</volume>, <fpage>133</fpage>&#x2013;<lpage>136</lpage>. doi: <pub-id pub-id-type="doi">10.3758/CABN.4.2.133</pub-id>, PMID: <pub-id pub-id-type="pmid">15460920</pub-id></citation></ref>
<ref id="ref37"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Motoki</surname> <given-names>K.</given-names></name> <name><surname>Marks</surname> <given-names>L. E.</given-names></name> <name><surname>Velasco</surname> <given-names>C.</given-names></name></person-group> (<year>2023</year>). <article-title>Reflections on cross-modal correspondences: current understanding and issues for future research</article-title>. <source>Multisens. Res.</source> <volume>37</volume>, <fpage>1</fpage>&#x2013;<lpage>23</lpage>. doi: <pub-id pub-id-type="doi">10.1163/22134808-bja10114</pub-id>, PMID: <pub-id pub-id-type="pmid">37963487</pub-id></citation></ref>
<ref id="ref38"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Palmer</surname> <given-names>S. E.</given-names></name> <name><surname>Langlois</surname> <given-names>T. A.</given-names></name> <name><surname>Schloss</surname> <given-names>K. B.</given-names></name></person-group> (<year>2016</year>). <article-title>Music-to-color associations of single-line piano melodies in non-synesthetes</article-title>. <source>Multisens. Res.</source> <volume>29</volume>, <fpage>157</fpage>&#x2013;<lpage>193</lpage>. doi: <pub-id pub-id-type="doi">10.1163/22134808-00002486</pub-id>, PMID: <pub-id pub-id-type="pmid">27311295</pub-id></citation></ref>
<ref id="ref39"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Palmer</surname> <given-names>S. E.</given-names></name> <name><surname>Schloss</surname> <given-names>K. B.</given-names></name> <name><surname>Xu</surname> <given-names>Z.</given-names></name> <name><surname>Prado-Leon</surname> <given-names>L. R.</given-names></name></person-group> (<year>2013</year>). <article-title>Music-color associations are mediated by emotion</article-title>. <source>Proc. Natl. Acad. Sci.</source> <volume>110</volume>, <fpage>8836</fpage>&#x2013;<lpage>8841</lpage>. doi: <pub-id pub-id-type="doi">10.1073/pnas.1212562110</pub-id>, PMID: <pub-id pub-id-type="pmid">23671106</pub-id></citation></ref>
<ref id="ref40"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Parise</surname> <given-names>C. V.</given-names></name> <name><surname>Spence</surname> <given-names>C.</given-names></name></person-group> (<year>2013</year>). <article-title>Audiovisual crossmodal correspondences and sound symbolism: a study using the implicit association test</article-title>. <source>Exp. Brain Res.</source> <volume>220</volume>, <fpage>319</fpage>&#x2013;<lpage>333</lpage>. doi: <pub-id pub-id-type="doi">10.1007/s00221-012-3140-6</pub-id>, PMID: <pub-id pub-id-type="pmid">22706551</pub-id></citation></ref>
<ref id="ref41"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Qi</surname> <given-names>Y.</given-names></name> <name><surname>Huang</surname> <given-names>F.</given-names></name> <name><surname>Li</surname> <given-names>Z.</given-names></name> <name><surname>Wan</surname> <given-names>X.</given-names></name></person-group> (<year>2020</year>). <article-title>Crossmodal correspondences in the sounds of Chinese instruments</article-title>. <source>Perception</source> <volume>49</volume>, <fpage>81</fpage>&#x2013;<lpage>97</lpage>. doi: <pub-id pub-id-type="doi">10.1177/0301006619888992</pub-id>, PMID: <pub-id pub-id-type="pmid">31747837</pub-id></citation></ref>
<ref id="ref42"><citation citation-type="book"><person-group person-group-type="author"><collab id="coll1">R Core Team</collab></person-group> (<year>2023</year>). <source>R: A language and environment for statistical computing</source>. <publisher-loc>Vienna, Austria</publisher-loc>: <publisher-name>R Foundation for Statistical Computing</publisher-name>.</citation></ref>
<ref id="ref43"><citation citation-type="confproc"><person-group person-group-type="author"><name><surname>Reuter</surname> <given-names>C.</given-names></name> <name><surname>Jewanski</surname> <given-names>J.</given-names></name> <name><surname>Saitis</surname> <given-names>C.</given-names></name> <name><surname>Czedik-Eysenberg</surname> <given-names>I.</given-names></name> <name><surname>Siddiq</surname> <given-names>S.</given-names></name> <name><surname>Kruchten</surname> <given-names>S.</given-names></name> <etal/></person-group>. (<year>2018</year>). <article-title>Colors and timbres: consistency and tendencies of color-timbre mappings in non-synesthetic individuals</article-title>. <conf-name>Conference of the Deutschen Gesellschaft f&#x00FC;r Musikpsychologie (DGM)</conf-name>. Available at: <ext-link xlink:href="https://www.researchgate.net/profile/Christoph-Reuter/publication/327682218_Colors_and_Timbres_-_Consistency_and_Tendencies_of_Color-Timbre_Mappings_in_non-synesthetic_Individuals/links/5b9eea4d92851ca9ed10d5ba/Colors-and-Timbres-Consistency-and-Tendencies-of-Color-Timbre-Mappings-in-non-synesthetic-Individuals.pdf" ext-link-type="uri">https://www.researchgate.net/profile/Christoph-Reuter/publication/327682218_Colors_and_Timbres_-_Consistency_and_Tendencies_of_Color-Timbre_Mappings_in_non-synesthetic_Individuals/links/5b9eea4d92851ca9ed10d5ba/Colors-and-Timbres-Consistency-and-Tendencies-of-Color-Timbre-Mappings-in-non-synesthetic-Individuals.pdf</ext-link></citation></ref>
<ref id="ref44"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Reymore</surname> <given-names>L.</given-names></name></person-group> (<year>2021</year>). <article-title>Variations in timbre qualia with register and dynamics in the oboe and French horn</article-title>. <source>Empiric. Musicol. Rev.</source> <volume>16</volume>, <fpage>231</fpage>&#x2013;<lpage>275</lpage>. doi: <pub-id pub-id-type="doi">10.18061/emr.v16i2.8005</pub-id></citation></ref>
<ref id="ref45"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Reymore</surname> <given-names>L.</given-names></name> <name><surname>Huron</surname> <given-names>D.</given-names></name></person-group> (<year>2020</year>). <article-title>Using auditory imagery tasks to map the cognitive linguistic dimensions of musical instrument timbre qualia</article-title>. <source>Psychomusicology</source> <volume>30</volume>, <fpage>124</fpage>&#x2013;<lpage>144</lpage>. doi: <pub-id pub-id-type="doi">10.1037/pmu0000263</pub-id></citation></ref>
<ref id="ref46"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Reymore</surname> <given-names>L.</given-names></name> <name><surname>Noble</surname> <given-names>J.</given-names></name> <name><surname>Saitis</surname> <given-names>C.</given-names></name> <name><surname>Traube</surname> <given-names>C.</given-names></name> <name><surname>Wallmark</surname> <given-names>Z.</given-names></name></person-group> (<year>2023</year>). <article-title>Timbre semantic associations vary both between and within instruments: an empirical study incorporating register and pitch height</article-title>. <source>Music Percept.</source> <volume>40</volume>, <fpage>253</fpage>&#x2013;<lpage>274</lpage>. doi: <pub-id pub-id-type="doi">10.1525/mp.2023.40.3.253</pub-id></citation></ref>
<ref id="ref47"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Rosi</surname> <given-names>V.</given-names></name> <name><surname>Houix</surname> <given-names>O.</given-names></name> <name><surname>Misdariis</surname> <given-names>N.</given-names></name> <name><surname>Susini</surname> <given-names>P.</given-names></name></person-group> (<year>2022</year>). <article-title>Investigating the shared meaning of metaphorical sound attributes: bright, warm, round, and rough</article-title>. <source>Music Percept.</source> <volume>39</volume>, <fpage>468</fpage>&#x2013;<lpage>483</lpage>. doi: <pub-id pub-id-type="doi">10.1525/mp.2022.39.5.468</pub-id></citation></ref>
<ref id="ref48"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Saitis</surname> <given-names>C.</given-names></name> <name><surname>Wallmark</surname> <given-names>Z.</given-names></name></person-group> (<year>2024</year>). <article-title>Timbral brightness perception investigated through multimodal interference</article-title>. <source>Atten. Percept. Psychophys.</source> <volume>86</volume>, <fpage>1835</fpage>&#x2013;<lpage>1845</lpage>. doi: <pub-id pub-id-type="doi">10.3758/s13414-024-02934-2</pub-id>, PMID: <pub-id pub-id-type="pmid">39090510</pub-id></citation></ref>
<ref id="ref49"><citation citation-type="book"><person-group person-group-type="author"><name><surname>Saitis</surname> <given-names>C.</given-names></name> <name><surname>Weinzierl</surname> <given-names>S.</given-names></name></person-group> (<year>2019</year>). &#x201C;<article-title>The semantics of timbre</article-title>&#x201D; in <source>Timbre: Acoustics, perception, cognition</source>. eds. <person-group person-group-type="editor"><name><surname>Siedenburg</surname> <given-names>K.</given-names></name> <name><surname>Saitis</surname> <given-names>C.</given-names></name> <name><surname>McAdams</surname> <given-names>S.</given-names></name> <name><surname>Popper</surname> <given-names>A. N.</given-names></name> <name><surname>Fay</surname> <given-names>R. R.</given-names></name></person-group> (<publisher-loc>Cham, Switzerland</publisher-loc>: <publisher-name>Springer International Publishing</publisher-name>), <fpage>119</fpage>&#x2013;<lpage>149</lpage>.</citation></ref>
<ref id="ref50"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Saitis</surname> <given-names>C.</given-names></name> <name><surname>Weinzierl</surname> <given-names>S.</given-names></name> <name><surname>von Kriegstein</surname> <given-names>K.</given-names></name> <name><surname>Ystad</surname> <given-names>S.</given-names></name> <name><surname>Cuskley</surname> <given-names>C.</given-names></name></person-group> (<year>2020</year>). <article-title>Timbre semantics through the lens of crossmodal correspondences: a new way of asking old questions</article-title>. <source>Acoust. Sci. Technol.</source> <volume>41</volume>, <fpage>365</fpage>&#x2013;<lpage>368</lpage>. doi: <pub-id pub-id-type="doi">10.1250/ast.41.365</pub-id></citation></ref>
<ref id="ref51"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Siedenburg</surname> <given-names>K.</given-names></name> <name><surname>McAdams</surname> <given-names>S.</given-names></name></person-group> (<year>2017</year>). <article-title>Four distinctions for the auditory &#x201C;wastebasket&#x201D; of timbre</article-title>. <source>Front. Psychol.</source> <volume>8</volume>:<fpage>1747</fpage>. doi: <pub-id pub-id-type="doi">10.3389/fpsyg.2017.01747</pub-id>, PMID: <pub-id pub-id-type="pmid">29046659</pub-id></citation></ref>
<ref id="ref52"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Spence</surname> <given-names>C.</given-names></name></person-group> (<year>2011</year>). <article-title>Crossmodal correspondences: A tutorial review</article-title>. <source>Atten. Percept. Psychophys.</source> <volume>73</volume>, <fpage>971</fpage>&#x2013;<lpage>995</lpage>. doi: <pub-id pub-id-type="doi">10.3758/s13414-010-0073-7</pub-id>, PMID: <pub-id pub-id-type="pmid">21264748</pub-id></citation></ref>
<ref id="ref53"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Spence</surname> <given-names>C.</given-names></name></person-group> (<year>2020a</year>). <article-title>Assessing the role of emotional mediation in explaining crossmodal correspondences involving musical stimuli</article-title>. <source>Multisens. Res.</source> <volume>33</volume>, <fpage>1</fpage>&#x2013;<lpage>29</lpage>. doi: <pub-id pub-id-type="doi">10.1163/22134808-20191469</pub-id>, PMID: <pub-id pub-id-type="pmid">31648195</pub-id></citation></ref>
<ref id="ref54"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Spence</surname> <given-names>C.</given-names></name></person-group> (<year>2020b</year>). <article-title>Simple and complex crossmodal correspondences involving audition</article-title>. <source>Acoust. Sci. Technol.</source> <volume>41</volume>, <fpage>6</fpage>&#x2013;<lpage>12</lpage>. doi: <pub-id pub-id-type="doi">10.1250/ast.41.6</pub-id></citation></ref>
<ref id="ref55"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Spence</surname> <given-names>C.</given-names></name> <name><surname>Di Stefano</surname> <given-names>N.</given-names></name></person-group> (<year>2022</year>). <article-title>Coloured hearing, colour music, colour organs, and the search for perceptually meaningful correspondences between colour and sound</article-title>. <source>i-Perception</source> <volume>13</volume>. doi: <pub-id pub-id-type="doi">10.1177/20416695221092802</pub-id>, PMID: <pub-id pub-id-type="pmid">35572076</pub-id></citation></ref>
<ref id="ref56"><citation citation-type="book"><person-group person-group-type="author"><name><surname>Spence</surname> <given-names>C.</given-names></name> <name><surname>Sathian</surname> <given-names>K.</given-names></name></person-group> (<year>2020</year>). &#x201C;<article-title>Audiovisual crossmodal correspondences: Behavioral consequences and neural underpinnings</article-title>&#x201D; in <source>Multisensory perception</source> (<publisher-loc>London, United Kingdom</publisher-loc>: <publisher-name>Academic Press</publisher-name>), <fpage>239</fpage>&#x2013;<lpage>258</lpage>.</citation></ref>
<ref id="ref57"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Walker</surname> <given-names>P.</given-names></name> <name><surname>Scallon</surname> <given-names>G.</given-names></name> <name><surname>Francis</surname> <given-names>B.</given-names></name></person-group> (<year>2017</year>). <article-title>Cross-sensory correspondences: heaviness is dark and low-pitched</article-title>. <source>Perception</source> <volume>46</volume>, <fpage>772</fpage>&#x2013;<lpage>792</lpage>. doi: <pub-id pub-id-type="doi">10.1177/0301006616684369</pub-id>, PMID: <pub-id pub-id-type="pmid">28622755</pub-id></citation></ref>
<ref id="ref58"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Walker</surname> <given-names>P.</given-names></name> <name><surname>Smith</surname> <given-names>S.</given-names></name></person-group> (<year>1984</year>). <article-title>Stroop interference based on the synaesthetic qualities of auditory pitch</article-title>. <source>Perception</source> <volume>13</volume>, <fpage>75</fpage>&#x2013;<lpage>81</lpage>. doi: <pub-id pub-id-type="doi">10.1068/p130075</pub-id>, PMID: <pub-id pub-id-type="pmid">6473055</pub-id></citation></ref>
<ref id="ref59"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Walker</surname> <given-names>L.</given-names></name> <name><surname>Walker</surname> <given-names>P.</given-names></name> <name><surname>Francis</surname> <given-names>B.</given-names></name></person-group> (<year>2012</year>). <article-title>A common scheme for cross-sensory correspondences across stimulus domains</article-title>. <source>Perception</source> <volume>41</volume>, <fpage>1186</fpage>&#x2013;<lpage>1192</lpage>. doi: <pub-id pub-id-type="doi">10.1068/p7149</pub-id>, PMID: <pub-id pub-id-type="pmid">23469700</pub-id></citation></ref>
<ref id="ref60"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wallmark</surname> <given-names>Z.</given-names></name></person-group> (<year>2019a</year>). <article-title>A corpus analysis of timbre semantics in orchestration treatises</article-title>. <source>Psychol. Music</source> <volume>47</volume>, <fpage>585</fpage>&#x2013;<lpage>605</lpage>. doi: <pub-id pub-id-type="doi">10.1177/0305735618768102</pub-id></citation></ref>
<ref id="ref61"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wallmark</surname> <given-names>Z.</given-names></name></person-group> (<year>2019b</year>). <article-title>Semantic crosstalk in timbre perception</article-title>. <source>Music Sci.</source> <volume>2</volume>, <fpage>1</fpage>&#x2013;<lpage>18</lpage>. doi: <pub-id pub-id-type="doi">10.1177/2059204319846617</pub-id>, PMID: <pub-id pub-id-type="pmid">39688133</pub-id></citation></ref>
<ref id="ref62"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wallmark</surname> <given-names>Z.</given-names></name> <name><surname>Allen</surname> <given-names>S. E.</given-names></name></person-group> (<year>2020</year>). <article-title>Preschoolers&#x2019; crossmodal mappings of timbre</article-title>. <source>Atten. Percept. Psychophys.</source> <volume>82</volume>, <fpage>2230</fpage>&#x2013;<lpage>2236</lpage>. doi: <pub-id pub-id-type="doi">10.3758/s13414-020-02015-0</pub-id>, PMID: <pub-id pub-id-type="pmid">32166645</pub-id></citation></ref>
<ref id="ref63"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wallmark</surname> <given-names>Z.</given-names></name> <name><surname>Nghiem</surname> <given-names>L.</given-names></name> <name><surname>Marks</surname> <given-names>L. E.</given-names></name></person-group> (<year>2021</year>). <article-title>Does timbre modulate visual perception? Exploring crossmodal interactions</article-title>. <source>Music Percept.</source> <volume>39</volume>, <fpage>1</fpage>&#x2013;<lpage>20</lpage>. doi: <pub-id pub-id-type="doi">10.1525/mp.2021.39.1.1</pub-id></citation></ref>
<ref id="ref64"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wang</surname> <given-names>Q. J.</given-names></name> <name><surname>Spence</surname> <given-names>C.</given-names></name></person-group> (<year>2017</year>). <article-title>The role of pitch and tempo in sound-temperature crossmodal correspondences</article-title>. <source>Multisens. Res.</source> <volume>30</volume>, <fpage>307</fpage>&#x2013;<lpage>320</lpage>. doi: <pub-id pub-id-type="doi">10.1163/22134808-00002564</pub-id>, PMID: <pub-id pub-id-type="pmid">31287077</pub-id></citation></ref>
<ref id="ref65"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ward</surname> <given-names>J.</given-names></name> <name><surname>Huckstep</surname> <given-names>B.</given-names></name> <name><surname>Tsakanikos</surname> <given-names>E.</given-names></name></person-group> (<year>2006</year>). <article-title>Sound-colour synaesthesia: to what extent does it use cross-modal mechanisms common to us all?</article-title> <source>Cortex</source> <volume>42</volume>, <fpage>264</fpage>&#x2013;<lpage>280</lpage>. doi: <pub-id pub-id-type="doi">10.1016/S0010-9452(08)70352-6</pub-id>, PMID: <pub-id pub-id-type="pmid">16683501</pub-id></citation></ref>
<ref id="ref66"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Warrier</surname> <given-names>C. M.</given-names></name> <name><surname>Zatorre</surname> <given-names>R. J.</given-names></name></person-group> (<year>2002</year>). <article-title>Influence of tonal context and timbral variation on perception of pitch</article-title>. <source>Percept. Psychophys.</source> <volume>64</volume>, <fpage>198</fpage>&#x2013;<lpage>207</lpage>. doi: <pub-id pub-id-type="doi">10.3758/BF03195786</pub-id>, PMID: <pub-id pub-id-type="pmid">12013375</pub-id></citation></ref>
<ref id="ref67"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Weinreich</surname> <given-names>A.</given-names></name> <name><surname>Gollwitzer</surname> <given-names>A.</given-names></name></person-group> (<year>2016</year>). <article-title>Automaticity and affective responses in valence transfer: insights from the crossmodal auditory-visual paradigm</article-title>. <source>Psychol. Music</source> <volume>44</volume>, <fpage>1304</fpage>&#x2013;<lpage>1317</lpage>. doi: <pub-id pub-id-type="doi">10.1177/0305735615626519</pub-id></citation></ref>
<ref id="ref68"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Whiteford</surname> <given-names>K. L.</given-names></name> <name><surname>Schloss</surname> <given-names>K. B.</given-names></name> <name><surname>Helwig</surname> <given-names>N. E.</given-names></name> <name><surname>Palmer</surname> <given-names>S. E.</given-names></name></person-group> (<year>2018</year>). <article-title>Color, music, and emotion: Bach to the blues</article-title>. <source>i-Perception</source> <volume>9</volume>, <fpage>1</fpage>&#x2013;<lpage>27</lpage>. doi: <pub-id pub-id-type="doi">10.1177/2041669518808535</pub-id>, PMID: <pub-id pub-id-type="pmid">30479734</pub-id></citation></ref>
<ref id="ref69"><citation citation-type="book"><person-group person-group-type="author"><name><surname>Wyszecki</surname> <given-names>G.</given-names></name> <name><surname>Stiles</surname> <given-names>W. S.</given-names></name></person-group> (<year>1982</year>). <source>Color science</source>, vol. <volume>8</volume>. <publisher-loc>New York</publisher-loc>: <publisher-name>Wiley</publisher-name>.</citation></ref>
</ref-list>
</back>
</article>