<?xml version="1.0" encoding="UTF-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Archiving and Interchange DTD v2.3 20070202//EN" "archivearticle.dtd">
<article article-type="data-paper" dtd-version="2.3" xml:lang="EN" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Sig. Proc.</journal-id>
<journal-title>Frontiers in Signal Processing</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Sig. Proc.</abbrev-journal-title>
<issn pub-type="epub">2673-8198</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="publisher-id">1380060</article-id>
<article-id pub-id-type="doi">10.3389/frsip.2024.1380060</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Signal Processing</subject>
<subj-group>
<subject>Data Report</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>A multi-loudspeaker binaural room impulse response dataset with high-resolution translational and rotational head coordinates in a listening room</article-title>
<alt-title alt-title-type="left-running-head">Qiao et al.</alt-title>
<alt-title alt-title-type="right-running-head">
<ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3389/frsip.2024.1380060">10.3389/frsip.2024.1380060</ext-link>
</alt-title>
</title-group>
<contrib-group>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Qiao</surname>
<given-names>Yue</given-names>
</name>
<xref ref-type="corresp" rid="c001">&#x2a;</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2376222/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Gonzales</surname>
<given-names>Ryan Miguel</given-names>
</name>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Choueiri</surname>
<given-names>Edgar</given-names>
</name>
<role content-type="https://credit.niso.org/contributor-roles/funding-acquisition/"/>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
<role content-type="https://credit.niso.org/contributor-roles/resources/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
</contrib-group>
<aff>
<institution>3D Audio and Applied Acoustics Laboratory</institution>, <institution>Princeton University</institution>, <addr-line>Princeton</addr-line>, <addr-line>NJ</addr-line>, <country>United States</country>
</aff>
<author-notes>
<fn fn-type="edited-by">
<p>
<bold>Edited by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/2287762/overview">Ante Jukic</ext-link>, Nvidia, United States</p>
</fn>
<fn fn-type="edited-by">
<p>
<bold>Reviewed by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1488370/overview">Emanu&#xeb;l Habets</ext-link>, International Audio Laboratories Erlangen (AudioLabs), Germany</p>
<p>
<ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1110624/overview">Lukas Asp&#xf6;ck</ext-link>, RWTH Aachen University, Germany</p>
</fn>
<corresp id="c001">&#x2a;Correspondence: Yue Qiao, <email>yqiao@princeton.edu</email>
</corresp>
</author-notes>
<pub-date pub-type="epub">
<day>29</day>
<month>04</month>
<year>2024</year>
</pub-date>
<pub-date pub-type="collection">
<year>2024</year>
</pub-date>
<volume>4</volume>
<elocation-id>1380060</elocation-id>
<history>
<date date-type="received">
<day>31</day>
<month>01</month>
<year>2024</year>
</date>
<date date-type="accepted">
<day>05</day>
<month>04</month>
<year>2024</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#xa9; 2024 Qiao, Gonzales and Choueiri.</copyright-statement>
<copyright-year>2024</copyright-year>
<copyright-holder>Qiao, Gonzales and Choueiri</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<kwd-group>
<kwd>binaural room impulse response</kwd>
<kwd>acoustic dataset</kwd>
<kwd>spatial audio</kwd>
<kwd>listener movement</kwd>
<kwd>acoustic measurement</kwd>
</kwd-group>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-at-acceptance</meta-name>
<meta-value>Audio and Acoustic Signal Processing</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec id="s1">
<title>1 Introduction</title>
<p>A binaural room impulse response (BRIR) describes the characteristics of acoustic wave interactions from a sound source in a room to the torso, head, and ears of a listener. The use of BRIRs has been ubiquitous in many audio applications. For example, in spatial audio reproduction with headphones, BRIRs are used as audio filters to simulate or reproduce an immersive and perceptually plausible sounding environment; in loudspeaker-based applications, the frequency-domain counterparts of BRIRs are equivalent to the acoustic transfer functions between the loudspeakers and the listener&#x2019;s ears, based on which audio filters are designed for tasks such as crosstalk cancellation (<xref ref-type="bibr" rid="B5">Cooper and Bauck, 1989</xref>; <xref ref-type="bibr" rid="B9">Gardner, 1998</xref>; <xref ref-type="bibr" rid="B4">Choueiri, 2018</xref>), room correction/loudspeaker equalization (<xref ref-type="bibr" rid="B11">Karjalainen et al., 1999</xref>; <xref ref-type="bibr" rid="B15">Lindfors et al., 2022</xref>), and personal sound zones (<xref ref-type="bibr" rid="B6">Druyvesteyn and Garas, 1997</xref>; <xref ref-type="bibr" rid="B2">Betlehem et al., 2015</xref>; <xref ref-type="bibr" rid="B17">Qiao and Choueiri, 2023a</xref>). In addition to audio reproduction and rendering, BRIRs have also played an important role in other audio-related tasks, such as sound source localization (<xref ref-type="bibr" rid="B19">Shinn-Cunningham et al., 2005</xref>), sound source separation (<xref ref-type="bibr" rid="B23">Yu et al., 2016</xref>), and audio-visual learning (<xref ref-type="bibr" rid="B22">Younes et al., 2023</xref>).</p>
<p>As implied by its name, a BRIR is dependent on both the listener&#x2019;s anthropometric features (e.g., ear size and shape) and the room&#x2019;s geometry and acoustic properties. Due to the complex acoustic interactions, such as sound reflections in the room and scattering off the listener, a BRIR varies with both the listener&#x2019;s position and orientation in the room. This is unlike room impulse response (RIR), which only depends on the position, or anechoic head-related impulse response (HRIR) which, in the far-field case, only depends on the orientation. Although there have been multiple HRIR and RIR datasets (<xref ref-type="bibr" rid="B20">Sridhar et al., 2017</xref>; <xref ref-type="bibr" rid="B3">Brinkmann et al., 2019</xref>; <xref ref-type="bibr" rid="B14">Koyama et al., 2021</xref>) available, and while it is possible to synthesize BRIRs from HRIRs and RIRs using methods such as the image source model (<xref ref-type="bibr" rid="B21">Wendt et al., 2014</xref>), synthesized BRIRs lose physical accuracy and can only maintain perceptual plausibility. While this is sufficient for some applications, such as headphone-based auralization, it is not appropriate for others, such as crosstalk cancellation and personal sound zones, where measured BRIRs are required. Moreover, the lack of a high-resolution BRIR dataset, compared to existing ones measured at sparse listener positions and orientations (<xref ref-type="bibr" rid="B10">Jeub et al., 2009</xref>; <xref ref-type="bibr" rid="B13">Kayser et al., 2009</xref>; <xref ref-type="bibr" rid="B7">Erbes et al., 2015</xref>), limits the study of BRIR modeling and interpolation at high frequencies and the development of machine learning-based audio applications.</p>
<p>In this paper, we introduce a dataset that contains BRIRs measured in an acoustically treated listening room using multiple loudspeakers and at high-resolution translational and rotational head coordinates. Although the dataset is only measured in one specific room, it is expected that it will be useful for a wide range of applications as it faithfully captures the spatial dependency of BRIRs on listener positions and orientations. For example, it can be directly applied to studies of BRIR modeling and to interpolation and applications that require multi-loudspeaker BRIRs, such as crosstalk cancellation and personal sound zones. It has been shown (<xref ref-type="bibr" rid="B18">Qiao and Choueiri, 2023b</xref>) that the spatial sampling resolution of the dataset is adequate for rendering personal sound zones with continuous listener movements within a certain frequency range. With proper interpolation between the measured BRIRs, the dataset can also be used to simulate binaural audio for headphone-based auralization with continuous listener movements. In addition, the dataset can be used for either data augmentation or performance evaluation in a wide range of machine learning-based tasks that require binaural audio with listener movements in multiple degrees of freedom.</p>
</sec>
<sec sec-type="methods" id="s2">
<title>2 Methods</title>
<sec id="s2-1">
<title>2.1 Data collection</title>
<p>We measured BRIRs in an irregularly shaped listening room of a near-shoebox shape (see <xref ref-type="fig" rid="F1">Figure 1</xref> for the exact dimensions). The room had a <italic>RT</italic>
<sub>60</sub> of 0.24&#xa0;s averaged in the range of 1,300 and 6,300&#xa0;Hz. Its floor was covered with carpet, and the walls and ceiling were partially covered with acoustic panels. <xref ref-type="fig" rid="F1">Figure 1</xref> also shows the setup and dimensions of the measurement system. A linear array of eight loudspeakers was used as a sound source; each was a Focal Shape 40 4-inch Flax woofer. The loudspeaker array layout was initially intended for sound field control applications, such as rendering personal sound zones. A Br&#xfc;el &#x26; Kj&#xe6;r Head and Torso Simulator (HATS, Type 4,100) was used as the mannequin listener, with its built-in microphones replaced with a pair of in-ear binaural microphones (Theoretica Applied Physics BACCH-BM Pro). The microphones were calibrated and free-field equalized before the measurement. A custom-made, computer-controlled mechanical translation platform was applied to enable translational movements, and a turntable (Outline ET250-3D) was mounted on top of the platform for rotational movements in the azimuth.</p>
<fig id="F1" position="float">
<label>FIGURE 1</label>
<caption>
<p>Illustrations of the measurement system <bold>(A)</bold>. Photograph of the measurement system. Note that the tweeter loudspeaker arrays in the photograph were not used in the data collection <bold>(B)</bold>. Schematic representation of the listening room. Left: side view. Right: top view <bold>(C)</bold>. Schematic diagram of the measurement system <bold>(D)</bold>. Dimensions of the loudspeaker array.</p>
</caption>
<graphic xlink:href="frsip-04-1380060-g001.tif"/>
</fig>
<p>The BRIR measurement grid has a range of [0.5, 1.0] m in the <italic>y</italic> direction (front/back) and [-0.5, 0.5] m in the <italic>x</italic> direction (left/right), with a 0.05-m spacing between adjacent grid points. The distances are relative to the center of the loudspeaker array. At each grid point, the BRIRs were measured at 37 different azimuth angles from the listener facing left to facing right, with a 5&#xb0; spacing between adjacent angles. In total, there were 68376 (&#x3d;11 <italic>y</italic>-translations &#xd7; 21 <italic>x</italic>-translations &#xd7; 37 azimuthal rotations &#xd7; 8 loudspeakers) BRIRs measured.</p>
<p>We measured BRIRs by playing back exponential sine sweep (ESS) signals from the loudspeakers and recording the signals received with binaural microphones. Each sine sweep signal had a length of 500&#xa0;ms at a 48&#xa0;kHz sampling rate and was generated using the synchronized ESS method (<xref ref-type="bibr" rid="B16">Novak et al., 2015</xref>), with a start frequency of 100&#xa0;Hz and an end frequency of 24&#xa0;kHz. Synchronized ESS is a variant of the traditional ESS method (<xref ref-type="bibr" rid="B8">Farina, 2000</xref>), with the advantage of correctly estimating higher harmonic frequency responses. All eight loudspeakers were triggered in series with no overlapping between the ESS signals.</p>
<p>The entire data collection process was split into multiple measurement sessions. For each session, we manually fixed the distance from the listener to the array in the <italic>y</italic> direction and automated the movements in the <italic>x</italic> direction and the azimuthal rotations. The measurement automation, signal generation, and data collection were implemented in Cycling &#x2019;74 Max 8. The BRIR post-processing was performed in MATLAB. Each session lasted for approximately 2&#xa0;h, and the entire data collection process took 9 days.</p>
</sec>
<sec id="s2-2">
<title>2.2 Data processing</title>
<p>The BRIRs were obtained by first deconvolving the recorded signals with the ESS signal in the frequency domain, with 32768-length FFT at a 48&#xa0;kHz sampling rate. A fourth-order highpass Butterworth filter with a cutoff frequency of 100&#xa0;Hz was then applied to the deconvolved signals to remove the low-frequency noise present during the measurement. Finally, the deconvolved signals were truncated to the first 16,384 samples (corresponding to 341.3&#xa0;ms) and globally normalized. No loudspeaker equalization was applied to the BRIRs as loudspeaker-specific information, such as directivity, is an integral part of the BRIR and was therefore difficult to compensate for. The processed BRIRs, together with the corresponding listener position and orientation coordinates, were saved as separate files corresponding to different <italic>y</italic> translations in the SOFA (spatially oriented format for acoustics, <xref ref-type="bibr" rid="B1">AES69-2022 (2022)</xref>) format, following the AES69-2022 (SOFA 2.1) standard. The dataset was generated using SOFA Toolbox for MATLAB/Octave version 2.2.0.</p>
</sec>
</sec>
<sec id="s3">
<title>3 Data visualization</title>
<p>We examine the dataset by visualizing 1) the time index of the BRIR onset, 2) the peak amplitude of the BRIR, and 3) the interaural time difference (ITD) of the BRIR as functions of the listener position. Both the onset and the ITD of the BRIRs were calculated in a thresholding approach (see <xref ref-type="bibr" rid="B12">Katz and Noisternig (2014)</xref> as an example). <xref ref-type="fig" rid="F2">Figure 2</xref> shows the onset, peak amplitude, and ITD of the BRIRs measured at the fourth loudspeaker (counting from the left) and with the listener facing forward. The colors in the figures were interpolated using the &#x201c;interp&#x201d; option of the MATLAB &#x201c;pcolor&#x201d; function. All three figures show clear spatial dependency on the listener&#x2019;s position, with the onset increasing and the peak amplitude decreasing as the listener moved away from the loudspeaker. The ITD was nearly zero when the listener was on-axis with the loudspeaker and increased as the listener moved to off-axis positions. Note that the maxima of the peak amplitude in <xref ref-type="fig" rid="F2">Figure 2</xref> are not on-axis with the loudspeaker, which was due to the occlusion effect of the listener&#x2019;s head and ear.</p>
<fig id="F2" position="float">
<label>FIGURE 2</label>
<caption>
<p>Three example figures showing the properties of the BRIRs measured at the fourth loudspeaker (counting from the left) and with the listener facing forward. The <italic>x</italic>-axis of the figures is aligned with the loudspeaker array. The dashed line in the figure indicates the center of the fourth loudspeaker. Top figure: onsets of left-ear BRIRs. Middle figure: peak amplitude of left-ear BRIRs. Bottom figure: calculated ITD of BRIRs.</p>
</caption>
<graphic xlink:href="frsip-04-1380060-g002.tif"/>
</fig>
</sec>
</body>
<back>
<sec sec-type="data-availability" id="s4">
<title>Data availability statement</title>
<p>The datasets presented in this study can be found in online repositories. The names of the repository/repositories and accession number(s) can be found in the article/Supplementary material <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.34770/6gc9-5787">https://doi.org/10.34770/6gc9-5787</ext-link>.</p>
</sec>
<sec id="s5">
<title>Author contributions</title>
<p>YQ: conceptualization, data curation, formal analysis, methodology, software, visualization, writing&#x2013;original draft, and writing&#x2013;review and editing. RG: formal analysis, investigation, software, validation, visualization, and writing&#x2013;review and editing. EC: funding acquisition, project administration, resources, supervision, and writing&#x2013;review and editing.</p>
</sec>
<sec sec-type="funding-information" id="s6">
<title>Funding</title>
<p>The author(s) declare financial support was received for the research, authorship, and/or publication of this article. This work was supported by a research grant from Masimo Corporation.</p>
</sec>
<ack>
<p>The authors wish to thank K. Tworek, R. Sridhar, and J. Tylka for their prior work in developing the control software program for the translation platform and the turntable. They also wish to thank L. Guadagnin for his support on the loudspeaker array.</p>
</ack>
<sec sec-type="COI-statement" id="s7">
<title>Conflict of interest</title>
<p>The authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec sec-type="disclaimer" id="s8">
<title>Publisher&#x2019;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors, and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<ref-list>
<title>References</title>
<ref id="B1">
<citation citation-type="book">
<collab>AES69-2022</collab> (<year>2022</year>). <source>AES standard for file exchange-spatial acoustic data file format</source>. <publisher-loc>Standard</publisher-loc>: <publisher-name>Audio Engineering Society</publisher-name>.</citation>
</ref>
<ref id="B2">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Betlehem</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Poletti</surname>
<given-names>M. A.</given-names>
</name>
<name>
<surname>Abhayapala</surname>
<given-names>T. D.</given-names>
</name>
</person-group> (<year>2015</year>). <article-title>Personal sound zones: delivering interface-free audio to multiple listeners</article-title>. <source>IEEE Signal Process. Mag.</source> <volume>32</volume>, <fpage>81</fpage>&#x2013;<lpage>91</lpage>. <pub-id pub-id-type="doi">10.1109/msp.2014.2360707</pub-id>
</citation>
</ref>
<ref id="B3">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Brinkmann</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Dinakaran</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Pelzer</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Grosche</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Voss</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Weinzierl</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>A cross-evaluated database of measured and simulated hrtfs including 3d head meshes, anthropometric features, and headphone impulse responses</article-title>. <source>J. Audio Eng. Soc.</source> <volume>67</volume>, <fpage>705</fpage>&#x2013;<lpage>718</lpage>. <pub-id pub-id-type="doi">10.17743/jaes.2019.0024</pub-id>
</citation>
</ref>
<ref id="B4">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Choueiri</surname>
<given-names>E.</given-names>
</name>
</person-group> (<year>2018</year>). &#x201c;<article-title>Binaural audio through loudspeakers</article-title>,&#x201d; in <source>Immersive sound: the art and science of binaural and multi-channel audio</source>. Editors <person-group person-group-type="editor">
<name>
<surname>Roginska</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Geluso</surname>
<given-names>P.</given-names>
</name>
</person-group> (<publisher-loc>New York, NY, USA</publisher-loc>. <publisher-name>Taylor and Francis</publisher-name>), <fpage>124</fpage>&#x2013;<lpage>179</lpage>. <comment>chap. 6</comment>.</citation>
</ref>
<ref id="B5">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Cooper</surname>
<given-names>D. H.</given-names>
</name>
<name>
<surname>Bauck</surname>
<given-names>J. L.</given-names>
</name>
</person-group> (<year>1989</year>). <article-title>Prospects for transaural recording</article-title>. <source>J. Audio Eng. Soc.</source> <volume>37</volume>, <fpage>3</fpage>&#x2013;<lpage>19</lpage>.</citation>
</ref>
<ref id="B6">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Druyvesteyn</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Garas</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>1997</year>). <article-title>Personal sound</article-title>. <source>J. Audio Eng. Soc.</source> <volume>45</volume>, <fpage>685</fpage>&#x2013;<lpage>701</lpage>.</citation>
</ref>
<ref id="B7">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Erbes</surname>
<given-names>V.</given-names>
</name>
<name>
<surname>Geier</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Weinzierl</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Spors</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2015</year>). &#x201c;<article-title>Database of single-channel and binaural room impulse responses of a 64-channel loudspeaker array</article-title>,&#x201d; in <source>Audio engineering society convention</source> (<publisher-loc>Warsaw, Poland</publisher-loc>. <publisher-name>Audio Engineering Society</publisher-name>), <volume>138</volume>.</citation>
</ref>
<ref id="B8">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Farina</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2000</year>). &#x201c;<article-title>Simultaneous measurement of impulse response and distortion with a swept-sine technique</article-title>,&#x201d; in <source>Audio engineering society convention</source> (<publisher-loc>Paris, France</publisher-loc>. <publisher-name>Audio Engineering Society</publisher-name>), <volume>108</volume>.</citation>
</ref>
<ref id="B9">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Gardner</surname>
<given-names>W. G.</given-names>
</name>
</person-group> (<year>1998</year>). <source>D audio using loudspeakers</source>. in <source>3</source>, <volume>444</volume>. <publisher-name>Springer Science and Business Media</publisher-name>.</citation>
</ref>
<ref id="B10">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Jeub</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Schafer</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Vary</surname>
<given-names>P.</given-names>
</name>
</person-group> (<year>2009</year>). &#x201c;<article-title>A binaural room impulse response database for the evaluation of dereverberation algorithms</article-title>,&#x201d; in <source>2009 16th international conference on digital signal processing</source> (<publisher-name>IEEE</publisher-name>), <fpage>1</fpage>&#x2013;<lpage>5</lpage>.</citation>
</ref>
<ref id="B11">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Karjalainen</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Piiril&#xe4;</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>J&#xe4;rvinen</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Huopaniemi</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>1999</year>). <article-title>Comparison of loudspeaker equalization methods based on dsp techniques</article-title>. <source>J. Audio Eng. Soc.</source> <volume>47</volume>, <fpage>14</fpage>&#x2013;<lpage>31</lpage>.</citation>
</ref>
<ref id="B12">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Katz</surname>
<given-names>B. F.</given-names>
</name>
<name>
<surname>Noisternig</surname>
<given-names>M.</given-names>
</name>
</person-group> (<year>2014</year>). <article-title>A comparative study of interaural time delay estimation methods</article-title>. <source>J. Acoust. Soc. Am.</source> <volume>135</volume>, <fpage>3530</fpage>&#x2013;<lpage>3540</lpage>. <pub-id pub-id-type="doi">10.1121/1.4875714</pub-id>
</citation>
</ref>
<ref id="B13">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kayser</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Ewert</surname>
<given-names>S. D.</given-names>
</name>
<name>
<surname>Anem&#xfc;ller</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Rohdenburg</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Hohmann</surname>
<given-names>V.</given-names>
</name>
<name>
<surname>Kollmeier</surname>
<given-names>B.</given-names>
</name>
</person-group> (<year>2009</year>). <article-title>Database of multichannel in-ear and behind-the-ear head-related and binaural room impulse responses</article-title>. <source>EURASIP J. Adv. signal Process.</source> <volume>2009</volume>, <fpage>298605</fpage>&#x2013;<lpage>298610</lpage>. <pub-id pub-id-type="doi">10.1155/2009/298605</pub-id>
</citation>
</ref>
<ref id="B14">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Koyama</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Nishida</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Kimura</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Abe</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Ueno</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Brunnstr&#xf6;m</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2021</year>). &#x201c;<article-title>Meshrir: a dataset of room impulse responses on meshed grid points for evaluating sound field analysis and synthesis methods</article-title>,&#x201d; in <source>2021 IEEE workshop on applications of signal processing to audio and acoustics (WASPAA)</source> (<publisher-name>IEEE</publisher-name>), <fpage>1</fpage>&#x2013;<lpage>5</lpage>.</citation>
</ref>
<ref id="B15">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lindfors</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Liski</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>V&#xe4;lim&#xe4;ki</surname>
<given-names>V.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Loudspeaker equalization for a moving listener</article-title>. <source>J. Audio Eng. Soc.</source> <volume>70</volume>, <fpage>722</fpage>&#x2013;<lpage>730</lpage>. <pub-id pub-id-type="doi">10.17743/jaes.2022.0020</pub-id>
</citation>
</ref>
<ref id="B16">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Novak</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Lotton</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Simon</surname>
<given-names>L.</given-names>
</name>
</person-group> (<year>2015</year>). <article-title>Synchronized swept-sine: theory, application, and implementation</article-title>. <source>J. Audio Eng. Soc.</source> <volume>63</volume>, <fpage>786</fpage>&#x2013;<lpage>798</lpage>. <pub-id pub-id-type="doi">10.17743/jaes.2015.0071</pub-id>
</citation>
</ref>
<ref id="B17">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Qiao</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Choueiri</surname>
<given-names>E.</given-names>
</name>
</person-group> (<year>2023a</year>). <article-title>The effects of individualized binaural room transfer functions for personal sound zones</article-title>. <source>J. Audio Eng. Soc.</source> <volume>71</volume>, <fpage>848</fpage>&#x2013;<lpage>858</lpage>. <pub-id pub-id-type="doi">10.17743/jaes.2022.0109</pub-id>
</citation>
</ref>
<ref id="B18">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Qiao</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Choueiri</surname>
<given-names>E.</given-names>
</name>
</person-group> (<year>2023b</year>). &#x201c;<article-title>Optimal spatial sampling of plant transfer functions for head-tracked personal sound zones</article-title>,&#x201d; in <source>Audio engineering society convention</source> (<publisher-loc>Espoo, Helsinki, Finland</publisher-loc>. <publisher-name>Audio Engineering Society</publisher-name>), <volume>154</volume>.</citation>
</ref>
<ref id="B19">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Shinn-Cunningham</surname>
<given-names>B. G.</given-names>
</name>
<name>
<surname>Kopco</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Martin</surname>
<given-names>T. J.</given-names>
</name>
</person-group> (<year>2005</year>). <article-title>Localizing nearby sound sources in a classroom: binaural room impulse responses</article-title>. <source>J. Acoust. Soc. Am.</source> <volume>117</volume>, <fpage>3100</fpage>&#x2013;<lpage>3115</lpage>. <pub-id pub-id-type="doi">10.1121/1.1872572</pub-id>
</citation>
</ref>
<ref id="B20">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Sridhar</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Tylka</surname>
<given-names>J. G.</given-names>
</name>
<name>
<surname>Choueiri</surname>
<given-names>E.</given-names>
</name>
</person-group> (<year>2017</year>). &#x201c;<article-title>A database of head-related transfer functions and morphological measurements</article-title>,&#x201d; in <source>Audio engineering society convention</source> (<publisher-loc>New York, NY, USA</publisher-loc>. <publisher-name>Audio Engineering Society</publisher-name>), <volume>143</volume>.</citation>
</ref>
<ref id="B21">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wendt</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Van De Par</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Ewert</surname>
<given-names>S. D.</given-names>
</name>
</person-group> (<year>2014</year>). <article-title>A computationally-efficient and perceptually-plausible algorithm for binaural room impulse response simulation</article-title>. <source>J. Audio Eng. Soc.</source> <volume>62</volume>, <fpage>748</fpage>&#x2013;<lpage>766</lpage>. <pub-id pub-id-type="doi">10.17743/jaes.2014.0042</pub-id>
</citation>
</ref>
<ref id="B22">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Younes</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Honerkamp</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Welschehold</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Valada</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Catch me if you hear me: audio-visual navigation in complex unmapped environments with moving sounds</article-title>. <source>IEEE Robotics Automation Lett.</source> <volume>8</volume>, <fpage>928</fpage>&#x2013;<lpage>935</lpage>. <pub-id pub-id-type="doi">10.1109/lra.2023.3234766</pub-id>
</citation>
</ref>
<ref id="B23">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yu</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Han</surname>
<given-names>P.</given-names>
</name>
</person-group> (<year>2016</year>). <article-title>Localization based stereo speech source separation using probabilistic time-frequency masking and deep neural networks</article-title>. <source>EURASIP J. Audio, Speech, Music Process.</source> <volume>2016</volume>, <fpage>7</fpage>&#x2013;<lpage>18</lpage>. <pub-id pub-id-type="doi">10.1186/s13636-016-0085-x</pub-id>
</citation>
</ref>
</ref-list>
</back>
</article>