<?xml version="1.0"?>
<dblpperson name="Tomoki Toda" pid="85/741" n="584">
<person key="homepages/85/741" mdate="2025-01-10">
<author pid="85/741">Tomoki Toda</author>
<url>https://orcid.org/0000-0001-8146-1279</url>
<url>https://www.wikidata.org/entity/Q130982065</url>
</person>
<r><article key="journals/csl/YasudaT26" mdate="2025-10-14">
<author orcid="0000-0002-2130-747X" pid="228/9342">Yusuke Yasuda</author>
<author pid="85/741">Tomoki Toda</author>
<title>Automatic design optimization of preference-based subjective evaluation with online learning in crowdsourcing environment.</title>
<pages>101888</pages>
<year>2026</year>
<volume>96</volume>
<journal>Comput. Speech Lang.</journal>
<ee>https://doi.org/10.1016/j.csl.2025.101888</ee>
<url>db/journals/csl/csl96.html#YasudaT26</url>
<stream>streams/journals/csl</stream>
</article>
</r>
<r><article key="journals/csl/MiSMHFT26" mdate="2026-06-25">
<author orcid="0009-0007-8554-0987" pid="387/9787">Jinyi Mi</author>
<author pid="89/3530">Xiaohan Shi</author>
<author pid="30/6718">Ding Ma</author>
<author pid="205/5074">Jiajun He</author>
<author pid="216/3624">Takuya Fujimura</author>
<author pid="85/741">Tomoki Toda</author>
<title>Robust speech emotion recognition under human speech noise.</title>
<year>2026</year>
<pages>101987</pages>
<volume>100</volume>
<journal>Comput. Speech Lang.</journal>
<ee type="oa">https://doi.org/10.1016/j.csl.2026.101987</ee>
<url>db/journals/csl/csl100.html#MiSMHFT26</url>
<stream>streams/journals/csl</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2602-12723" mdate="2026-03-27">
<author pid="271/4266">Bence Mark Halpern</author>
<author pid="317/5237">Thomas Tienkamp</author>
<author pid="329/8124">Defne Abur</author>
<author pid="85/741">Tomoki Toda</author>
<title>Towards explainable reference-free speech intelligibility evaluation of people with pathological speech.</title>
<year>2026</year>
<month>February</month>
<volume>abs/2602.12723</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2602.12723</ee>
<url>db/journals/corr/corr2602.html#abs-2602-12723</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2603-08097" mdate="2026-04-09">
<author pid="271/4266">Bence Mark Halpern</author>
<author pid="317/5237">Thomas Tienkamp</author>
<author pid="329/8124">Defne Abur</author>
<author pid="85/741">Tomoki Toda</author>
<title>PathBench: Speech Intelligibility Benchmark for Automatic Pathological Speech Assessment.</title>
<year>2026</year>
<month>March</month>
<volume>abs/2603.08097</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2603.08097</ee>
<url>db/journals/corr/corr2603.html#abs-2603-08097</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2606-01905" mdate="2026-07-10">
<author pid="30/6718">Ding Ma</author>
<author pid="387/9787">Jinyi Mi</author>
<author pid="372/9374">Fengji Li</author>
<author pid="304/2864">Lester Phillip Violeta</author>
<author pid="205/5074">Jiajun He</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="85/741">Tomoki Toda</author>
<title>Advancing Electrolaryngeal Speech Enhancement Through Speech-Text Representation Learning.</title>
<year>2026</year>
<month>June</month>
<volume>abs/2606.01905</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2606.01905</ee>
<url>db/journals/corr/corr2606.html#abs-2606-01905</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2606-06200" mdate="2026-07-05">
<author pid="387/9787">Jinyi Mi</author>
<author pid="30/6718">Ding Ma</author>
<author pid="85/741">Tomoki Toda</author>
<title>Learning Emotion-discriminative Representations for Zero-Shot Cross-lingual Speech Emotion Recognition.</title>
<year>2026</year>
<month>June</month>
<volume>abs/2606.06200</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2606.06200</ee>
<url>db/journals/corr/corr2606.html#abs-2606-06200</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2606-19792" mdate="2026-07-08">
<author pid="332/1313">Masato Murata</author>
<author pid="27/9637">Koichi Miyazaki</author>
<author pid="44/9232">Tomoki Koriyama</author>
<author pid="85/741">Tomoki Toda</author>
<title>Exploring Pre-training Benefits on Phoneme Addition through Fine-tuning in Speech Synthesis.</title>
<year>2026</year>
<month>June</month>
<volume>abs/2606.19792</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2606.19792</ee>
<url>db/journals/corr/corr2606.html#abs-2606-19792</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2606-31105" mdate="2026-07-10">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="85/741">Tomoki Toda</author>
<title>Attacking UTMOS: Probing the Robustness of a Speech Quality Assessment Model.</title>
<year>2026</year>
<month>June</month>
<volume>abs/2606.31105</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2606.31105</ee>
<url>db/journals/corr/corr2606.html#abs-2606-31105</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article key="journals/access/EshghiT25" mdate="2025-06-11">
<author orcid="0000-0003-3878-6363" pid="42/3681">Mohammad Eshghi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Predicting Fundamental Frequency Patterns in Electrolaryngeal Speech Using Automated Phoneme Extraction.</title>
<pages>73831-73847</pages>
<year>2025</year>
<volume>13</volume>
<journal>IEEE Access</journal>
<ee type="oa">https://doi.org/10.1109/ACCESS.2025.3564648</ee>
<url>db/journals/access/access13.html#EshghiT25</url>
<stream>streams/journals/access</stream>
</article>
</r>
<r><article key="journals/access/OguraOOCTK25" mdate="2025-08-09">
<author orcid="0000-0002-4507-0008" pid="243/3880">Tadashi Ogura</author>
<author orcid="0000-0001-9913-4647" pid="132/9091">Takuma Okamoto</author>
<author orcid="0009-0009-2961-2821" pid="34/8763">Yamato Ohtani</author>
<author pid="03/7184">Erica Cooper</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author orcid="0000-0002-0914-5092" pid="32/4341">Hisashi Kawai</author>
<title>Phoneme-Level Duration Controllable Neural Text-to-Speech With Phoneme Embedding Skip Connection and Modified Gaussian Duration Modeling.</title>
<pages>118369-118380</pages>
<year>2025</year>
<volume>13</volume>
<journal>IEEE Access</journal>
<ee type="oa">https://doi.org/10.1109/ACCESS.2025.3585135</ee>
<url>db/journals/access/access13.html#OguraOOCTK25</url>
<stream>streams/journals/access</stream>
</article>
</r>
<r><article key="journals/access/YamashitaOTOTTK25" mdate="2026-01-09">
<author orcid="0009-0001-8553-9317" pid="368/3515">Haruki Yamashita</author>
<author orcid="0000-0001-9913-4647" pid="132/9091">Takuma Okamoto</author>
<author orcid="0000-0002-9808-0250" pid="22/8758">Ryoichi Takashima</author>
<author orcid="0009-0009-2961-2821" pid="34/8763">Yamato Ohtani</author>
<author orcid="0000-0001-5005-7679" pid="79/4485">Tetsuya Takiguchi</author>
<author pid="85/741">Tomoki Toda</author>
<author orcid="0000-0002-0914-5092" pid="32/4341">Hisashi Kawai</author>
<title>Sequence-to-Sequence Voice Conversion With Weighted Guided Attention.</title>
<pages>216583-216595</pages>
<year>2025</year>
<volume>13</volume>
<journal>IEEE Access</journal>
<ee type="oa">https://doi.org/10.1109/ACCESS.2025.3647153</ee>
<url>db/journals/access/access13.html#YamashitaOTOTTK25</url>
<stream>streams/journals/access</stream>
</article>
</r>
<r><article key="journals/csl/HuYT25" mdate="2025-05-01">
<author orcid="0000-0002-6112-9380" pid="156/0822">Cheng-Hung Hu</author>
<author pid="228/9342">Yusuke Yasuda</author>
<author pid="85/741">Tomoki Toda</author>
<title>E2EPref: An end-to-end preference-based framework for speech quality assessment to alleviate bias in direct assessment scores.</title>
<pages>101799</pages>
<year>2025</year>
<volume>93</volume>
<journal>Comput. Speech Lang.</journal>
<ee>https://doi.org/10.1016/j.csl.2025.101799</ee>
<url>db/journals/csl/csl93.html#HuYT25</url>
<stream>streams/journals/csl</stream>
</article>
</r>
<r><article key="journals/jstsp/HalpernTRSVWAT25" mdate="2026-01-01">
<author orcid="0000-0001-8787-359X" pid="271/4266">Bence Mark Halpern</author>
<author orcid="0000-0002-8374-6654" pid="317/5237">Thomas B. Tienkamp</author>
<author orcid="0000-0002-1986-8470" pid="317/5018">Teja Rebernik</author>
<author orcid="0000-0001-6321-7635" pid="34/4715">Rob J. J. H. van Son</author>
<author orcid="0000-0001-8964-5427" pid="329/8363">Sebastiaan A. H. J. de Visscher</author>
<author pid="227/2349">Max J. H. Witjes</author>
<author orcid="0000-0002-0410-8487" pid="329/8124">Defne Abur</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>XPPG-PCA: Reference-Free Automatic Speech Severity Evaluation With Principal Components.</title>
<pages>783-795</pages>
<year>2025</year>
<month>July</month>
<volume>19</volume>
<journal>IEEE J. Sel. Top. Signal Process.</journal>
<number>5</number>
<ee>https://doi.org/10.1109/JSTSP.2025.3617859</ee>
<url>db/journals/jstsp/jstsp19.html#HalpernTRSVWAT25</url>
<stream>streams/journals/jstsp</stream>
</article>
</r>
<r><article key="journals/jstsp/VioletaHMYKT25" mdate="2026-03-10">
<author orcid="0000-0001-9118-2994" pid="304/2864">Lester Phillip Violeta</author>
<author orcid="0000-0003-3172-3335" pid="225/7821">Wen-Chin Huang</author>
<author orcid="0009-0002-6564-4571" pid="30/6718">Ding Ma</author>
<author pid="81/3672">Ryuichi Yamamoto</author>
<author orcid="0000-0001-5047-4165" pid="70/8239">Kazuhiro Kobayashi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Resolving Domain Mismatches in Electrolaryngeal Speech Enhancement With Linguistic Intermediates.</title>
<pages>827-839</pages>
<year>2025</year>
<month>July</month>
<volume>19</volume>
<journal>IEEE J. Sel. Top. Signal Process.</journal>
<number>5</number>
<ee type="oa">https://doi.org/10.1109/JSTSP.2025.3584195</ee>
<url>db/journals/jstsp/jstsp19.html#VioletaHMYKT25</url>
<stream>streams/journals/jstsp</stream>
</article>
</r>
<r><inproceedings key="conf/apsipa/HattoriHTT25" mdate="2026-03-16">
<author pid="429/5419">Kimihiro Hattori</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="38/3483">Kazuya Takeda</author>
<author pid="85/741">Tomoki Toda</author>
<title>An Evaluation of Supervised Virtual Microphone Estimators in Reverberant Sound Fields.</title>
<pages>125-130</pages>
<year>2025</year>
<booktitle>APSIPA</booktitle>
<ee>https://doi.org/10.1109/APSIPAASC65261.2025.11249347</ee>
<crossref>conf/apsipa/2025</crossref>
<url>db/conf/apsipa/apsipa2025.html#HattoriHTT25</url>
</inproceedings>
</r>
<r><inproceedings key="conf/apsipa/MiyajiSHT25" mdate="2026-03-16">
<author pid="429/5349">Hikari Miyaji</author>
<author pid="429/5300">Keito Sawada</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="85/741">Tomoki Toda</author>
<title>Designing a Music Difficulty Measure for Controllable Automatic Piano Rearrangement.</title>
<pages>246-251</pages>
<year>2025</year>
<booktitle>APSIPA</booktitle>
<ee>https://doi.org/10.1109/APSIPAASC65261.2025.11249163</ee>
<crossref>conf/apsipa/2025</crossref>
<url>db/conf/apsipa/apsipa2025.html#MiyajiSHT25</url>
</inproceedings>
</r>
<r><inproceedings key="conf/apsipa/SawadaHT25" mdate="2026-03-16">
<author pid="429/5300">Keito Sawada</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="85/741">Tomoki Toda</author>
<title>Hierarchical Symbolic Music Generation with Variational Autoencoder-Based Bar-Wise Feature Sequences.</title>
<pages>299-304</pages>
<year>2025</year>
<booktitle>APSIPA</booktitle>
<ee>https://doi.org/10.1109/APSIPAASC65261.2025.11249414</ee>
<crossref>conf/apsipa/2025</crossref>
<url>db/conf/apsipa/apsipa2025.html#SawadaHT25</url>
</inproceedings>
</r>
<r><inproceedings key="conf/apsipa/KanekoHT25" mdate="2026-03-16">
<author pid="95/1951">Masataka Kaneko</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="85/741">Tomoki Toda</author>
<title>Estimating Speaker's Seating Position from Monaural Speech in a Simulated Vehicle Interior Sound Field.</title>
<pages>625-629</pages>
<year>2025</year>
<booktitle>APSIPA</booktitle>
<ee>https://doi.org/10.1109/APSIPAASC65261.2025.11249106</ee>
<crossref>conf/apsipa/2025</crossref>
<url>db/conf/apsipa/apsipa2025.html#KanekoHT25</url>
</inproceedings>
</r>
<r><inproceedings key="conf/apsipa/NiwaKT25" mdate="2026-03-16">
<author pid="429/5304">Kiseki Niwa</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="85/741">Tomoki Toda</author>
<title>Investigation of the Effectiveness of Converted Speech Auditory Feedback in Low-Latency Real-Time Voice Conversion.</title>
<pages>753-758</pages>
<year>2025</year>
<booktitle>APSIPA</booktitle>
<ee>https://doi.org/10.1109/APSIPAASC65261.2025.11249417</ee>
<crossref>conf/apsipa/2025</crossref>
<url>db/conf/apsipa/apsipa2025.html#NiwaKT25</url>
</inproceedings>
</r>
<r><inproceedings key="conf/apsipa/NakataYHT25" mdate="2026-03-16">
<author pid="429/5477">Yuuto Nakata</author>
<author pid="200/0472">Daiki Yoshioka</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="85/741">Tomoki Toda</author>
<title>Disfluency Disentanglement Enhancement in Spoken-Text-Style Transfer for Spontaneous Speech Synthesis.</title>
<pages>1098-1103</pages>
<year>2025</year>
<booktitle>APSIPA</booktitle>
<ee>https://doi.org/10.1109/APSIPAASC65261.2025.11248977</ee>
<crossref>conf/apsipa/2025</crossref>
<url>db/conf/apsipa/apsipa2025.html#NakataYHT25</url>
</inproceedings>
</r>
<r><inproceedings key="conf/apsipa/YoonT25" mdate="2026-03-16">
<author pid="429/5371">Dohyun Yoon</author>
<author pid="85/741">Tomoki Toda</author>
<title>Neural Semi-Fragile Watermarking for Proactive Deepfake Speech Detection.</title>
<pages>2092-2097</pages>
<year>2025</year>
<booktitle>APSIPA</booktitle>
<ee>https://doi.org/10.1109/APSIPAASC65261.2025.11249352</ee>
<crossref>conf/apsipa/2025</crossref>
<url>db/conf/apsipa/apsipa2025.html#YoonT25</url>
</inproceedings>
</r>
<r><inproceedings key="conf/apsipa/TangLCLTL25" mdate="2026-03-16">
<author pid="429/5417">Shaoqi Tang</author>
<author pid="284/4048">Zeyan Liu</author>
<author pid="88/1450">Liping Chen</author>
<author pid="35/4621">Kong Aik Lee</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="70/5210">Zhenhua Ling</author>
<title>A Preliminary Study on Sectional Voice Anonymization and Detection.</title>
<pages>2229-2234</pages>
<year>2025</year>
<booktitle>APSIPA</booktitle>
<ee>https://doi.org/10.1109/APSIPAASC65261.2025.11249094</ee>
<crossref>conf/apsipa/2025</crossref>
<url>db/conf/apsipa/apsipa2025.html#TangLCLTL25</url>
</inproceedings>
</r>
<r><inproceedings key="conf/apsipa/ChenLL0DT025" mdate="2026-05-07">
<author pid="88/1450">Liping Chen</author>
<author pid="35/4621">Kong-Aik Lee</author>
<author pid="70/5210">Zhen-Hua Ling</author>
<author orcid="0000-0001-8246-0606" pid="10/5630-37">Xin Wang 0037</author>
<author pid="158/4101">Rohan Kumar Das</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="36/4118">Haizhou Li 0001</author>
<title>Speaker Privacy and Security in the Big Data Era: Protection and Defense Against Deepfake.</title>
<pages>2570-2575</pages>
<year>2025</year>
<booktitle>APSIPA</booktitle>
<ee>https://doi.org/10.1109/APSIPAASC65261.2025.11249028</ee>
<crossref>conf/apsipa/2025</crossref>
<url>db/conf/apsipa/apsipa2025.html#ChenLL0DT025</url>
</inproceedings>
</r>
<r><inproceedings key="conf/asru/CooperOOTK25" mdate="2026-04-14">
<author pid="03/7184">Erica Cooper</author>
<author pid="132/9091">Takuma Okamoto</author>
<author pid="34/8763">Yamato Ohtani</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="32/4341">Hisashi Kawai</author>
<title>Layer-wise Analysis for Quality of Multilingual Synthesized Speech.</title>
<booktitle>ASRU</booktitle>
<year>2025</year>
<pages>1-7</pages>
<crossref>conf/asru/2025</crossref>
<ee>https://doi.org/10.1109/ASRU65441.2025.11434769</ee>
<url>db/conf/asru/asru2025.html#CooperOOTK25</url>
<stream>streams/conf/asru</stream>
</inproceedings>
</r>
<r><inproceedings key="conf/asru/HeSMT25" mdate="2026-04-14">
<author pid="205/5074">Jiajun He</author>
<author pid="70/2093">Naoki Sawada</author>
<author pid="27/9637">Koichi Miyazaki</author>
<author pid="85/741">Tomoki Toda</author>
<title>PARCO: Phoneme-Augmented Robust Contextual ASR via Contrastive Entity Disambiguation.</title>
<booktitle>ASRU</booktitle>
<year>2025</year>
<pages>1-7</pages>
<crossref>conf/asru/2025</crossref>
<ee>https://doi.org/10.1109/ASRU65441.2025.11434772</ee>
<url>db/conf/asru/asru2025.html#HeSMT25</url>
<stream>streams/conf/asru</stream>
</inproceedings>
</r>
<r><inproceedings key="conf/asru/HuangWLWTHCQT25" mdate="2026-04-14">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="39/721">Hui Wang</author>
<author pid="15/2288">Cheng Liu</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="166/6471">Andros Tjandra</author>
<author pid="160/9923">Wei-Ning Hsu</author>
<author pid="03/7184">Erica Cooper</author>
<author pid="20/4298">Yong Qin</author>
<author pid="85/741">Tomoki Toda</author>
<title>The AudioMOS Challenge 2025.</title>
<booktitle>ASRU</booktitle>
<year>2025</year>
<pages>1-8</pages>
<crossref>conf/asru/2025</crossref>
<ee>https://doi.org/10.1109/ASRU65441.2025.11434660</ee>
<url>db/conf/asru/asru2025.html#HuangWLWTHCQT25</url>
<stream>streams/conf/asru</stream>
</inproceedings>
</r>
<r><inproceedings key="conf/asru/OhtaniOTK25" mdate="2026-04-14">
<author pid="34/8763">Yamato Ohtani</author>
<author pid="132/9091">Takuma Okamoto</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="32/4341">Hisashi Kawai</author>
<title>Voice Factor Control Using FIR-Based Fast Neural Vocoder for Speech Generation Applications.</title>
<booktitle>ASRU</booktitle>
<year>2025</year>
<pages>1-4</pages>
<crossref>conf/asru/2025</crossref>
<ee>https://doi.org/10.1109/ASRU65441.2025.11434739</ee>
<url>db/conf/asru/asru2025.html#OhtaniOTK25</url>
<stream>streams/conf/asru</stream>
</inproceedings>
</r>
<r><inproceedings key="conf/eusipco/VioletaHT25" mdate="2026-01-06">
<author pid="304/2864">Lester Phillip Violeta</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="85/741">Tomoki Toda</author>
<title>Serenade: A Singing Style Conversion Framework Based on Audio Infilling.</title>
<pages>411-415</pages>
<year>2025</year>
<booktitle>EUSIPCO</booktitle>
<ee>https://ieeexplore.ieee.org/document/11226227</ee>
<crossref>conf/eusipco/2025</crossref>
<url>db/conf/eusipco/eusipco2025.html#VioletaHT25</url>
<stream>streams/conf/eusipco</stream>
</inproceedings>
</r>
<r><inproceedings key="conf/eusipco/OgitaYHT25" mdate="2026-01-06">
<author pid="422/2585">Kenichi Ogita</author>
<author pid="290/1755">Reo Yoneyama</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="85/741">Tomoki Toda</author>
<title>VAE-SiFiGAN: Source-Filter HiFi-GAN Based on Variational Autoencoder Representations with Enhanced Pitch Controllability.</title>
<pages>531-535</pages>
<year>2025</year>
<booktitle>EUSIPCO</booktitle>
<ee>https://ieeexplore.ieee.org/document/11226579</ee>
<crossref>conf/eusipco/2025</crossref>
<url>db/conf/eusipco/eusipco2025.html#OgitaYHT25</url>
<stream>streams/conf/eusipco</stream>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/FujimuraKT25" mdate="2025-07-03">
<author pid="216/3624">Takuya Fujimura</author>
<author pid="294/8499">Ibuki Kuroyanagi</author>
<author pid="85/741">Tomoki Toda</author>
<title>Improvements of Discriminative Feature Space Training for Anomalous Sound Detection in Unlabeled Conditions.</title>
<pages>1-5</pages>
<year>2025</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP49660.2025.10890020</ee>
<crossref>conf/icassp/2025</crossref>
<url>db/conf/icassp/icassp2025.html#FujimuraKT25</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/HashizumeT25" mdate="2025-07-02">
<author pid="334/0168">Yuka Hashizume</author>
<author pid="85/741">Tomoki Toda</author>
<title>Investigation of perceptual music similarity focusing on each instrumental part.</title>
<pages>1-5</pages>
<year>2025</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP49660.2025.10887810</ee>
<crossref>conf/icassp/2025</crossref>
<url>db/conf/icassp/icassp2025.html#HashizumeT25</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/NishizawaYHT25" mdate="2025-07-02">
<author pid="409/3900">Kaito Nishizawa</author>
<author pid="81/3672">Ryuichi Yamamoto</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="85/741">Tomoki Toda</author>
<title>Investigating Factors Related to the Naturalness of Synthesized Unison Singing.</title>
<pages>1-5</pages>
<year>2025</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP49660.2025.10889744</ee>
<crossref>conf/icassp/2025</crossref>
<url>db/conf/icassp/icassp2025.html#NishizawaYHT25</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/OguraOOCTK25" mdate="2025-07-02">
<author pid="243/3880">Tadashi Ogura</author>
<author pid="132/9091">Takuma Okamoto</author>
<author pid="34/8763">Yamato Ohtani</author>
<author pid="03/7184">Erica Cooper</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="32/4341">Hisashi Kawai</author>
<title>Mora-Level Prosody Prediction for Text-to-Speech Using Japanese BERT Without Accentual Labels.</title>
<pages>1-5</pages>
<year>2025</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP49660.2025.10887607</ee>
<crossref>conf/icassp/2025</crossref>
<url>db/conf/icassp/icassp2025.html#OguraOOCTK25</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/HalpernTRS0AT25" mdate="2025-12-07">
<author pid="271/4266">Bence Mark Halpern</author>
<author orcid="0000-0002-8374-6654" pid="317/5237">Thomas Tienkamp</author>
<author pid="317/5018">Teja Rebernik</author>
<author pid="34/4715">Rob J. J. H. van Son</author>
<author orcid="0000-0003-0434-1526" pid="35/2985">Martijn Wieling 0001</author>
<author pid="329/8124">Defne Abur</author>
<author pid="85/741">Tomoki Toda</author>
<title>Relationship between objective and subjective perceptual measures of speech in individuals with head and neck cancer.</title>
<year>2025</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2025-1127</ee>
<crossref>conf/interspeech/2025</crossref>
<url>db/conf/interspeech/interspeech2025.html#HalpernTRS0AT25</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/HeMT25" mdate="2025-11-20">
<author pid="205/5074">Jiajun He</author>
<author pid="387/9787">Jinyi Mi</author>
<author pid="85/741">Tomoki Toda</author>
<title>GIA-MIC: Multimodal Emotion Recognition with Gated Interactive Attention and Modality-Invariant Learning Constraints.</title>
<year>2025</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2025-2696</ee>
<crossref>conf/interspeech/2025</crossref>
<url>db/conf/interspeech/interspeech2025.html#HeMT25</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/HeSMT25" mdate="2025-11-20">
<author pid="205/5074">Jiajun He</author>
<author pid="70/2093">Naoki Sawada</author>
<author pid="27/9637">Koichi Miyazaki</author>
<author pid="85/741">Tomoki Toda</author>
<title>CMT-LLM: Contextual Multi-Talker ASR Utilizing Large Language Models.</title>
<year>2025</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2025-943</ee>
<crossref>conf/interspeech/2025</crossref>
<url>db/conf/interspeech/interspeech2025.html#HeSMT25</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/HuYYT25" mdate="2025-11-20">
<author pid="156/0822">Cheng-Hung Hu</author>
<author pid="228/9342">Yusuke Yasuda</author>
<author pid="162/7852">Akifumi Yoshimoto</author>
<author pid="85/741">Tomoki Toda</author>
<title>Unifying Listener Scoring Scales: Comparison Learning Framework for Speech Quality Assessment and Continuous Speech Emotion Recognition.</title>
<year>2025</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2025-1435</ee>
<crossref>conf/interspeech/2025</crossref>
<url>db/conf/interspeech/interspeech2025.html#HuYYT25</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/HuangCT25" mdate="2025-11-20">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="03/7184">Erica Cooper</author>
<author pid="85/741">Tomoki Toda</author>
<title>SHEET: A Multi-purpose Open-source Speech Human Evaluation Estimation Toolkit.</title>
<year>2025</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2025-1977</ee>
<crossref>conf/interspeech/2025</crossref>
<url>db/conf/interspeech/interspeech2025.html#HuangCT25</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/MurataMKT25" mdate="2025-11-20">
<author pid="332/1313">Masato Murata</author>
<author pid="27/9637">Koichi Miyazaki</author>
<author pid="44/9232">Tomoki Koriyama</author>
<author pid="85/741">Tomoki Toda</author>
<title>Eigenvoice Synthesis based on Model Editing for Speaker Generation.</title>
<year>2025</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2025-277</ee>
<crossref>conf/interspeech/2025</crossref>
<url>db/conf/interspeech/interspeech2025.html#MurataMKT25</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/OguraOOCTK25" mdate="2025-11-20">
<author pid="243/3880">Tadashi Ogura</author>
<author pid="132/9091">Takuma Okamoto</author>
<author pid="34/8763">Yamato Ohtani</author>
<author pid="03/7184">Erica Cooper</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="32/4341">Hisashi Kawai</author>
<title>GST-BERT-TTS: Prosody Prediction Without Accentual Labels For Multi-Speaker TTS Using BERT With Global Style Tokens.</title>
<year>2025</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2025-1098</ee>
<crossref>conf/interspeech/2025</crossref>
<url>db/conf/interspeech/interspeech2025.html#OguraOOCTK25</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/Shi0T25" mdate="2026-03-24">
<author orcid="0009-0001-3721-4964" pid="89/3530">Xiaohan Shi</author>
<author pid="63/9121-1">Xingfeng Li 0001</author>
<author pid="85/741">Tomoki Toda</author>
<title>Who, When, and What: Leveraging the &#34;Three Ws&#34; Concept for Emotion Recognition in Conversation.</title>
<year>2025</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2025-1433</ee>
<crossref>conf/interspeech/2025</crossref>
<url>db/conf/interspeech/interspeech2025.html#Shi0T25</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/Shi0T25a" mdate="2026-03-24">
<author orcid="0009-0001-3721-4964" pid="89/3530">Xiaohan Shi</author>
<author pid="63/9121-1">Xingfeng Li 0001</author>
<author pid="85/741">Tomoki Toda</author>
<title>Speaker-Aware Multi-Task Learning for Speech Emotion Recognition.</title>
<year>2025</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2025-1439</ee>
<crossref>conf/interspeech/2025</crossref>
<url>db/conf/interspeech/interspeech2025.html#Shi0T25a</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/ShiM0T25" mdate="2026-03-24">
<author orcid="0009-0001-3721-4964" pid="89/3530">Xiaohan Shi</author>
<author pid="387/9787">Jinyi Mi</author>
<author pid="63/9121-1">Xingfeng Li 0001</author>
<author pid="85/741">Tomoki Toda</author>
<title>Advancing Emotion Recognition via Ensemble Learning: Integrating Speech, Context, and Text Representations.</title>
<year>2025</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2025-1445</ee>
<crossref>conf/interspeech/2025</crossref>
<url>db/conf/interspeech/interspeech2025.html#ShiM0T25</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/YoneyamaKTYT25" mdate="2025-11-20">
<author pid="290/1755">Reo Yoneyama</author>
<author pid="192/7186">Masaya Kawamura</author>
<author pid="67/8451">Ryo Terashima</author>
<author pid="81/3672">Ryuichi Yamamoto</author>
<author pid="85/741">Tomoki Toda</author>
<title>Comparative Analysis of Fast and High-Fidelity Neural Vocoders for Low-Latency Streaming Synthesis in Resource-Constrained Environments.</title>
<year>2025</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2025-1819</ee>
<crossref>conf/interspeech/2025</crossref>
<url>db/conf/interspeech/interspeech2025.html#YoneyamaKTYT25</url>
</inproceedings>
</r>
<r><article publtype="informal" key="journals/corr/abs-2502-02138" mdate="2025-03-10">
<author pid="334/0168">Yuka Hashizume</author>
<author pid="85/741">Tomoki Toda</author>
<title>Investigation of perceptual music similarity focusing on each instrumental part.</title>
<year>2025</year>
<month>February</month>
<volume>abs/2502.02138</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2502.02138</ee>
<url>db/journals/corr/corr2502.html#abs-2502-02138</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2503-10435" mdate="2025-05-01">
<author orcid="0000-0003-4200-9129" pid="207/9559">Kevin Wilkinghoff</author>
<author pid="216/3624">Takuya Fujimura</author>
<author pid="140/2806">Keisuke Imoto</author>
<author pid="36/4575">Jonathan Le Roux</author>
<author orcid="0000-0001-6856-8928" pid="39/4898">Zheng-Hua Tan</author>
<author pid="85/741">Tomoki Toda</author>
<title>Handling Domain Shifts for Anomalous Sound Detection: A Review of DCASE-Related Work.</title>
<year>2025</year>
<month>March</month>
<volume>abs/2503.10435</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2503.10435</ee>
<url>db/journals/corr/corr2503.html#abs-2503-10435</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2503-12388" mdate="2025-04-13">
<author pid="304/2864">Lester Phillip Violeta</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="85/741">Tomoki Toda</author>
<title>Serenade: A Singing Style Conversion Framework Based On Audio Infilling.</title>
<year>2025</year>
<month>March</month>
<volume>abs/2503.12388</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2503.12388</ee>
<url>db/journals/corr/corr2503.html#abs-2503-12388</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2503-17281" mdate="2025-04-15">
<author pid="334/0168">Yuka Hashizume</author>
<author pid="53/2189-63">Li Li 0063</author>
<author pid="357/2546">Atsushi Miyashita</author>
<author pid="85/741">Tomoki Toda</author>
<title>Learning disentangled representations for instrument-based music similarity.</title>
<year>2025</year>
<month>March</month>
<volume>abs/2503.17281</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2503.17281</ee>
<url>db/journals/corr/corr2503.html#abs-2503-17281</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2503-18486" mdate="2025-04-19">
<author pid="398/1356">Takehiro Imamura</author>
<author pid="334/0168">Yuka Hashizume</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="85/741">Tomoki Toda</author>
<title>Music Similarity Representation Learning Focusing on Individual Instruments with Source Separation and Human Preference.</title>
<year>2025</year>
<month>March</month>
<volume>abs/2503.18486</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2503.18486</ee>
<url>db/journals/corr/corr2503.html#abs-2503-18486</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2505-15061" mdate="2025-06-25">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="03/7184">Erica Cooper</author>
<author pid="85/741">Tomoki Toda</author>
<title>SHEET: A Multi-purpose Open-source Speech Human Evaluation Estimation Toolkit.</title>
<year>2025</year>
<month>May</month>
<volume>abs/2505.15061</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2505.15061</ee>
<url>db/journals/corr/corr2505.html#abs-2505-15061</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2505-18980" mdate="2025-06-26">
<author pid="294/8499">Ibuki Kuroyanagi</author>
<author pid="216/3624">Takuya Fujimura</author>
<author pid="38/3483">Kazuya Takeda</author>
<author pid="85/741">Tomoki Toda</author>
<title>Improving Anomalous Sound Detection through Pseudo-anomalous Set Selection and Pseudo-label Utilization under Unlabeled Conditions.</title>
<year>2025</year>
<month>May</month>
<volume>abs/2505.18980</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2505.18980</ee>
<url>db/journals/corr/corr2505.html#abs-2505-18980</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2505-18982" mdate="2025-06-26">
<author pid="294/8499">Ibuki Kuroyanagi</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="38/3483">Kazuya Takeda</author>
<author pid="85/741">Tomoki Toda</author>
<title>Serial-OE: Anomalous sound detection based on serial method with outlier exposure capable of using small amounts of anomalous data for training.</title>
<year>2025</year>
<month>May</month>
<volume>abs/2505.18982</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2505.18982</ee>
<url>db/journals/corr/corr2505.html#abs-2505-18982</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2506-00865" mdate="2025-07-06">
<author pid="205/5074">Jiajun He</author>
<author pid="387/9787">Jinyi Mi</author>
<author pid="85/741">Tomoki Toda</author>
<title>GIA-MIC: Multimodal Emotion Recognition with Gated Interactive Attention and Modality-Invariant Learning Constraints.</title>
<year>2025</year>
<month>June</month>
<volume>abs/2506.00865</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2506.00865</ee>
<url>db/journals/corr/corr2506.html#abs-2506-00865</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2506-03554" mdate="2025-07-06">
<author pid="290/1755">Reo Yoneyama</author>
<author pid="192/7186">Masaya Kawamura</author>
<author pid="67/8451">Ryo Terashima</author>
<author pid="81/3672">Ryuichi Yamamoto</author>
<author pid="85/741">Tomoki Toda</author>
<title>Comparative Analysis of Fast and High-Fidelity Neural Vocoders for Low-Latency Streaming Synthesis in Resource-Constrained Environments.</title>
<year>2025</year>
<month>June</month>
<volume>abs/2506.03554</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2506.03554</ee>
<url>db/journals/corr/corr2506.html#abs-2506-03554</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2506-11064" mdate="2025-07-14">
<author pid="205/5074">Jiajun He</author>
<author pid="85/741">Tomoki Toda</author>
<title>PMF-CEC: Phoneme-augmented Multimodal Fusion for Context-aware ASR Error Correction with Error-specific Selective Decoding.</title>
<year>2025</year>
<month>June</month>
<volume>abs/2506.11064</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2506.11064</ee>
<url>db/journals/corr/corr2506.html#abs-2506-11064</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2506-12059" mdate="2025-07-14">
<author pid="205/5074">Jiajun He</author>
<author pid="70/2093">Naoki Sawada</author>
<author pid="27/9637">Koichi Miyazaki</author>
<author pid="85/741">Tomoki Toda</author>
<title>CMT-LLM: Contextual Multi-Talker ASR Utilizing Large Language Models.</title>
<year>2025</year>
<month>June</month>
<volume>abs/2506.12059</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2506.12059</ee>
<url>db/journals/corr/corr2506.html#abs-2506-12059</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2507-01611" mdate="2025-08-22">
<author pid="342/6470">Shaowen Chen</author>
<author pid="85/741">Tomoki Toda</author>
<title>QHARMA-GAN: Quasi-Harmonic Neural Vocoder based on Autoregressive Moving Average Model.</title>
<year>2025</year>
<month>July</month>
<volume>abs/2507.01611</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2507.01611</ee>
<url>db/journals/corr/corr2507.html#abs-2507-01611</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2507-03377" mdate="2025-08-10">
<author pid="332/1313">Masato Murata</author>
<author pid="27/9637">Koichi Miyazaki</author>
<author pid="44/9232">Tomoki Koriyama</author>
<author pid="85/741">Tomoki Toda</author>
<title>Eigenvoice Synthesis based on Model Editing for Speaker Generation.</title>
<year>2025</year>
<month>July</month>
<volume>abs/2507.03377</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2507.03377</ee>
<url>db/journals/corr/corr2507.html#abs-2507-03377</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2507-13626" mdate="2025-08-22">
<author pid="156/0822">Cheng-Hung Hu</author>
<author pid="228/9342">Yusuke Yasuda</author>
<author pid="162/7852">Akifumi Yoshimoto</author>
<author pid="85/741">Tomoki Toda</author>
<title>Unifying Listener Scoring Scales: Comparison Learning Framework for Speech Quality Assessment and Continuous Speech Emotion Recognition.</title>
<year>2025</year>
<month>July</month>
<volume>abs/2507.13626</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2507.13626</ee>
<url>db/journals/corr/corr2507.html#abs-2507-13626</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2507-21426" mdate="2026-02-01">
<author orcid="0000-0001-8787-359X" pid="271/4266">Bence Mark Halpern</author>
<author pid="317/5237">Thomas Tienkamp</author>
<author pid="317/5018">Teja Rebernik</author>
<author pid="34/4715">Rob J. J. H. van Son</author>
<author pid="35/2985">Martijn Wieling 0001</author>
<author orcid="0000-0002-0410-8487" pid="329/8124">Defne Abur</author>
<author pid="85/741">Tomoki Toda</author>
<title>Relationship between objective and subjective perceptual measures of speech in individuals with head and neck cancer.</title>
<year>2025</year>
<month>July</month>
<volume>abs/2507.21426</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2507.21426</ee>
<url>db/journals/corr/corr2507.html#abs-2507-21426</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2509-01336" mdate="2025-10-08">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="39/721">Hui Wang</author>
<author pid="15/2288">Cheng Liu</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="166/6471">Andros Tjandra</author>
<author pid="160/9923">Wei-Ning Hsu</author>
<author pid="03/7184">Erica Cooper</author>
<author pid="20/4298">Yong Qin</author>
<author pid="85/741">Tomoki Toda</author>
<title>The AudioMOS Challenge 2025.</title>
<year>2025</year>
<month>September</month>
<volume>abs/2509.01336</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2509.01336</ee>
<url>db/journals/corr/corr2509.html#abs-2509-01336</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2509-04357" mdate="2025-10-12">
<author pid="205/5074">Jiajun He</author>
<author pid="70/2093">Naoki Sawada</author>
<author pid="27/9637">Koichi Miyazaki</author>
<author pid="85/741">Tomoki Toda</author>
<title>PARCO: Phoneme-Augmented Robust Contextual ASR via Contrastive Entity Disambiguation.</title>
<year>2025</year>
<month>September</month>
<volume>abs/2509.04357</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2509.04357</ee>
<url>db/journals/corr/corr2509.html#abs-2509-04357</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2509-04830" mdate="2025-10-21">
<author pid="03/7184">Erica Cooper</author>
<author pid="132/9091">Takuma Okamoto</author>
<author pid="34/8763">Yamato Ohtani</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="32/4341">Hisashi Kawai</author>
<title>Layer-wise Analysis for Quality of Multilingual Synthesized Speech.</title>
<year>2025</year>
<month>September</month>
<volume>abs/2509.04830</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2509.04830</ee>
<url>db/journals/corr/corr2509.html#abs-2509-04830</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2509-15629" mdate="2025-10-18">
<author pid="304/2864">Lester Phillip Violeta</author>
<author pid="237/9484">Xueyao Zhang</author>
<author pid="229/3529">Jiatong Shi</author>
<author pid="228/9342">Yusuke Yasuda</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="29/8054-1">Zhizheng Wu 0001</author>
<author pid="85/741">Tomoki Toda</author>
<title>The Singing Voice Conversion Challenge 2025: From Singer Identity Conversion To Singing Style Conversion.</title>
<year>2025</year>
<month>September</month>
<volume>abs/2509.15629</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2509.15629</ee>
<url>db/journals/corr/corr2509.html#abs-2509-15629</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2509-18706" mdate="2025-10-18">
<author pid="205/5074">Jiajun He</author>
<author pid="89/3530">Xiaohan Shi</author>
<author pid="156/0822">Cheng-Hung Hu</author>
<author pid="387/9787">Jinyi Mi</author>
<author pid="63/9121-1">Xingfeng Li 0001</author>
<author pid="85/741">Tomoki Toda</author>
<title>M4SER: Multimodal, Multirepresentation, Multitask, and Multistrategy Learning for Speech Emotion Recognition.</title>
<year>2025</year>
<month>September</month>
<volume>abs/2509.18706</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2509.18706</ee>
<url>db/journals/corr/corr2509.html#abs-2509-18706</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2510-00639" mdate="2025-11-08">
<author pid="271/4266">Bence Mark Halpern</author>
<author pid="85/741">Tomoki Toda</author>
<title>Reference-free automatic speech severity evaluation using acoustic unit language modelling.</title>
<year>2025</year>
<month>October</month>
<volume>abs/2510.00639</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2510.00639</ee>
<url>db/journals/corr/corr2510.html#abs-2510-00639</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2510-00657" mdate="2025-11-08">
<author pid="271/4266">Bence Mark Halpern</author>
<author pid="317/5237">Thomas B. Tienkamp</author>
<author pid="317/5018">Teja Rebernik</author>
<author pid="34/4715">Rob J. J. H. van Son</author>
<author pid="329/8363">Sebastiaan A. H. J. de Visscher</author>
<author pid="227/2349">Max J. H. Witjes</author>
<author pid="329/8124">Defne Abur</author>
<author pid="85/741">Tomoki Toda</author>
<title>XPPG-PCA: Reference-free automatic speech severity evaluation with principal components.</title>
<year>2025</year>
<month>October</month>
<volume>abs/2510.00657</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2510.00657</ee>
<url>db/journals/corr/corr2510.html#abs-2510-00657</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2511-20106" mdate="2026-03-24">
<author pid="63/9121-1">Xingfeng Li 0001</author>
<author orcid="0009-0001-3721-4964" pid="89/3530">Xiaohan Shi</author>
<author pid="83/5144">Junjie Li</author>
<author pid="172/4495">Yongwei Li</author>
<author pid="73/1580">Masashi Unoki</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="18/2260">Masato Akagi</author>
<title>EM2LDL: A Multilingual Speech Corpus for Mixed Emotion Recognition through Label Distribution Learning.</title>
<year>2025</year>
<month>November</month>
<volume>abs/2511.20106</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2511.20106</ee>
<url>db/journals/corr/corr2511.html#abs-2511-20106</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article key="journals/access/YamashitaOTOTTK24" mdate="2024-10-03">
<author orcid="0009-0001-8553-9317" pid="368/3515">Haruki Yamashita</author>
<author orcid="0000-0001-9913-4647" pid="132/9091">Takuma Okamoto</author>
<author orcid="0000-0002-9808-0250" pid="22/8758">Ryoichi Takashima</author>
<author orcid="0009-0009-2961-2821" pid="34/8763">Yamato Ohtani</author>
<author orcid="0000-0001-5005-7679" pid="79/4485">Tetsuya Takiguchi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author orcid="0000-0002-0914-5092" pid="32/4341">Hisashi Kawai</author>
<title>Fast Neural Speech Waveform Generative Models With Fully-Connected Layer-Based Upsampling.</title>
<pages>31409-31421</pages>
<year>2024</year>
<volume>12</volume>
<journal>IEEE Access</journal>
<ee type="oa">https://doi.org/10.1109/ACCESS.2024.3366707</ee>
<url>db/journals/access/access12.html#YamashitaOTOTTK24</url>
</article>
</r>
<r><article key="journals/access/EshghiT24" mdate="2024-05-04">
<author orcid="0000-0003-3878-6363" pid="42/3681">Mohammad Eshghi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>An Investigation of Fundamental Frequency Pattern Prediction for Japanese Electrolaryngeal Speech Enhancement Based on Frame-Wise Phoneme Representations.</title>
<pages>50137-50153</pages>
<year>2024</year>
<volume>12</volume>
<journal>IEEE Access</journal>
<ee type="oa">https://doi.org/10.1109/ACCESS.2024.3384973</ee>
<url>db/journals/access/access12.html#EshghiT24</url>
</article>
</r>
<r><article key="journals/spl/HuangWT24" mdate="2024-12-09">
<author orcid="0000-0003-3172-3335" pid="225/7821">Wen-Chin Huang</author>
<author orcid="0000-0003-4390-1354" pid="188/5943">Yi-Chiao Wu</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Multi-Speaker Text-to-Speech Training With Speaker Anonymized Data.</title>
<pages>2995-2999</pages>
<year>2024</year>
<volume>31</volume>
<journal>IEEE Signal Process. Lett.</journal>
<ee type="oa">https://doi.org/10.1109/LSP.2024.3482701</ee>
<url>db/journals/spl/spl31.html#HuangWT24</url>
<stream>streams/journals/spl</stream>
</article>
</r>
<r><article key="journals/taslp/WangLT24" mdate="2026-03-11">
<author orcid="0009-0003-0770-9936" pid="06/2293-169">Rui Wang 0169</author>
<author orcid="0000-0002-3121-7857" pid="53/2189-63">Li Li 0063</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Dual-Channel Target Speaker Extraction Based on Conditional Variational Autoencoder and Directional Information.</title>
<pages>1968-1979</pages>
<year>2024</year>
<volume>32</volume>
<journal>IEEE ACM Trans. Audio Speech Lang. Process.</journal>
<ee type="oa">https://doi.org/10.1109/TASLP.2024.3376154</ee>
<url>db/journals/taslp/taslp32.html#WangLT24</url>
</article>
</r>
<r><article key="journals/taslp/VioletaMHT24" mdate="2024-06-18">
<author orcid="0000-0001-9118-2994" pid="304/2864">Lester Phillip Violeta</author>
<author orcid="0009-0002-6564-4571" pid="30/6718">Ding Ma</author>
<author orcid="0000-0003-3172-3335" pid="225/7821">Wen-Chin Huang</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Pretraining and Adaptation Techniques for Electrolaryngeal Speech Recognition.</title>
<pages>2777-2789</pages>
<year>2024</year>
<volume>32</volume>
<journal>IEEE ACM Trans. Audio Speech Lang. Process.</journal>
<ee type="oa">https://doi.org/10.1109/TASLP.2024.3402557</ee>
<url>db/journals/taslp/taslp32.html#VioletaMHT24</url>
</article>
</r>
<r><article key="journals/taslp/LuanWT24" mdate="2024-12-09">
<author orcid="0009-0005-5489-0787" pid="331/6661">Shuming Luan</author>
<author orcid="0000-0001-8846-3699" pid="158/4215">Yukoh Wakabayashi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Unequally Spaced Sound Field Interpolation for Rotation-Robust Beamforming.</title>
<pages>3185-3199</pages>
<year>2024</year>
<volume>32</volume>
<journal>IEEE ACM Trans. Audio Speech Lang. Process.</journal>
<ee type="oa">https://doi.org/10.1109/TASLP.2024.3410879</ee>
<url>db/journals/taslp/taslp32.html#LuanWT24</url>
</article>
</r>
<r><inproceedings key="conf/apsipa/ImamuraHT24" mdate="2025-02-26">
<author pid="398/1356">Takehiro Imamura</author>
<author pid="334/0168">Yuka Hashizume</author>
<author pid="85/741">Tomoki Toda</author>
<title>Multi-Task Learning Approaches for Music Similarity Representation Learning Based on Individual Instrument Sounds.</title>
<pages>1-6</pages>
<year>2024</year>
<booktitle>APSIPA</booktitle>
<ee>https://doi.org/10.1109/APSIPAASC63619.2025.10849029</ee>
<crossref>conf/apsipa/2024</crossref>
<url>db/conf/apsipa/apsipa2024.html#ImamuraHT24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/apsipa/MiKT24" mdate="2025-02-26">
<author pid="387/9787">Jinyi Mi</author>
<author pid="47/4994">Sehun Kim</author>
<author pid="85/741">Tomoki Toda</author>
<title>Improved Architecture for High-resolution Piano Transcription to Efficiently Capture Acoustic Characteristics of Music Signals.</title>
<pages>1-6</pages>
<year>2024</year>
<booktitle>APSIPA</booktitle>
<ee>https://doi.org/10.1109/APSIPAASC63619.2025.10848780</ee>
<crossref>conf/apsipa/2024</crossref>
<url>db/conf/apsipa/apsipa2024.html#MiKT24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/apsipa/MiSMHFT24" mdate="2026-03-24">
<author pid="387/9787">Jinyi Mi</author>
<author orcid="0009-0001-3721-4964" pid="89/3530">Xiaohan Shi</author>
<author pid="30/6718">Ding Ma</author>
<author pid="205/5074">Jiajun He</author>
<author pid="216/3624">Takuya Fujimura</author>
<author pid="85/741">Tomoki Toda</author>
<title>Two-stage Framework for Robust Speech Emotion Recognition Using Target Speaker Extraction in Human Speech Noise Conditions.</title>
<pages>1-6</pages>
<year>2024</year>
<booktitle>APSIPA</booktitle>
<ee>https://doi.org/10.1109/APSIPAASC63619.2025.10848943</ee>
<crossref>conf/apsipa/2024</crossref>
<url>db/conf/apsipa/apsipa2024.html#MiSMHFT24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/apsipa/ShiGHMLT24" mdate="2026-03-24">
<author orcid="0009-0001-3721-4964" pid="89/3530">Xiaohan Shi</author>
<author pid="76/2452-40">Yuan Gao 0040</author>
<author pid="205/5074">Jiajun He</author>
<author pid="387/9787">Jinyi Mi</author>
<author orcid="0000-0002-8958-0341" pid="63/9121-1">Xingfeng Li 0001</author>
<author pid="85/741">Tomoki Toda</author>
<title>A Study on Multimodal Fusion and Layer Adapter in Emotion Recognition.</title>
<pages>1-6</pages>
<year>2024</year>
<booktitle>APSIPA</booktitle>
<ee>https://doi.org/10.1109/APSIPAASC63619.2025.10848773</ee>
<crossref>conf/apsipa/2024</crossref>
<url>db/conf/apsipa/apsipa2024.html#ShiGHMLT24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/apsipa/YangHT24" mdate="2025-02-26">
<author pid="176/5718">Zekun Yang</author>
<author pid="205/5074">Jiajun He</author>
<author pid="85/741">Tomoki Toda</author>
<title>Multi-Modal Video Summarization Based on Two-Stage Fusion of Audio, Visual, and Recognized Text Information.</title>
<pages>1-6</pages>
<year>2024</year>
<booktitle>APSIPA</booktitle>
<ee>https://doi.org/10.1109/APSIPAASC63619.2025.10849046</ee>
<crossref>conf/apsipa/2024</crossref>
<url>db/conf/apsipa/apsipa2024.html#YangHT24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/embc/LiSMZZW0LCTN24" mdate="2025-08-06">
<author pid="372/9374">Fengji Li</author>
<author pid="99/6100">Fei Shen</author>
<author pid="30/6718">Ding Ma</author>
<author pid="313/2304">Shaochuan Zhang</author>
<author pid="00/5012">Jie Zhou</author>
<author pid="58/6810-111">Li Wang 0111</author>
<author pid="20/4226-2">Fan Fan 0002</author>
<author pid="43/656">Tao Liu</author>
<author pid="02/1438">Xiaohong Chen</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="36/2169">Haijun Niu</author>
<title>Mandarin Speech Reconstruction from Tongue Motion Ultrasound Images based on Generative Adversarial Networks.</title>
<pages>1-4</pages>
<year>2024</year>
<booktitle>EMBC</booktitle>
<ee>https://doi.org/10.1109/EMBC53108.2024.10781847</ee>
<crossref>conf/embc/2024</crossref>
<url>db/conf/embc/embc2024.html#LiSMZZW0LCTN24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/embc/MaCLXKT24" mdate="2025-01-08">
<author pid="30/6718">Ding Ma</author>
<author pid="323/5236">Yeonjong Choi</author>
<author pid="372/9374">Fengji Li</author>
<author pid="04/5576">Chao Xie</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="85/741">Tomoki Toda</author>
<title>Robust Sequence-to-sequence Voice Conversion for Electrolaryngeal Speech Enhancement in Noisy and Reverberant Conditions.</title>
<pages>1-4</pages>
<year>2024</year>
<booktitle>EMBC</booktitle>
<ee>https://doi.org/10.1109/EMBC53108.2024.10781979</ee>
<crossref>conf/embc/2024</crossref>
<url>db/conf/embc/embc2024.html#MaCLXKT24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/eusipco/FujimuraIT24" mdate="2024-11-06">
<author pid="216/3624">Takuya Fujimura</author>
<author pid="140/2806">Keisuke Imoto</author>
<author pid="85/741">Tomoki Toda</author>
<title>Discriminative Neighborhood Smoothing for Generative Anomalous Sound Detection.</title>
<pages>156-160</pages>
<year>2024</year>
<booktitle>EUSIPCO</booktitle>
<ee>https://ieeexplore.ieee.org/document/10715201</ee>
<crossref>conf/eusipco/2024</crossref>
<url>db/conf/eusipco/eusipco2024.html#FujimuraIT24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/eusipco/WangT24" mdate="2024-11-06">
<author pid="145/6290">Jiachen Wang</author>
<author pid="85/741">Tomoki Toda</author>
<title>Unsupervised Training of Neural Network-Based Virtual Microphone Estimator.</title>
<pages>256-260</pages>
<year>2024</year>
<booktitle>EUSIPCO</booktitle>
<ee>https://ieeexplore.ieee.org/document/10715335</ee>
<crossref>conf/eusipco/2024</crossref>
<url>db/conf/eusipco/eusipco2024.html#WangT24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/KomatsuFTT24" mdate="2024-08-05">
<author pid="136/5113">Tatsuya Komatsu</author>
<author pid="50/6352">Yusuke Fujita</author>
<author pid="38/3483">Kazuya Takeda</author>
<author pid="85/741">Tomoki Toda</author>
<title>Audio Difference Learning for Audio Captioning.</title>
<pages>1456-1460</pages>
<year>2024</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP48485.2024.10448085</ee>
<crossref>conf/icassp/2024</crossref>
<url>db/conf/icassp/icassp2024.html#KomatsuFTT24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/OhtaniOTK24" mdate="2024-08-06">
<author pid="34/8763">Yamato Ohtani</author>
<author pid="132/9091">Takuma Okamoto</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="32/4341">Hisashi Kawai</author>
<title>FIRNet: Fundamental Frequency Controllable Fast Neural Vocoder With Trainable Finite Impulse Response Filter.</title>
<pages>10871-10875</pages>
<year>2024</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP48485.2024.10446960</ee>
<crossref>conf/icassp/2024</crossref>
<url>db/conf/icassp/icassp2024.html#OhtaniOTK24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/VioletaHMYKT24" mdate="2024-08-07">
<author pid="304/2864">Lester Phillip Violeta</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="30/6718">Ding Ma</author>
<author pid="81/3672">Ryuichi Yamamoto</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="85/741">Tomoki Toda</author>
<title>Electrolaryngeal Speech Intelligibility Enhancement through Robust Linguistic Encoders.</title>
<pages>10961-10965</pages>
<year>2024</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP48485.2024.10447197</ee>
<crossref>conf/icassp/2024</crossref>
<url>db/conf/icassp/icassp2024.html#VioletaHMYKT24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/HeS0T24" mdate="2026-03-24">
<author pid="205/5074">Jiajun He</author>
<author orcid="0009-0001-3721-4964" pid="89/3530">Xiaohan Shi</author>
<author orcid="0000-0002-8958-0341" pid="63/9121-1">Xingfeng Li 0001</author>
<author pid="85/741">Tomoki Toda</author>
<title>MF-AED-AEC: Speech Emotion Recognition by Leveraging Multimodal Fusion, Asr Error Detection, and Asr Error Correction.</title>
<pages>11066-11070</pages>
<year>2024</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP48485.2024.10446548</ee>
<crossref>conf/icassp/2024</crossref>
<url>db/conf/icassp/icassp2024.html#HeS0T24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/OkamotoOTK24" mdate="2024-08-07">
<author pid="132/9091">Takuma Okamoto</author>
<author pid="34/8763">Yamato Ohtani</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="32/4341">Hisashi Kawai</author>
<title>Convnext-TTS And Convnext-VC: Convnext-Based Fast End-To-End Sequence-To-Sequence Text-To-Speech And Voice Conversion.</title>
<pages>12456-12460</pages>
<year>2024</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP48485.2024.10446890</ee>
<crossref>conf/icassp/2024</crossref>
<url>db/conf/icassp/icassp2024.html#OkamotoOTK24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/ChenT24" mdate="2025-05-20">
<author pid="342/6470">Shaowen Chen</author>
<author pid="85/741">Tomoki Toda</author>
<title>QHM-GAN: Neural Vocoder based on Quasi-Harmonic Modeling.</title>
<year>2024</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2024-2371</ee>
<crossref>conf/interspeech/2024</crossref>
<url>db/conf/interspeech/interspeech2024.html#ChenT24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/FengYT24" mdate="2025-05-20">
<author pid="174/4737">Jingyi Feng</author>
<author pid="228/9342">Yusuke Yasuda</author>
<author pid="85/741">Tomoki Toda</author>
<title>Exploring the Robustness of Text-to-Speech Synthesis Based on Diffusion Probabilistic Models to Heavily Noisy Transcriptions.</title>
<year>2024</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2024-2337</ee>
<crossref>conf/interspeech/2024</crossref>
<url>db/conf/interspeech/interspeech2024.html#FengYT24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/HalpernTHVRVW0A24" mdate="2026-02-01">
<author orcid="0000-0001-8787-359X" pid="271/4266">Bence Mark Halpern</author>
<author orcid="0000-0002-8374-6654" pid="317/5237">Thomas Tienkamp</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="304/2864">Lester Phillip Violeta</author>
<author orcid="0000-0002-1986-8470" pid="317/5018">Teja Rebernik</author>
<author orcid="0000-0001-8964-5427" pid="329/8363">Sebastiaan A. H. J. de Visscher</author>
<author pid="227/2349">Max J. H. Witjes</author>
<author orcid="0000-0003-0434-1526" pid="35/2985">Martijn Wieling 0001</author>
<author orcid="0000-0002-0410-8487" pid="329/8124">Defne Abur</author>
<author pid="85/741">Tomoki Toda</author>
<title>Quantifying the effect of speech pathology on automatic and human speaker verification.</title>
<year>2024</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2024-1400</ee>
<crossref>conf/interspeech/2024</crossref>
<url>db/conf/interspeech/interspeech2024.html#HalpernTHVRVW0A24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/HeT24" mdate="2025-05-20">
<author pid="205/5074">Jiajun He</author>
<author pid="85/741">Tomoki Toda</author>
<title>2DP-2MRC: 2-Dimensional Pointer-based Machine Reading Comprehension Method for Multimodal Moment Retrieval.</title>
<year>2024</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2024-1633</ee>
<crossref>conf/interspeech/2024</crossref>
<url>db/conf/interspeech/interspeech2024.html#HeT24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/HuYT24" mdate="2025-05-20">
<author pid="156/0822">Cheng-Hung Hu</author>
<author pid="228/9342">Yusuke Yasuda</author>
<author pid="85/741">Tomoki Toda</author>
<title>Embedding Learning for Preference-based Speech Quality Assessment.</title>
<year>2024</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2024-1243</ee>
<crossref>conf/interspeech/2024</crossref>
<url>db/conf/interspeech/interspeech2024.html#HuYT24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/OkamotoOSTK24" mdate="2025-05-20">
<author pid="132/9091">Takuma Okamoto</author>
<author pid="34/8763">Yamato Ohtani</author>
<author pid="14/5586">Sota Shimizu</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="32/4341">Hisashi Kawai</author>
<title>Challenge of Singing Voice Synthesis Using Only Text-To-Speech Corpus With FIRNet Source-Filter Neural Vocoder.</title>
<year>2024</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2024-2504</ee>
<crossref>conf/interspeech/2024</crossref>
<url>db/conf/interspeech/interspeech2024.html#OkamotoOSTK24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/Shi0T24" mdate="2026-03-24">
<author orcid="0009-0001-3721-4964" pid="89/3530">Xiaohan Shi</author>
<author orcid="0000-0002-8958-0341" pid="63/9121-1">Xingfeng Li 0001</author>
<author pid="85/741">Tomoki Toda</author>
<title>Multimodal Fusion of Music Theory-Inspired and Self-Supervised Representations for Improved Emotion Recognition.</title>
<year>2024</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2024-2350</ee>
<crossref>conf/interspeech/2024</crossref>
<url>db/conf/interspeech/interspeech2024.html#Shi0T24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/ZangS0YHTXZGTD24" mdate="2025-05-31">
<author pid="348/9712">Yongyi Zang</author>
<author pid="229/3529">Jiatong Shi</author>
<author orcid="0000-0002-4649-278X" pid="26/3166-1">You Zhang 0001</author>
<author pid="81/3672">Ryuichi Yamamoto</author>
<author pid="378/4662">Jionghao Han</author>
<author pid="358/9271">Yuxun Tang</author>
<author pid="386/5432">Shengyuan Xu</author>
<author pid="77/8083">Wenxiao Zhao</author>
<author pid="95/5907">Jing Guo</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="04/6716">Zhiyao Duan</author>
<title>CtrSVDD: A Benchmark Dataset and Baseline Analysis for Controlled Singing Voice Deepfake Detection.</title>
<year>2024</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2024-2242</ee>
<crossref>conf/interspeech/2024</crossref>
<url>db/conf/interspeech/interspeech2024.html#ZangS0YHTXZGTD24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mmasia/HalpernT24" mdate="2026-03-05">
<author orcid="0000-0001-8787-359X" pid="271/4266">Bence Mark Halpern</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Reference-free automatic speech severity evaluation using acoustic unit language modelling.</title>
<pages>1:1-1:5</pages>
<year>2024</year>
<booktitle>MMAsia Workshops</booktitle>
<ee>https://doi.org/10.1145/3700410.3702114</ee>
<crossref>conf/mmasia/2024w</crossref>
<url>db/conf/mmasia/mmasia2024w.html#HalpernT24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/slt/ZhangZSYTD24" mdate="2025-03-03">
<author orcid="0000-0002-4649-278X" pid="26/3166-1">You Zhang 0001</author>
<author pid="348/9712">Yongyi Zang</author>
<author pid="229/3529">Jiatong Shi</author>
<author pid="81/3672">Ryuichi Yamamoto</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="04/6716">Zhiyao Duan</author>
<title>SVDD 2024: The Inaugural Singing Voice Deepfake Detection Challenge.</title>
<pages>782-787</pages>
<year>2024</year>
<booktitle>SLT</booktitle>
<ee>https://doi.org/10.1109/SLT61566.2024.10832284</ee>
<crossref>conf/slt/2024</crossref>
<url>db/conf/slt/slt2024.html#ZhangZSYTD24</url>
<stream>streams/conf/slt</stream>
</inproceedings>
</r>
<r><inproceedings key="conf/slt/HuangFCZTWYT24" mdate="2026-02-01">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="160/0591">Szu-Wei Fu</author>
<author pid="03/7184">Erica Cooper</author>
<author orcid="0000-0001-7319-8263" pid="199/7652">Ryandhimas E. Zezario</author>
<author pid="85/741">Tomoki Toda</author>
<author orcid="0000-0003-3599-5071" pid="28/5019">Hsin-Min Wang</author>
<author pid="87/3979">Junichi Yamagishi</author>
<author orcid="0000-0001-6956-0418" pid="66/7146-1">Yu Tsao 0001</author>
<title>The Voicemos Challenge 2024: Beyond Speech Quality Prediction.</title>
<pages>803-810</pages>
<year>2024</year>
<booktitle>SLT</booktitle>
<ee>https://doi.org/10.1109/SLT61566.2024.10832295</ee>
<crossref>conf/slt/2024</crossref>
<url>db/conf/slt/slt2024.html#HuangFCZTWYT24</url>
<stream>streams/conf/slt</stream>
</inproceedings>
</r>
<r><article publtype="informal" key="journals/corr/abs-2401-13260" mdate="2026-03-24">
<author pid="205/5074">Jiajun He</author>
<author orcid="0009-0001-3721-4964" pid="89/3530">Xiaohan Shi</author>
<author orcid="0000-0002-8958-0341" pid="63/9121-1">Xingfeng Li 0001</author>
<author pid="85/741">Tomoki Toda</author>
<title>MF-AED-AEC: Speech Emotion Recognition by Leveraging Multimodal Fusion, ASR Error Detection, and ASR Error Correction.</title>
<year>2024</year>
<volume>abs/2401.13260</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2401.13260</ee>
<url>db/journals/corr/corr2401.html#abs-2401-13260</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2403-06100" mdate="2024-04-04">
<author pid="228/9342">Yusuke Yasuda</author>
<author pid="85/741">Tomoki Toda</author>
<title>Automatic design optimization of preference-based subjective evaluation with online learning in crowdsourcing environment.</title>
<year>2024</year>
<volume>abs/2403.06100</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2403.06100</ee>
<url>db/journals/corr/corr2403.html#abs-2403-06100</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2404-06682" mdate="2024-05-16">
<author pid="334/0168">Yuka Hashizume</author>
<author pid="53/2189-63">Li Li 0063</author>
<author pid="357/2546">Atsushi Miyashita</author>
<author pid="85/741">Tomoki Toda</author>
<title>Learning Multidimensional Disentangled Representations of Instrumental Sounds for Musical Similarity Assessment.</title>
<year>2024</year>
<volume>abs/2404.06682</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2404.06682</ee>
<url>db/journals/corr/corr2404.html#abs-2404-06682</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2405-05244" mdate="2024-06-24">
<author pid="26/3166-1">You Zhang 0001</author>
<author pid="348/9712">Yongyi Zang</author>
<author pid="229/3529">Jiatong Shi</author>
<author pid="81/3672">Ryuichi Yamamoto</author>
<author pid="378/4662">Jionghao Han</author>
<author pid="358/9271">Yuxun Tang</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="04/6716">Zhiyao Duan</author>
<title>SVDD Challenge 2024: A Singing Voice Deepfake Detection Challenge Evaluation Plan.</title>
<year>2024</year>
<volume>abs/2405.05244</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2405.05244</ee>
<url>db/journals/corr/corr2405.html#abs-2405-05244</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2405-11767" mdate="2024-06-24">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="85/741">Tomoki Toda</author>
<title>Multi-speaker Text-to-speech Training with Speaker Anonymized Data.</title>
<year>2024</year>
<volume>abs/2405.11767</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2405.11767</ee>
<url>db/journals/corr/corr2405.html#abs-2405-11767</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2406-02438" mdate="2024-07-24">
<author pid="348/9712">Yongyi Zang</author>
<author pid="229/3529">Jiatong Shi</author>
<author pid="26/3166-1">You Zhang 0001</author>
<author pid="81/3672">Ryuichi Yamamoto</author>
<author pid="378/4662">Jionghao Han</author>
<author pid="358/9271">Yuxun Tang</author>
<author pid="386/5432">Shengyuan Xu</author>
<author pid="77/8083">Wenxiao Zhao</author>
<author pid="95/5907">Jing Guo</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="04/6716">Zhiyao Duan</author>
<title>CtrSVDD: A Benchmark Dataset and Baseline Analysis for Controlled Singing Voice Deepfake Detection.</title>
<year>2024</year>
<volume>abs/2406.02438</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2406.02438</ee>
<url>db/journals/corr/corr2406.html#abs-2406-02438</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2406-06201" mdate="2024-07-13">
<author pid="205/5074">Jiajun He</author>
<author pid="85/741">Tomoki Toda</author>
<title>2DP-2MRC: 2-Dimensional Pointer-based Machine Reading Comprehension Method for Multimodal Moment Retrieval.</title>
<year>2024</year>
<volume>abs/2406.06201</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2406.06201</ee>
<url>db/journals/corr/corr2406.html#abs-2406-06201</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2406-06208" mdate="2026-02-01">
<author orcid="0000-0001-8787-359X" pid="271/4266">Bence Mark Halpern</author>
<author orcid="0000-0002-8374-6654" pid="317/5237">Thomas Tienkamp</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="304/2864">Lester Phillip Violeta</author>
<author orcid="0000-0002-1986-8470" pid="317/5018">Teja Rebernik</author>
<author orcid="0000-0001-8964-5427" pid="329/8363">Sebastiaan A. H. J. de Visscher</author>
<author pid="227/2349">Max J. H. Witjes</author>
<author orcid="0000-0003-0434-1526" pid="35/2985">Martijn Wieling 0001</author>
<author orcid="0000-0002-0410-8487" pid="329/8124">Defne Abur</author>
<author pid="85/741">Tomoki Toda</author>
<title>Quantifying the effect of speech pathology on automatic and human speaker verification.</title>
<year>2024</year>
<volume>abs/2406.06208</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2406.06208</ee>
<url>db/journals/corr/corr2406.html#abs-2406-06208</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2408-16132" mdate="2024-09-28">
<author pid="26/3166-1">You Zhang 0001</author>
<author pid="348/9712">Yongyi Zang</author>
<author pid="229/3529">Jiatong Shi</author>
<author pid="81/3672">Ryuichi Yamamoto</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="04/6716">Zhiyao Duan</author>
<title>SVDD 2024: The Inaugural Singing Voice Deepfake Detection Challenge.</title>
<year>2024</year>
<volume>abs/2408.16132</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2408.16132</ee>
<url>db/journals/corr/corr2408.html#abs-2408-16132</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2409-07001" mdate="2026-02-01">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="160/0591">Szu-Wei Fu</author>
<author pid="03/7184">Erica Cooper</author>
<author pid="199/7652">Ryandhimas E. Zezario</author>
<author pid="85/741">Tomoki Toda</author>
<author orcid="0000-0003-3599-5071" pid="28/5019">Hsin-Min Wang</author>
<author pid="87/3979">Junichi Yamagishi</author>
<author orcid="0000-0001-6956-0418" pid="66/7146-1">Yu Tsao 0001</author>
<title>The VoiceMOS Challenge 2024: Beyond Speech Quality Prediction.</title>
<year>2024</year>
<volume>abs/2409.07001</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2409.07001</ee>
<url>db/journals/corr/corr2409.html#abs-2409-07001</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2409-09332" mdate="2024-10-21">
<author pid="216/3624">Takuya Fujimura</author>
<author pid="294/8499">Ibuki Kuroyanagi</author>
<author pid="85/741">Tomoki Toda</author>
<title>Improvements of Discriminative Feature Space Training for Anomalous Sound Detection in Unlabeled Conditions.</title>
<year>2024</year>
<volume>abs/2409.09332</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2409.09332</ee>
<url>db/journals/corr/corr2409.html#abs-2409-09332</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2409-19585" mdate="2026-03-24">
<author pid="387/9787">Jinyi Mi</author>
<author orcid="0009-0001-3721-4964" pid="89/3530">Xiaohan Shi</author>
<author pid="30/6718">Ding Ma</author>
<author pid="205/5074">Jiajun He</author>
<author pid="216/3624">Takuya Fujimura</author>
<author pid="85/741">Tomoki Toda</author>
<title>Two-stage Framework for Robust Speech Emotion Recognition Using Target Speaker Extraction in Human Speech Noise Conditions.</title>
<year>2024</year>
<volume>abs/2409.19585</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2409.19585</ee>
<url>db/journals/corr/corr2409.html#abs-2409-19585</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2409-19614" mdate="2024-10-17">
<author pid="387/9787">Jinyi Mi</author>
<author pid="47/4994">Sehun Kim</author>
<author pid="85/741">Tomoki Toda</author>
<title>Improved Architecture for High-resolution Piano Transcription to Efficiently Capture Acoustic Characteristics of Music Signals.</title>
<year>2024</year>
<volume>abs/2409.19614</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2409.19614</ee>
<url>db/journals/corr/corr2409.html#abs-2409-19614</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2411-03715" mdate="2025-01-01">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="03/7184">Erica Cooper</author>
<author pid="85/741">Tomoki Toda</author>
<title>MOS-Bench: Benchmarking Generalization Abilities of Subjective Speech Quality Assessment Models.</title>
<year>2024</year>
<volume>abs/2411.03715</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2411.03715</ee>
<url>db/journals/corr/corr2411.html#abs-2411-03715</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2411-06807" mdate="2025-01-01">
<author pid="290/1755">Reo Yoneyama</author>
<author pid="357/2546">Atsushi Miyashita</author>
<author pid="81/3672">Ryuichi Yamamoto</author>
<author pid="85/741">Tomoki Toda</author>
<title>Wavehax: Aliasing-Free Neural Waveform Synthesis Based on 2D Convolution and Harmonic Prior for Reliable Complex Spectrogram Estimation.</title>
<year>2024</year>
<volume>abs/2411.06807</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2411.06807</ee>
<url>db/journals/corr/corr2411.html#abs-2411-06807</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article key="journals/taslp/MatsubaraOTTTK23" mdate="2023-07-27">
<author orcid="0000-0002-2935-668X" pid="157/7171">Keisuke Matsubara</author>
<author orcid="0000-0001-9913-4647" pid="132/9091">Takuma Okamoto</author>
<author orcid="0000-0002-9808-0250" pid="22/8758">Ryoichi Takashima</author>
<author orcid="0000-0001-5005-7679" pid="79/4485">Tetsuya Takiguchi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="32/4341">Hisashi Kawai</author>
<title>Harmonic-Net: Fundamental Frequency and Speech Rate Controllable Fast Neural Vocoder.</title>
<pages>1902-1915</pages>
<year>2023</year>
<volume>31</volume>
<journal>IEEE ACM Trans. Audio Speech Lang. Process.</journal>
<ee type="oa">https://doi.org/10.1109/TASLP.2023.3275032</ee>
<url>db/journals/taslp/taslp31.html#MatsubaraOTTTK23</url>
</article>
</r>
<r><article key="journals/taslp/YoneyamaWT23" mdate="2023-11-09">
<author orcid="0000-0001-9686-4783" pid="290/1755">Reo Yoneyama</author>
<author orcid="0000-0003-4390-1354" pid="188/5943">Yi-Chiao Wu</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>High-Fidelity and Pitch-Controllable Neural Vocoder Based on Unified Source-Filter Networks.</title>
<pages>3717-3729</pages>
<year>2023</year>
<volume>31</volume>
<journal>IEEE ACM Trans. Audio Speech Lang. Process.</journal>
<ee type="oa">https://doi.org/10.1109/TASLP.2023.3313410</ee>
<url>db/journals/taslp/taslp31.html#YoneyamaWT23</url>
</article>
</r>
<r><article key="journals/taslp/XieT23" mdate="2025-01-19">
<author orcid="0000-0002-6237-9816" pid="04/5576">Chao Xie</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Noisy-to-Noisy Voice Conversion Under Variations of Noisy Condition.</title>
<pages>3871-3882</pages>
<year>2023</year>
<volume>31</volume>
<journal>IEEE ACM Trans. Audio Speech Lang. Process.</journal>
<ee type="oa">https://doi.org/10.1109/TASLP.2023.3313426</ee>
<ee>https://www.wikidata.org/entity/Q130982067</ee>
<url>db/journals/taslp/taslp31.html#XieT23</url>
</article>
</r>
<r><inproceedings key="conf/apsipa/HuangT23" mdate="2023-12-02">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="85/741">Tomoki Toda</author>
<title>Evaluating Methods for Ground-Truth-Free Foreign Accent Conversion.</title>
<pages>1161-1166</pages>
<year>2023</year>
<booktitle>APSIPA ASC</booktitle>
<ee>https://doi.org/10.1109/APSIPAASC58517.2023.10317592</ee>
<crossref>conf/apsipa/2023</crossref>
<url>db/conf/apsipa/apsipa2023.html#HuangT23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/apsipa/VioletaT23" mdate="2023-12-02">
<author pid="304/2864">Lester Phillip Violeta</author>
<author pid="85/741">Tomoki Toda</author>
<title>An Analysis of Personalized Speech Recognition System Development for the Deaf and Hard-of-Hearing.</title>
<pages>1862-1867</pages>
<year>2023</year>
<booktitle>APSIPA ASC</booktitle>
<ee>https://doi.org/10.1109/APSIPAASC58517.2023.10317318</ee>
<crossref>conf/apsipa/2023</crossref>
<url>db/conf/apsipa/apsipa2023.html#VioletaT23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/asru/CooperHTWTY23" mdate="2026-02-01">
<author pid="03/7184">Erica Cooper</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author orcid="0000-0001-6956-0418" pid="66/7146-1">Yu Tsao 0001</author>
<author orcid="0000-0003-3599-5071" pid="28/5019">Hsin-Min Wang</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="87/3979">Junichi Yamagishi</author>
<title>The Voicemos Challenge 2023: Zero-Shot Subjective Speech Quality Prediction for Multiple Domains.</title>
<pages>1-7</pages>
<year>2023</year>
<booktitle>ASRU</booktitle>
<ee>https://doi.org/10.1109/ASRU57964.2023.10389763</ee>
<crossref>conf/asru/2023</crossref>
<url>db/conf/asru/asru2023.html#CooperHTWTY23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/asru/HalpernHVST23" mdate="2026-02-01">
<author orcid="0000-0001-8787-359X" pid="271/4266">Bence Mark Halpern</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="304/2864">Lester Phillip Violeta</author>
<author orcid="0000-0001-6321-7635" pid="34/4715">R. J. J. H. van Son</author>
<author pid="85/741">Tomoki Toda</author>
<title>Improving Severity Preservation of Healthy-to-Pathological Voice Conversion With Global Style Tokens.</title>
<pages>1-7</pages>
<year>2023</year>
<booktitle>ASRU</booktitle>
<ee>https://doi.org/10.1109/ASRU57964.2023.10389707</ee>
<crossref>conf/asru/2023</crossref>
<url>db/conf/asru/asru2023.html#HalpernHVST23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/asru/HeYT23" mdate="2024-10-06">
<author pid="205/5074">Jiajun He</author>
<author orcid="0000-0003-3518-8229" pid="176/5718">Zekun Yang</author>
<author pid="85/741">Tomoki Toda</author>
<title>ED-CEC: Improving Rare word Recognition Using ASR Postprocessing Based on Error Detection and Context-Aware Error Correction.</title>
<pages>1-6</pages>
<year>2023</year>
<booktitle>ASRU</booktitle>
<ee>https://doi.org/10.1109/ASRU57964.2023.10389661</ee>
<crossref>conf/asru/2023</crossref>
<url>db/conf/asru/asru2023.html#HeYT23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/asru/HuangVLST23" mdate="2024-02-13">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="304/2864">Lester Phillip Violeta</author>
<author pid="226/2031">Songxiang Liu</author>
<author pid="229/3529">Jiatong Shi</author>
<author pid="85/741">Tomoki Toda</author>
<title>The Singing Voice Conversion Challenge 2023.</title>
<pages>1-8</pages>
<year>2023</year>
<booktitle>ASRU</booktitle>
<ee>https://doi.org/10.1109/ASRU57964.2023.10389671</ee>
<crossref>conf/asru/2023</crossref>
<url>db/conf/asru/asru2023.html#HuangVLST23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/asru/OkamotoYOTK23" mdate="2024-02-13">
<author pid="132/9091">Takuma Okamoto</author>
<author pid="368/3515">Haruki Yamashita</author>
<author pid="34/8763">Yamato Ohtani</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="32/4341">Hisashi Kawai</author>
<title>WaveNeXt: ConvNeXt-Based Fast Neural Vocoder Without ISTFT layer.</title>
<pages>1-8</pages>
<year>2023</year>
<booktitle>ASRU</booktitle>
<ee>https://doi.org/10.1109/ASRU57964.2023.10389765</ee>
<crossref>conf/asru/2023</crossref>
<url>db/conf/asru/asru2023.html#OkamotoYOTK23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/asru/YamamotoYVHT23" mdate="2024-02-13">
<author pid="81/3672">Ryuichi Yamamoto</author>
<author pid="290/1755">Reo Yoneyama</author>
<author pid="304/2864">Lester Phillip Violeta</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="85/741">Tomoki Toda</author>
<title>A Comparative Study of Voice Conversion Models With Large-Scale Speech and Singing Data: The T13 Systems for the Singing Voice Conversion Challenge 2023.</title>
<pages>1-6</pages>
<year>2023</year>
<booktitle>ASRU</booktitle>
<ee>https://doi.org/10.1109/ASRU57964.2023.10389779</ee>
<crossref>conf/asru/2023</crossref>
<url>db/conf/asru/asru2023.html#YamamotoYVHT23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/eusipco/LuanWT23" mdate="2023-11-06">
<author pid="331/6661">Shuming Luan</author>
<author pid="158/4215">Yukoh Wakabayashi</author>
<author pid="85/741">Tomoki Toda</author>
<title>Sound Field Interpolation with Unsupervised Calibration for Freely Spaced Circular Microphone Array in Rotation-Robust Beamforming.</title>
<pages>21-25</pages>
<year>2023</year>
<booktitle>EUSIPCO</booktitle>
<ee>https://doi.org/10.23919/EUSIPCO58844.2023.10289792</ee>
<crossref>conf/eusipco/2023</crossref>
<url>db/conf/eusipco/eusipco2023.html#LuanWT23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/FujimuraT23" mdate="2023-11-05">
<author pid="216/3624">Takuya Fujimura</author>
<author pid="85/741">Tomoki Toda</author>
<title>Analysis Of Noisy-Target Training For Dnn-Based Speech Enhancement.</title>
<pages>1-5</pages>
<year>2023</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP49357.2023.10096481</ee>
<crossref>conf/icassp/2023</crossref>
<url>db/conf/icassp/icassp2023.html#FujimuraT23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/KobayashiHT23" mdate="2023-11-05">
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="85/741">Tomoki Toda</author>
<title>Low-Latency Electrolaryngeal Speech Enhancement Based on Fastspeech2-Based Voice Conversion and Self-Supervised Speech Representation.</title>
<pages>1-5</pages>
<year>2023</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP49357.2023.10096442</ee>
<crossref>conf/icassp/2023</crossref>
<url>db/conf/icassp/icassp2023.html#KobayashiHT23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/MiyashitaT23" mdate="2023-11-05">
<author pid="357/2546">Atsushi Miyashita</author>
<author pid="85/741">Tomoki Toda</author>
<title>Representation of Vocal Tract Length Transformation Based on Group Theory.</title>
<pages>1-5</pages>
<year>2023</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP49357.2023.10095239</ee>
<crossref>conf/icassp/2023</crossref>
<url>db/conf/icassp/icassp2023.html#MiyashitaT23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/VioletaMHT23" mdate="2023-11-05">
<author pid="304/2864">Lester Phillip Violeta</author>
<author pid="30/6718">Ding Ma</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="85/741">Tomoki Toda</author>
<title>Intermediate Fine-Tuning Using Imperfect Synthetic Speech for Improving Electrolaryngeal Speech Recognition.</title>
<pages>1-5</pages>
<year>2023</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP49357.2023.10095931</ee>
<crossref>conf/icassp/2023</crossref>
<url>db/conf/icassp/icassp2023.html#VioletaMHT23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/YamamotoYT23" mdate="2023-11-05">
<author pid="81/3672">Ryuichi Yamamoto</author>
<author pid="290/1755">Reo Yoneyama</author>
<author pid="85/741">Tomoki Toda</author>
<title>NNSVS: A Neural Network-Based Singing Voice Synthesis Toolkit.</title>
<pages>1-5</pages>
<year>2023</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP49357.2023.10096239</ee>
<crossref>conf/icassp/2023</crossref>
<url>db/conf/icassp/icassp2023.html#YamamotoYT23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/YasudaT23" mdate="2023-11-05">
<author pid="228/9342">Yusuke Yasuda</author>
<author pid="85/741">Tomoki Toda</author>
<title>Text-To-Speech Synthesis Based on Latent Variable Conversion Using Diffusion Probabilistic Model and Variational Autoencoder.</title>
<pages>1-5</pages>
<year>2023</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP49357.2023.10094298</ee>
<crossref>conf/icassp/2023</crossref>
<url>db/conf/icassp/icassp2023.html#YasudaT23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/YoneyamaWT23" mdate="2023-11-05">
<author pid="290/1755">Reo Yoneyama</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="85/741">Tomoki Toda</author>
<title>Source-Filter HiFi-GAN: Fast and Pitch Controllable High-Fidelity Neural Vocoder.</title>
<pages>1-5</pages>
<year>2023</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP49357.2023.10095298</ee>
<crossref>conf/icassp/2023</crossref>
<url>db/conf/icassp/icassp2023.html#YoneyamaWT23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/HuYT23" mdate="2024-06-14">
<author pid="156/0822">Cheng-Hung Hu</author>
<author pid="228/9342">Yusuke Yasuda</author>
<author pid="85/741">Tomoki Toda</author>
<title>Preference-based training framework for automatic speech quality assessment using deep neural network.</title>
<pages>546-550</pages>
<year>2023</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2023-589</ee>
<crossref>conf/interspeech/2023</crossref>
<url>db/conf/interspeech/interspeech2023.html#HuYT23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/Shi0T23" mdate="2026-03-24">
<author orcid="0009-0001-3721-4964" pid="89/3530">Xiaohan Shi</author>
<author orcid="0000-0002-8958-0341" pid="63/9121-1">Xingfeng Li 0001</author>
<author pid="85/741">Tomoki Toda</author>
<title>Emotion Awareness in Multi-utterance Turn for Improving Emotion Prediction in Multi-Speaker Conversation.</title>
<pages>765-769</pages>
<year>2023</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2023-1236</ee>
<crossref>conf/interspeech/2023</crossref>
<url>db/conf/interspeech/interspeech2023.html#Shi0T23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/OkamotoTK23" mdate="2024-06-14">
<author pid="132/9091">Takuma Okamoto</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="32/4341">Hisashi Kawai</author>
<title>E2E-S2S-VC: End-To-End Sequence-To-Sequence Voice Conversion.</title>
<pages>2043-2047</pages>
<year>2023</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2023-2518</ee>
<crossref>conf/interspeech/2023</crossref>
<url>db/conf/interspeech/interspeech2023.html#OkamotoTK23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/ChoiXT23" mdate="2024-06-14">
<author pid="323/5236">Yeonjong Choi</author>
<author pid="04/5576">Chao Xie</author>
<author pid="85/741">Tomoki Toda</author>
<title>Reverberation-Controllable Voice Conversion Using Reverberation Time Estimator.</title>
<pages>2103-2107</pages>
<year>2023</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2023-1356</ee>
<crossref>conf/interspeech/2023</crossref>
<url>db/conf/interspeech/interspeech2023.html#ChoiXT23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/YasudaT23" mdate="2024-06-14">
<author pid="228/9342">Yusuke Yasuda</author>
<author pid="85/741">Tomoki Toda</author>
<title>Analysis of Mean Opinion Scores in Subjective Evaluation of Synthetic Speech Based on Tail Probabilities.</title>
<pages>5491-5495</pages>
<year>2023</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2023-1285</ee>
<crossref>conf/interspeech/2023</crossref>
<url>db/conf/interspeech/interspeech2023.html#YasudaT23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/ismir/KimTT23" mdate="2025-02-18">
<author pid="47/4994">Sehun Kim</author>
<author pid="38/3483">Kazuya Takeda</author>
<author pid="85/741">Tomoki Toda</author>
<title>Sequence-to-Sequence Network Training Methods for Automatic Guitar Transcription With Tokenized Outputs.</title>
<pages>524-531</pages>
<year>2023</year>
<booktitle>ISMIR</booktitle>
<ee type="oa">https://doi.org/10.5281/zenodo.10265341</ee>
<crossref>conf/ismir/2023</crossref>
<url>db/conf/ismir/ismir2023.html#KimTT23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mrac/TianHSH0GTXH23" mdate="2025-09-12">
<author orcid="0009-0000-0865-9422" pid="313/9922">Jingguang Tian</author>
<author orcid="0009-0000-6348-0409" pid="02/4151">Desheng Hu</author>
<author orcid="0000-0002-1917-4479" pid="89/3530">Xiaohan Shi</author>
<author orcid="0009-0006-6489-5220" pid="205/5074">Jiajun He</author>
<author orcid="0000-0002-8958-0341" pid="63/9121-1">Xingfeng Li 0001</author>
<author orcid="0000-0002-2147-1835" pid="76/2452-40">Yuan Gao 0040</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author orcid="0009-0003-2771-1398" pid="277/3578">Xinkang Xu</author>
<author orcid="0009-0009-1433-9324" pid="04/1555">Xinhui Hu</author>
<title>Semi-supervised Multimodal Emotion Recognition with Consensus Decision-making and Label Correction.</title>
<pages>67-73</pages>
<year>2023</year>
<booktitle>MRAC@MM</booktitle>
<ee>https://doi.org/10.1145/3607865.3613182</ee>
<crossref>conf/mrac/2023</crossref>
<url>db/conf/mrac/mrac2023.html#TianHSH0GTXH23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/waspaa/MiyashitaT23" mdate="2023-09-23">
<author pid="357/2546">Atsushi Miyashita</author>
<author pid="85/741">Tomoki Toda</author>
<title>Differentiable Representation of Warping Based on Lie Group Theory.</title>
<pages>1-5</pages>
<year>2023</year>
<booktitle>WASPAA</booktitle>
<ee>https://doi.org/10.1109/WASPAA58266.2023.10248099</ee>
<crossref>conf/waspaa/2023</crossref>
<url>db/conf/waspaa/waspaa2023.html#MiyashitaT23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/waspaa/WangT23" mdate="2026-03-12">
<author pid="06/2293-169">Rui Wang 0169</author>
<author pid="85/741">Tomoki Toda</author>
<title>Directional Target Speaker Extraction under Noisy Underdetermined Conditions through Conditional Variational Autoencoder with Global Style Tokens.</title>
<pages>1-5</pages>
<year>2023</year>
<booktitle>WASPAA</booktitle>
<ee>https://doi.org/10.1109/WASPAA58266.2023.10248146</ee>
<crossref>conf/waspaa/2023</crossref>
<url>db/conf/waspaa/waspaa2023.html#WangT23</url>
</inproceedings>
</r>
<r><article publtype="informal" key="journals/corr/abs-2306-13953" mdate="2023-06-27">
<author pid="304/2864">Lester Phillip Violeta</author>
<author pid="85/741">Tomoki Toda</author>
<title>An Analysis of Personalized Speech Recognition System Development for the Deaf and Hard-of-Hearing.</title>
<year>2023</year>
<volume>abs/2306.13953</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2306.13953</ee>
<url>db/journals/corr/corr2306.html#abs-2306-13953</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2306-14422" mdate="2023-06-27">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="304/2864">Lester Phillip Violeta</author>
<author pid="226/2031">Songxiang Liu</author>
<author pid="229/3529">Jiatong Shi</author>
<author pid="228/9342">Yusuke Yasuda</author>
<author pid="85/741">Tomoki Toda</author>
<title>The Singing Voice Conversion Challenge 2023.</title>
<year>2023</year>
<volume>abs/2306.14422</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2306.14422</ee>
<url>db/journals/corr/corr2306.html#abs-2306-14422</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2309-02133" mdate="2023-09-11">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="85/741">Tomoki Toda</author>
<title>Evaluating Methods for Ground-Truth-Free Foreign Accent Conversion.</title>
<year>2023</year>
<volume>abs/2309.02133</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2309.02133</ee>
<url>db/journals/corr/corr2309.html#abs-2309-02133</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2309-07598" mdate="2023-09-19">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="85/741">Tomoki Toda</author>
<title>AAS-VC: On the Generalization Ability of Automatic Alignment Search based Non-autoregressive Sequence-to-sequence Voice Conversion.</title>
<year>2023</year>
<volume>abs/2309.07598</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2309.07598</ee>
<url>db/journals/corr/corr2309.html#abs-2309-07598</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2309-08141" mdate="2023-09-26">
<author pid="136/5113">Tatsuya Komatsu</author>
<author pid="50/6352">Yusuke Fujita</author>
<author pid="38/3483">Kazuya Takeda</author>
<author pid="85/741">Tomoki Toda</author>
<title>Audio Difference Learning for Audio Captioning.</title>
<year>2023</year>
<volume>abs/2309.08141</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2309.08141</ee>
<url>db/journals/corr/corr2309.html#abs-2309-08141</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2309-09627" mdate="2023-09-22">
<author pid="304/2864">Lester Phillip Violeta</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="30/6718">Ding Ma</author>
<author pid="81/3672">Ryuichi Yamamoto</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="85/741">Tomoki Toda</author>
<title>Electrolaryngeal Speech Intelligibility Enhancement Through Robust Linguistic Encoders.</title>
<year>2023</year>
<volume>abs/2309.09627</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2309.09627</ee>
<url>db/journals/corr/corr2309.html#abs-2309-09627</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2310-02570" mdate="2023-10-19">
<author pid="271/4266">Bence Mark Halpern</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="304/2864">Lester Phillip Violeta</author>
<author pid="34/4715">R. J. J. H. van Son</author>
<author pid="85/741">Tomoki Toda</author>
<title>Improving severity preservation of healthy-to-pathological voice conversion with global style tokens.</title>
<year>2023</year>
<volume>abs/2310.02570</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2310.02570</ee>
<url>db/journals/corr/corr2310.html#abs-2310-02570</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2310-05129" mdate="2023-10-20">
<author pid="205/5074">Jiajun He</author>
<author pid="176/5718">Zekun Yang</author>
<author pid="85/741">Tomoki Toda</author>
<title>ed-cec: improving rare word recognition using asr postprocessing based on error detection and context-aware error correction.</title>
<year>2023</year>
<volume>abs/2310.05129</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2310.05129</ee>
<url>db/journals/corr/corr2310.html#abs-2310-05129</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2310-05203" mdate="2023-10-26">
<author pid="81/3672">Ryuichi Yamamoto</author>
<author pid="290/1755">Reo Yoneyama</author>
<author pid="304/2864">Lester Phillip Violeta</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="85/741">Tomoki Toda</author>
<title>A Comparative Study of Voice Conversion Models with Large-Scale Speech and Singing Data: The T13 Systems for the Singing Voice Conversion Challenge 2023.</title>
<year>2023</year>
<volume>abs/2310.05203</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2310.05203</ee>
<url>db/journals/corr/corr2310.html#abs-2310-05203</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2311-07093" mdate="2026-03-24">
<author orcid="0009-0001-3721-4964" pid="89/3530">Xiaohan Shi</author>
<author pid="205/5074">Jiajun He</author>
<author orcid="0000-0002-8958-0341" pid="63/9121-1">Xingfeng Li 0001</author>
<author pid="85/741">Tomoki Toda</author>
<title>On the Effectiveness of ASR Representations in Real-world Noisy Speech Emotion Recognition.</title>
<year>2023</year>
<volume>abs/2311.07093</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2311.07093</ee>
<url>db/journals/corr/corr2311.html#abs-2311-07093</url>
</article>
</r>
<r><article key="journals/jstsp/HuangYHT22" mdate="2022-11-13">
<author orcid="0000-0003-3172-3335" pid="225/7821">Wen-Chin Huang</author>
<author orcid="0000-0001-5503-9410" pid="246/0774">Shu-Wen Yang</author>
<author orcid="0000-0001-8782-4093" pid="82/8616">Tomoki Hayashi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>A Comparative Study of Self-Supervised Speech Representation Based Voice Conversion.</title>
<pages>1308-1318</pages>
<year>2022</year>
<volume>16</volume>
<journal>IEEE J. Sel. Top. Signal Process.</journal>
<number>6</number>
<ee>https://doi.org/10.1109/JSTSP.2022.3193761</ee>
<url>db/journals/jstsp/jstsp16.html#HuangYHT22</url>
</article>
</r>
<r><article key="journals/jstsp/YasudaT22" mdate="2025-01-19">
<author orcid="0000-0002-2130-747X" pid="228/9342">Yusuke Yasuda</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Investigation of Japanese PnG BERT Language Model in Text-to-Speech Synthesis for Pitch Accent Language.</title>
<pages>1319-1328</pages>
<year>2022</year>
<volume>16</volume>
<journal>IEEE J. Sel. Top. Signal Process.</journal>
<number>6</number>
<ee type="oa">https://doi.org/10.1109/JSTSP.2022.3190672</ee>
<ee>https://www.wikidata.org/entity/Q125590291</ee>
<url>db/journals/jstsp/jstsp16.html#YasudaT22</url>
</article>
</r>
<r><article key="journals/speech/OkamotoMTSK22" mdate="2022-05-13">
<author orcid="0000-0001-9913-4647" pid="132/9091">Takuma Okamoto</author>
<author orcid="0000-0002-2935-668X" pid="157/7171">Keisuke Matsubara</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="02/2978">Yoshinori Shiga</author>
<author pid="32/4341">Hisashi Kawai</author>
<title>Neural speech-rate conversion with multispeaker WaveNet vocoder.</title>
<pages>1-12</pages>
<year>2022</year>
<volume>138</volume>
<journal>Speech Commun.</journal>
<ee type="oa">https://doi.org/10.1016/j.specom.2022.01.003</ee>
<url>db/journals/speech/speech138.html#OkamotoMTSK22</url>
</article>
</r>
<r><inproceedings key="conf/eusipco/KimHT22" mdate="2022-10-25">
<author pid="47/4994">Sehun Kim</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="85/741">Tomoki Toda</author>
<title>Note-level Automatic Guitar Transcription Using Attention Mechanism.</title>
<pages>229-233</pages>
<year>2022</year>
<booktitle>EUSIPCO</booktitle>
<ee>https://ieeexplore.ieee.org/document/9909659</ee>
<crossref>conf/eusipco/2022</crossref>
<url>db/conf/eusipco/eusipco2022.html#KimHT22</url>
</inproceedings>
</r>
<r><inproceedings key="conf/eusipco/KuroyanagiHTT22" mdate="2022-10-25">
<author pid="294/8499">Ibuki Kuroyanagi</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="38/3483">Kazuya Takeda</author>
<author pid="85/741">Tomoki Toda</author>
<title>Improvement of Serial Approach to Anomalous Sound Detection by Incorporating Two Binary Cross-Entropies for Outlier Exposure.</title>
<pages>294-298</pages>
<year>2022</year>
<booktitle>EUSIPCO</booktitle>
<ee>https://ieeexplore.ieee.org/document/9909621</ee>
<crossref>conf/eusipco/2022</crossref>
<url>db/conf/eusipco/eusipco2022.html#KuroyanagiHTT22</url>
</inproceedings>
</r>
<r><inproceedings key="conf/eusipco/LuanWT22" mdate="2022-10-25">
<author pid="331/6661">Shuming Luan</author>
<author pid="158/4215">Yukoh Wakabayashi</author>
<author pid="85/741">Tomoki Toda</author>
<title>Modified Sound Field Interpolation Method for Rotation-robust Beamforming with Unequally Spaced Circular Microphone Array.</title>
<pages>344-348</pages>
<year>2022</year>
<booktitle>EUSIPCO</booktitle>
<ee>https://ieeexplore.ieee.org/document/9909954</ee>
<crossref>conf/eusipco/2022</crossref>
<url>db/conf/eusipco/eusipco2022.html#LuanWT22</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/HuangCYT22" mdate="2022-06-07">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="03/7184">Erica Cooper</author>
<author pid="87/3979">Junichi Yamagishi</author>
<author pid="85/741">Tomoki Toda</author>
<title>LDNet: Unified Listener Dependent Modeling in MOS Prediction for Synthetic Speech.</title>
<pages>896-900</pages>
<year>2022</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP43922.2022.9747222</ee>
<crossref>conf/icassp/2022</crossref>
<url>db/conf/icassp/icassp2022.html#HuangCYT22</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/HuangYHLWT22" mdate="2022-12-07">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="246/0774">Shu-Wen Yang</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="81/8056">Hung-Yi Lee</author>
<author orcid="0000-0002-5970-8631" pid="39/3245-1">Shinji Watanabe 0001</author>
<author pid="85/741">Tomoki Toda</author>
<title>S3PRL-VC: Open-Source Voice Conversion Framework with Self-Supervised Speech Representations.</title>
<pages>6552-6556</pages>
<year>2022</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP43922.2022.9746430</ee>
<crossref>conf/icassp/2022</crossref>
<url>db/conf/icassp/icassp2022.html#HuangYHLWT22</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/HuangHVST22" mdate="2026-02-01">
<author pid="225/7821">Wen-Chin Huang</author>
<author orcid="0000-0001-8787-359X" pid="271/4266">Bence Mark Halpern</author>
<author pid="304/2864">Lester Phillip Violeta</author>
<author orcid="0000-0003-0693-8852" pid="87/1365">Odette Scharenborg</author>
<author pid="85/741">Tomoki Toda</author>
<title>Towards Identity Preserving Normal to Dysarthric Voice Conversion.</title>
<pages>6672-6676</pages>
<year>2022</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP43922.2022.9747550</ee>
<crossref>conf/icassp/2022</crossref>
<url>db/conf/icassp/icassp2022.html#HuangHVST22</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/XieWTHT22" mdate="2023-06-26">
<author pid="04/5576">Chao Xie</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author orcid="0000-0003-2792-8418" pid="158/4210">Patrick Lumban Tobing</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="85/741">Tomoki Toda</author>
<title>Direct Noisy Speech Modeling for Noisy-To-Noisy Voice Conversion.</title>
<pages>6787-6791</pages>
<year>2022</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP43922.2022.9747894</ee>
<crossref>conf/icassp/2022</crossref>
<url>db/conf/icassp/icassp2022.html#XieWTHT22</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/HayashiKT22" mdate="2022-06-07">
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="85/741">Tomoki Toda</author>
<title>An Investigation of Streaming Non-Autoregressive sequence-to-sequence Voice Conversion.</title>
<pages>6802-6806</pages>
<year>2022</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP43922.2022.9747558</ee>
<crossref>conf/icassp/2022</crossref>
<url>db/conf/icassp/icassp2022.html#HayashiKT22</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/CooperHTY22" mdate="2022-06-07">
<author pid="03/7184">Erica Cooper</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="87/3979">Junichi Yamagishi</author>
<title>Generalization Ability of MOS Prediction Networks.</title>
<pages>8442-8446</pages>
<year>2022</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP43922.2022.9746395</ee>
<crossref>conf/icassp/2022</crossref>
<url>db/conf/icassp/icassp2022.html#CooperHTY22</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/VioletaHT22" mdate="2023-06-21">
<author pid="304/2864">Lester Phillip Violeta</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="85/741">Tomoki Toda</author>
<title>Investigating Self-supervised Pretraining Frameworks for Pathological Speech Recognition.</title>
<pages>41-45</pages>
<year>2022</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2022-10043</ee>
<crossref>conf/interspeech/2022</crossref>
<url>db/conf/interspeech/interspeech2022.html#VioletaHT22</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/YoneyamaWT22" mdate="2023-06-21">
<author pid="290/1755">Reo Yoneyama</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="85/741">Tomoki Toda</author>
<title>Unified Source-Filter GAN with Harmonic-plus-Noise Source Excitation Generation.</title>
<pages>848-852</pages>
<year>2022</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2022-11130</ee>
<crossref>conf/interspeech/2022</crossref>
<url>db/conf/interspeech/interspeech2022.html#YoneyamaWT22</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/HuangC0WTY22" mdate="2026-02-01">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="03/7184">Erica Cooper</author>
<author orcid="0000-0001-6956-0418" pid="66/7146-1">Yu Tsao 0001</author>
<author orcid="0000-0003-3599-5071" pid="28/5019">Hsin-Min Wang</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="87/3979">Junichi Yamagishi</author>
<title>The VoiceMOS Challenge 2022.</title>
<pages>4536-4540</pages>
<year>2022</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2022-970</ee>
<crossref>conf/interspeech/2022</crossref>
<url>db/conf/interspeech/interspeech2022.html#HuangC0WTY22</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/YoshiokaYMOT22" mdate="2023-06-21">
<author pid="200/0472">Daiki Yoshioka</author>
<author pid="228/9342">Yusuke Yasuda</author>
<author pid="265/6508">Noriyuki Matsunaga</author>
<author pid="34/8763">Yamato Ohtani</author>
<author pid="85/741">Tomoki Toda</author>
<title>Spoken-Text-Style Transfer with Conditional Variational Autoencoder and Content Word Storage.</title>
<pages>4576-4580</pages>
<year>2022</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2022-10118</ee>
<crossref>conf/interspeech/2022</crossref>
<url>db/conf/interspeech/interspeech2022.html#YoshiokaYMOT22</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/ChoiXT22" mdate="2023-06-21">
<author pid="323/5236">Yeonjong Choi</author>
<author pid="04/5576">Chao Xie</author>
<author pid="85/741">Tomoki Toda</author>
<title>An Evaluation of Three-Stage Voice Conversion Framework for Noisy and Reverberant Conditions.</title>
<pages>4910-4914</pages>
<year>2022</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2022-10158</ee>
<crossref>conf/interspeech/2022</crossref>
<url>db/conf/interspeech/interspeech2022.html#ChoiXT22</url>
</inproceedings>
</r>
<r><inproceedings key="conf/slt/MaVKT22" mdate="2023-02-06">
<author pid="30/6718">Ding Ma</author>
<author pid="304/2864">Lester Phillip Violeta</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="85/741">Tomoki Toda</author>
<title>Two-Stage Training Method for Japanese Electrolaryngeal Speech Enhancement Based on Sequence-to-Sequence Voice Conversion.</title>
<pages>949-954</pages>
<year>2022</year>
<booktitle>SLT</booktitle>
<ee>https://doi.org/10.1109/SLT54892.2023.10023033</ee>
<crossref>conf/slt/2022</crossref>
<url>db/conf/slt/slt2022.html#MaVKT22</url>
</inproceedings>
</r>
<r><article publtype="informal" key="journals/corr/abs-2203-11389" mdate="2026-02-01">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="03/7184">Erica Cooper</author>
<author orcid="0000-0001-6956-0418" pid="66/7146-1">Yu Tsao 0001</author>
<author orcid="0000-0003-3599-5071" pid="28/5019">Hsin-Min Wang</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="87/3979">Junichi Yamagishi</author>
<title>The VoiceMOS Challenge 2022.</title>
<year>2022</year>
<volume>abs/2203.11389</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2203.11389</ee>
<url>db/journals/corr/corr2203.html#abs-2203-11389</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2203-15431" mdate="2022-04-04">
<author pid="304/2864">Lester Phillip Violeta</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="85/741">Tomoki Toda</author>
<title>Investigating Self-supervised Pretraining Frameworks for Pathological Speech Recognition.</title>
<year>2022</year>
<volume>abs/2203.15431</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2203.15431</ee>
<url>db/journals/corr/corr2203.html#abs-2203-15431</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2205-06053" mdate="2022-05-17">
<author pid="290/1755">Reo Yoneyama</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="85/741">Tomoki Toda</author>
<title>Unified Source-Filter GAN with Harmonic-plus-Noise Source Excitation Generation.</title>
<year>2022</year>
<volume>abs/2205.06053</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2205.06053</ee>
<url>db/journals/corr/corr2205.html#abs-2205-06053</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2206-05929" mdate="2022-06-20">
<author pid="294/8499">Ibuki Kuroyanagi</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="38/3483">Kazuya Takeda</author>
<author pid="85/741">Tomoki Toda</author>
<title>Improvement of Serial Approach to Anomalous Sound Detection by Incorporating Two Binary Cross-Entropies for Outlier Exposure.</title>
<year>2022</year>
<volume>abs/2206.05929</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2206.05929</ee>
<url>db/journals/corr/corr2206.html#abs-2206-05929</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2206-15155" mdate="2022-07-04">
<author pid="323/5236">Yeonjong Choi</author>
<author pid="04/5576">Chao Xie</author>
<author pid="85/741">Tomoki Toda</author>
<title>An Evaluation of Three-Stage Voice Conversion Framework for Noisy and Reverberant Conditions.</title>
<year>2022</year>
<volume>abs/2206.15155</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2206.15155</ee>
<url>db/journals/corr/corr2206.html#abs-2206-15155</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2207-04356" mdate="2022-07-13">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="246/0774">Shu-Wen Yang</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="85/741">Tomoki Toda</author>
<title>A Comparative Study of Self-supervised Speech Representation Based Voice Conversion.</title>
<year>2022</year>
<volume>abs/2207.04356</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2207.04356</ee>
<url>db/journals/corr/corr2207.html#abs-2207-04356</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2207-05913" mdate="2022-07-19">
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="158/4210">Patrick Lumban Tobing</author>
<author pid="265/6506">Kazuki Yasuhara</author>
<author pid="265/6508">Noriyuki Matsunaga</author>
<author pid="34/8763">Yamato Ohtani</author>
<author pid="85/741">Tomoki Toda</author>
<title>A Cyclical Approach to Synthetic and Natural Speech Mismatch Refinement of Neural Post-filter for Low-cost Text-to-speech System.</title>
<year>2022</year>
<volume>abs/2207.05913</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2207.05913</ee>
<url>db/journals/corr/corr2207.html#abs-2207-05913</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2210-10314" mdate="2022-10-24">
<author pid="30/6718">Ding Ma</author>
<author pid="304/2864">Lester Phillip Violeta</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="85/741">Tomoki Toda</author>
<title>Two-stage training method for Japanese electrolaryngeal speech enhancement based on sequence-to-sequence voice conversion.</title>
<year>2022</year>
<volume>abs/2210.10314</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2210.10314</ee>
<url>db/journals/corr/corr2210.html#abs-2210-10314</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2210-15533" mdate="2022-11-02">
<author pid="290/1755">Reo Yoneyama</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="85/741">Tomoki Toda</author>
<title>Source-Filter HiFi-GAN: Fast and Pitch Controllable High-Fidelity Neural Vocoder.</title>
<year>2022</year>
<volume>abs/2210.15533</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2210.15533</ee>
<url>db/journals/corr/corr2210.html#abs-2210-15533</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2210-15987" mdate="2022-11-03">
<author pid="81/3672">Ryuichi Yamamoto</author>
<author pid="290/1755">Reo Yoneyama</author>
<author pid="85/741">Tomoki Toda</author>
<title>NNSVS: A Neural Network-Based Singing Voice Synthesis Toolkit.</title>
<year>2022</year>
<volume>abs/2210.15987</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2210.15987</ee>
<url>db/journals/corr/corr2210.html#abs-2210-15987</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2211-01079" mdate="2022-11-04">
<author pid="304/2864">Lester Phillip Violeta</author>
<author pid="30/6718">Ding Ma</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="85/741">Tomoki Toda</author>
<title>Intermediate Fine-Tuning Using Imperfect Synthetic Speech for Improving Electrolaryngeal Speech Recognition.</title>
<year>2022</year>
<volume>abs/2211.01079</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2211.01079</ee>
<url>db/journals/corr/corr2211.html#abs-2211-01079</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2211-01198" mdate="2022-11-04">
<author pid="216/3624">Takuya Fujimura</author>
<author pid="85/741">Tomoki Toda</author>
<title>Analysis of Noisy-target Training for DNN-based speech enhancement.</title>
<year>2022</year>
<volume>abs/2211.01198</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2211.01198</ee>
<url>db/journals/corr/corr2211.html#abs-2211-01198</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2211-07863" mdate="2022-11-23">
<author pid="334/0168">Yuka Hashizume</author>
<author pid="53/2189-63">Li Li 0063</author>
<author pid="85/741">Tomoki Toda</author>
<title>Music Similarity Calculation of Individual Instrumental Sounds Using Metric Learning.</title>
<year>2022</year>
<volume>abs/2211.07863</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2211.07863</ee>
<url>db/journals/corr/corr2211.html#abs-2211-07863</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2212-08321" mdate="2023-01-02">
<author pid="228/9342">Yusuke Yasuda</author>
<author pid="85/741">Tomoki Toda</author>
<title>Investigation of Japanese PnG BERT language model in text-to-speech synthesis for pitch accent language.</title>
<year>2022</year>
<volume>abs/2212.08321</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2212.08321</ee>
<url>db/journals/corr/corr2212.html#abs-2212-08321</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2212-08329" mdate="2023-01-02">
<author pid="228/9342">Yusuke Yasuda</author>
<author pid="85/741">Tomoki Toda</author>
<title>Text-to-speech synthesis based on latent variable conversion using diffusion probabilistic model and variational autoencoder.</title>
<year>2022</year>
<volume>abs/2212.08329</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2212.08329</ee>
<url>db/journals/corr/corr2212.html#abs-2212-08329</url>
</article>
</r>
<r><article key="journals/access/MatsubaraOTTTSK21" mdate="2021-09-16">
<author orcid="0000-0002-2935-668X" pid="157/7171">Keisuke Matsubara</author>
<author orcid="0000-0001-9913-4647" pid="132/9091">Takuma Okamoto</author>
<author orcid="0000-0002-9808-0250" pid="22/8758">Ryoichi Takashima</author>
<author orcid="0000-0001-5005-7679" pid="79/4485">Tetsuya Takiguchi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="02/2978">Yoshinori Shiga</author>
<author pid="32/4341">Hisashi Kawai</author>
<title>Full-Band LPCNet: A Real-Time Neural Vocoder for 48 kHz Audio With a CPU.</title>
<pages>94923-94933</pages>
<year>2021</year>
<volume>9</volume>
<journal>IEEE Access</journal>
<ee type="oa">https://doi.org/10.1109/ACCESS.2021.3089565</ee>
<url>db/journals/access/access9.html#MatsubaraOTTTSK21</url>
</article>
</r>
<r><article key="journals/taslp/KameokaHTKHT21" mdate="2021-02-11">
<author orcid="0000-0003-3102-0162" pid="97/941">Hirokazu Kameoka</author>
<author orcid="0000-0003-3172-3335" pid="225/7821">Wen-Chin Huang</author>
<author pid="140/2727">Kou Tanaka</author>
<author pid="119/1623">Takuhiro Kaneko</author>
<author pid="158/4110">Nobukatsu Hojo</author>
<author pid="85/741">Tomoki Toda</author>
<title>Many-to-Many Voice Transformer Network.</title>
<pages>656-670</pages>
<year>2021</year>
<volume>29</volume>
<journal>IEEE ACM Trans. Audio Speech Lang. Process.</journal>
<ee>https://doi.org/10.1109/TASLP.2020.3047262</ee>
<url>db/journals/taslp/taslp29.html#KameokaHTKHT21</url>
</article>
</r>
<r><article key="journals/taslp/HuangHWKT21" mdate="2023-07-27">
<author orcid="0000-0003-3172-3335" pid="225/7821">Wen-Chin Huang</author>
<author orcid="0000-0001-8782-4093" pid="82/8616">Tomoki Hayashi</author>
<author orcid="0000-0003-4390-1354" pid="188/5943">Yi-Chiao Wu</author>
<author orcid="0000-0003-3102-0162" pid="97/941">Hirokazu Kameoka</author>
<author pid="85/741">Tomoki Toda</author>
<title>Pretraining Techniques for Sequence-to-Sequence Voice Conversion.</title>
<pages>745-755</pages>
<year>2021</year>
<volume>29</volume>
<journal>IEEE ACM Trans. Audio Speech Lang. Process.</journal>
<ee type="oa">https://doi.org/10.1109/TASLP.2021.3049336</ee>
<url>db/journals/taslp/taslp29.html#HuangHWKT21</url>
</article>
</r>
<r><article key="journals/taslp/WuHOKT21" mdate="2025-01-19">
<author orcid="0000-0003-4390-1354" pid="188/5943">Yi-Chiao Wu</author>
<author orcid="0000-0001-8782-4093" pid="82/8616">Tomoki Hayashi</author>
<author pid="132/9091">Takuma Okamoto</author>
<author pid="32/4341">Hisashi Kawai</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Quasi-Periodic Parallel WaveGAN: A Non-Autoregressive Raw Waveform Generative Model With Pitch-Dependent Dilated Convolution Neural Network.</title>
<pages>792-806</pages>
<year>2021</year>
<volume>29</volume>
<journal>IEEE ACM Trans. Audio Speech Lang. Process.</journal>
<ee type="oa">https://doi.org/10.1109/TASLP.2021.3051765</ee>
<ee>https://www.wikidata.org/entity/Q131163702</ee>
<url>db/journals/taslp/taslp29.html#WuHOKT21</url>
</article>
</r>
<r><article key="journals/taslp/WuHTKT21" mdate="2025-01-19">
<author orcid="0000-0003-4390-1354" pid="188/5943">Yi-Chiao Wu</author>
<author orcid="0000-0001-8782-4093" pid="82/8616">Tomoki Hayashi</author>
<author orcid="0000-0003-2792-8418" pid="158/4210">Patrick Lumban Tobing</author>
<author orcid="0000-0001-5047-4165" pid="70/8239">Kazuhiro Kobayashi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Quasi-Periodic WaveNet: An Autoregressive Raw Waveform Generative Model With Pitch-Dependent Dilated Convolution Neural Network.</title>
<pages>1134-1148</pages>
<year>2021</year>
<volume>29</volume>
<journal>IEEE ACM Trans. Audio Speech Lang. Process.</journal>
<ee type="oa">https://doi.org/10.1109/TASLP.2021.3061245</ee>
<ee>https://www.wikidata.org/entity/Q131163716</ee>
<url>db/journals/taslp/taslp29.html#WuHTKT21</url>
</article>
</r>
<r><inproceedings key="conf/apsipa/QianNWKZT21" mdate="2025-08-07">
<author pid="193/6214">Zhaopeng Qian</author>
<author pid="36/2169">Haijun Niu</author>
<author pid="58/6810-111">Li Wang 0111</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="313/2304">Shaochuan Zhang</author>
<author pid="85/741">Tomoki Toda</author>
<title>Mandarin Electro-Laryngeal Speech Enhancement based on Statistical Voice Conversion and Manual Tone Control.</title>
<pages>546-552</pages>
<year>2021</year>
<booktitle>APSIPA ASC</booktitle>
<ee>https://ieeexplore.ieee.org/document/9689354</ee>
<crossref>conf/apsipa/2021</crossref>
<url>db/conf/apsipa/apsipa2021.html#QianNWKZT21</url>
</inproceedings>
</r>
<r><inproceedings key="conf/apsipa/XieWTHT21" mdate="2022-02-09">
<author pid="04/5576">Chao Xie</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="158/4210">Patrick Lumban Tobing</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="85/741">Tomoki Toda</author>
<title>Noisy-to-Noisy Voice Conversion Framework with Denoising Model.</title>
<pages>814-820</pages>
<year>2021</year>
<booktitle>APSIPA ASC</booktitle>
<ee>https://ieeexplore.ieee.org/document/9689325</ee>
<crossref>conf/apsipa/2021</crossref>
<url>db/conf/apsipa/apsipa2021.html#XieWTHT21</url>
</inproceedings>
</r>
<r><inproceedings key="conf/apsipa/MaHT21" mdate="2022-02-09">
<author pid="30/6718">Ding Ma</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="85/741">Tomoki Toda</author>
<title>Investigation of Text-to-Speech-based Synthetic Parallel Data for Sequence-to-Sequence Non-Parallel Voice Conversion.</title>
<pages>870-877</pages>
<year>2021</year>
<booktitle>APSIPA ASC</booktitle>
<ee>https://ieeexplore.ieee.org/document/9689675</ee>
<crossref>conf/apsipa/2021</crossref>
<url>db/conf/apsipa/apsipa2021.html#MaHT21</url>
</inproceedings>
</r>
<r><inproceedings key="conf/apsipa/LiouHYTPTTW21" mdate="2022-02-09">
<author pid="301/7716">Yi-Syuan Liou</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="155/7903">Ming-Chi Yen</author>
<author pid="301/7776">Shu-Wei Tsai</author>
<author pid="214/2303">Yu-Huai Peng</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="66/7146-1">Yu Tsao 0001</author>
<author pid="28/5019">Hsin-Min Wang</author>
<title>Time Alignment using Lip Images for Frame-based Electrolaryngeal Voice Conversion.</title>
<pages>1234-1238</pages>
<year>2021</year>
<booktitle>APSIPA ASC</booktitle>
<ee>https://ieeexplore.ieee.org/document/9689296</ee>
<crossref>conf/apsipa/2021</crossref>
<url>db/conf/apsipa/apsipa2021.html#LiouHYTPTTW21</url>
</inproceedings>
</r>
<r><inproceedings key="conf/asru/OkamotoTK21" mdate="2022-02-09">
<author pid="132/9091">Takuma Okamoto</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="32/4341">Hisashi Kawai</author>
<title>Multi-Stream HiFi-GAN with Data-Driven Waveform Decomposition.</title>
<pages>610-617</pages>
<year>2021</year>
<booktitle>ASRU</booktitle>
<ee>https://doi.org/10.1109/ASRU51503.2021.9688194</ee>
<crossref>conf/asru/2021</crossref>
<url>db/conf/asru/asru2021.html#OkamotoTK21</url>
</inproceedings>
</r>
<r><inproceedings key="conf/asru/HuangHLWT21" mdate="2023-03-21">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="69/8079">Xinjian Li</author>
<author orcid="0000-0002-5970-8631" pid="39/3245-1">Shinji Watanabe 0001</author>
<author pid="85/741">Tomoki Toda</author>
<title>On Prosody Modeling for ASR+TTS Based Voice Conversion.</title>
<pages>642-649</pages>
<year>2021</year>
<booktitle>ASRU</booktitle>
<ee>https://doi.org/10.1109/ASRU51503.2021.9688010</ee>
<crossref>conf/asru/2021</crossref>
<url>db/conf/asru/asru2021.html#HuangHLWT21</url>
</inproceedings>
</r>
<r><inproceedings key="conf/asru/YenHKPTTTJW21" mdate="2026-02-01">
<author pid="155/7903">Ming-Chi Yen</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="214/2303">Yu-Huai Peng</author>
<author pid="301/7776">Shu-Wei Tsai</author>
<author orcid="0000-0001-6956-0418" pid="66/7146-1">Yu Tsao 0001</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="81/829">Jyh-Shing Roger Jang</author>
<author orcid="0000-0003-3599-5071" pid="28/5019">Hsin-Min Wang</author>
<title>Mandarin Electrolaryngeal Speech Voice Conversion with Sequence-to-Sequence Modeling.</title>
<pages>650-657</pages>
<year>2021</year>
<booktitle>ASRU</booktitle>
<ee>https://doi.org/10.1109/ASRU51503.2021.9687908</ee>
<crossref>conf/asru/2021</crossref>
<url>db/conf/asru/asru2021.html#YenHKPTTTJW21</url>
</inproceedings>
</r>
<r><inproceedings key="conf/asru/ChiangWYTWHT21" mdate="2026-04-07">
<author pid="241/2983">Hsin-Tien Chiang</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="43/3340">Cheng Yu</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="28/5019">Hsin-Min Wang</author>
<author pid="51/6805">Yih-Chun Hu</author>
<author orcid="0000-0001-6956-0418" pid="66/7146-1">Yu Tsao 0001</author>
<title>HASA-Net: A Non-Intrusive Hearing-Aid Speech Assessment Network.</title>
<pages>907-913</pages>
<year>2021</year>
<booktitle>ASRU</booktitle>
<ee>https://doi.org/10.1109/ASRU51503.2021.9687972</ee>
<crossref>conf/asru/2021</crossref>
<url>db/conf/asru/asru2021.html#ChiangWYTWHT21</url>
</inproceedings>
</r>
<r><inproceedings key="conf/dcase/KuroyanagiHAYTT21" mdate="2021-12-21">
<author pid="294/8499">Ibuki Kuroyanagi</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="168/2958">Yusuke Adachi</author>
<author pid="173/6383">Takenori Yoshimura</author>
<author pid="38/3483">Kazuya Takeda</author>
<author pid="85/741">Tomoki Toda</author>
<title>An Ensemble Approach to Anomalous Sound Detection Based on Conformer-Based Autoencoder and Binary Classifier Incorporated with Metric Learning.</title>
<pages>110-114</pages>
<year>2021</year>
<booktitle>DCASE</booktitle>
<ee type="oa">http://dcase.community/documents/workshop2021/proceedings/DCASE2021Workshop_Kuroyanagi_40.pdf</ee>
<crossref>conf/dcase/2021</crossref>
<url>db/conf/dcase/dcase2021.html#KuroyanagiHAYTT21</url>
</inproceedings>
</r>
<r><inproceedings key="conf/eusipco/KuroyanagiHTT21" mdate="2021-12-09">
<author pid="294/8499">Ibuki Kuroyanagi</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="38/3483">Kazuya Takeda</author>
<author pid="85/741">Tomoki Toda</author>
<title>Anomalous Sound Detection Using a Binary Classification Model and Class Centroids.</title>
<pages>1995-1999</pages>
<year>2021</year>
<booktitle>EUSIPCO</booktitle>
<ee>https://doi.org/10.23919/EUSIPCO54536.2021.9616198</ee>
<crossref>conf/eusipco/2021</crossref>
<url>db/conf/eusipco/eusipco2021.html#KuroyanagiHTT21</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/KobayashiHWTHT21" mdate="2023-03-21">
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author orcid="0000-0003-2792-8418" pid="158/4210">Patrick Lumban Tobing</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="85/741">Tomoki Toda</author>
<title>Crank: An Open-Source Software for Nonparallel Voice Conversion Based on Vector-Quantized Variational Autoencoder.</title>
<pages>5934-5938</pages>
<year>2021</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP39728.2021.9413959</ee>
<crossref>conf/icassp/2021</crossref>
<url>db/conf/icassp/icassp2021.html#KobayashiHWTHT21</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/OkamotoTSK21" mdate="2021-07-09">
<author pid="132/9091">Takuma Okamoto</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="02/2978">Yoshinori Shiga</author>
<author pid="32/4341">Hisashi Kawai</author>
<title>Noise Level Limited Sub-Modeling for Diffusion Probabilistic Vocoders.</title>
<pages>6029-6033</pages>
<year>2021</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP39728.2021.9415087</ee>
<crossref>conf/icassp/2021</crossref>
<url>db/conf/icassp/icassp2021.html#OkamotoTSK21</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/AndoMSMAIT21" mdate="2026-02-12">
<author pid="173/6654">Atsushi Ando</author>
<author pid="08/10650">Ryo Masumura</author>
<author pid="55/6900-2">Hiroshi Sato 0002</author>
<author pid="175/8811">Takafumi Moriya</author>
<author pid="260/4522">Takanori Ashihara</author>
<author pid="67/8052">Yusuke Ijima</author>
<author pid="85/741">Tomoki Toda</author>
<title>Speech Emotion Recognition Based on Listener Adaptive Models.</title>
<pages>6274-6278</pages>
<year>2021</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP39728.2021.9414698</ee>
<crossref>conf/icassp/2021</crossref>
<url>db/conf/icassp/icassp2021.html#AndoMSMAIT21</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/MatsubaraOTTTSK21" mdate="2021-07-09">
<author pid="157/7171">Keisuke Matsubara</author>
<author pid="132/9091">Takuma Okamoto</author>
<author pid="22/8758">Ryoichi Takashima</author>
<author pid="79/4485">Tetsuya Takiguchi</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="02/2978">Yoshinori Shiga</author>
<author pid="32/4341">Hisashi Kawai</author>
<title>High-Intelligibility Speech Synthesis for Dysarthric Speakers with LPCNet-Based TTS and CycleVAE-Based VC.</title>
<pages>7058-7062</pages>
<year>2021</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP39728.2021.9414136</ee>
<crossref>conf/icassp/2021</crossref>
<url>db/conf/icassp/icassp2021.html#MatsubaraOTTTSK21</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/HayashiHKT21" mdate="2021-07-09">
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="85/741">Tomoki Toda</author>
<title>Non-Autoregressive Sequence-To-Sequence Voice Conversion.</title>
<pages>7068-7072</pages>
<year>2021</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP39728.2021.9413973</ee>
<crossref>conf/icassp/2021</crossref>
<url>db/conf/icassp/icassp2021.html#HayashiHKT21</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/HuangWLCWT21" mdate="2025-02-12">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="160/2718">Chia-Hua Wu</author>
<author pid="234/6347">Shang-Bao Luo</author>
<author pid="35/6313-2">Kuan-Yu Chen 0002</author>
<author orcid="0000-0003-3599-5071" pid="28/5019">Hsin-Min Wang</author>
<author pid="85/741">Tomoki Toda</author>
<title>Speech Recognition by Simply Fine-Tuning Bert.</title>
<pages>7343-7347</pages>
<year>2021</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP39728.2021.9413668</ee>
<crossref>conf/icassp/2021</crossref>
<url>db/conf/icassp/icassp2021.html#HuangWLCWT21</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/HuangKPLTWT21" mdate="2026-02-01">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="214/2303">Yu-Huai Peng</author>
<author pid="196/4607">Ching-Feng Liu</author>
<author orcid="0000-0001-6956-0418" pid="66/7146-1">Yu Tsao 0001</author>
<author orcid="0000-0003-3599-5071" pid="28/5019">Hsin-Min Wang</author>
<author pid="85/741">Tomoki Toda</author>
<title>A Preliminary Study of a Two-Stage Paradigm for Preserving Speaker Identity in Dysarthric Voice Conversion.</title>
<pages>1329-1333</pages>
<year>2021</year>
<booktitle>Interspeech</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2021-208</ee>
<crossref>conf/interspeech/2021</crossref>
<url>db/conf/interspeech/interspeech2021.html#HuangKPLTWT21</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/YoneyamaWT21" mdate="2023-06-21">
<author pid="290/1755">Reo Yoneyama</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="85/741">Tomoki Toda</author>
<title>Unified Source-Filter GAN: Unified Source-Filter Network Based On Factorization of Quasi-Periodic Parallel WaveGAN.</title>
<pages>2187-2191</pages>
<year>2021</year>
<booktitle>Interspeech</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2021-517</ee>
<crossref>conf/interspeech/2021</crossref>
<url>db/conf/interspeech/interspeech2021.html#YoneyamaWT21</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/TobingT21" mdate="2023-06-26">
<author orcid="0000-0003-2792-8418" pid="158/4210">Patrick Lumban Tobing</author>
<author pid="85/741">Tomoki Toda</author>
<title>High-Fidelity and Low-Latency Universal Neural Vocoder Based on Multiband WaveRNN with Data-Driven Linear Prediction for Discrete Waveform Modeling.</title>
<pages>2217-2221</pages>
<year>2021</year>
<booktitle>Interspeech</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2021-1984</ee>
<crossref>conf/interspeech/2021</crossref>
<url>db/conf/interspeech/interspeech2021.html#TobingT21</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/WuHLPHTWT21" mdate="2026-02-01">
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="156/0822">Cheng-Hung Hu</author>
<author pid="13/8052">Hung-Shin Lee</author>
<author pid="214/2303">Yu-Huai Peng</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author orcid="0000-0001-6956-0418" pid="66/7146-1">Yu Tsao 0001</author>
<author orcid="0000-0003-3599-5071" pid="28/5019">Hsin-Min Wang</author>
<author pid="85/741">Tomoki Toda</author>
<title>Relational Data Selection for Data Augmentation of Speaker-Dependent Multi-Band MelGAN Vocoder.</title>
<pages>3630-3634</pages>
<year>2021</year>
<booktitle>Interspeech</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2021-806</ee>
<crossref>conf/interspeech/2021</crossref>
<url>db/conf/interspeech/interspeech2021.html#WuHLPHTWT21</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mlsp/SekiTT21" mdate="2021-11-23">
<author pid="194/1307">Shogo Seki</author>
<author pid="306/9710">Haruka Taga</author>
<author pid="85/741">Tomoki Toda</author>
<title>Singing Fundamental Frequency Contour Generation Using Generalized Command-Response Model and Score-Conditional Variational Autoencoder.</title>
<pages>1-3</pages>
<year>2021</year>
<booktitle>MLSP</booktitle>
<ee>https://doi.org/10.1109/MLSP52302.2021.9596428</ee>
<crossref>conf/mlsp/2021</crossref>
<url>db/conf/mlsp/mlsp2021.html#SekiTT21</url>
</inproceedings>
</r>
<r><inproceedings key="conf/ssw/TobingT21" mdate="2024-08-01">
<author pid="158/4210">Patrick Lumban Tobing</author>
<author pid="85/741">Tomoki Toda</author>
<title>Low-latency real-time non-parallel voice conversion based on cyclic variational autoencoder and multiband WaveRNN with data-driven linear prediction.</title>
<pages>142-147</pages>
<year>2021</year>
<booktitle>SSW</booktitle>
<ee type="oa">https://doi.org/10.21437/SSW.2021-25</ee>
<crossref>conf/ssw/2021</crossref>
<url>db/conf/ssw/ssw2021.html#TobingT21</url>
</inproceedings>
</r>
<r><article publtype="informal" key="journals/corr/abs-2102-00291" mdate="2025-02-12">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="160/2718">Chia-Hua Wu</author>
<author pid="234/6347">Shang-Bao Luo</author>
<author pid="35/6313-2">Kuan-Yu Chen 0002</author>
<author pid="28/5019">Hsin-Min Wang</author>
<author pid="85/741">Tomoki Toda</author>
<title>Speech Recognition by Simply Fine-tuning BERT.</title>
<year>2021</year>
<volume>abs/2102.00291</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2102.00291</ee>
<url>db/journals/corr/corr2102.html#abs-2102-00291</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2103-02858" mdate="2021-03-16">
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="158/4210">Patrick Lumban Tobing</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="85/741">Tomoki Toda</author>
<title>crank: An Open-Source Software for Nonparallel Voice Conversion Based on Vector-Quantized Variational Autoencoder.</title>
<year>2021</year>
<volume>abs/2103.02858</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2103.02858</ee>
<url>db/journals/corr/corr2103.html#abs-2103-02858</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2104-03009" mdate="2026-07-16">
<author pid="156/0822">Cheng-Hung Hu</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="214/2303">Yu-Huai Peng</author>
<author pid="49/8346-6">Yu-Wen Chen 0006</author>
<author pid="289/7336">Pin-Jui Ku</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="66/7146-1">Yu Tsao 0001</author>
<author pid="28/5019">Hsin-Min Wang</author>
<title>The AS-NU System for the M2VoC Challenge.</title>
<year>2021</year>
<volume>abs/2104.03009</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2104.03009</ee>
<url>db/journals/corr/corr2104.html#abs-2104-03009</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2104-04668" mdate="2021-04-19">
<author pid="290/1755">Reo Yoneyama</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="85/741">Tomoki Toda</author>
<title>Unified Source-Filter GAN: Unified Source-filter Network Based On Factorization of Quasi-Periodic Parallel WaveGAN.</title>
<year>2021</year>
<volume>abs/2104.04668</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2104.04668</ee>
<url>db/journals/corr/corr2104.html#abs-2104-04668</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2104-06793" mdate="2021-04-19">
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="85/741">Tomoki Toda</author>
<title>Non-autoregressive sequence-to-sequence voice conversion.</title>
<year>2021</year>
<volume>abs/2104.06793</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2104.06793</ee>
<url>db/journals/corr/corr2104.html#abs-2104-06793</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2105-09856" mdate="2021-05-31">
<author pid="158/4210">Patrick Lumban Tobing</author>
<author pid="85/741">Tomoki Toda</author>
<title>High-Fidelity and Low-Latency Universal Neural Vocoder based on Multiband WaveRNN with Data-Driven Linear Prediction for Discrete Waveform Modeling.</title>
<year>2021</year>
<volume>abs/2105.09856</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2105.09856</ee>
<url>db/journals/corr/corr2105.html#abs-2105-09856</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2105-09858" mdate="2021-05-31">
<author pid="158/4210">Patrick Lumban Tobing</author>
<author pid="85/741">Tomoki Toda</author>
<title>Low-Latency Real-Time Non-Parallel Voice Conversion based on Cyclic Variational Autoencoder and Multiband WaveRNN with Data-Driven Linear Prediction.</title>
<year>2021</year>
<volume>abs/2105.09858</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2105.09858</ee>
<url>db/journals/corr/corr2105.html#abs-2105-09858</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2106-01415" mdate="2021-06-10">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="214/2303">Yu-Huai Peng</author>
<author pid="196/4607">Ching-Feng Liu</author>
<author pid="66/7146-1">Yu Tsao 0001</author>
<author pid="28/5019">Hsin-Min Wang</author>
<author pid="85/741">Tomoki Toda</author>
<title>A Preliminary Study of a Two-Stage Paradigm for Preserving Speaker Identity in Dysarthric Voice Conversion.</title>
<year>2021</year>
<volume>abs/2106.01415</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2106.01415</ee>
<url>db/journals/corr/corr2106.html#abs-2106-01415</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2106-06151" mdate="2021-06-15">
<author pid="294/8499">Ibuki Kuroyanagi</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="38/3483">Kazuya Takeda</author>
<author pid="85/741">Tomoki Toda</author>
<title>Anomalous Sound Detection Using a Binary Classification Model and Class Centroids.</title>
<year>2021</year>
<volume>abs/2106.06151</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2106.06151</ee>
<url>db/journals/corr/corr2106.html#abs-2106-06151</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2107-09477" mdate="2021-07-29">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="69/8079">Xinjian Li</author>
<author pid="39/3245-1">Shinji Watanabe 0001</author>
<author pid="85/741">Tomoki Toda</author>
<title>On Prosody Modeling for ASR+TTS based Voice Conversion.</title>
<year>2021</year>
<volume>abs/2107.09477</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2107.09477</ee>
<url>db/journals/corr/corr2107.html#abs-2107-09477</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2109-03551" mdate="2022-10-12">
<author pid="301/7716">Yi-Syuan Liou</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="155/7903">Ming-Chi Yen</author>
<author pid="301/7776">Shu-Wei Tsai</author>
<author pid="214/2303">Yu-Huai Peng</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="66/7146-1">Yu Tsao 0001</author>
<author pid="28/5019">Hsin-Min Wang</author>
<title>Time Alignment using Lip Images for Frame-based Electrolaryngeal Voice Conversion.</title>
<year>2021</year>
<volume>abs/2109.03551</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2109.03551</ee>
<url>db/journals/corr/corr2109.html#abs-2109-03551</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2109-10608" mdate="2021-09-27">
<author pid="04/5576">Chao Xie</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="158/4210">Patrick Lumban Tobing</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="85/741">Tomoki Toda</author>
<title>Noisy-to-Noisy Voice Conversion Framework with Denoising Model.</title>
<year>2021</year>
<volume>abs/2109.10608</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2109.10608</ee>
<url>db/journals/corr/corr2109.html#abs-2109-10608</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2110-06280" mdate="2021-10-22">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="246/0774">Shu-Wen Yang</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="81/8056">Hung-Yi Lee</author>
<author pid="39/3245-1">Shinji Watanabe 0001</author>
<author pid="85/741">Tomoki Toda</author>
<title>S3PRL-VC: Open-source Voice Conversion Framework with Self-supervised Speech Representations.</title>
<year>2021</year>
<volume>abs/2110.06280</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2110.06280</ee>
<url>db/journals/corr/corr2110.html#abs-2110-06280</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2110-08213" mdate="2021-10-22">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="271/4266">Bence Mark Halpern</author>
<author pid="304/2864">Lester Phillip Violeta</author>
<author pid="87/1365">Odette Scharenborg</author>
<author pid="85/741">Tomoki Toda</author>
<title>Towards Identity Preserving Normal to Dysarthric Voice Conversion.</title>
<year>2021</year>
<volume>abs/2110.08213</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2110.08213</ee>
<url>db/journals/corr/corr2110.html#abs-2110-08213</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2110-09103" mdate="2021-10-22">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="03/7184">Erica Cooper</author>
<author pid="87/3979">Junichi Yamagishi</author>
<author pid="85/741">Tomoki Toda</author>
<title>LDNet: Unified Listener Dependent Modeling in MOS Prediction for Synthetic Speech.</title>
<year>2021</year>
<volume>abs/2110.09103</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2110.09103</ee>
<url>db/journals/corr/corr2110.html#abs-2110-09103</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2111-05691" mdate="2021-11-16">
<author pid="241/2983">Hsin-Tien Chiang</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="43/3340">Cheng Yu</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="28/5019">Hsin-Min Wang</author>
<author pid="51/6805">Yih-Chun Hu</author>
<author pid="66/7146-1">Yu Tsao 0001</author>
<title>HASA-net: A non-intrusive hearing-aid speech assessment network.</title>
<year>2021</year>
<volume>abs/2111.05691</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2111.05691</ee>
<url>db/journals/corr/corr2111.html#abs-2111-05691</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2111-07116" mdate="2021-11-16">
<author pid="04/5576">Chao Xie</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="158/4210">Patrick Lumban Tobing</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="85/741">Tomoki Toda</author>
<title>Direct Noisy Speech Modeling for Noisy-to-Noisy Voice Conversion.</title>
<year>2021</year>
<volume>abs/2111.07116</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2111.07116</ee>
<url>db/journals/corr/corr2111.html#abs-2111-07116</url>
</article>
</r>
<r><article key="journals/access/WuTKHT20" mdate="2021-04-09">
<author orcid="0000-0003-4390-1354" pid="188/5943">Yi-Chiao Wu</author>
<author orcid="0000-0003-2792-8418" pid="158/4210">Patrick Lumban Tobing</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Non-Parallel Voice Conversion System With WaveNet Vocoder and Collapsed Speech Suppression.</title>
<pages>62094-62106</pages>
<year>2020</year>
<volume>8</volume>
<journal>IEEE Access</journal>
<ee type="oa">https://doi.org/10.1109/ACCESS.2020.2984007</ee>
<url>db/journals/access/access8.html#WuTKHT20</url>
</article>
</r>
<r><article key="journals/taslp/AndoMKKAT20" mdate="2021-04-09">
<author orcid="0000-0002-3971-0654" pid="173/6654">Atsushi Ando</author>
<author pid="08/10650">Ryo Masumura</author>
<author pid="212/6389">Hosana Kamiyama</author>
<author pid="09/3769">Satoshi Kobashikawa</author>
<author pid="175/7791">Yushi Aono</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Customer Satisfaction Estimation in Contact Center Calls Based on a Hierarchical Multi-Task Model.</title>
<pages>715-728</pages>
<year>2020</year>
<volume>28</volume>
<journal>IEEE ACM Trans. Audio Speech Lang. Process.</journal>
<ee>https://doi.org/10.1109/TASLP.2020.2966857</ee>
<url>db/journals/taslp/taslp28.html#AndoMKKAT20</url>
</article>
</r>
<r><inproceedings key="conf/apsipa/NakataniTTT20" mdate="2021-02-11">
<author pid="278/9159">Hikaru Nakatani</author>
<author pid="158/4210">Patrick Lumban Tobing</author>
<author pid="38/3483">Kazuya Takeda</author>
<author pid="85/741">Tomoki Toda</author>
<title>Cross-Lingual Voice Conversion using a Cyclic Variational Auto-encoder and a WaveNet Vocoder.</title>
<pages>520-526</pages>
<year>2020</year>
<booktitle>APSIPA</booktitle>
<ee>https://ieeexplore.ieee.org/document/9306335</ee>
<crossref>conf/apsipa/2020</crossref>
<url>db/conf/apsipa/apsipa2020.html#NakataniTTT20</url>
</inproceedings>
</r>
<r><inproceedings key="conf/apsipa/EshghiKTKT20" mdate="2021-02-11">
<author pid="42/3681">Mohammad Eshghi</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="140/2727">Kou Tanaka</author>
<author pid="97/941">Hirokazu Kameoka</author>
<author pid="85/741">Tomoki Toda</author>
<title>Phoneme Embeddings on Predicting Fundamental Frequency Pattern for Electrolaryngeal Speech.</title>
<pages>572-577</pages>
<year>2020</year>
<booktitle>APSIPA</booktitle>
<ee>https://ieeexplore.ieee.org/document/9306228</ee>
<crossref>conf/apsipa/2020</crossref>
<url>db/conf/apsipa/apsipa2020.html#EshghiKTKT20</url>
</inproceedings>
</r>
<r><inproceedings key="conf/blizzard/0006HTYDKLT20" mdate="2024-09-18">
<author pid="51/4138-6">Yi Zhao 0006</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="153/0728">Xiaohai Tian</author>
<author pid="87/3979">Junichi Yamagishi</author>
<author pid="158/4101">Rohan Kumar Das</author>
<author pid="94/4754">Tomi Kinnunen</author>
<author pid="70/5210">Zhen-Hua Ling</author>
<author pid="85/741">Tomoki Toda</author>
<title>Voice Conversion Challenge 2020 -- Intra-lingual semi-parallel and cross-lingual voice conversion --.</title>
<year>2020</year>
<booktitle>Blizzard Challenge / Voice Conversion Challenge</booktitle>
<ee type="oa">https://doi.org/10.21437/VCCBC.2020-14</ee>
<crossref>conf/blizzard/2020</crossref>
<url>db/conf/blizzard/blizzard2020.html#0006HTYDKLT20</url>
</inproceedings>
</r>
<r><inproceedings key="conf/blizzard/DasKHLY0TT20" mdate="2024-09-18">
<author pid="158/4101">Rohan Kumar Das</author>
<author pid="94/4754">Tomi Kinnunen</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="70/5210">Zhen-Hua Ling</author>
<author pid="87/3979">Junichi Yamagishi</author>
<author pid="51/4138-6">Yi Zhao 0006</author>
<author pid="153/0728">Xiaohai Tian</author>
<author pid="85/741">Tomoki Toda</author>
<title>Predictions of Subjective Ratings and Spoofing Assessments of Voice Conversion Challenge 2020 Submissions.</title>
<year>2020</year>
<booktitle>Blizzard Challenge / Voice Conversion Challenge</booktitle>
<ee type="oa">https://doi.org/10.21437/VCCBC.2020-15</ee>
<crossref>conf/blizzard/2020</crossref>
<url>db/conf/blizzard/blizzard2020.html#DasKHLY0TT20</url>
</inproceedings>
</r>
<r><inproceedings key="conf/blizzard/HuangH0T20" mdate="2024-09-18">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="39/3245-1">Shinji Watanabe 0001</author>
<author pid="85/741">Tomoki Toda</author>
<title>The Sequence-to-Sequence Baseline for the Voice Conversion Challenge 2020: Cascading ASR and TTS.</title>
<year>2020</year>
<booktitle>Blizzard Challenge / Voice Conversion Challenge</booktitle>
<ee type="oa">https://doi.org/10.21437/VCCBC.2020-24</ee>
<crossref>conf/blizzard/2020</crossref>
<url>db/conf/blizzard/blizzard2020.html#HuangH0T20</url>
</inproceedings>
</r>
<r><inproceedings key="conf/blizzard/HuangTWKT20" mdate="2024-09-18">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="158/4210">Patrick Lumban Tobing</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="85/741">Tomoki Toda</author>
<title>The NU Voice Conversion System for the Voice Conversion Challenge 2020: On the Effectiveness of Sequence-to-sequence Models and Autoregressive Neural Vocoders.</title>
<year>2020</year>
<booktitle>Blizzard Challenge / Voice Conversion Challenge</booktitle>
<ee type="oa">https://doi.org/10.21437/VCCBC.2020-25</ee>
<crossref>conf/blizzard/2020</crossref>
<url>db/conf/blizzard/blizzard2020.html#HuangTWKT20</url>
</inproceedings>
</r>
<r><inproceedings key="conf/blizzard/TobingWT20" mdate="2024-09-18">
<author pid="158/4210">Patrick Lumban Tobing</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="85/741">Tomoki Toda</author>
<title>Baseline System of Voice Conversion Challenge 2020 with Cyclic Variational Autoencoder and Parallel WaveGAN.</title>
<year>2020</year>
<booktitle>Blizzard Challenge / Voice Conversion Challenge</booktitle>
<ee type="oa">https://doi.org/10.21437/VCCBC.2020-23</ee>
<crossref>conf/blizzard/2020</crossref>
<url>db/conf/blizzard/blizzard2020.html#TobingWT20</url>
</inproceedings>
</r>
<r><inproceedings key="conf/dcase/MiyazakiKHWTT20" mdate="2021-12-21">
<author pid="27/9637">Koichi Miyazaki</author>
<author pid="136/5113">Tatsuya Komatsu</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="39/3245-1">Shinji Watanabe 0001</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="38/3483">Kazuya Takeda</author>
<title>Conformer-Based Sound Event Detection with Semi-Supervised Learning and Data Augmentation.</title>
<pages>100-104</pages>
<year>2020</year>
<booktitle>DCASE</booktitle>
<ee type="oa">http://dcase.community/documents/workshop2020/proceedings/DCASE2020Workshop_Miyazaki_92.pdf</ee>
<crossref>conf/dcase/2020</crossref>
<url>db/conf/dcase/dcase2020.html#MiyazakiKHWTT20</url>
</inproceedings>
</r>
<r><inproceedings key="conf/eusipco/KobayashiT20" mdate="2021-01-08">
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="85/741">Tomoki Toda</author>
<title>Implementation of low-latency electrolaryngeal speech enhancement based on multi-task CLDNN.</title>
<pages>396-400</pages>
<year>2020</year>
<booktitle>EUSIPCO</booktitle>
<ee>https://doi.org/10.23919/Eusipco47968.2020.9287721</ee>
<crossref>conf/eusipco/2020</crossref>
<url>db/conf/eusipco/eusipco2020.html#KobayashiT20</url>
</inproceedings>
</r>
<r><inproceedings key="conf/eusipco/TakadaSTT20" mdate="2021-03-15">
<author pid="237/0010">Moe Takada</author>
<author pid="194/1307">Shogo Seki</author>
<author pid="158/4210">Patrick Lumban Tobing</author>
<author pid="85/741">Tomoki Toda</author>
<title>Semi-Supervised Enhancement and Suppression of Self-Produced Speech Using Correspondence between Air- and Body-Conducted Signals.</title>
<pages>456-460</pages>
<year>2020</year>
<booktitle>EUSIPCO</booktitle>
<ee>https://doi.org/10.23919/Eusipco47968.2020.9287512</ee>
<crossref>conf/eusipco/2020</crossref>
<url>db/conf/eusipco/eusipco2020.html#TakadaSTT20</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/MiyazakiKH0TT20" mdate="2021-04-09">
<author pid="27/9637">Koichi Miyazaki</author>
<author pid="136/5113">Tatsuya Komatsu</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author orcid="0000-0002-5970-8631" pid="39/3245-1">Shinji Watanabe 0001</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="38/3483">Kazuya Takeda</author>
<title>Weakly-Supervised Sound Event Detection with Self-Attention.</title>
<pages>66-70</pages>
<year>2020</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP40776.2020.9053609</ee>
<crossref>conf/icassp/2020</crossref>
<url>db/conf/icassp/icassp2020.html#MiyazakiKH0TT20</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/OkamotoTSK20" mdate="2021-04-09">
<author pid="132/9091">Takuma Okamoto</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="02/2978">Yoshinori Shiga</author>
<author pid="32/4341">Hisashi Kawai</author>
<title>Transformer-Based Text-to-Speech with Weighted Forced Attention.</title>
<pages>6729-6733</pages>
<year>2020</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP40776.2020.9053915</ee>
<crossref>conf/icassp/2020</crossref>
<url>db/conf/icassp/icassp2020.html#OkamotoTSK20</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/TobingWHKT20" mdate="2023-03-21">
<author orcid="0000-0003-2792-8418" pid="158/4210">Patrick Lumban Tobing</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Efficient Shallow Wavenet Vocoder Using Multiple Samples Output Based on Laplacian Distribution and Linear Prediction.</title>
<pages>7204-7208</pages>
<year>2020</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP40776.2020.9053991</ee>
<crossref>conf/icassp/2020</crossref>
<url>db/conf/icassp/icassp2020.html#TobingWHKT20</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/HayashiYIY0TTZT20" mdate="2025-04-01">
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="81/3672">Ryuichi Yamamoto</author>
<author pid="214/2329">Katsuki Inoue</author>
<author pid="173/6383">Takenori Yoshimura</author>
<author orcid="0000-0002-5970-8631" pid="39/3245-1">Shinji Watanabe 0001</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="38/3483">Kazuya Takeda</author>
<author pid="50/671-33">Yu Zhang 0033</author>
<author orcid="0000-0001-5631-0639" pid="96/10484-3">Xu Tan 0003</author>
<title>Espnet-TTS: Unified, Reproducible, and Integratable Open Source End-to-End Text-to-Speech Toolkit.</title>
<pages>7654-7658</pages>
<year>2020</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP40776.2020.9053512</ee>
<crossref>conf/icassp/2020</crossref>
<url>db/conf/icassp/icassp2020.html#HayashiYIY0TTZT20</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/WuHOKT20" mdate="2021-04-09">
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="132/9091">Takuma Okamoto</author>
<author pid="32/4341">Hisashi Kawai</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Quasi-Periodic Parallel WaveGAN Vocoder: A Non-Autoregressive Pitch-Dependent Dilated Convolution Model for Parametric Speech Generation.</title>
<pages>3535-3539</pages>
<year>2020</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2020-1070</ee>
<crossref>conf/interspeech/2020</crossref>
<url>db/conf/interspeech/interspeech2020.html#WuHOKT20</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/WuTYMOT20" mdate="2023-06-26">
<author pid="188/5943">Yi-Chiao Wu</author>
<author orcid="0000-0003-2792-8418" pid="158/4210">Patrick Lumban Tobing</author>
<author pid="265/6506">Kazuki Yasuhara</author>
<author pid="265/6508">Noriyuki Matsunaga</author>
<author pid="34/8763">Yamato Ohtani</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>A Cyclical Post-Filtering Approach to Mismatch Refinement of Neural Vocoder for Text-to-Speech Systems.</title>
<pages>3540-3544</pages>
<year>2020</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2020-1072</ee>
<crossref>conf/interspeech/2020</crossref>
<url>db/conf/interspeech/interspeech2020.html#WuTYMOT20</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/SekiTT20" mdate="2021-04-09">
<author pid="194/1307">Shogo Seki</author>
<author pid="237/0010">Moe Takada</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Semi-Supervised Self-Produced Speech Enhancement and Suppression Based on Joint Source Modeling of Air- and Body-Conducted Signals Using Variational Autoencoder.</title>
<pages>4039-4043</pages>
<year>2020</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2020-2055</ee>
<crossref>conf/interspeech/2020</crossref>
<url>db/conf/interspeech/interspeech2020.html#SekiTT20</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/HikosakaSHKTBT20" mdate="2021-04-09">
<author pid="277/3528">Shu Hikosaka</author>
<author pid="194/1307">Shogo Seki</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="38/3483">Kazuya Takeda</author>
<author pid="95/6268">Hideki Banno</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Intelligibility Enhancement Based on Speech Waveform Modification Using Hearing Impairment.</title>
<pages>4059-4063</pages>
<year>2020</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2020-2062</ee>
<crossref>conf/interspeech/2020</crossref>
<url>db/conf/interspeech/interspeech2020.html#HikosakaSHKTBT20</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/HuangHWKT20" mdate="2021-04-09">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="97/941">Hirokazu Kameoka</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Voice Transformer Network: Sequence-to-Sequence Voice Conversion Using Transformer with Text-to-Speech Pretraining.</title>
<pages>4676-4680</pages>
<year>2020</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2020-1066</ee>
<crossref>conf/interspeech/2020</crossref>
<url>db/conf/interspeech/interspeech2020.html#HuangHWKT20</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/TobingHWKT20" mdate="2023-06-26">
<author orcid="0000-0003-2792-8418" pid="158/4210">Patrick Lumban Tobing</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Cyclic Spectral Modeling for Unsupervised Unit Discovery into Voice Conversion with Excitation and Waveform Modeling.</title>
<pages>4861-4865</pages>
<year>2020</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2020-2559</ee>
<crossref>conf/interspeech/2020</crossref>
<url>db/conf/interspeech/interspeech2020.html#TobingHWKT20</url>
</inproceedings>
</r>
<r><proceedings key="conf/blizzard/2020" mdate="2026-01-13">
<editor pid="87/3979">Junichi Yamagishi</editor>
<editor pid="70/5210">Zhenhua Ling</editor>
<editor pid="158/4101">Rohan Kumar Das</editor>
<editor pid="68/2005">Simon King 0001</editor>
<editor pid="94/4754">Tomi Kinnunen</editor>
<editor pid="85/741">Tomoki Toda</editor>
<editor pid="225/7821">Wen-Chin Huang</editor>
<editor pid="267/2864-24">Xiao Zhou 0024</editor>
<editor pid="153/0728">Xiaohai Tian</editor>
<editor pid="51/4138-6">Yi Zhao 0006</editor>
<title>Joint Workshop for the Blizzard Challenge and Voice Conversion Challenge 2020, Shanghai, China, October 30, 2020</title>
<booktitle>Blizzard Challenge / Voice Conversion Challenge</booktitle>
<publisher>ISCA</publisher>
<year>2020</year>
<ee type="oa">https://doi.org/10.21437/VCCBC.2020</ee>
<url>db/conf/blizzard/blizzard2020.html</url>
</proceedings>
</r>
<r><article publtype="informal" key="journals/corr/abs-2003-11750" mdate="2020-04-02">
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="158/4210">Patrick Lumban Tobing</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="85/741">Tomoki Toda</author>
<title>Non-parallel Voice Conversion System with WaveNet Vocoder and Collapsed Speech Suppression.</title>
<year>2020</year>
<volume>abs/2003.11750</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2003.11750</ee>
<url>db/journals/corr/corr2003.html#abs-2003-11750</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2005-08445" mdate="2020-05-22">
<author pid="97/941">Hirokazu Kameoka</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="140/2727">Kou Tanaka</author>
<author pid="119/1623">Takuhiro Kaneko</author>
<author pid="158/4110">Nobukatsu Hojo</author>
<author pid="85/741">Tomoki Toda</author>
<title>Many-to-Many Voice Transformer Network.</title>
<year>2020</year>
<volume>abs/2005.08445</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2005.08445</ee>
<url>db/journals/corr/corr2005.html#abs-2005-08445</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2005-08654" mdate="2020-05-22">
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="132/9091">Takuma Okamoto</author>
<author pid="32/4341">Hisashi Kawai</author>
<author pid="85/741">Tomoki Toda</author>
<title>Quasi-Periodic Parallel WaveGAN Vocoder: A Non-autoregressive Pitch-dependent Dilated Convolution Model for Parametric Speech Generation.</title>
<year>2020</year>
<volume>abs/2005.08654</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2005.08654</ee>
<url>db/journals/corr/corr2005.html#abs-2005-08654</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2005-08659" mdate="2020-05-22">
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="158/4210">Patrick Lumban Tobing</author>
<author pid="265/6506">Kazuki Yasuhara</author>
<author pid="265/6508">Noriyuki Matsunaga</author>
<author pid="34/8763">Yamato Ohtani</author>
<author pid="85/741">Tomoki Toda</author>
<title>A Cyclical Post-filtering Approach to Mismatch Refinement of Neural Vocoder for Text-to-speech Systems.</title>
<year>2020</year>
<volume>abs/2005.08659</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2005.08659</ee>
<url>db/journals/corr/corr2005.html#abs-2005-08659</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2007-05663" mdate="2020-07-22">
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="158/4210">Patrick Lumban Tobing</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="85/741">Tomoki Toda</author>
<title>Quasi-Periodic WaveNet: An Autoregressive Raw Waveform Generative Model with Pitch-dependent Dilated Convolution Neural Network.</title>
<year>2020</year>
<volume>abs/2007.05663</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2007.05663</ee>
<url>db/journals/corr/corr2007.html#abs-2007-05663</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2007-12955" mdate="2020-07-29">
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="132/9091">Takuma Okamoto</author>
<author pid="32/4341">Hisashi Kawai</author>
<author pid="85/741">Tomoki Toda</author>
<title>Quasi-Periodic Parallel WaveGAN: A Non-autoregressive Raw Waveform Generative Model with Pitch-dependent Dilated Convolution Neural Network.</title>
<year>2020</year>
<volume>abs/2007.12955</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2007.12955</ee>
<url>db/journals/corr/corr2007.html#abs-2007-12955</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2008-03088" mdate="2020-08-17">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="97/941">Hirokazu Kameoka</author>
<author pid="85/741">Tomoki Toda</author>
<title>Pretraining Techniques for Sequence-to-Sequence Voice Conversion.</title>
<year>2020</year>
<volume>abs/2008.03088</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2008.03088</ee>
<url>db/journals/corr/corr2008.html#abs-2008-03088</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2008-12527" mdate="2020-09-16">
<author pid="51/4138-6">Yi Zhao 0006</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="153/0728">Xiaohai Tian</author>
<author pid="87/3979">Junichi Yamagishi</author>
<author pid="158/4101">Rohan Kumar Das</author>
<author pid="94/4754">Tomi Kinnunen</author>
<author pid="70/5210">Zhen-Hua Ling</author>
<author pid="85/741">Tomoki Toda</author>
<title>Voice Conversion Challenge 2020: Intra-lingual semi-parallel and cross-lingual voice conversion.</title>
<year>2020</year>
<volume>abs/2008.12527</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2008.12527</ee>
<url>db/journals/corr/corr2008.html#abs-2008-12527</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2009-03554" mdate="2020-09-18">
<author pid="158/4101">Rohan Kumar Das</author>
<author pid="94/4754">Tomi Kinnunen</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="70/5210">Zhen-Hua Ling</author>
<author pid="87/3979">Junichi Yamagishi</author>
<author pid="51/4138-6">Yi Zhao 0006</author>
<author pid="153/0728">Xiaohai Tian</author>
<author pid="85/741">Tomoki Toda</author>
<title>Predictions of Subjective Ratings and Spoofing Assessments of Voice Conversion Challenge 2020 Submissions.</title>
<year>2020</year>
<volume>abs/2009.03554</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2009.03554</ee>
<url>db/journals/corr/corr2009.html#abs-2009-03554</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2010-02434" mdate="2020-10-13">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="39/3245-1">Shinji Watanabe 0001</author>
<author pid="85/741">Tomoki Toda</author>
<title>The Sequence-to-Sequence Baseline for the Voice Conversion Challenge 2020: Cascading ASR and TTS.</title>
<year>2020</year>
<volume>abs/2010.02434</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2010.02434</ee>
<url>db/journals/corr/corr2010.html#abs-2010-02434</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2010-04429" mdate="2020-10-13">
<author pid="158/4210">Patrick Lumban Tobing</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="85/741">Tomoki Toda</author>
<title>Baseline System of Voice Conversion Challenge 2020 with Cyclic Variational Autoencoder and Parallel WaveGAN.</title>
<year>2020</year>
<volume>abs/2010.04429</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2010.04429</ee>
<url>db/journals/corr/corr2010.html#abs-2010-04429</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2010-04446" mdate="2020-10-13">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="158/4210">Patrick Lumban Tobing</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="85/741">Tomoki Toda</author>
<title>The NU Voice Conversion System for the Voice Conversion Challenge 2020: On the Effectiveness of Sequence-to-sequence Models and Autoregressive Neural Vocoders.</title>
<year>2020</year>
<volume>abs/2010.04446</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2010.04446</ee>
<url>db/journals/corr/corr2010.html#abs-2010-04446</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2010-12231" mdate="2020-10-27">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="85/741">Tomoki Toda</author>
<title>Any-to-One Sequence-to-Sequence Voice Conversion using Self-Supervised Discrete Speech Representations.</title>
<year>2020</year>
<volume>abs/2010.12231</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2010.12231</ee>
<url>db/journals/corr/corr2010.html#abs-2010-12231</url>
</article>
</r>
<r><article key="journals/access/SekiKLTT19" mdate="2021-04-09">
<author orcid="0000-0001-8284-188X" pid="194/1307">Shogo Seki</author>
<author pid="97/941">Hirokazu Kameoka</author>
<author pid="53/2189-63">Li Li 0063</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="38/3483">Kazuya Takeda</author>
<title>Underdetermined Source Separation Based on Generalized Multichannel Variational Autoencoder.</title>
<pages>168104-168115</pages>
<year>2019</year>
<volume>7</volume>
<journal>IEEE Access</journal>
<ee type="oa">https://doi.org/10.1109/ACCESS.2019.2954120</ee>
<url>db/journals/access/access7.html#SekiKLTT19</url>
</article>
</r>
<r><article key="journals/access/TobingWHKT19" mdate="2021-04-09">
<author orcid="0000-0003-2792-8418" pid="158/4210">Patrick Lumban Tobing</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Voice Conversion With CycleRNN-Based Spectral Mapping and Finely Tuned WaveNet Vocoder.</title>
<pages>171114-171125</pages>
<year>2019</year>
<volume>7</volume>
<journal>IEEE Access</journal>
<ee type="oa">https://doi.org/10.1109/ACCESS.2019.2955978</ee>
<url>db/journals/access/access7.html#TobingWHKT19</url>
</article>
</r>
<r><article key="journals/spm/VijayanLT19" mdate="2021-04-09">
<author orcid="0000-0001-7281-1329" pid="148/9726">Karthika Vijayan</author>
<author orcid="0000-0001-9158-9401" pid="36/4118">Haizhou Li 0001</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Speech-to-Singing Voice Conversion: The Challenges and Strategies for Improving Vocal Conversion Processes.</title>
<pages>95-102</pages>
<year>2019</year>
<volume>36</volume>
<journal>IEEE Signal Process. Mag.</journal>
<number>1</number>
<ee>https://doi.org/10.1109/MSP.2018.2875195</ee>
<url>db/journals/spm/spm36.html#VijayanLT19</url>
</article>
</r>
<r><inproceedings key="conf/asru/TobingHT19" mdate="2021-04-09">
<author pid="158/4210">Patrick Lumban Tobing</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Investigation of Shallow Wavenet Vocoder with Laplacian Distribution Output.</title>
<pages>176-183</pages>
<year>2019</year>
<booktitle>ASRU</booktitle>
<ee>https://doi.org/10.1109/ASRU46091.2019.9003800</ee>
<crossref>conf/asru/2019</crossref>
<url>db/conf/asru/asru2019.html#TobingHT19</url>
</inproceedings>
</r>
<r><inproceedings key="conf/asru/OkamotoTSK19" mdate="2021-04-09">
<author pid="132/9091">Takuma Okamoto</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="02/2978">Yoshinori Shiga</author>
<author pid="32/4341">Hisashi Kawai</author>
<title>Tacotron-Based Acoustic Model Using Phoneme Alignment for Practical Neural Text-to-Speech Systems.</title>
<pages>214-221</pages>
<year>2019</year>
<booktitle>ASRU</booktitle>
<ee>https://doi.org/10.1109/ASRU46091.2019.9003956</ee>
<crossref>conf/asru/2019</crossref>
<url>db/conf/asru/asru2019.html#OkamotoTSK19</url>
</inproceedings>
</r>
<r><inproceedings key="conf/assets/AhmadiKT19" mdate="2021-04-09">
<author pid="52/9232">Farzaneh Ahmadi</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Development of a Real-time Bionic Voice Generation System based on Statistical Excitation Prediction.</title>
<pages>655-657</pages>
<year>2019</year>
<booktitle>ASSETS</booktitle>
<ee>https://doi.org/10.1145/3308561.3354591</ee>
<crossref>conf/assets/2019</crossref>
<url>db/conf/assets/assets2019.html#AhmadiKT19</url>
</inproceedings>
</r>
<r><inproceedings key="conf/eusipco/HuangWHTHKTTW19" mdate="2023-03-21">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="85/2787">Hsin-Te Hwang</author>
<author orcid="0000-0003-2792-8418" pid="158/4210">Patrick Lumban Tobing</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author orcid="0000-0001-6956-0418" pid="66/7146-1">Yu Tsao 0001</author>
<author orcid="0000-0003-3599-5071" pid="28/5019">Hsin-Min Wang</author>
<title>Refined WaveNet Vocoder for Variational Autoencoder Based Voice Conversion.</title>
<pages>1-5</pages>
<year>2019</year>
<booktitle>EUSIPCO</booktitle>
<ee>https://doi.org/10.23919/EUSIPCO.2019.8902651</ee>
<crossref>conf/eusipco/2019</crossref>
<url>db/conf/eusipco/eusipco2019.html#HuangWHTHKTTW19</url>
</inproceedings>
</r>
<r><inproceedings key="conf/eusipco/SekiKLTT19" mdate="2020-01-14">
<author pid="194/1307">Shogo Seki</author>
<author pid="97/941">Hirokazu Kameoka</author>
<author pid="53/2189-63">Li Li 0063</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="38/3483">Kazuya Takeda</author>
<title>Generalized Multichannel Variational Autoencoder for Underdetermined Source Separation.</title>
<pages>1-5</pages>
<year>2019</year>
<booktitle>EUSIPCO</booktitle>
<ee>https://doi.org/10.23919/EUSIPCO.2019.8903054</ee>
<crossref>conf/eusipco/2019</crossref>
<url>db/conf/eusipco/eusipco2019.html#SekiKLTT19</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/KomatsuHKTT19" mdate="2019-06-30">
<author pid="136/5113">Tatsuya Komatsu</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="180/2712">Reishi Kondo</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="38/3483">Kazuya Takeda</author>
<title>Scene-dependent Anomalous Acoustic-event Detection Based on Conditional Wavenet and I-vector.</title>
<pages>870-874</pages>
<year>2019</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2019.8683068</ee>
<crossref>conf/icassp/2019</crossref>
<url>db/conf/icassp/icassp2019.html#KomatsuHKTT19</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/TobingWHKT19" mdate="2023-03-21">
<author orcid="0000-0003-2792-8418" pid="158/4210">Patrick Lumban Tobing</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Voice Conversion with Cyclic Recurrent Neural Network and Fine-tuned Wavenet Vocoder.</title>
<pages>6815-6819</pages>
<year>2019</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2019.8682156</ee>
<crossref>conf/icassp/2019</crossref>
<url>db/conf/icassp/icassp2019.html#TobingWHKT19</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/OkamotoTSK19" mdate="2021-04-09">
<author pid="132/9091">Takuma Okamoto</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="02/2978">Yoshinori Shiga</author>
<author pid="32/4341">Hisashi Kawai</author>
<title>Investigations of Real-time Gaussian Fftnet and Parallel Wavenet Neural Vocoders with Simple Acoustic Features.</title>
<pages>7020-7024</pages>
<year>2019</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2019.8682320</ee>
<crossref>conf/icassp/2019</crossref>
<url>db/conf/icassp/icassp2019.html#OkamotoTSK19</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/WuHTKT19" mdate="2023-06-26">
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author orcid="0000-0003-2792-8418" pid="158/4210">Patrick Lumban Tobing</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Quasi-Periodic WaveNet Vocoder: A Pitch Dependent Dilated Convolution Model for Parametric Speech Generation.</title>
<pages>196-200</pages>
<year>2019</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2019-1232</ee>
<crossref>conf/interspeech/2019</crossref>
<url>db/conf/interspeech/interspeech2019.html#WuHTKT19</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/TobingWHKT19" mdate="2023-06-26">
<author orcid="0000-0003-2792-8418" pid="158/4210">Patrick Lumban Tobing</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Non-Parallel Voice Conversion with Cyclic Variational Autoencoder.</title>
<pages>674-678</pages>
<year>2019</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2019-2307</ee>
<crossref>conf/interspeech/2019</crossref>
<url>db/conf/interspeech/interspeech2019.html#TobingWHKT19</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/KuritaKTT19" mdate="2021-04-09">
<author pid="131/8830">Yusuke Kurita</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="38/3483">Kazuya Takeda</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Robustness of Statistical Voice Conversion Based on Direct Waveform Modification Against Background Sounds.</title>
<pages>684-688</pages>
<year>2019</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2019-2206</ee>
<crossref>conf/interspeech/2019</crossref>
<url>db/conf/interspeech/interspeech2019.html#KuritaKTT19</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/HuangWLTHKT0W19" mdate="2023-06-26">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="234/6665">Chen-Chou Lo</author>
<author orcid="0000-0003-2792-8418" pid="158/4210">Patrick Lumban Tobing</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author orcid="0000-0001-6956-0418" pid="66/7146-1">Yu Tsao 0001</author>
<author orcid="0000-0003-3599-5071" pid="28/5019">Hsin-Min Wang</author>
<title>Investigation of F0 Conditioning and Fully Convolutional Networks in Variational Autoencoder Based Voice Conversion.</title>
<pages>709-713</pages>
<year>2019</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2019-1774</ee>
<crossref>conf/interspeech/2019</crossref>
<url>db/conf/interspeech/interspeech2019.html#HuangWLTHKT0W19</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/OkamotoTSK19" mdate="2021-04-09">
<author pid="132/9091">Takuma Okamoto</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="02/2978">Yoshinori Shiga</author>
<author pid="32/4341">Hisashi Kawai</author>
<title>Real-Time Neural Text-to-Speech with Sequence-to-Sequence Acoustic Model and WaveGlow or Single Gaussian WaveRNN Vocoders.</title>
<pages>1308-1312</pages>
<year>2019</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2019-1288</ee>
<crossref>conf/interspeech/2019</crossref>
<url>db/conf/interspeech/interspeech2019.html#OkamotoTSK19</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/HayashiWTTTL19" mdate="2024-10-06">
<author pid="82/8616">Tomoki Hayashi</author>
<author orcid="0000-0002-5970-8631" pid="39/3245-1">Shinji Watanabe 0001</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="38/3483">Kazuya Takeda</author>
<author pid="160/4302">Shubham Toshniwal</author>
<author orcid="0000-0003-4962-946X" pid="51/2464">Karen Livescu</author>
<title>Pre-Trained Text Embeddings for Enhanced Text-to-Speech Synthesis.</title>
<pages>4430-4434</pages>
<year>2019</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2019-3177</ee>
<crossref>conf/interspeech/2019</crossref>
<url>db/conf/interspeech/interspeech2019.html#HayashiWTTTL19</url>
</inproceedings>
</r>
<r><inproceedings key="conf/ismir/LiTMKM19" mdate="2020-03-12">
<author pid="53/2189-63">Li Li 0063</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="214/2225">Kazuho Morikawa</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="31/6801">Shoji Makino</author>
<title>Improving Singing Aid System for Laryngectomees With Statistical Voice Conversion and VAE-SPACE.</title>
<pages>784-790</pages>
<year>2019</year>
<booktitle>ISMIR</booktitle>
<ee type="oa">http://archives.ismir.net/ismir2019/paper/000096.pdf</ee>
<crossref>conf/ismir/2019</crossref>
<url>db/conf/ismir/ismir2019.html#LiTMKM19</url>
</inproceedings>
</r>
<r><inproceedings key="conf/ssw/HuangWKPHT0WT19" mdate="2024-07-31">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="214/2303">Yu-Huai Peng</author>
<author pid="85/2787">Hsin-Te Hwang</author>
<author pid="158/4210">Patrick Lumban Tobing</author>
<author pid="66/7146-1">Yu Tsao 0001</author>
<author pid="28/5019">Hsin-Min Wang</author>
<author pid="85/741">Tomoki Toda</author>
<title>Generalization of Spectrum Differential based Direct Waveform Modification for Voice Conversion.</title>
<pages>57-62</pages>
<year>2019</year>
<booktitle>SSW</booktitle>
<ee type="oa">https://doi.org/10.21437/SSW.2019-11</ee>
<crossref>conf/ssw/2019</crossref>
<url>db/conf/ssw/ssw2019.html#HuangWKPHT0WT19</url>
</inproceedings>
</r>
<r><inproceedings key="conf/ssw/WuTHKT19" mdate="2024-07-31">
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="158/4210">Patrick Lumban Tobing</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="85/741">Tomoki Toda</author>
<title>Statistical Voice Conversion with Quasi-periodic WaveNet Vocoder.</title>
<pages>63-68</pages>
<year>2019</year>
<booktitle>SSW</booktitle>
<ee type="oa">https://doi.org/10.21437/SSW.2019-12</ee>
<crossref>conf/ssw/2019</crossref>
<url>db/conf/ssw/ssw2019.html#WuTHKT19</url>
</inproceedings>
</r>
<r><inproceedings key="conf/ssw/EshghiTKKT19" mdate="2024-07-31">
<author pid="42/3681">Mohammad Eshghi</author>
<author pid="140/2727">Kou Tanaka</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="97/941">Hirokazu Kameoka</author>
<author pid="85/741">Tomoki Toda</author>
<title>An Investigation of Features for Fundamental Frequency Pattern Prediction in Electrolaryngeal Speech Enhancement.</title>
<pages>251-256</pages>
<year>2019</year>
<booktitle>SSW</booktitle>
<ee type="oa">https://doi.org/10.21437/SSW.2019-45</ee>
<crossref>conf/ssw/2019</crossref>
<url>db/conf/ssw/ssw2019.html#EshghiTKKT19</url>
</inproceedings>
</r>
<r><article publtype="informal" key="journals/corr/abs-1905-00615" mdate="2019-12-20">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="234/6665">Chen-Chou Lo</author>
<author pid="158/4210">Patrick Lumban Tobing</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="66/7146-1">Yu Tsao 0001</author>
<author pid="28/5019">Hsin-Min Wang</author>
<title>Investigation of F0 conditioning and Fully Convolutional Networks in Variational Autoencoder based Voice Conversion.</title>
<year>2019</year>
<volume>abs/1905.00615</volume>
<journal>CoRR</journal>
<ee type="oa">http://arxiv.org/abs/1905.00615</ee>
<url>db/journals/corr/corr1905.html#abs-1905-00615</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-1907-00797" mdate="2019-07-08">
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="158/4210">Patrick Lumban Tobing</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="85/741">Tomoki Toda</author>
<title>Quasi-Periodic WaveNet Vocoder: A Pitch Dependent Dilated Convolution Model for Parametric Speech Generation.</title>
<year>2019</year>
<volume>abs/1907.00797</volume>
<journal>CoRR</journal>
<ee type="oa">http://arxiv.org/abs/1907.00797</ee>
<url>db/journals/corr/corr1907.html#abs-1907-00797</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-1907-08940" mdate="2019-07-30">
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="158/4210">Patrick Lumban Tobing</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="85/741">Tomoki Toda</author>
<title>Statistical Voice Conversion with Quasi-Periodic WaveNet Vocoder.</title>
<year>2019</year>
<volume>abs/1907.08940</volume>
<journal>CoRR</journal>
<ee type="oa">http://arxiv.org/abs/1907.08940</ee>
<url>db/journals/corr/corr1907.html#abs-1907-08940</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-1907-10185" mdate="2019-07-30">
<author pid="158/4210">Patrick Lumban Tobing</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="85/741">Tomoki Toda</author>
<title>Non-Parallel Voice Conversion with Cyclic Variational Autoencoder.</title>
<year>2019</year>
<volume>abs/1907.10185</volume>
<journal>CoRR</journal>
<ee type="oa">http://arxiv.org/abs/1907.10185</ee>
<url>db/journals/corr/corr1907.html#abs-1907-10185</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-1910-10909" mdate="2024-08-06">
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="81/3672">Ryuichi Yamamoto</author>
<author pid="214/2329">Katsuki Inoue</author>
<author pid="173/6383">Takenori Yoshimura</author>
<author pid="39/3245-1">Shinji Watanabe 0001</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="38/3483">Kazuya Takeda</author>
<author pid="50/671-33">Yu Zhang 0033</author>
<author pid="96/10484-3">Xu Tan 0003</author>
<title>ESPnet-TTS: Unified, Reproducible, and Integratable Open Source End-to-End Text-to-Speech Toolkit.</title>
<year>2019</year>
<volume>abs/1910.10909</volume>
<journal>CoRR</journal>
<ee type="oa">http://arxiv.org/abs/1910.10909</ee>
<url>db/journals/corr/corr1910.html#abs-1910-10909</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-1911-01601" mdate="2020-10-26">
<author pid="10/5630-37">Xin Wang 0037</author>
<author pid="87/3979">Junichi Yamagishi</author>
<author pid="89/6736">Massimiliano Todisco</author>
<author pid="21/10462">H&#233;ctor Delgado</author>
<author pid="49/10647">Andreas Nautsch</author>
<author pid="84/366">Nicholas W. D. Evans</author>
<author pid="93/9671">Md. Sahidullah</author>
<author pid="201/7499">Ville Vestman</author>
<author pid="94/4754">Tomi Kinnunen</author>
<author pid="35/4621">Kong Aik Lee</author>
<author pid="158/4213">Lauri Juvela</author>
<author pid="99/2726">Paavo Alku</author>
<author pid="214/2303">Yu-Huai Peng</author>
<author pid="85/2787">Hsin-Te Hwang</author>
<author pid="66/7146-1">Yu Tsao 0001</author>
<author pid="28/5019">Hsin-Min Wang</author>
<author pid="47/10649">S&#233;bastien Le Maguer</author>
<author pid="63/4934">Markus Becker</author>
<author pid="h/FergusHenderson">Fergus Henderson</author>
<author pid="23/6618">Rob Clark</author>
<author pid="50/671-33">Yu Zhang 0033</author>
<author pid="86/5728">Quan Wang</author>
<author pid="217/2520">Ye Jia</author>
<author pid="252/5717">Kai Onuma</author>
<author pid="252/4963">Koji Mushika</author>
<author pid="218/9267">Takashi Kaneda</author>
<author pid="02/393-6">Yuan Jiang 0006</author>
<author pid="93/6344">Li-Juan Liu</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="140/2727">Kou Tanaka</author>
<author pid="97/941">Hirokazu Kameoka</author>
<author pid="64/9231">Ingmar Steiner</author>
<author pid="53/1360">Driss Matrouf</author>
<author pid="72/3130">Jean-Fran&#231;ois Bonastre</author>
<author pid="173/6404">Avashna Govender</author>
<author pid="180/2841">Srikanth Ronanki</author>
<author pid="223/5831">Jing-Xuan Zhang</author>
<author pid="70/5210">Zhen-Hua Ling</author>
<title>The ASVspoof 2019 database.</title>
<year>2019</year>
<volume>abs/1911.01601</volume>
<journal>CoRR</journal>
<ee type="oa">http://arxiv.org/abs/1911.01601</ee>
<url>db/journals/corr/corr1911.html#abs-1911-01601</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-1912-06813" mdate="2020-01-07">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="97/941">Hirokazu Kameoka</author>
<author pid="85/741">Tomoki Toda</author>
<title>Voice Transformer Network: Sequence-to-Sequence Voice Conversion Using Transformer with Text-to-Speech Pretraining.</title>
<year>2019</year>
<volume>abs/1912.06813</volume>
<journal>CoRR</journal>
<ee type="oa">http://arxiv.org/abs/1912.06813</ee>
<url>db/journals/corr/corr1912.html#abs-1912-06813</url>
</article>
</r>
<r><article key="journals/ieicet/HayashiNKTT18" mdate="2021-04-09">
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="28/2662">Masafumi Nishida</author>
<author pid="07/6964">Norihide Kitaoka</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="38/3483">Kazuya Takeda</author>
<title>Daily Activity Recognition with Large-Scaled Real-Life Recording Datasets Based on Deep Neural Network Using Multi-Modal Signals.</title>
<pages>199-210</pages>
<year>2018</year>
<volume>101-A</volume>
<journal>IEICE Trans. Fundam. Electron. Commun. Comput. Sci.</journal>
<number>1</number>
<ee>https://doi.org/10.1587/transfun.E101.A.199</ee>
<ee>http://search.ieice.org/bin/summary.php?id=e101-a_1_199</ee>
<url>db/journals/ieicet/ieicet101a.html#HayashiNKTT18</url>
</article>
</r>
<r><article key="journals/ieicet/SekiTT18" mdate="2021-04-09">
<author pid="194/1307">Shogo Seki</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="38/3483">Kazuya Takeda</author>
<title>Stereophonic Music Separation Based on Non-Negative Tensor Factorization with Cepstral Distance Regularization.</title>
<pages>1057-1064</pages>
<year>2018</year>
<volume>101-A</volume>
<journal>IEICE Trans. Fundam. Electron. Commun. Comput. Sci.</journal>
<number>7</number>
<ee>https://doi.org/10.1587/transfun.E101.A.1057</ee>
<ee>http://search.ieice.org/bin/summary.php?id=e101-a_7_1057</ee>
<url>db/journals/ieicet/ieicet101a.html#SekiTT18</url>
</article>
</r>
<r><article key="journals/mt/KanoTSNTN18" mdate="2021-04-09">
<author orcid="0000-0001-9693-3785" pid="140/2697">Takatomo Kano</author>
<author pid="124/9221">Shinnosuke Takamichi</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>An end-to-end model for cross-lingual transformation of paralinguistic information.</title>
<pages>353-368</pages>
<year>2018</year>
<volume>32</volume>
<journal>Mach. Transl.</journal>
<number>4</number>
<ee>https://doi.org/10.1007/s10590-018-9217-7</ee>
<url>db/journals/mt/mt32.html#KanoTSNTN18</url>
</article>
</r>
<r><article key="journals/speech/KobayashiTN18" mdate="2021-04-09">
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Intra-gender statistical singing voice conversion with direct waveform modification using log-spectral differential.</title>
<pages>211-220</pages>
<year>2018</year>
<volume>99</volume>
<journal>Speech Commun.</journal>
<ee type="oa">https://doi.org/10.1016/j.specom.2018.03.011</ee>
<url>db/journals/speech/speech99.html#KobayashiTN18</url>
</article>
</r>
<r><inproceedings key="conf/apsipa/TakadaST18" mdate="2021-04-09">
<author pid="237/0010">Moe Takada</author>
<author pid="194/1307">Shogo Seki</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Self-Produced Speech Enhancement and Suppression Method using Air- and Body-Conductive Microphones.</title>
<pages>1240-1245</pages>
<year>2018</year>
<booktitle>APSIPA</booktitle>
<ee>https://doi.org/10.23919/APSIPA.2018.8659663</ee>
<crossref>conf/apsipa/2018</crossref>
<url>db/conf/apsipa/apsipa2018.html#TakadaST18</url>
</inproceedings>
</r>
<r><inproceedings key="conf/educon/SeiyaIOTODT18" mdate="2024-05-07">
<author pid="220/9281">Shunya Seiya</author>
<author pid="220/9537">Ryuya Ito</author>
<author pid="220/9460">Kosuke Okamoto</author>
<author pid="200/0164">Ukyo Tanikawa</author>
<author pid="36/2324">Shigeki Ohira</author>
<author orcid="0000-0003-0603-8790" pid="00/2175">Daisuke Deguchi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Development of &#34;KamiRepo&#34; system with automatic student identification to handle handwritten assignments on LMS.</title>
<pages>835-842</pages>
<year>2018</year>
<booktitle>EDUCON</booktitle>
<ee>https://doi.org/10.1109/EDUCON.2018.8363317</ee>
<crossref>conf/educon/2018</crossref>
<url>db/conf/educon/educon2018.html#SeiyaIOTODT18</url>
</inproceedings>
</r>
<r><inproceedings key="conf/eusipco/MiyazakiHTT18" mdate="2021-04-09">
<author pid="27/9637">Koichi Miyazaki</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="38/3483">Kazuya Takeda</author>
<title>Connectionist Temporal Classification-based Sound Event Encoder for Converting Sound Events into Onomatopoeic Representations.</title>
<pages>852-856</pages>
<year>2018</year>
<booktitle>EUSIPCO</booktitle>
<ee>https://doi.org/10.23919/EUSIPCO.2018.8553374</ee>
<crossref>conf/eusipco/2018</crossref>
<url>db/conf/eusipco/eusipco2018.html#MiyazakiHTT18</url>
</inproceedings>
</r>
<r><inproceedings key="conf/eusipco/KobayashiT18" mdate="2021-04-09">
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Electrolaryngeal Speech Enhancement with Statistical Voice Conversion based on CLDNN.</title>
<pages>2115-2119</pages>
<year>2018</year>
<booktitle>EUSIPCO</booktitle>
<ee>https://doi.org/10.23919/EUSIPCO.2018.8553154</ee>
<crossref>conf/eusipco/2018</crossref>
<url>db/conf/eusipco/eusipco2018.html#KobayashiT18</url>
</inproceedings>
</r>
<r><inproceedings key="conf/eusipco/HayashiKKTT18" mdate="2021-04-09">
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="136/5113">Tatsuya Komatsu</author>
<author pid="180/2712">Reishi Kondo</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="38/3483">Kazuya Takeda</author>
<title>Anomalous Sound Event Detection Based on WaveNet.</title>
<pages>2494-2498</pages>
<year>2018</year>
<booktitle>EUSIPCO</booktitle>
<ee>https://doi.org/10.23919/EUSIPCO.2018.8553423</ee>
<crossref>conf/eusipco/2018</crossref>
<url>db/conf/eusipco/eusipco2018.html#HayashiKKTT18</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/OkamotoTTSK18" mdate="2021-04-09">
<author pid="132/9091">Takuma Okamoto</author>
<author pid="35/8056">Kentaro Tachibana</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="02/2978">Yoshinori Shiga</author>
<author pid="32/4341">Hisashi Kawai</author>
<title>An Investigation of Subband Wavenet Vocoder Covering Entire Audible Frequency Range with Limited Acoustic Features.</title>
<pages>5654-5658</pages>
<year>2018</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2018.8462237</ee>
<crossref>conf/icassp/2018</crossref>
<url>db/conf/icassp/icassp2018.html#OkamotoTTSK18</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/TachibanaTSK18" mdate="2021-04-09">
<author pid="35/8056">Kentaro Tachibana</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="02/2978">Yoshinori Shiga</author>
<author pid="32/4341">Hisashi Kawai</author>
<title>An Investigation of Noise Shaping with Perceptual Weighting for Wavenet-Based Speech Generation.</title>
<pages>5664-5668</pages>
<year>2018</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2018.8461332</ee>
<crossref>conf/icassp/2018</crossref>
<url>db/conf/icassp/icassp2018.html#TachibanaTSK18</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/HayashiWTT18" mdate="2021-04-09">
<author pid="82/8616">Tomoki Hayashi</author>
<author orcid="0000-0002-5970-8631" pid="39/3245-1">Shinji Watanabe 0001</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="38/3483">Kazuya Takeda</author>
<title>Multi-Head Decoder for End-to-End Speech Recognition.</title>
<pages>801-805</pages>
<year>2018</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2018-1655</ee>
<crossref>conf/interspeech/2018</crossref>
<url>db/conf/interspeech/interspeech2018.html#HayashiWTT18</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/WuKHTT18" mdate="2023-03-21">
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author orcid="0000-0003-2792-8418" pid="158/4210">Patrick Lumban Tobing</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Collapsed Speech Segment Detection and Suppression for WaveNet Vocoder.</title>
<pages>1988-1992</pages>
<year>2018</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2018-1210</ee>
<crossref>conf/interspeech/2018</crossref>
<url>db/conf/interspeech/interspeech2018.html#WuKHTT18</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/KawaharaSMBTI18" mdate="2021-04-09">
<author pid="84/3249">Hideki Kawahara</author>
<author pid="11/9231">Ken-Ichi Sakakibara</author>
<author pid="75/6649">Masanori Morise</author>
<author pid="95/6268">Hideki Banno</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="72/3903">Toshio Irino</author>
<title>Frequency Domain Variants of Velvet Noise and Their Application to Speech Processing and Synthesis.</title>
<pages>2027-2031</pages>
<year>2018</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2018-43</ee>
<crossref>conf/interspeech/2018</crossref>
<url>db/conf/interspeech/interspeech2018.html#KawaharaSMBTI18</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/TamuraHEHT18" mdate="2021-04-09">
<author pid="32/4043">Satoshi Tamura</author>
<author pid="226/1799">Kento Horio</author>
<author pid="226/1934">Hajime Endo</author>
<author pid="24/64">Satoru Hayamizu</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Audio-visual Voice Conversion Using Deep Canonical Correlation Analysis for Deep Bottleneck Features.</title>
<pages>2469-2473</pages>
<year>2018</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2018-2286</ee>
<crossref>conf/interspeech/2018</crossref>
<url>db/conf/interspeech/interspeech2018.html#TamuraHEHT18</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/AhmadiT18" mdate="2021-04-09">
<author pid="52/9232">Farzaneh Ahmadi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Designing a Pneumatic Bionic Voice Prosthesis - A Statistical Approach for Source Excitation Generation.</title>
<pages>3142-3146</pages>
<year>2018</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2018-1043</ee>
<crossref>conf/interspeech/2018</crossref>
<url>db/conf/interspeech/interspeech2018.html#AhmadiT18</url>
</inproceedings>
</r>
<r><inproceedings key="conf/odyssey/KinnunenLYTSVL18" mdate="2021-02-10">
<author pid="94/4754">Tomi Kinnunen</author>
<author pid="124/9038">Jaime Lorenzo-Trueba</author>
<author pid="87/3979">Junichi Yamagishi</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="17/7825">Daisuke Saito</author>
<author pid="64/2286">Fernando Villavicencio</author>
<author pid="70/5210">Zhen-Hua Ling</author>
<title>A Spoofing Benchmark for the 2018 Voice Conversion Challenge: Leveraging from Spoofing Countermeasures for Speech Artifact Assessment.</title>
<pages>187-194</pages>
<year>2018</year>
<booktitle>Odyssey</booktitle>
<ee type="oa">https://doi.org/10.21437/Odyssey.2018-27</ee>
<crossref>conf/odyssey/2018</crossref>
<url>db/conf/odyssey/odyssey2018.html#KinnunenLYTSVL18</url>
</inproceedings>
</r>
<r><inproceedings key="conf/odyssey/Lorenzo-TruebaY18" mdate="2021-02-10">
<author pid="124/9038">Jaime Lorenzo-Trueba</author>
<author pid="87/3979">Junichi Yamagishi</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="17/7825">Daisuke Saito</author>
<author pid="64/2286">Fernando Villavicencio</author>
<author pid="94/4754">Tomi Kinnunen</author>
<author pid="70/5210">Zhen-Hua Ling</author>
<title>The Voice Conversion Challenge 2018: Promoting Development of Parallel and Nonparallel Methods.</title>
<pages>195-202</pages>
<year>2018</year>
<booktitle>Odyssey</booktitle>
<ee type="oa">https://doi.org/10.21437/Odyssey.2018-28</ee>
<crossref>conf/odyssey/2018</crossref>
<url>db/conf/odyssey/odyssey2018.html#Lorenzo-TruebaY18</url>
</inproceedings>
</r>
<r><inproceedings key="conf/odyssey/KobayashiT18" mdate="2021-02-10">
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="85/741">Tomoki Toda</author>
<title>sprocket: Open-Source Voice Conversion Software.</title>
<pages>203-210</pages>
<year>2018</year>
<booktitle>Odyssey</booktitle>
<ee type="oa">https://doi.org/10.21437/Odyssey.2018-29</ee>
<crossref>conf/odyssey/2018</crossref>
<url>db/conf/odyssey/odyssey2018.html#KobayashiT18</url>
</inproceedings>
</r>
<r><inproceedings key="conf/odyssey/WuTHKT18" mdate="2021-02-10">
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="158/4210">Patrick Lumban Tobing</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="85/741">Tomoki Toda</author>
<title>The NU Non-Parallel Voice Conversion System for the Voice Conversion Challenge 2018.</title>
<pages>211-218</pages>
<year>2018</year>
<booktitle>Odyssey</booktitle>
<ee type="oa">https://doi.org/10.21437/Odyssey.2018-30</ee>
<crossref>conf/odyssey/2018</crossref>
<url>db/conf/odyssey/odyssey2018.html#WuTHKT18</url>
</inproceedings>
</r>
<r><inproceedings key="conf/odyssey/TobingWHKT18" mdate="2021-02-10">
<author pid="158/4210">Patrick Lumban Tobing</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="85/741">Tomoki Toda</author>
<title>NU Voice Conversion System for the Voice Conversion Challenge 2018.</title>
<pages>219-226</pages>
<year>2018</year>
<booktitle>Odyssey</booktitle>
<ee type="oa">https://doi.org/10.21437/Odyssey.2018-31</ee>
<crossref>conf/odyssey/2018</crossref>
<url>db/conf/odyssey/odyssey2018.html#TobingWHKT18</url>
</inproceedings>
</r>
<r><inproceedings key="conf/slt/TobingHWKT18" mdate="2023-03-21">
<author orcid="0000-0003-2792-8418" pid="158/4210">Patrick Lumban Tobing</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>An Evaluation of Deep Spectral Mappings and WaveNet Vocoder for Voice Conversion.</title>
<pages>297-303</pages>
<year>2018</year>
<booktitle>SLT</booktitle>
<ee>https://doi.org/10.1109/SLT.2018.8639608</ee>
<crossref>conf/slt/2018</crossref>
<url>db/conf/slt/slt2018.html#TobingHWKT18</url>
</inproceedings>
</r>
<r><inproceedings key="conf/slt/OkamotoTSK18" mdate="2021-04-09">
<author pid="132/9091">Takuma Okamoto</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="02/2978">Yoshinori Shiga</author>
<author pid="32/4341">Hisashi Kawai</author>
<title>Improving FFTNet Vocoder with Noise Shaping and Subband Approaches.</title>
<pages>304-311</pages>
<year>2018</year>
<booktitle>SLT</booktitle>
<ee>https://doi.org/10.1109/SLT.2018.8639687</ee>
<crossref>conf/slt/2018</crossref>
<url>db/conf/slt/slt2018.html#OkamotoTSK18</url>
</inproceedings>
</r>
<r><inproceedings key="conf/slt/HayashiWZTHAT18" mdate="2022-08-11">
<author pid="82/8616">Tomoki Hayashi</author>
<author orcid="0000-0002-5970-8631" pid="39/3245-1">Shinji Watanabe 0001</author>
<author pid="50/671-33">Yu Zhang 0033</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="46/3941">Takaaki Hori</author>
<author pid="56/7987">Ram&#243;n Fernandez Astudillo</author>
<author pid="38/3483">Kazuya Takeda</author>
<title>Back-Translation-Style Data Augmentation for end-to-end ASR.</title>
<pages>426-433</pages>
<year>2018</year>
<booktitle>SLT</booktitle>
<ee>https://doi.org/10.1109/SLT.2018.8639619</ee>
<crossref>conf/slt/2018</crossref>
<url>db/conf/slt/slt2018.html#HayashiWZTHAT18</url>
</inproceedings>
</r>
<r><article publtype="informal" key="journals/corr/abs-1804-04262" mdate="2018-08-13">
<author pid="124/9038">Jaime Lorenzo-Trueba</author>
<author pid="87/3979">Junichi Yamagishi</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="17/7825">Daisuke Saito</author>
<author pid="64/2286">Fernando Villavicencio</author>
<author pid="94/4754">Tomi Kinnunen</author>
<author pid="70/5210">Zhen-Hua Ling</author>
<title>The Voice Conversion Challenge 2018: Promoting Development of Parallel and Nonparallel Methods.</title>
<year>2018</year>
<volume>abs/1804.04262</volume>
<journal>CoRR</journal>
<ee type="oa">http://arxiv.org/abs/1804.04262</ee>
<url>db/journals/corr/corr1804.html#abs-1804-04262</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-1804-08050" mdate="2020-06-30">
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="39/3245-1">Shinji Watanabe 0001</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="38/3483">Kazuya Takeda</author>
<title>Multi-Head Decoder for End-to-End Speech Recognition.</title>
<year>2018</year>
<volume>abs/1804.08050</volume>
<journal>CoRR</journal>
<ee type="oa">http://arxiv.org/abs/1804.08050</ee>
<url>db/journals/corr/corr1804.html#abs-1804-08050</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-1804-08438" mdate="2018-08-13">
<author pid="94/4754">Tomi Kinnunen</author>
<author pid="124/9038">Jaime Lorenzo-Trueba</author>
<author pid="87/3979">Junichi Yamagishi</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="17/7825">Daisuke Saito</author>
<author pid="64/2286">Fernando Villavicencio</author>
<author pid="70/5210">Zhen-Hua Ling</author>
<title>A Spoofing Benchmark for the 2018 Voice Conversion Challenge: Leveraging from Spoofing Countermeasures for Speech Artifact Assessment.</title>
<year>2018</year>
<volume>abs/1804.08438</volume>
<journal>CoRR</journal>
<ee type="oa">http://arxiv.org/abs/1804.08438</ee>
<url>db/journals/corr/corr1804.html#abs-1804-08438</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-1804-11055" mdate="2018-08-13">
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="158/4210">Patrick Lumban Tobing</author>
<author pid="85/741">Tomoki Toda</author>
<title>Collapsed speech segment detection and suppression for WaveNet vocoder.</title>
<year>2018</year>
<volume>abs/1804.11055</volume>
<journal>CoRR</journal>
<ee type="oa">http://arxiv.org/abs/1804.11055</ee>
<url>db/journals/corr/corr1804.html#abs-1804-11055</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-1806-06812" mdate="2018-08-13">
<author pid="84/3249">Hideki Kawahara</author>
<author pid="11/9231">Ken-Ichi Sakakibara</author>
<author pid="75/6649">Masanori Morise</author>
<author pid="95/6268">Hideki Banno</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="72/3903">Toshio Irino</author>
<title>Frequency domain variants of velvet noise and their application to speech processing and synthesis: with appendices.</title>
<year>2018</year>
<volume>abs/1806.06812</volume>
<journal>CoRR</journal>
<ee type="oa">http://arxiv.org/abs/1806.06812</ee>
<url>db/journals/corr/corr1806.html#abs-1806-06812</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-1807-10893" mdate="2022-08-11">
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="39/3245-1">Shinji Watanabe 0001</author>
<author pid="50/671-33">Yu Zhang 0033</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="46/3941">Takaaki Hori</author>
<author pid="56/7987">Ram&#243;n Fernandez Astudillo</author>
<author pid="38/3483">Kazuya Takeda</author>
<title>Back-Translation-Style Data Augmentation for End-to-End ASR.</title>
<year>2018</year>
<volume>abs/1807.10893</volume>
<journal>CoRR</journal>
<ee type="oa">http://arxiv.org/abs/1807.10893</ee>
<url>db/journals/corr/corr1807.html#abs-1807-10893</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-1810-00223" mdate="2020-01-14">
<author pid="194/1307">Shogo Seki</author>
<author pid="97/941">Hirokazu Kameoka</author>
<author pid="53/2189-63">Li Li 0063</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="38/3483">Kazuya Takeda</author>
<title>Generalized Multichannel Variational Autoencoder for Underdetermined Source Separation.</title>
<year>2018</year>
<volume>abs/1810.00223</volume>
<journal>CoRR</journal>
<ee type="oa">http://arxiv.org/abs/1810.00223</ee>
<url>db/journals/corr/corr1810.html#abs-1810-00223</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-1811-11078" mdate="2019-12-20">
<author pid="225/7821">Wen-Chin Huang</author>
<author pid="188/5943">Yi-Chiao Wu</author>
<author pid="85/2787">Hsin-Te Hwang</author>
<author pid="158/4210">Patrick Lumban Tobing</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="66/7146-1">Yu Tsao 0001</author>
<author pid="28/5019">Hsin-Min Wang</author>
<title>Refined WaveNet Vocoder for Variational Autoencoder Based Voice Conversion.</title>
<year>2018</year>
<volume>abs/1811.11078</volume>
<journal>CoRR</journal>
<ee type="oa">http://arxiv.org/abs/1811.11078</ee>
<url>db/journals/corr/corr1811.html#abs-1811-11078</url>
</article>
</r>
<r><article key="journals/ieicet/TanakaT017" mdate="2021-04-09">
<author pid="140/2727">Kou Tanaka</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>A Vibration Control Method of an Electrolarynx Based on Statistical <i>F</i><sub>0</sub> Pattern Prediction.</title>
<pages>2165-2173</pages>
<year>2017</year>
<volume>100-D</volume>
<journal>IEICE Trans. Inf. Syst.</journal>
<number>9</number>
<ee>https://doi.org/10.1587/transinf.2016EDP7485</ee>
<ee>http://search.ieice.org/bin/summary.php?id=e100-d_9_2165</ee>
<url>db/journals/ieicet/ieicet100d.html#TanakaT017</url>
</article>
</r>
<r><article key="journals/taslp/DoTNSN17" mdate="2021-04-09">
<author pid="166/6489">Quoc Truong Do</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Preserving Word-Level Emphasis in Speech-to-Speech Translation.</title>
<pages>544-556</pages>
<year>2017</year>
<volume>25</volume>
<journal>IEEE ACM Trans. Audio Speech Lang. Process.</journal>
<number>3</number>
<ee>https://doi.org/10.1109/TASLP.2016.2643280</ee>
<url>db/journals/taslp/taslp25.html#DoTNSN17</url>
</article>
</r>
<r><article key="journals/taslp/HayashiWTHRT17" mdate="2021-04-09">
<author pid="82/8616">Tomoki Hayashi</author>
<author orcid="0000-0002-5970-8631" pid="39/3245-1">Shinji Watanabe 0001</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="46/3941">Takaaki Hori</author>
<author pid="36/4575">Jonathan Le Roux</author>
<author pid="38/3483">Kazuya Takeda</author>
<title>Duration-Controlled LSTM for Polyphonic Sound Event Detection.</title>
<pages>2059-2070</pages>
<year>2017</year>
<volume>25</volume>
<journal>IEEE ACM Trans. Audio Speech Lang. Process.</journal>
<number>11</number>
<ee type="oa">https://doi.org/10.1109/TASLP.2017.2740002</ee>
<url>db/journals/taslp/taslp25.html#HayashiWTHRT17</url>
</article>
</r>
<r><article key="journals/taslp/TobingKT17" mdate="2023-03-21">
<author orcid="0000-0003-2792-8418" pid="158/4210">Patrick Lumban Tobing</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Articulatory Controllable Speech Modification Based on Statistical Inversion and Production Mappings.</title>
<pages>2337-2350</pages>
<year>2017</year>
<volume>25</volume>
<journal>IEEE ACM Trans. Audio Speech Lang. Process.</journal>
<number>12</number>
<ee>https://doi.org/10.1109/TASLP.2017.2753583</ee>
<url>db/journals/taslp/taslp25.html#TobingKT17</url>
</article>
</r>
<r><inproceedings key="conf/apsipa/MorikawaT17" mdate="2021-04-09">
<author pid="214/2225">Kazuho Morikawa</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Electrolaryngeal speech modification towards singing aid system for laryngectomees.</title>
<pages>610-613</pages>
<year>2017</year>
<booktitle>APSIPA</booktitle>
<ee>https://doi.org/10.1109/APSIPA.2017.8282097</ee>
<crossref>conf/apsipa/2017</crossref>
<url>db/conf/apsipa/apsipa2017.html#MorikawaT17</url>
</inproceedings>
</r>
<r><inproceedings key="conf/apsipa/TobingKT17" mdate="2023-03-21">
<author orcid="0000-0003-2792-8418" pid="158/4210">Patrick Lumban Tobing</author>
<author pid="97/941">Hirokazu Kameoka</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Deep acoustic-to-articulatory inversion mapping with latent trajectory modeling.</title>
<pages>1274-1277</pages>
<year>2017</year>
<booktitle>APSIPA</booktitle>
<ee>https://doi.org/10.1109/APSIPA.2017.8282219</ee>
<crossref>conf/apsipa/2017</crossref>
<url>db/conf/apsipa/apsipa2017.html#TobingKT17</url>
</inproceedings>
</r>
<r><inproceedings key="conf/apsipa/TamamoriHTT17" mdate="2026-06-21">
<author orcid="0009-0000-8893-0058" pid="01/8760">Akira Tamamori</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="38/3483">Kazuya Takeda</author>
<title>An investigation of recurrent neural network for daily activity recognition using multi-modal signals.</title>
<pages>1334-1340</pages>
<year>2017</year>
<booktitle>APSIPA</booktitle>
<ee>https://doi.org/10.1109/APSIPA.2017.8282239</ee>
<crossref>conf/apsipa/2017</crossref>
<url>db/conf/apsipa/apsipa2017.html#TamamoriHTT17</url>
</inproceedings>
</r>
<r><inproceedings key="conf/apsipa/KuboKTNS017" mdate="2021-04-09">
<author pid="214/2335">Kazutaka Kubo</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>An investigation of how to design control parameters for statistical voice timbre control.</title>
<pages>1520-1523</pages>
<year>2017</year>
<booktitle>APSIPA</booktitle>
<ee>https://doi.org/10.1109/APSIPA.2017.8282283</ee>
<crossref>conf/apsipa/2017</crossref>
<url>db/conf/apsipa/apsipa2017.html#KuboKTNS017</url>
</inproceedings>
</r>
<r><inproceedings key="conf/apsipa/KawaharaSMBT17" mdate="2021-04-09">
<author pid="84/3249">Hideki Kawahara</author>
<author pid="11/9231">Ken-Ichi Sakakibara</author>
<author pid="75/6649">Masanori Morise</author>
<author pid="95/6268">Hideki Banno</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Accurate estimation of f0 and aperiodicity based on periodicity detector residuals and deviations of phase derivatives.</title>
<pages>1556-1564</pages>
<year>2017</year>
<booktitle>APSIPA</booktitle>
<ee>https://doi.org/10.1109/APSIPA.2017.8282290</ee>
<crossref>conf/apsipa/2017</crossref>
<url>db/conf/apsipa/apsipa2017.html#KawaharaSMBT17</url>
</inproceedings>
</r>
<r><inproceedings key="conf/asru/OkamotoTTSK17" mdate="2021-04-09">
<author pid="132/9091">Takuma Okamoto</author>
<author pid="35/8056">Kentaro Tachibana</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="02/2978">Yoshinori Shiga</author>
<author pid="32/4341">Hisashi Kawai</author>
<title>Subband wavenet with overlapped single-sideband filterbanks.</title>
<pages>698-704</pages>
<year>2017</year>
<booktitle>ASRU</booktitle>
<ee>https://doi.org/10.1109/ASRU.2017.8269005</ee>
<crossref>conf/asru/2017</crossref>
<url>db/conf/asru/asru2017.html#OkamotoTTSK17</url>
</inproceedings>
</r>
<r><inproceedings key="conf/asru/HayashiTKTT17" mdate="2026-06-21">
<author pid="82/8616">Tomoki Hayashi</author>
<author orcid="0009-0000-8893-0058" pid="01/8760">Akira Tamamori</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="38/3483">Kazuya Takeda</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>An investigation of multi-speaker training for wavenet vocoder.</title>
<pages>712-718</pages>
<year>2017</year>
<booktitle>ASRU</booktitle>
<ee>https://doi.org/10.1109/ASRU.2017.8269007</ee>
<crossref>conf/asru/2017</crossref>
<url>db/conf/asru/asru2017.html#HayashiTKTT17</url>
</inproceedings>
</r>
<r><inproceedings key="conf/eusipco/SekiTT17" mdate="2021-04-09">
<author pid="194/1307">Shogo Seki</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="38/3483">Kazuya Takeda</author>
<title>Stereophonic music separation based on non-negative tensor factorization with cepstrum regularization.</title>
<pages>981-985</pages>
<year>2017</year>
<booktitle>EUSIPCO</booktitle>
<ee>https://doi.org/10.23919/EUSIPCO.2017.8081354</ee>
<crossref>conf/eusipco/2017</crossref>
<url>db/conf/eusipco/eusipco2017.html#SekiTT17</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/HayashiWTHRT17" mdate="2021-04-09">
<author pid="82/8616">Tomoki Hayashi</author>
<author orcid="0000-0002-5970-8631" pid="39/3245-1">Shinji Watanabe 0001</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="46/3941">Takaaki Hori</author>
<author pid="36/4575">Jonathan Le Roux</author>
<author pid="38/3483">Kazuya Takeda</author>
<title>BLSTM-HMM hybrid system combined with sound activity detection network for polyphonic Sound Event Detection.</title>
<pages>766-770</pages>
<year>2017</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2017.7952259</ee>
<crossref>conf/icassp/2017</crossref>
<url>db/conf/icassp/icassp2017.html#HayashiWTHRT17</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/TajiriKT17" mdate="2021-04-09">
<author pid="173/6532">Yusuke Tajiri</author>
<author pid="97/941">Hirokazu Kameoka</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>A noise suppression method for body-conducted soft speech based on non-negative tensor factorization of air- and body-conducted signals.</title>
<pages>4960-4964</pages>
<year>2017</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2017.7953100</ee>
<crossref>conf/icassp/2017</crossref>
<url>db/conf/icassp/icassp2017.html#TajiriKT17</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/KawaharaSMBT17" mdate="2023-08-06">
<author pid="84/3249">Hideki Kawahara</author>
<author pid="11/9231">Ken-Ichi Sakakibara</author>
<author pid="75/6649">Masanori Morise</author>
<author pid="95/6268">Hideki Banno</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>A Modulation Property of Time-Frequency Derivatives of Filtered Phase and its Application to Aperiodicity and f<sub>o</sub> Estimation.</title>
<pages>424-428</pages>
<year>2017</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2017-436</ee>
<crossref>conf/interspeech/2017</crossref>
<url>db/conf/interspeech/interspeech2017.html#KawaharaSMBT17</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/TanakaKT017" mdate="2023-08-06">
<author pid="140/2727">Kou Tanaka</author>
<author pid="97/941">Hirokazu Kameoka</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Physically Constrained Statistical F<sub>0</sub> Prediction for Electrolaryngeal Speech Enhancement.</title>
<pages>1069-1073</pages>
<year>2017</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2017-688</ee>
<crossref>conf/interspeech/2017</crossref>
<url>db/conf/interspeech/interspeech2017.html#TanakaKT017</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/TamamoriHKTT17" mdate="2026-06-21">
<author orcid="0009-0000-8893-0058" pid="01/8760">Akira Tamamori</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="38/3483">Kazuya Takeda</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Speaker-Dependent WaveNet Vocoder.</title>
<pages>1118-1122</pages>
<year>2017</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2017-314</ee>
<crossref>conf/interspeech/2017</crossref>
<url>db/conf/interspeech/interspeech2017.html#TamamoriHKTT17</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/KobayashiHTT17" mdate="2026-06-21">
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author orcid="0009-0000-8893-0058" pid="01/8760">Akira Tamamori</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Statistical Voice Conversion with WaveNet-Based Waveform Generation.</title>
<pages>1138-1142</pages>
<year>2017</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2017-986</ee>
<crossref>conf/interspeech/2017</crossref>
<url>db/conf/interspeech/interspeech2017.html#KobayashiHTT17</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/KawaharaSMBTI17" mdate="2023-08-06">
<author pid="84/3249">Hideki Kawahara</author>
<author pid="11/9231">Ken-Ichi Sakakibara</author>
<author pid="75/6649">Masanori Morise</author>
<author pid="95/6268">Hideki Banno</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="72/3903">Toshio Irino</author>
<title>A New Cosine Series Antialiasing Function and its Application to Aliasing-Free Glottal Source Models for Speech and Singing Synthesis.</title>
<pages>1358-1362</pages>
<year>2017</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2017-15</ee>
<crossref>conf/interspeech/2017</crossref>
<url>db/conf/interspeech/interspeech2017.html#KawaharaSMBTI17</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/LiKTM17" mdate="2023-08-06">
<author pid="53/2189-63">Li Li 0063</author>
<author pid="97/941">Hirokazu Kameoka</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author orcid="0000-0003-1934-640X" pid="31/6801">Shoji Makino</author>
<title>Speech Enhancement Using Non-Negative Spectrogram Models with Mel-Generalized Cepstral Regularization.</title>
<pages>1998-2002</pages>
<year>2017</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2017-1492</ee>
<crossref>conf/interspeech/2017</crossref>
<url>db/conf/interspeech/interspeech2017.html#LiKTM17</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mlsp/SekiKTT17" mdate="2021-04-09">
<author pid="194/1307">Shogo Seki</author>
<author pid="97/941">Hirokazu Kameoka</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="38/3483">Kazuya Takeda</author>
<title>Missing component restoration for masked speech signals based on time-domain spectrogram factorization.</title>
<pages>1-6</pages>
<year>2017</year>
<booktitle>MLSP</booktitle>
<ee>https://doi.org/10.1109/MLSP.2017.8168125</ee>
<crossref>conf/mlsp/2017</crossref>
<url>db/conf/mlsp/mlsp2017.html#SekiKTT17</url>
</inproceedings>
</r>
<r><article publtype="informal" key="journals/corr/KawaharaSBMTI17" mdate="2018-08-13">
<author pid="84/3249">Hideki Kawahara</author>
<author pid="11/9231">Ken-Ichi Sakakibara</author>
<author pid="95/6268">Hideki Banno</author>
<author pid="75/6649">Masanori Morise</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="72/3903">Toshio Irino</author>
<title>A new cosine series antialiasing function and its application to aliasing-free glottal source models for speech and singing synthesis.</title>
<year>2017</year>
<volume>abs/1702.06724</volume>
<journal>CoRR</journal>
<ee type="oa">http://arxiv.org/abs/1702.06724</ee>
<url>db/journals/corr/corr1702.html#KawaharaSBMTI17</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/KawaharaSMBT17" mdate="2018-08-13">
<author pid="84/3249">Hideki Kawahara</author>
<author pid="11/9231">Ken-Ichi Sakakibara</author>
<author pid="75/6649">Masanori Morise</author>
<author pid="95/6268">Hideki Banno</author>
<author pid="85/741">Tomoki Toda</author>
<title>A modulation property of time-frequency derivatives of filtered phase and its application to aperiodicity and fo estimation.</title>
<year>2017</year>
<volume>abs/1706.02964</volume>
<journal>CoRR</journal>
<ee type="oa">http://arxiv.org/abs/1706.02964</ee>
<url>db/journals/corr/corr1706.html#KawaharaSMBT17</url>
</article>
</r>
<r><article key="journals/ieicet/MakiTSNN16" mdate="2021-04-09">
<author pid="166/6698">Hayato Maki</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Enhancing Event-Related Potentials Based on Maximum a Posteriori Estimation with a Spatial Correlation Prior.</title>
<pages>1437-1446</pages>
<year>2016</year>
<volume>99-D</volume>
<journal>IEICE Trans. Inf. Syst.</journal>
<number>6</number>
<ee>https://doi.org/10.1587/transinf.2015CBP0008</ee>
<ee>http://search.ieice.org/bin/summary.php?id=e99-d_6_1437</ee>
<url>db/journals/ieicet/ieicet99d.html#MakiTSNN16</url>
</article>
</r>
<r><article key="journals/ieicet/TakamichiTNSN16" mdate="2021-04-09">
<author pid="124/9221">Shinnosuke Takamichi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>A Statistical Sample-Based Approach to GMM-Based Voice Conversion Using Tied-Covariance Acoustic Models.</title>
<pages>2490-2498</pages>
<year>2016</year>
<volume>99-D</volume>
<journal>IEICE Trans. Inf. Syst.</journal>
<number>10</number>
<ee>https://doi.org/10.1587/transinf.2016SLP0020</ee>
<ee>http://search.ieice.org/bin/summary.php?id=e99-d_10_2490</ee>
<url>db/journals/ieicet/ieicet99d.html#TakamichiTNSN16</url>
</article>
</r>
<r><article key="journals/ieicet/KobayashiTNGN16" mdate="2022-10-02">
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="49/5220">Tomoyasu Nakano</author>
<author orcid="0000-0003-1167-0977" pid="85/1753">Masataka Goto</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Improvements of Voice Timbre Control Based on Perceived Age in Singing Voice Conversion.</title>
<pages>2767-2777</pages>
<year>2016</year>
<volume>99-D</volume>
<journal>IEICE Trans. Inf. Syst.</journal>
<number>11</number>
<ee>https://doi.org/10.1587/transinf.2016EDP7234</ee>
<ee>http://search.ieice.org/bin/summary.php?id=e99-d_11_2767</ee>
<url>db/journals/ieicet/ieicet99d.html#KobayashiTNGN16</url>
</article>
</r>
<r><article key="journals/ieicet/OshimaTTNSN16" mdate="2021-04-09">
<author pid="130/8033">Yuji Oshima</author>
<author pid="124/9221">Shinnosuke Takamichi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Non-Native Text-to-Speech Preserving Speaker Individuality Based on Partial Correction of Prosodic and Phonetic Characteristics.</title>
<pages>3132-3139</pages>
<year>2016</year>
<volume>99-D</volume>
<journal>IEICE Trans. Inf. Syst.</journal>
<number>12</number>
<ee>https://doi.org/10.1587/transinf.2016EDP7231</ee>
<ee>http://search.ieice.org/bin/summary.php?id=e99-d_12_3132</ee>
<url>db/journals/ieicet/ieicet99d.html#OshimaTTNSN16</url>
</article>
</r>
<r><article key="journals/speech/HiraokaNSTN16" mdate="2021-04-09">
<author pid="31/977">Takuya Hiraoka</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Learning cooperative persuasive dialogue policies using framing.</title>
<pages>83-96</pages>
<year>2016</year>
<volume>84</volume>
<journal>Speech Commun.</journal>
<ee>https://doi.org/10.1016/j.specom.2016.09.002</ee>
<url>db/journals/speech/speech84.html#HiraokaNSTN16</url>
</article>
</r>
<r><article key="journals/taslp/TakamichiTBNSN16" mdate="2021-04-09">
<author pid="124/9221">Shinnosuke Takamichi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="b/AlanWBlack">Alan W. Black</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Postfilters to Modify the Modulation Spectrum for Statistical Parametric Speech Synthesis.</title>
<pages>755-767</pages>
<year>2016</year>
<volume>24</volume>
<journal>IEEE ACM Trans. Audio Speech Lang. Process.</journal>
<number>4</number>
<ee>https://doi.org/10.1109/TASLP.2016.2522655</ee>
<url>db/journals/taslp/taslp24.html#TakamichiTBNSN16</url>
</article>
</r>
<r><article key="journals/taslp/WuLDKKLSSTWY16" mdate="2024-01-29">
<author pid="29/8054-1">Zhizheng Wu 0001</author>
<author orcid="0000-0002-7665-9632" pid="15/4307">Phillip L. De Leon</author>
<author pid="21/2706">Cenk Demiroglu</author>
<author orcid="0000-0002-2873-4140" pid="149/9670">Ali Khodabakhsh 0001</author>
<author pid="68/2005">Simon King 0001</author>
<author pid="70/5210">Zhen-Hua Ling</author>
<author pid="17/7825">Daisuke Saito</author>
<author pid="124/9083">Bryan Stewart</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="30/1879">Mirjam Wester</author>
<author pid="87/3979">Junichi Yamagishi</author>
<title>Anti-Spoofing for Text-Independent Speaker Verification: An Initial Database, Comparison of Countermeasures, and Human Performance.</title>
<pages>768-783</pages>
<year>2016</year>
<volume>24</volume>
<journal>IEEE ACM Trans. Audio Speech Lang. Process.</journal>
<number>4</number>
<ee type="oa">https://doi.org/10.1109/TASLP.2016.2526653</ee>
<url>db/journals/taslp/taslp24.html#WuLDKKLSSTWY16</url>
</article>
</r>
<r><article key="journals/tiis/TanakaSNTNIN16" mdate="2022-06-23">
<author orcid="0000-0002-0548-6252" pid="66/2476">Hiroki Tanaka</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="160/4279">Hideki Negoro</author>
<author pid="160/4287">Hidemi Iwasaka</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Teaching Social Communication Skills Through Human-Agent Interaction.</title>
<pages>18:1-18:26</pages>
<year>2016</year>
<volume>6</volume>
<journal>ACM Trans. Interact. Intell. Syst.</journal>
<number>2</number>
<ee>https://doi.org/10.1145/2937757</ee>
<url>db/journals/tiis/tiis6.html#TanakaSNTNIN16</url>
</article>
</r>
<r><inproceedings key="conf/dcase/HayashiWTHRT16" mdate="2021-12-22">
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="39/3245-1">Shinji Watanabe 0001</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="46/3941">Takaaki Hori</author>
<author pid="36/4575">Jonathan Le Roux</author>
<author pid="38/3483">Kazuya Takeda</author>
<title>Bidirectional LSTM-HMM Hybrid System for Polyphonic Sound Event Detection.</title>
<pages>35-39</pages>
<year>2016</year>
<booktitle>DCASE</booktitle>
<ee type="oa">http://dcase.community/documents/workshop2016/proceedings/Hayashi-DCASE2016workshop.pdf</ee>
<crossref>conf/dcase/2016</crossref>
<url>db/conf/dcase/dcase2016.html#HayashiWTHRT16</url>
</inproceedings>
</r>
<r><inproceedings key="conf/embc/MakiTSNN16" mdate="2021-04-09">
<author pid="166/6698">Hayato Maki</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Removing noise from event-related potentials using a probabilistic generative model with grouped covariance matrices.</title>
<pages>3728-3731</pages>
<year>2016</year>
<booktitle>EMBC</booktitle>
<ee>https://doi.org/10.1109/EMBC.2016.7591538</ee>
<ee>https://www.wikidata.org/entity/Q38923507</ee>
<crossref>conf/embc/2016</crossref>
<url>db/conf/embc/embc2016.html#MakiTSNN16</url>
</inproceedings>
</r>
<r><inproceedings key="conf/eusipco/TanakaTNN16" mdate="2021-04-09">
<author pid="140/2727">Kou Tanaka</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Real-time vibration control of an electrolarynx based on statistical F0 contour prediction.</title>
<pages>1333-1337</pages>
<year>2016</year>
<booktitle>EUSIPCO</booktitle>
<ee>https://doi.org/10.1109/EUSIPCO.2016.7760465</ee>
<crossref>conf/eusipco/2016</crossref>
<url>db/conf/eusipco/eusipco2016.html#TanakaTNN16</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/YamaneKTNGN16" mdate="2024-10-06">
<author pid="180/2666">Soichi Yamane</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="49/5220">Tomoyasu Nakano</author>
<author orcid="0000-0003-1167-0977" pid="85/1753">Masataka Goto</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>An estimation method of voice timbre evaluation values using feature extraction with Gaussian mixture model based on reference singer.</title>
<pages>5265-5269</pages>
<year>2016</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2016.7472682</ee>
<crossref>conf/icassp/2016</crossref>
<url>db/conf/icassp/icassp2016.html#YamaneKTNGN16</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/TanakaKTN16" mdate="2021-04-09">
<author pid="140/2727">Kou Tanaka</author>
<author pid="97/941">Hirokazu Kameoka</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Statistical F0 prediction for electrolaryngeal speech enhancement considering generative process of F0 contours within product of experts framework.</title>
<pages>5665-5669</pages>
<year>2016</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2016.7472762</ee>
<crossref>conf/icassp/2016</crossref>
<url>db/conf/icassp/icassp2016.html#TanakaKTN16</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/KobayashiTN16" mdate="2021-04-09">
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Implementation of F0 transformation for statistical singing voice conversion based on direct waveform modification.</title>
<pages>5670-5674</pages>
<year>2016</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2016.7472763</ee>
<crossref>conf/icassp/2016</crossref>
<url>db/conf/icassp/icassp2016.html#KobayashiTN16</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/TajiriTN16" mdate="2021-04-09">
<author pid="173/6532">Yusuke Tajiri</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Noise suppression method for body-conducted soft speech enhancement based on external noise monitoring.</title>
<pages>5935-5939</pages>
<year>2016</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2016.7472816</ee>
<crossref>conf/icassp/2016</crossref>
<url>db/conf/icassp/icassp2016.html#TajiriTN16</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/TobingTKN16" mdate="2023-03-21">
<author orcid="0000-0003-2792-8418" pid="158/4210">Patrick Lumban Tobing</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="97/941">Hirokazu Kameoka</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Acoustic-to-Articulatory Inversion Mapping Based on Latent Trajectory Gaussian Mixture Model.</title>
<pages>953-957</pages>
<year>2016</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2016-1196</ee>
<crossref>conf/interspeech/2016</crossref>
<url>db/conf/interspeech/interspeech2016.html#TobingTKN16</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/TodaCSVWWY16" mdate="2023-02-18">
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="85/9874">Ling-Hui Chen</author>
<author pid="17/7825">Daisuke Saito</author>
<author pid="64/2286">Fernando Villavicencio</author>
<author pid="30/1879">Mirjam Wester</author>
<author pid="29/8054-1">Zhizheng Wu 0001</author>
<author pid="87/3979">Junichi Yamagishi</author>
<title>The Voice Conversion Challenge 2016.</title>
<pages>1632-1636</pages>
<year>2016</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2016-1066</ee>
<crossref>conf/interspeech/2016</crossref>
<url>db/conf/interspeech/interspeech2016.html#TodaCSVWWY16</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/KobayashiTNT16" mdate="2021-04-09">
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="124/9221">Shinnosuke Takamichi</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>The NU-NAIST Voice Conversion System for the Voice Conversion Challenge 2016.</title>
<pages>1667-1671</pages>
<year>2016</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2016-970</ee>
<crossref>conf/interspeech/2016</crossref>
<url>db/conf/interspeech/interspeech2016.html#KobayashiTNT16</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/TachibanaTSK16" mdate="2021-04-09">
<author pid="35/8056">Kentaro Tachibana</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="02/2978">Yoshinori Shiga</author>
<author pid="32/4341">Hisashi Kawai</author>
<title>Model Integration for HMM- and DNN-Based Speech Synthesis Using Product-of-Experts Framework.</title>
<pages>2288-2292</pages>
<year>2016</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2016-1006</ee>
<crossref>conf/interspeech/2016</crossref>
<url>db/conf/interspeech/interspeech2016.html#TachibanaTSK16</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/DoTNSN16" mdate="2021-04-09">
<author pid="166/6489">Quoc Truong Do</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>A Hybrid System for Continuous Word-Level Emphasis Modeling Based on HMM State Clustering and Adaptive Training.</title>
<pages>3196-3200</pages>
<year>2016</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2016-930</ee>
<crossref>conf/interspeech/2016</crossref>
<url>db/conf/interspeech/interspeech2016.html#DoTNSN16</url>
</inproceedings>
</r>
<r><inproceedings key="conf/iwsds/HiraokaNYT016" mdate="2021-04-09">
<author pid="31/977">Takuya Hiraoka</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="59/8883">Koichiro Yoshino</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Active Learning for Example-Based Dialog Systems.</title>
<pages>67-78</pages>
<year>2016</year>
<booktitle>IWSDS</booktitle>
<ee>https://doi.org/10.1007/978-981-10-2585-3_5</ee>
<crossref>conf/iwsds/2016</crossref>
<url>db/conf/iwsds/iwsds2016.html#HiraokaNYT016</url>
</inproceedings>
</r>
<r><inproceedings key="conf/slt/KobayashiTN16" mdate="2021-04-09">
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>F0 transformation techniques for statistical voice conversion with direct waveform modification with spectral differential.</title>
<pages>693-700</pages>
<year>2016</year>
<booktitle>SLT</booktitle>
<ee>https://doi.org/10.1109/SLT.2016.7846338</ee>
<crossref>conf/slt/2016</crossref>
<url>db/conf/slt/slt2016.html#KobayashiTN16</url>
</inproceedings>
</r>
<r><inproceedings key="conf/ssw/TajiriT16" mdate="2021-02-01">
<author pid="173/6532">Yusuke Tajiri</author>
<author pid="85/741">Tomoki Toda</author>
<title>Nonaudible murmur enhancement based on statistical voice conversion and noise suppression with external noise monitoring.</title>
<pages>52-58</pages>
<year>2016</year>
<booktitle>SSW</booktitle>
<ee type="oa">https://doi.org/10.21437/SSW.2016-9</ee>
<crossref>conf/ssw/2016</crossref>
<url>db/conf/ssw/ssw2016.html#TajiriT16</url>
</inproceedings>
</r>
<r><article key="journals/ieicet/TanakaSNTN15" mdate="2022-06-23">
<author orcid="0000-0002-0548-6252" pid="66/2476">Hiroki Tanaka</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>NOCOA+: Multimodal Computer-Based Training for Social and Communication Skills.</title>
<pages>1536-1544</pages>
<year>2015</year>
<volume>98-D</volume>
<journal>IEICE Trans. Inf. Syst.</journal>
<number>8</number>
<ee>https://doi.org/10.1587/transinf.2014EDP7400</ee>
<ee>http://search.ieice.org/bin/summary.php?id=e98-d_8_1536</ee>
<url>db/journals/ieicet/ieicet98d.html#TanakaSNTN15</url>
</article>
</r>
<r><article key="journals/tacl/ArthurNSTN15" mdate="2024-06-19">
<author pid="150/5396">Philip Arthur</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Semantic Parsing of Ambiguous Input through Paraphrasing and Verification.</title>
<pages>571-584</pages>
<year>2015</year>
<volume>3</volume>
<journal>Trans. Assoc. Comput. Linguistics</journal>
<ee type="oa">https://doi.org/10.1162/tacl_a_00159</ee>
<url>db/journals/tacl/tacl3.html#ArthurNSTN15</url>
</article>
</r>
<r><inproceedings key="books/sp/15/MizukamiNSTN15" mdate="2021-10-14">
<author pid="169/3173">Masahiro Mizukami</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Linguistic Individuality Transformation for Spoken Language.</title>
<pages>129-143</pages>
<year>2015</year>
<booktitle>IWSDS</booktitle>
<ee>https://doi.org/10.1007/978-3-319-19291-8_13</ee>
<crossref>books/sp/2015LKJK</crossref>
<url>db/books/collections/LKJK2015.html#MizukamiNSTN15</url>
</inproceedings>
</r>
<r><inproceedings key="books/sp/15/KotoSNTAN15" mdate="2021-10-14">
<author pid="160/0019">Fajri Koto</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="15/6057">Mirna Adriani</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>A Study on Natural Expressive Speech: Automatic Memorable Spoken Quote Detection.</title>
<pages>145-152</pages>
<year>2015</year>
<booktitle>IWSDS</booktitle>
<ee>https://doi.org/10.1007/978-3-319-19291-8_14</ee>
<crossref>books/sp/2015LKJK</crossref>
<url>db/books/collections/LKJK2015.html#KotoSNTAN15</url>
</inproceedings>
</r>
<r><inproceedings key="books/sp/15/HiraokaNSTN15" mdate="2021-10-14">
<author pid="31/977">Takuya Hiraoka</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Evaluation of a Fully Automatic Cooperative Persuasive Dialogue System.</title>
<pages>153-167</pages>
<year>2015</year>
<booktitle>IWSDS</booktitle>
<ee>https://doi.org/10.1007/978-3-319-19291-8_15</ee>
<crossref>books/sp/2015LKJK</crossref>
<url>db/books/collections/LKJK2015.html#HiraokaNSTN15</url>
</inproceedings>
</r>
<r><inproceedings key="books/sp/15/SasakuraSNTN15" mdate="2021-10-14">
<author pid="159/9937">Takafumi Sasakura</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Unknown Word Detection Based on Event-Related Brain Desynchronization Responses.</title>
<pages>169-175</pages>
<year>2015</year>
<booktitle>IWSDS</booktitle>
<ee>https://doi.org/10.1007/978-3-319-19291-8_16</ee>
<crossref>books/sp/2015LKJK</crossref>
<url>db/books/collections/LKJK2015.html#SasakuraSNTN15</url>
</inproceedings>
</r>
<r><inproceedings key="books/sp/15/TsunomoriNSTN15" mdate="2021-10-14">
<author pid="195/4988">Yuiko Tsunomori</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>An Analysis Towards Dialogue-Based Deception Detection.</title>
<pages>177-187</pages>
<year>2015</year>
<booktitle>IWSDS</booktitle>
<ee>https://doi.org/10.1007/978-3-319-19291-8_17</ee>
<crossref>books/sp/2015LKJK</crossref>
<url>db/books/collections/LKJK2015.html#TsunomoriNSTN15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/acl/OdaNSTN15" mdate="2021-08-06">
<author pid="148/4505">Yusuke Oda</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Syntax-based Simultaneous Translation through Prediction of Unseen Syntactic Constituents.</title>
<pages>198-207</pages>
<year>2015</year>
<booktitle>ACL (1)</booktitle>
<ee type="oa">https://doi.org/10.3115/v1/p15-1020</ee>
<ee type="oa">https://aclanthology.org/P15-1020/</ee>
<crossref>conf/acl/2015-1</crossref>
<url>db/conf/acl/acl2015-1.html#OdaNSTN15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/acl/MiuraNSTN15" mdate="2021-08-06">
<author pid="166/1748">Akiva Miura</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Improving Pivot Translation by Remembering the Pivot.</title>
<pages>573-577</pages>
<year>2015</year>
<booktitle>ACL (2)</booktitle>
<ee type="oa">https://doi.org/10.3115/v1/p15-2094</ee>
<ee type="oa">https://aclanthology.org/P15-2094/</ee>
<crossref>conf/acl/2015-2</crossref>
<url>db/conf/acl/acl2015-2.html#MiuraNSTN15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/apsipa/KawaharaSBMTI15" mdate="2021-04-09">
<author pid="84/3249">Hideki Kawahara</author>
<author pid="11/9231">Ken-Ichi Sakakibara</author>
<author pid="95/6268">Hideki Banno</author>
<author pid="75/6649">Masanori Morise</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="72/3903">Toshio Irino</author>
<title>Aliasing-free implementation of discrete-time glottal source models and their applications to speech synthesis and F0 extractor evaluation.</title>
<pages>520-529</pages>
<year>2015</year>
<booktitle>APSIPA</booktitle>
<ee>https://doi.org/10.1109/APSIPA.2015.7415325</ee>
<crossref>conf/apsipa/2015</crossref>
<url>db/conf/apsipa/apsipa2015.html#KawaharaSBMTI15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/asru/SaktiINTPN15" mdate="2025-05-31">
<author pid="71/3717">Sakriani Sakti</author>
<author pid="175/8915">Faiz Ilham</author>
<author pid="03/8155">Graham Neubig</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author orcid="0000-0002-5016-3700" pid="15/2527">Ayu Purwarianti</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Incremental sentence compression using LSTM recurrent networks.</title>
<pages>252-258</pages>
<year>2015</year>
<booktitle>ASRU</booktitle>
<ee>https://doi.org/10.1109/ASRU.2015.7404802</ee>
<crossref>conf/asru/2015</crossref>
<url>db/conf/asru/asru2015.html#SaktiINTPN15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/asru/DoHSNTN15" mdate="2026-07-18">
<author pid="166/6489">Quoc Truong Do</author>
<author orcid="0000-0001-9841-5025" pid="120/6962">Michael Heck</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>The NAIST ASR system for the 2015 Multi-Genre Broadcast challenge: On combination of deep learning systems using a rank-score function.</title>
<pages>654-659</pages>
<year>2015</year>
<booktitle>ASRU</booktitle>
<ee>https://doi.org/10.1109/ASRU.2015.7404858</ee>
<crossref>conf/asru/2015</crossref>
<url>db/conf/asru/asru2015.html#DoHSNTN15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/asru/LubisSNYTN15" mdate="2021-04-09">
<author pid="160/8114">Nurul Lubis</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="59/8883">Koichiro Yoshino</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>A study of social-affective communication: Automatic prediction of emotion triggers and responses in television talk shows.</title>
<pages>777-783</pages>
<year>2015</year>
<booktitle>ASRU</booktitle>
<ee>https://doi.org/10.1109/ASRU.2015.7404867</ee>
<crossref>conf/asru/2015</crossref>
<url>db/conf/asru/asru2015.html#LubisSNYTN15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/asru/MizukamiKNNYSTN15" mdate="2021-04-09">
<author pid="169/3173">Masahiro Mizukami</author>
<author pid="175/8818">Hideaki Kizuki</author>
<author pid="43/4382">Toshio Nomura</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="59/8883">Koichiro Yoshino</author>
<author pid="71/3717">Sakriani Sakti</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Adaptive selection from multiple response candidates in example-based dialogue.</title>
<pages>784-790</pages>
<year>2015</year>
<booktitle>ASRU</booktitle>
<ee>https://doi.org/10.1109/ASRU.2015.7404868</ee>
<crossref>conf/asru/2015</crossref>
<url>db/conf/asru/asru2015.html#MizukamiKNNYSTN15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/assets/TanakaTNSN15" mdate="2021-04-09">
<author pid="140/2727">Kou Tanaka</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>An Enhanced Electrolarynx with Automatic Fundamental Frequency Control based on Statistical Prediction.</title>
<pages>435-436</pages>
<year>2015</year>
<booktitle>ASSETS</booktitle>
<ee>https://doi.org/10.1145/2700648.2811340</ee>
<crossref>conf/assets/2015</crossref>
<url>db/conf/assets/assets2015.html#TanakaTNSN15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/blizzard/TakamichiKTT015" mdate="2024-09-27">
<author pid="124/9221">Shinnosuke Takamichi</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="140/2727">Kou Tanaka</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>The NAIST Text-to-Speech System for the Blizzard Challenge 2015.</title>
<year>2015</year>
<booktitle>Blizzard Challenge</booktitle>
<ee type="oa">https://doi.org/10.21437/Blizzard.2015-7</ee>
<crossref>conf/blizzard/2015</crossref>
<url>db/conf/blizzard/blizzard2015.html#TakamichiKTT015</url>
</inproceedings>
</r>
<r><inproceedings key="conf/embc/MakiTSNN15" mdate="2021-04-09">
<author pid="166/6698">Hayato Maki</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>An evaluation of EEG ocular artifact removal with a multi-channel wiener filter based on probabilistic generative model.</title>
<pages>2775-2778</pages>
<year>2015</year>
<booktitle>EMBC</booktitle>
<ee>https://doi.org/10.1109/EMBC.2015.7318967</ee>
<ee>https://www.wikidata.org/entity/Q40138572</ee>
<crossref>conf/embc/2015</crossref>
<url>db/conf/embc/embc2015.html#MakiTSNN15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/MakiTSNN15" mdate="2021-04-09">
<author pid="166/6698">Hayato Maki</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>EEG signal enhancement using multi-channel wiener filter with a spatial correlation prior.</title>
<pages>2639-2643</pages>
<year>2015</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2015.7178449</ee>
<crossref>conf/icassp/2015</crossref>
<url>db/conf/icassp/icassp2015.html#MakiTSNN15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/TakamichiTBN15" mdate="2021-04-09">
<author pid="124/9221">Shinnosuke Takamichi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="b/AlanWBlack">Alan W. Black</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Parameter generation algorithm considering Modulation Spectrum for HMM-based speech synthesis.</title>
<pages>4210-4214</pages>
<year>2015</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2015.7178764</ee>
<crossref>conf/icassp/2015</crossref>
<url>db/conf/icassp/icassp2015.html#TakamichiTBN15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/WuKDYSTK15" mdate="2024-01-29">
<author pid="29/8054-1">Zhizheng Wu 0001</author>
<author orcid="0000-0002-2873-4140" pid="149/9670">Ali Khodabakhsh 0001</author>
<author pid="21/2706">Cenk Demiroglu</author>
<author pid="87/3979">Junichi Yamagishi</author>
<author pid="17/7825">Daisuke Saito</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="68/2005">Simon King 0001</author>
<title>SAS: A speaker verification spoofing database containing diverse attacks.</title>
<pages>4440-4444</pages>
<year>2015</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2015.7178810</ee>
<crossref>conf/icassp/2015</crossref>
<url>db/conf/icassp/icassp2015.html#WuKDYSTK15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/TjandraSNTAN15" mdate="2021-04-09">
<author pid="166/6471">Andros Tjandra</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="15/6057">Mirna Adriani</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Combination of two-dimensional cochleogram and spectrogram features for deep learning-based ASR.</title>
<pages>4525-4529</pages>
<year>2015</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2015.7178827</ee>
<crossref>conf/icassp/2015</crossref>
<url>db/conf/icassp/icassp2015.html#TjandraSNTAN15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/TakamichiTBN15a" mdate="2021-04-09">
<author pid="124/9221">Shinnosuke Takamichi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="b/AlanWBlack">Alan W. Black</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Modulation spectrum-constrained trajectory training algorithm for GMM-based Voice Conversion.</title>
<pages>4859-4863</pages>
<year>2015</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2015.7178894</ee>
<crossref>conf/icassp/2015</crossref>
<url>db/conf/icassp/icassp2015.html#TakamichiTBN15a</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/OshimaTTNSN15" mdate="2023-06-23">
<author pid="130/8033">Yuji Oshima</author>
<author pid="124/9221">Shinnosuke Takamichi</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Non-native speech synthesis preserving speaker individuality based on partial correction of prosodic and phonetic characteristics.</title>
<pages>299-303</pages>
<year>2015</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2015-121</ee>
<crossref>conf/interspeech/2015</crossref>
<url>db/conf/interspeech/interspeech2015.html#OshimaTTNSN15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/TakamichiTBN15" mdate="2023-06-23">
<author pid="124/9221">Shinnosuke Takamichi</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="b/AlanWBlack">Alan W. Black</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Modulation spectrum-constrained trajectory training algorithm for HMM-based speech synthesis.</title>
<pages>1206-1210</pages>
<year>2015</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2015-305</ee>
<crossref>conf/interspeech/2015</crossref>
<url>db/conf/interspeech/interspeech2015.html#TakamichiTBN15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/MienoNSTN15" mdate="2023-06-23">
<author pid="173/6560">Takashi Mieno</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Speed or accuracy? a study in evaluation of simultaneous speech translation.</title>
<pages>2267-2271</pages>
<year>2015</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2015-498</ee>
<crossref>conf/interspeech/2015</crossref>
<url>db/conf/interspeech/interspeech2015.html#MienoNSTN15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/NguyenNSSTN15" mdate="2023-06-23">
<author pid="03/7981">The Tung Nguyen</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="97/8317">Hiroyuki Shindo</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>A latent variable model for joint pause prediction and dependency parsing.</title>
<pages>2719-2723</pages>
<year>2015</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2015-573</ee>
<crossref>conf/interspeech/2015</crossref>
<url>db/conf/interspeech/interspeech2015.html#NguyenNSSTN15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/KobayashiTNSN15" mdate="2023-06-23">
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Statistical singing voice conversion based on direct waveform modification with global variance.</title>
<pages>2754-2758</pages>
<year>2015</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2015-580</ee>
<crossref>conf/interspeech/2015</crossref>
<url>db/conf/interspeech/interspeech2015.html#KobayashiTNSN15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/TajiriTTNSN15" mdate="2023-06-23">
<author pid="173/6532">Yusuke Tajiri</author>
<author pid="140/2727">Kou Tanaka</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Non-audible murmur enhancement based on statistical conversion using air- and body-conductive microphones in noisy environments.</title>
<pages>2769-2773</pages>
<year>2015</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2015-583</ee>
<crossref>conf/interspeech/2015</crossref>
<url>db/conf/interspeech/interspeech2015.html#TajiriTTNSN15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/TobingKTNSN15" mdate="2023-06-23">
<author pid="158/4210">Patrick Lumban Tobing</author>
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Articulatory controllable speech modification based on Gaussian mixture models with direct waveform modification using spectrum differential.</title>
<pages>3350-3354</pages>
<year>2015</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2015-138</ee>
<crossref>conf/interspeech/2015</crossref>
<url>db/conf/interspeech/interspeech2015.html#TobingKTNSN15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/DoTSNTN15" mdate="2023-06-23">
<author pid="166/6489">Quoc Truong Do</author>
<author pid="124/9221">Shinnosuke Takamichi</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Preserving word-level emphasis in speech-to-speech translation using linear regression HSMMs.</title>
<pages>3665-3669</pages>
<year>2015</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2015-727</ee>
<crossref>conf/interspeech/2015</crossref>
<url>db/conf/interspeech/interspeech2015.html#DoTSNTN15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/iui/TanakaSNTNIN15" mdate="2022-06-23">
<author orcid="0000-0002-0548-6252" pid="66/2476">Hiroki Tanaka</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="160/4279">Hideki Negoro</author>
<author pid="160/4287">Hidemi Iwasaka</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Automated Social Skills Trainer.</title>
<pages>17-27</pages>
<year>2015</year>
<booktitle>IUI</booktitle>
<ee>https://doi.org/10.1145/2678025.2701368</ee>
<crossref>conf/iui/2015</crossref>
<url>db/conf/iui/iui2015.html#TanakaSNTNIN15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/iwslt/DoSNTN15" mdate="2024-08-01">
<author pid="166/6489">Quoc Truong Do</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Improving translation of emphasis with pause prediction in speech-to-speech translation systems.</title>
<year>2015</year>
<booktitle>IWSLT</booktitle>
<ee type="oa">https://aclanthology.org/2015.iwslt-papers.12</ee>
<crossref>conf/iwslt/2015</crossref>
<url>db/conf/iwslt/iwslt2015.html#DoSNTN15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/kbse/OdaFNHSTN15" mdate="2023-03-24">
<author pid="148/4505">Yusuke Oda</author>
<author pid="16/1027">Hiroyuki Fudaba</author>
<author pid="03/8155">Graham Neubig</author>
<author orcid="0000-0003-0708-5222" pid="57/4288">Hideaki Hata</author>
<author pid="71/3717">Sakriani Sakti</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Learning to Generate Pseudo-Code from Source Code Using Statistical Machine Translation (T).</title>
<pages>574-584</pages>
<year>2015</year>
<booktitle>ASE</booktitle>
<ee>https://doi.org/10.1109/ASE.2015.36</ee>
<ee>https://doi.ieeecomputersociety.org/10.1109/ASE.2015.36</ee>
<ee>https://dl.acm.org/citation.cfm?id=3343959</ee>
<crossref>conf/kbse/2015</crossref>
<url>db/conf/kbse/ase2015.html#OdaFNHSTN15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/kbse/FudabaOANHSTN15" mdate="2023-09-30">
<author pid="16/1027">Hiroyuki Fudaba</author>
<author pid="148/4505">Yusuke Oda</author>
<author orcid="0000-0002-1624-497X" pid="151/8457">Koichi Akabe</author>
<author pid="03/8155">Graham Neubig</author>
<author orcid="0000-0003-0708-5222" pid="57/4288">Hideaki Hata</author>
<author pid="71/3717">Sakriani Sakti</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Pseudogen: A Tool to Automatically Generate Pseudo-Code from Source Code.</title>
<pages>824-829</pages>
<year>2015</year>
<booktitle>ASE</booktitle>
<ee>https://doi.org/10.1109/ASE.2015.107</ee>
<ee>https://doi.ieeecomputersociety.org/10.1109/ASE.2015.107</ee>
<ee>https://dl.acm.org/citation.cfm?id=3343995</ee>
<crossref>conf/kbse/2015</crossref>
<url>db/conf/kbse/ase2015.html#FudabaOANHSTN15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/naacl/OdaNSTN15" mdate="2021-08-06">
<author pid="148/4505">Yusuke Oda</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Ckylark: A More Robust PCFG-LA Parser.</title>
<pages>41-45</pages>
<year>2015</year>
<booktitle>HLT-NAACL</booktitle>
<ee type="oa">https://doi.org/10.3115/v1/n15-3009</ee>
<ee type="oa">https://aclanthology.org/N15-3009/</ee>
<crossref>conf/naacl/2015</crossref>
<url>db/conf/naacl/naacl2015.html#OdaNSTN15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/ococosda/LubisSNTN15" mdate="2021-04-09">
<author pid="160/8114">Nurul Lubis</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Construction and analysis of social-affective interaction corpus in English and Indonesian.</title>
<pages>202-206</pages>
<year>2015</year>
<booktitle>O-COCOSDA/CASLRE</booktitle>
<ee>https://doi.org/10.1109/ICSDA.2015.7357892</ee>
<crossref>conf/ococosda/2015</crossref>
<url>db/conf/ococosda/ococosda2015.html#LubisSNTN15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/wmt/SugiyamaMNYSTN15" mdate="2021-08-06">
<author pid="154/5283">Kyoshiro Sugiyama</author>
<author pid="169/3173">Masahiro Mizukami</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="59/8883">Koichiro Yoshino</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>An Investigation of Machine Translation Evaluation Metrics in Cross-lingual Question Answering.</title>
<pages>442-449</pages>
<year>2015</year>
<booktitle>WMT@EMNLP</booktitle>
<ee type="oa">https://doi.org/10.18653/v1/w15-3057</ee>
<ee type="oa">https://aclanthology.org/W15-3057/</ee>
<crossref>conf/wmt/2015</crossref>
<url>db/conf/wmt/wmt2015.html#SugiyamaMNYSTN15</url>
</inproceedings>
</r>
<r><article key="journals/ieicet/KobayashiTDNGNSN14" mdate="2024-10-06">
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="68/4056">Hironori Doi</author>
<author pid="49/5220">Tomoyasu Nakano</author>
<author orcid="0000-0003-1167-0977" pid="85/1753">Masataka Goto</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Voice Timbre Control Based on Perceived Age in Singing Voice Conversion.</title>
<pages>1419-1428</pages>
<year>2014</year>
<volume>97-D</volume>
<journal>IEICE Trans. Inf. Syst.</journal>
<number>6</number>
<ee>https://doi.org/10.1587/transinf.E97.D.1419</ee>
<ee>http://search.ieice.org/bin/summary.php?id=e97-d_6_1419</ee>
<url>db/journals/ieicet/ieicet97d.html#KobayashiTDNGNSN14</url>
</article>
</r>
<r><article key="journals/ieicet/TanakaTNSN14" mdate="2021-04-09">
<author pid="140/2727">Kou Tanaka</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>A Hybrid Approach to Electrolaryngeal Speech Enhancement Based on Noise Reduction and Statistical Excitation Generation.</title>
<pages>1429-1437</pages>
<year>2014</year>
<volume>97-D</volume>
<journal>IEICE Trans. Inf. Syst.</journal>
<number>6</number>
<ee>https://doi.org/10.1587/transinf.E97.D.1429</ee>
<ee>http://search.ieice.org/bin/summary.php?id=e97-d_6_1429</ee>
<url>db/journals/ieicet/ieicet97d.html#TanakaTNSN14</url>
</article>
</r>
<r><article key="journals/ieicet/KuboSNTN14" mdate="2021-04-09">
<author pid="124/9006">Keigo Kubo</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Structured Adaptive Regularization of Weight Vectors for a Robust Grapheme-to-Phoneme Conversion Model.</title>
<pages>1468-1476</pages>
<year>2014</year>
<volume>97-D</volume>
<journal>IEICE Trans. Inf. Syst.</journal>
<number>6</number>
<ee>https://doi.org/10.1587/transinf.E97.D.1468</ee>
<ee>http://search.ieice.org/bin/summary.php?id=e97-d_6_1468</ee>
<url>db/journals/ieicet/ieicet97d.html#KuboSNTN14</url>
</article>
</r>
<r><article key="journals/ieicet/NioSNTN14" mdate="2021-04-09">
<author pid="146/8261">Lasguido Nio</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Utilizing Human-to-Human Conversation Examples for a Multi Domain Chat-Oriented Dialog System.</title>
<pages>1497-1505</pages>
<year>2014</year>
<volume>97-D</volume>
<journal>IEICE Trans. Inf. Syst.</journal>
<number>6</number>
<ee>https://doi.org/10.1587/transinf.E97.D.1497</ee>
<ee>http://search.ieice.org/bin/summary.php?id=e97-d_6_1497</ee>
<url>db/journals/ieicet/ieicet97d.html#NioSNTN14</url>
</article>
</r>
<r><article key="journals/jstsp/TakamichiTSSNN14" mdate="2021-04-09">
<author pid="124/9221">Shinnosuke Takamichi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="02/2978">Yoshinori Shiga</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Parameter Generation Methods With Rich Context Models for High-Quality and Flexible Text-To-Speech Synthesis.</title>
<pages>239-250</pages>
<year>2014</year>
<volume>8</volume>
<journal>IEEE J. Sel. Top. Signal Process.</journal>
<number>2</number>
<ee>https://doi.org/10.1109/JSTSP.2013.2288599</ee>
<url>db/journals/jstsp/jstsp8.html#TakamichiTSSNN14</url>
</article>
</r>
<r><article key="journals/taslp/DoiTNSS14" mdate="2021-04-09">
<author pid="68/4056">Hironori Doi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="31/8013">Keigo Nakamura</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Alaryngeal Speech Enhancement Based on One-to-Many Eigenvoice Conversion.</title>
<pages>172-183</pages>
<year>2014</year>
<volume>22</volume>
<journal>IEEE ACM Trans. Audio Speech Lang. Process.</journal>
<number>1</number>
<ee>https://doi.org/10.1109/TASLP.2013.2286917</ee>
<url>db/journals/taslp/taslp22.html#DoiTNSS14</url>
</article>
</r>
<r><inproceedings key="conf/acl/OdaNSTN14" mdate="2021-08-06">
<author pid="148/4505">Yusuke Oda</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Optimizing Segmentation Strategies for Simultaneous Speech Translation.</title>
<pages>551-556</pages>
<year>2014</year>
<booktitle>ACL (2)</booktitle>
<ee type="oa">https://doi.org/10.3115/v1/p14-2090</ee>
<ee type="oa">https://aclanthology.org/P14-2090/</ee>
<crossref>conf/acl/2014-2</crossref>
<url>db/conf/acl/acl2014-2.html#OdaNSTN14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/acl-clpsych/TanakaSNTN14" mdate="2020-03-20">
<author pid="66/2476">Hiroki Tanaka</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Linguistic and Acoustic Features for Automatic Identification of Autism Spectrum Disorders in Children's Narrative.</title>
<pages>88-96</pages>
<year>2014</year>
<booktitle>CLPsych@ACL</booktitle>
<ee>https://doi.org/10.3115/v1/W14-3211</ee>
<crossref>conf/acl-clpsych/2014</crossref>
<url>db/conf/acl-clpsych/acl-clpsych2014.html#TanakaSNTN14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/apsipa/KawaharaMTBNI14" mdate="2021-04-09">
<author pid="84/3249">Hideki Kawahara</author>
<author pid="75/6649">Masanori Morise</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="95/6268">Hideki Banno</author>
<author pid="75/2915">Ryuichi Nisimura</author>
<author pid="72/3903">Toshio Irino</author>
<title>Excitation source design for high-quality speech manipulation systems based on a temporally static group delay representation of periodic signals.</title>
<pages>1-10</pages>
<year>2014</year>
<booktitle>APSIPA</booktitle>
<ee>https://doi.org/10.1109/APSIPA.2014.7041594</ee>
<crossref>conf/apsipa/2014</crossref>
<url>db/conf/apsipa/apsipa2014.html#KawaharaMTBNI14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/apsipa/KobayashiTNGNSN14" mdate="2024-10-06">
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="49/5220">Tomoyasu Nakano</author>
<author orcid="0000-0003-1167-0977" pid="85/1753">Masataka Goto</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Gender-dependent spectrum differential models for perceived age control based on direct waveform modification in singing voice conversion.</title>
<pages>1-4</pages>
<year>2014</year>
<booktitle>APSIPA</booktitle>
<ee>https://doi.org/10.1109/APSIPA.2014.7041590</ee>
<crossref>conf/apsipa/2014</crossref>
<url>db/conf/apsipa/apsipa2014.html#KobayashiTNGNSN14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/apsipa/KotoSNTAN14" mdate="2021-04-09">
<author pid="160/0019">Fajri Koto</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="15/6057">Mirna Adriani</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>The use of semantic and acoustic features for open-domain TED talk summarization.</title>
<pages>1-4</pages>
<year>2014</year>
<booktitle>APSIPA</booktitle>
<ee>https://doi.org/10.1109/APSIPA.2014.7041625</ee>
<crossref>conf/apsipa/2014</crossref>
<url>db/conf/apsipa/apsipa2014.html#KotoSNTAN14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/apsipa/NioSNTN14" mdate="2021-04-09">
<author pid="146/8261">Lasguido Nio</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Recursive neural network paraphrase identification for example-based dialog retrieval.</title>
<pages>1-4</pages>
<year>2014</year>
<booktitle>APSIPA</booktitle>
<ee>https://doi.org/10.1109/APSIPA.2014.7041777</ee>
<crossref>conf/apsipa/2014</crossref>
<url>db/conf/apsipa/apsipa2014.html#NioSNTN14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/apsipa/SaktiOSNTN14" mdate="2021-04-09">
<author pid="71/3717">Sakriani Sakti</author>
<author pid="159/9930">Yu Odagaki</author>
<author pid="159/9937">Takafumi Sasakura</author>
<author pid="03/8155">Graham Neubig</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>An event-related brain potential study on the impact of speech recognition errors.</title>
<pages>1-4</pages>
<year>2014</year>
<booktitle>APSIPA</booktitle>
<ee>https://doi.org/10.1109/APSIPA.2014.7041620</ee>
<crossref>conf/apsipa/2014</crossref>
<url>db/conf/apsipa/apsipa2014.html#SaktiOSNTN14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/apsipa/TakamichiTBN14" mdate="2021-04-09">
<author pid="124/9221">Shinnosuke Takamichi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="b/AlanWBlack">Alan W. Black</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Modulation spectrum-based post-filter for GMM-based Voice Conversion.</title>
<pages>1-4</pages>
<year>2014</year>
<booktitle>APSIPA</booktitle>
<ee>https://doi.org/10.1109/APSIPA.2014.7041540</ee>
<crossref>conf/apsipa/2014</crossref>
<url>db/conf/apsipa/apsipa2014.html#TakamichiTBN14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/apsipa/TanakaTNSN14" mdate="2021-04-09">
<author pid="140/2727">Kou Tanaka</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>An inter-speaker evaluation through simulation of electrolarynx control based on statistical F0 prediction.</title>
<pages>1-4</pages>
<year>2014</year>
<booktitle>APSIPA</booktitle>
<ee>https://doi.org/10.1109/APSIPA.2014.7041593</ee>
<crossref>conf/apsipa/2014</crossref>
<url>db/conf/apsipa/apsipa2014.html#TanakaTNSN14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/apsipa/TsurutaTTNSN14" mdate="2021-04-09">
<author pid="160/0028">Sakura Tsuruta</author>
<author pid="140/2727">Kou Tanaka</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>An evaluation of target speech for a nonaudible murmur enhancement system in noisy environments.</title>
<pages>1-4</pages>
<year>2014</year>
<booktitle>APSIPA</booktitle>
<ee>https://doi.org/10.1109/APSIPA.2014.7041618</ee>
<crossref>conf/apsipa/2014</crossref>
<url>db/conf/apsipa/apsipa2014.html#TsurutaTTNSN14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/apsipa/YoshidaHNSTN14" mdate="2021-04-09">
<author pid="159/9947">Riki Yoshida</author>
<author pid="31/977">Takuya Hiraoka</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Unnecessary utterance detection for avoiding digressions in discussion.</title>
<pages>1-4</pages>
<year>2014</year>
<booktitle>APSIPA</booktitle>
<ee>https://doi.org/10.1109/APSIPA.2014.7041572</ee>
<crossref>conf/apsipa/2014</crossref>
<url>db/conf/apsipa/apsipa2014.html#YoshidaHNSTN14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/coling/AkabeNSTN14" mdate="2021-08-06">
<author pid="151/8457">Koichi Akabe</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Discriminative Language Models as a Tool for Machine Translation Error Analysis.</title>
<pages>1124-1132</pages>
<year>2014</year>
<booktitle>COLING</booktitle>
<ee type="oa">https://aclanthology.org/C14-1106/</ee>
<crossref>conf/coling/2014</crossref>
<url>db/conf/coling/coling2014.html#AkabeNSTN14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/coling/HiraokaNSTN14" mdate="2021-08-06">
<author pid="31/977">Takuya Hiraoka</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Reinforcement Learning of Cooperative Persuasive Dialogue Policies using Framing.</title>
<pages>1706-1717</pages>
<year>2014</year>
<booktitle>COLING</booktitle>
<ee type="oa">https://aclanthology.org/C14-1161/</ee>
<crossref>conf/coling/2014</crossref>
<url>db/conf/coling/coling2014.html#HiraokaNSTN14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/eacl/VuNSTN14" mdate="2021-08-06">
<author pid="117/4989">Hoa Trong Vu</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Acquiring a Dictionary of Emotion-Provoking Events.</title>
<pages>128-132</pages>
<year>2014</year>
<booktitle>EACL</booktitle>
<ee type="oa">https://doi.org/10.3115/v1/e14-4025</ee>
<ee type="oa">https://aclanthology.org/E14-4025/</ee>
<crossref>conf/eacl/2014</crossref>
<url>db/conf/eacl/eacl2014.html#VuNSTN14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/globalsip/TakamichiTBN14" mdate="2021-04-09">
<author pid="124/9221">Shinnosuke Takamichi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="b/AlanWBlack">Alan W. Black</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Modified post-filter to recover modulation spectrum for HMM-based speech synthesis.</title>
<pages>547-551</pages>
<year>2014</year>
<booktitle>GlobalSIP</booktitle>
<ee>https://doi.org/10.1109/GlobalSIP.2014.7032177</ee>
<crossref>conf/globalsip/2014</crossref>
<url>db/conf/globalsip/globalsip2014.html#TakamichiTBN14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/globalsip/Toda14" mdate="2021-04-09">
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Augmented speech production based on real-time statistical voice conversion.</title>
<pages>592-596</pages>
<year>2014</year>
<booktitle>GlobalSIP</booktitle>
<ee>https://doi.org/10.1109/GlobalSIP.2014.7032186</ee>
<crossref>conf/globalsip/2014</crossref>
<url>db/conf/globalsip/globalsip2014.html#Toda14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/TakamichiTNSN14" mdate="2021-04-09">
<author pid="124/9221">Shinnosuke Takamichi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>A postfilter to modify the modulation spectrum in HMM-based speech synthesis.</title>
<pages>290-294</pages>
<year>2014</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2014.6853604</ee>
<crossref>conf/icassp/2014</crossref>
<url>db/conf/icassp/icassp2014.html#TakamichiTNSN14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/KuboSNTN14" mdate="2021-04-09">
<author pid="124/9006">Keigo Kubo</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Narrow Adaptive Regularization of weights for grapheme-to-phoneme conversion.</title>
<pages>2589-2593</pages>
<year>2014</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2014.6854068</ee>
<crossref>conf/icassp/2014</crossref>
<url>db/conf/icassp/icassp2014.html#KuboSNTN14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/TanakaTNSN14" mdate="2021-04-09">
<author pid="140/2727">Kou Tanaka</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>An evaluation of excitation feature prediction in a hybrid approach to electrolaryngeal speech enhancement.</title>
<pages>4488-4492</pages>
<year>2014</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2014.6854451</ee>
<crossref>conf/icassp/2014</crossref>
<url>db/conf/icassp/icassp2014.html#TanakaTNSN14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/KobayashiTNGNSN14" mdate="2024-10-06">
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="49/5220">Tomoyasu Nakano</author>
<author orcid="0000-0003-1167-0977" pid="85/1753">Masataka Goto</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Regression approaches to perceptual age control in singing voice conversion.</title>
<pages>7904-7908</pages>
<year>2014</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2014.6855139</ee>
<crossref>conf/icassp/2014</crossref>
<url>db/conf/icassp/icassp2014.html#KobayashiTNGNSN14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/TanakaTNSN14" mdate="2023-06-23">
<author pid="140/2727">Kou Tanaka</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Direct F<sub>0</sub> control of an electrolarynx based on statistical excitation feature prediction and its evaluation through simulation.</title>
<pages>31-35</pages>
<year>2014</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2014-7</ee>
<crossref>conf/interspeech/2014</crossref>
<url>db/conf/interspeech/interspeech2014.html#TanakaTNSN14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/JinboTTNSN14" mdate="2023-06-23">
<author pid="158/4125">Nozomi Jinbo</author>
<author pid="124/9221">Shinnosuke Takamichi</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>A hearing impairment simulation method using audiogram-based approximation of auditory charatecteristics.</title>
<pages>490-494</pages>
<year>2014</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2014-122</ee>
<crossref>conf/interspeech/2014</crossref>
<url>db/conf/interspeech/interspeech2014.html#JinboTTNSN14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/KuboSNTN14" mdate="2023-06-23">
<author pid="124/9006">Keigo Kubo</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Structured soft margin confidence weighted learning for grapheme-to-phoneme conversion.</title>
<pages>1263-1267</pages>
<year>2014</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2014-316</ee>
<crossref>conf/interspeech/2014</crossref>
<url>db/conf/interspeech/interspeech2014.html#KuboSNTN14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/MatsumiyaSNTN14" mdate="2023-06-23">
<author pid="146/3942">Sho Matsumiya</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Data-driven generation of text balloons based on linguistic and acoustic features of a comics-anime corpus.</title>
<pages>1801-1805</pages>
<year>2014</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2014-410</ee>
<crossref>conf/interspeech/2014</crossref>
<url>db/conf/interspeech/interspeech2014.html#MatsumiyaSNTN14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/KawaharaMTBNI14" mdate="2023-06-23">
<author pid="84/3249">Hideki Kawahara</author>
<author pid="75/6649">Masanori Morise</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="95/6268">Hideki Banno</author>
<author pid="75/2915">Ryuichi Nisimura</author>
<author pid="72/3903">Toshio Irino</author>
<title>Excitation source analysis for high-quality speech manipulation systems based on an interference-free representation of group delay with minimum phase response compensation.</title>
<pages>2243-2247</pages>
<year>2014</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2014-247</ee>
<crossref>conf/interspeech/2014</crossref>
<url>db/conf/interspeech/interspeech2014.html#KawaharaMTBNI14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/TobingTNSNP14" mdate="2023-06-23">
<author pid="158/4210">Patrick Lumban Tobing</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<author pid="15/2527">Ayu Purwarianti</author>
<title>Articulatory controllable speech modification based on statistical feature mapping with Gaussian mixture models.</title>
<pages>2298-2302</pages>
<year>2014</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2014-185</ee>
<crossref>conf/interspeech/2014</crossref>
<url>db/conf/interspeech/interspeech2014.html#TobingTNSNP14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/KobayashiTNSN14" mdate="2023-06-23">
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Statistical singing voice conversion with direct waveform modification based on the spectrum differential.</title>
<pages>2514-2518</pages>
<year>2014</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2014-539</ee>
<crossref>conf/interspeech/2014</crossref>
<url>db/conf/interspeech/interspeech2014.html#KobayashiTNSN14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/iwsds/LubisSNTPN14" mdate="2021-04-28">
<author pid="160/8114">Nurul Lubis</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="15/2527">Ayu Purwarianti</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Emotion and Its Triggers in Human Spoken Dialogue: Recognition and Analysis.</title>
<pages>103-110</pages>
<year>2014</year>
<booktitle>IWSDS</booktitle>
<ee>https://doi.org/10.1007/978-3-319-21834-2_10</ee>
<crossref>conf/iwsds/2014</crossref>
<url>db/conf/iwsds/iwsds2014.html#LubisSNTPN14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/iwsds/HiraokaNSTN14" mdate="2021-04-28">
<author pid="31/977">Takuya Hiraoka</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Construction and Analysis of a Persuasive Dialogue Corpus.</title>
<pages>125-138</pages>
<year>2014</year>
<booktitle>IWSDS</booktitle>
<ee>https://doi.org/10.1007/978-3-319-21834-2_12</ee>
<crossref>conf/iwsds/2014</crossref>
<url>db/conf/iwsds/iwsds2014.html#HiraokaNSTN14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/lrec/ShimizuNSTN14" mdate="2019-08-19">
<author pid="44/5142">Hiroaki Shimizu</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Collection of a Simultaneous Translation Corpus for Comparative Analysis.</title>
<pages>670-673</pages>
<year>2014</year>
<booktitle>LREC</booktitle>
<ee type="oa">http://www.lrec-conf.org/proceedings/lrec2014/summaries/162.html</ee>
<crossref>conf/lrec/2014</crossref>
<url>db/conf/lrec/lrec2014.html#ShimizuNSTN14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/lrec/SaktiKMNTNAI14" mdate="2019-08-19">
<author pid="71/3717">Sakriani Sakti</author>
<author pid="124/9006">Keigo Kubo</author>
<author pid="146/3942">Sho Matsumiya</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<author pid="84/4862">Fumihiro Adachi</author>
<author pid="36/6317">Ryosuke Isotani</author>
<title>Towards Multilingual Conversations in the Medical Domain: Development of Multilingual Medical Data and A Network-based ASR System.</title>
<pages>2639-2643</pages>
<year>2014</year>
<booktitle>LREC</booktitle>
<ee type="oa">http://www.lrec-conf.org/proceedings/lrec2014/summaries/709.html</ee>
<crossref>conf/lrec/2014</crossref>
<url>db/conf/lrec/lrec2014.html#SaktiKMNTNAI14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/ococosda/DoNSTN14" mdate="2021-04-09">
<author pid="166/6489">Quoc Truong Do</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Collection and analysis of a Japanese-English emphasized speech corpora.</title>
<pages>1-5</pages>
<year>2014</year>
<booktitle>O-COCOSDA</booktitle>
<ee>https://doi.org/10.1109/ICSDA.2014.7051424</ee>
<crossref>conf/ococosda/2014</crossref>
<url>db/conf/ococosda/ococosda2014.html#DoNSTN14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/ococosda/KotoSNTAN14" mdate="2021-04-09">
<author pid="160/0019">Fajri Koto</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="15/6057">Mirna Adriani</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Memorable spoken quote corpora of TED public speaking.</title>
<pages>1-4</pages>
<year>2014</year>
<booktitle>O-COCOSDA</booktitle>
<ee>https://doi.org/10.1109/ICSDA.2014.7051435</ee>
<crossref>conf/ococosda/2014</crossref>
<url>db/conf/ococosda/ococosda2014.html#KotoSNTAN14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/ococosda/MizukamiNSTN14" mdate="2021-04-09">
<author pid="169/3173">Masahiro Mizukami</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Building a free, general-domain paraphrase database for Japanese.</title>
<pages>1-4</pages>
<year>2014</year>
<booktitle>O-COCOSDA</booktitle>
<ee>https://doi.org/10.1109/ICSDA.2014.7051433</ee>
<crossref>conf/ococosda/2014</crossref>
<url>db/conf/ococosda/ococosda2014.html#MizukamiNSTN14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/ococosda/NioSNTN14" mdate="2021-04-09">
<author pid="146/8261">Lasguido Nio</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Conversation dialog corpora from television and movie scripts.</title>
<pages>1-4</pages>
<year>2014</year>
<booktitle>O-COCOSDA</booktitle>
<ee>https://doi.org/10.1109/ICSDA.2014.7051436</ee>
<crossref>conf/ococosda/2014</crossref>
<url>db/conf/ococosda/ococosda2014.html#NioSNTN14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/slt/NioSNTN14" mdate="2021-04-09">
<author pid="146/8261">Lasguido Nio</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Improving the robustness of example-based dialog retrieval using recursive neural network paraphrase identification.</title>
<pages>306-311</pages>
<year>2014</year>
<booktitle>SLT</booktitle>
<ee>https://doi.org/10.1109/SLT.2014.7078592</ee>
<crossref>conf/slt/2014</crossref>
<url>db/conf/slt/slt2014.html#NioSNTN14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/ssst/HatakoshiNSTN14" mdate="2021-08-06">
<author pid="04/11293">Yuto Hatakoshi</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Rule-based Syntactic Preprocessing for Syntax-based Machine Translation.</title>
<pages>34-42</pages>
<year>2014</year>
<booktitle>SSST@EMNLP</booktitle>
<ee type="oa">https://aclanthology.org/W14-4004/</ee>
<ee>https://doi.org/10.3115/v1/W14-4004</ee>
<crossref>conf/ssst/2014</crossref>
<url>db/conf/ssst/ssst2014.html#HatakoshiNSTN14</url>
</inproceedings>
</r>
<r><article key="journals/pieee/TokudaNTZYO13" mdate="2024-10-06">
<author pid="25/369">Keiichi Tokuda</author>
<author pid="90/302">Yoshihiko Nankaku</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author orcid="0000-0002-8959-5471" pid="42/7014">Heiga Zen</author>
<author pid="87/3979">Junichi Yamagishi</author>
<author pid="22/7077">Keiichiro Oura</author>
<title>Speech Synthesis Based on Hidden Markov Models.</title>
<pages>1234-1252</pages>
<year>2013</year>
<volume>101</volume>
<journal>Proc. IEEE</journal>
<number>5</number>
<ee>https://doi.org/10.1109/JPROC.2013.2251852</ee>
<url>db/journals/pieee/pieee101.html#TokudaNTZYO13</url>
</article>
</r>
<r><inproceedings key="conf/acl/NeubigSTNMII13" mdate="2021-08-06">
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<author pid="11/4619">Yuji Matsumoto 0001</author>
<author pid="36/6317">Ryosuke Isotani</author>
<author pid="213/0653">Yukichi Ikeda</author>
<title>Towards High-Reliability Speech Translation in the Medical Domain.</title>
<pages>22-29</pages>
<year>2013</year>
<booktitle>NLPHealthcare@IJCNLP</booktitle>
<ee type="oa">https://aclanthology.org/W13-4604/</ee>
<crossref>conf/ijcnlp/2013healthcare</crossref>
<url>db/conf/acl/healthcare2013.html#NeubigSTNMII13</url>
</inproceedings>
</r>
<r><inproceedings key="conf/asru/HiraokaYNSTN13" mdate="2021-04-09">
<author pid="31/977">Takuya Hiraoka</author>
<author pid="139/5461">Yuki Yamauchi</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Dialogue management for leading the conversation in persuasive dialogue systems.</title>
<pages>114-119</pages>
<year>2013</year>
<booktitle>ASRU</booktitle>
<ee>https://doi.org/10.1109/ASRU.2013.6707715</ee>
<crossref>conf/asru/2013</crossref>
<url>db/conf/asru/asru2013.html#HiraokaYNSTN13</url>
</inproceedings>
</r>
<r><inproceedings key="conf/clef/ArthurNSTN13" mdate="2023-03-10">
<author pid="150/5396">Philip Arthur</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Inter-Sentence Features and Thresholded Minimum Error Rate Training: NAIST at CLEF 2013 QA4MRE.</title>
<year>2013</year>
<booktitle>CLEF (Working Notes)</booktitle>
<ee type="oa">https://ceur-ws.org/Vol-1179/CLEF2013wn-QA4MRE-ArthurEt2013.pdf</ee>
<crossref>conf/clef/2013w</crossref>
<url>db/conf/clef/clef2013w.html#ArthurNSTN13</url>
</inproceedings>
</r>
<r><inproceedings key="conf/clef/ArthurNSTN13a" mdate="2023-03-10">
<author pid="150/5396">Philip Arthur</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>NAIST at the CLEF 2013 QA4MRE Pilot Task.</title>
<year>2013</year>
<booktitle>CLEF (Working Notes)</booktitle>
<ee type="oa">https://ceur-ws.org/Vol-1179/CLEF2013wn-QA4MRE-ArthurEt2013b.pdf</ee>
<crossref>conf/clef/2013w</crossref>
<url>db/conf/clef/clef2013w.html#ArthurNSTN13a</url>
</inproceedings>
</r>
<r><inproceedings key="conf/coginfocom/TanakaSNT013" mdate="2022-12-06">
<author orcid="0000-0002-0548-6252" pid="66/2476">Hiroki Tanaka</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Modality and contextual differences in computer based non-verbal communication training.</title>
<pages>127-132</pages>
<year>2013</year>
<booktitle>CogInfoCom</booktitle>
<ee>https://doi.org/10.1109/CogInfoCom.2013.6719227</ee>
<crossref>conf/coginfocom/2013</crossref>
<url>db/conf/coginfocom/coginfocom2013.html#TanakaSNT013</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/KawaharaMTNI13" mdate="2023-06-23">
<author pid="84/3249">Hideki Kawahara</author>
<author pid="75/6649">Masanori Morise</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="75/2915">Ryuichi Nisimura</author>
<author pid="72/3903">Toshio Irino</author>
<title>Beyond bandlimited sampling of speech spectral envelope imposed by the harmonic structure of voiced sounds.</title>
<pages>34-38</pages>
<year>2013</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2013-8</ee>
<crossref>conf/interspeech/2013</crossref>
<url>db/conf/interspeech/interspeech2013.html#KawaharaMTNI13</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/TakamichiTSSNN13" mdate="2023-06-23">
<author pid="124/9221">Shinnosuke Takamichi</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="02/2978">Yoshinori Shiga</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Improvements to HMM-based speech synthesis based on parameter generation with rich context models.</title>
<pages>364-368</pages>
<year>2013</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2013-101</ee>
<crossref>conf/interspeech/2013</crossref>
<url>db/conf/interspeech/interspeech2013.html#TakamichiTSSNN13</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/KobayashiDTNGNSN13" mdate="2024-10-06">
<author pid="70/8239">Kazuhiro Kobayashi</author>
<author pid="68/4056">Hironori Doi</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="49/5220">Tomoyasu Nakano</author>
<author orcid="0000-0003-1167-0977" pid="85/1753">Masataka Goto</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>An investigation of acoustic features for singing voice conversion based on perceptual age.</title>
<pages>1057-1061</pages>
<year>2013</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2013-118</ee>
<crossref>conf/interspeech/2013</crossref>
<url>db/conf/interspeech/interspeech2013.html#KobayashiDTNGNSN13</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/DoiTNGN13" mdate="2024-10-06">
<author pid="68/4056">Hironori Doi</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="49/5220">Tomoyasu Nakano</author>
<author orcid="0000-0003-1167-0977" pid="85/1753">Masataka Goto</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Evaluation of a singing voice conversion method based on many-to-many eigenvoice conversion.</title>
<pages>1067-1071</pages>
<year>2013</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2013-120</ee>
<crossref>conf/interspeech/2013</crossref>
<url>db/conf/interspeech/interspeech2013.html#DoiTNGN13</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/KuboSNTN13" mdate="2023-06-23">
<author pid="124/9006">Keigo Kubo</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Grapheme-to-phoneme conversion based on adaptive regularization of weight vectors.</title>
<pages>1946-1950</pages>
<year>2013</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2013-464</ee>
<crossref>conf/interspeech/2013</crossref>
<url>db/conf/interspeech/interspeech2013.html#KuboSNTN13</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/KanoTSNTN13" mdate="2023-06-23">
<author pid="140/2697">Takatomo Kano</author>
<author pid="124/9221">Shinnosuke Takamichi</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Generalizing continuous-space translation of paralinguistic information.</title>
<pages>2614-2618</pages>
<year>2013</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2013-602</ee>
<crossref>conf/interspeech/2013</crossref>
<url>db/conf/interspeech/interspeech2013.html#KanoTSNTN13</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/OhgushiNSTN13" mdate="2023-06-23">
<author pid="140/2857">Masaya Ohgushi</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>An empirical comparison of joint optimization techniques for speech translation.</title>
<pages>2619-2623</pages>
<year>2013</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2013-603</ee>
<crossref>conf/interspeech/2013</crossref>
<url>db/conf/interspeech/interspeech2013.html#OhgushiNSTN13</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/TanakaTNSN13" mdate="2023-06-23">
<author pid="140/2727">Kou Tanaka</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>A hybrid approach to electrolaryngeal speech enhancement based on spectral subtraction and statistical voice conversion.</title>
<pages>3067-3071</pages>
<year>2013</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2013-669</ee>
<crossref>conf/interspeech/2013</crossref>
<url>db/conf/interspeech/interspeech2013.html#TanakaTNSN13</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/MoriguchiTSSNSN13" mdate="2026-02-13">
<author pid="140/2835">Takuto Moriguchi</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="140/2888">Motoaki Sano</author>
<author pid="55/6900-2">Hiroshi Sato 0002</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>A digital signal processor implementation of silent/electrolaryngeal speech enhancement based on real-time statistical voice conversion.</title>
<pages>3072-3076</pages>
<year>2013</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2013-670</ee>
<crossref>conf/interspeech/2013</crossref>
<url>db/conf/interspeech/interspeech2013.html#MoriguchiTSSNSN13</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/FujitaNSTN13" mdate="2023-06-23">
<author pid="12/9858">Tomoki Fujita</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Simple, lexicalized choice of translation timing for simultaneous speech translation.</title>
<pages>3487-3491</pages>
<year>2013</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2013-615</ee>
<crossref>conf/interspeech/2013</crossref>
<url>db/conf/interspeech/interspeech2013.html#FujitaNSTN13</url>
</inproceedings>
</r>
<r><inproceedings key="conf/iwslt/SaktiKNTN13" mdate="2024-08-01">
<author pid="71/3717">Sakriani Sakti</author>
<author pid="124/9006">Keigo Kubo</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>The NAIST English speech recognition system for IWSLT 2013.</title>
<year>2013</year>
<booktitle>IWSLT (Evaluation Campaign)</booktitle>
<ee type="oa">https://aclanthology.org/2013.iwslt-evaluation.23</ee>
<crossref>conf/iwslt/2013eval</crossref>
<url>db/conf/iwslt/iwslt2013eval.html#SaktiKNTN13</url>
</inproceedings>
</r>
<r><inproceedings key="conf/iwslt/ShimizuNST013" mdate="2024-08-01">
<author pid="44/5142">Hiroaki Shimizu</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Constructing a speech translation system using simultaneous interpretation data.</title>
<year>2013</year>
<booktitle>IWSLT</booktitle>
<ee type="oa">https://aclanthology.org/2013.iwslt-papers.3</ee>
<crossref>conf/iwslt/2013</crossref>
<url>db/conf/iwslt/iwslt2013.html#ShimizuNST013</url>
</inproceedings>
</r>
<r><inproceedings key="conf/ssw/InukaiTNSN13" mdate="2024-07-31">
<author pid="178/0956">Tatsuo Inukai</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Investigation of intra-speaker spectral parameter variation and its prediction towards improvement of spectral conversion metric.</title>
<pages>89-94</pages>
<year>2013</year>
<booktitle>SSW</booktitle>
<ee type="oa">https://www.isca-archive.org/ssw_2013/inukai13_ssw.html</ee>
<crossref>conf/ssw/2013</crossref>
<url>db/conf/ssw/ssw2013.html#InukaiTNSN13</url>
</inproceedings>
</r>
<r><article key="journals/jirs/NakamuraSNITOO12" mdate="2026-05-07">
<author pid="63/4486">Tomoaki Nakamura</author>
<author orcid="0000-0002-0261-0510" pid="77/2654">Komei Sugiura</author>
<author pid="09/1646">Takayuki Nagai</author>
<author pid="22/2156">Naoto Iwahashi</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="24/2307">Hiroyuki Okada</author>
<author pid="21/4085">Takashi Omori</author>
<title>Learning Novel Objects for Extended Mobile Manipulation.</title>
<pages>187-204</pages>
<year>2012</year>
<volume>66</volume>
<journal>J. Intell. Robotic Syst.</journal>
<number>1-2</number>
<ee>https://doi.org/10.1007/s10846-011-9605-1</ee>
<url>db/journals/jirs/jirs66.html#NakamuraSNITOO12</url>
</article>
</r>
<r><article key="journals/speech/NakamuraTSS12" mdate="2021-04-09">
<author pid="31/8013">Keigo Nakamura</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Speaking-aid systems using GMM-based voice conversion for electrolaryngeal speech.</title>
<pages>134-146</pages>
<year>2012</year>
<volume>54</volume>
<journal>Speech Commun.</journal>
<number>1</number>
<ee>https://doi.org/10.1016/j.specom.2011.07.007</ee>
<url>db/journals/speech/speech54.html#NakamuraTSS12</url>
</article>
</r>
<r><article key="journals/taslp/TodaNS12" mdate="2021-04-09">
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="54/9262">Mikihiro Nakagiri</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Statistical Voice Conversion Techniques for Body-Conducted Unvoiced Speech Enhancement.</title>
<pages>2505-2517</pages>
<year>2012</year>
<volume>20</volume>
<journal>IEEE Trans. Speech Audio Process.</journal>
<number>9</number>
<ee>https://doi.org/10.1109/TASL.2012.2205241</ee>
<url>db/journals/taslp/taslp20.html#TodaNS12</url>
</article>
</r>
<r><inproceedings key="conf/apsipa/DoiTNGN12" mdate="2021-08-08">
<author pid="68/4056">Hironori Doi</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="49/5220">Tomoyasu Nakano</author>
<author pid="85/1753">Masataka Goto</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Singing voice conversion method based on many-to-many eigenvoice conversion and training data generation using a singing-to-singing synthesis system.</title>
<pages>1-6</pages>
<year>2012</year>
<booktitle>APSIPA</booktitle>
<ee>https://ieeexplore.ieee.org/document/6411800/</ee>
<crossref>conf/apsipa/2012</crossref>
<url>db/conf/apsipa/apsipa2012.html#DoiTNGN12</url>
</inproceedings>
</r>
<r><inproceedings key="conf/coginfocom/TanakaSNT0N12" mdate="2023-01-10">
<author orcid="0000-0002-0548-6252" pid="66/2476">Hiroki Tanaka</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="42/87">Nick Campbell 0001</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Non-verbal cognitive skills and autistic conditions: An analysis and training tool.</title>
<pages>41-46</pages>
<year>2012</year>
<booktitle>CogInfoCom</booktitle>
<ee>https://doi.org/10.1109/CogInfoCom.2012.6422034</ee>
<crossref>conf/coginfocom/2012</crossref>
<url>db/conf/coginfocom/coginfocom2012.html#TanakaSNT0N12</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/YamamotoTDSS12" mdate="2021-04-09">
<author pid="120/6772">Kenzo Yamamoto</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="68/4056">Hironori Doi</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Statistical approach to voice quality control in esophageal speech enhancement.</title>
<pages>4497-4500</pages>
<year>2012</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2012.6287949</ee>
<crossref>conf/icassp/2012</crossref>
<url>db/conf/icassp/icassp2012.html#YamamotoTDSS12</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/TodaMB12" mdate="2023-06-23">
<author pid="85/741">Tomoki Toda</author>
<author pid="40/9234">Takashi Muramatsu</author>
<author pid="95/6268">Hideki Banno</author>
<title>Implementation of Computationally Efficient Real-Time Voice Conversion.</title>
<year>2012</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2012-34</ee>
<crossref>conf/interspeech/2012</crossref>
<url>db/conf/interspeech/interspeech2012.html#TodaMB12</url>
<pages>94-97</pages>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/TakamichiTSKSN12" mdate="2023-06-23">
<author pid="124/9221">Shinnosuke Takamichi</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="02/2978">Yoshinori Shiga</author>
<author pid="32/4341">Hisashi Kawai</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>An Evaluation of Parameter Generation Methods with Rich Context Models in HMM-Based Speech Synthesis.</title>
<year>2012</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2012-358</ee>
<crossref>conf/interspeech/2012</crossref>
<url>db/conf/interspeech/interspeech2012.html#TakamichiTSKSN12</url>
<pages>1139-1142</pages>
</inproceedings>
</r>
<r><inproceedings key="conf/isspit/ItoiMTSS12" mdate="2023-03-24">
<author pid="140/1975">Miyuki Itoi</author>
<author pid="53/10650">Ryoichi Miyazaki</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Blind speech extraction for Non-Audible Murmur speech with speaker's movement noise.</title>
<pages>320-325</pages>
<year>2012</year>
<booktitle>ISSPIT</booktitle>
<ee>https://doi.org/10.1109/ISSPIT.2012.6621308</ee>
<ee>https://doi.ieeecomputersociety.org/10.1109/ISSPIT.2012.6621308</ee>
<crossref>conf/isspit/2012</crossref>
<url>db/conf/isspit/isspit2012.html#ItoiMTSS12</url>
</inproceedings>
</r>
<r><inproceedings key="conf/iwsds/NioSNTAN12" mdate="2021-04-29">
<author pid="146/8261">Lasguido Nio</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="15/6057">Mirna Adriani</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Developing Non-goal Dialog System Based on Examples of Drama Television.</title>
<pages>355-361</pages>
<year>2012</year>
<booktitle>IWSDS</booktitle>
<ee>https://doi.org/10.1007/978-1-4614-8280-2_32</ee>
<crossref>conf/iwsds/2012</crossref>
<url>db/conf/iwsds/iwsds2012.html#NioSNTAN12</url>
</inproceedings>
</r>
<r><inproceedings key="conf/iwslt/NeubigDOKKSTN12" mdate="2024-08-01">
<author pid="03/8155">Graham Neubig</author>
<author pid="58/3217">Kevin Duh</author>
<author pid="140/2857">Masaya Ogushi</author>
<author pid="140/2697">Takatomo Kano</author>
<author pid="97/9767">Tetsuo Kiso</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>The NAIST machine translation system for IWSLT2012.</title>
<pages>54-60</pages>
<year>2012</year>
<booktitle>IWSLT</booktitle>
<ee type="oa">https://www.isca-archive.org/iwslt_2012/neubig12_iwslt.html</ee>
<crossref>conf/iwslt/2012</crossref>
<url>db/conf/iwslt/iwslt2012.html#NeubigDOKKSTN12</url>
</inproceedings>
</r>
<r><inproceedings key="conf/iwslt/SaamMKHSKSSNTNW12" mdate="2024-08-01">
<author pid="138/0300">Christian Saam</author>
<author pid="60/7921">Christian Mohr</author>
<author pid="70/11337">Kevin Kilgour</author>
<author pid="120/6962">Michael Heck</author>
<author pid="79/9103">Matthias Sperber</author>
<author pid="124/9006">Keigo Kubo</author>
<author pid="15/1807">Sebastian St&#252;ker</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<author pid="08/2456">Alex Waibel</author>
<title>The 2012 KIT and KIT-NAIST English ASR systems for the IWSLT evaluation.</title>
<pages>87-90</pages>
<year>2012</year>
<booktitle>IWSLT</booktitle>
<ee type="oa">https://www.isca-archive.org/iwslt_2012/saam12_iwslt.html</ee>
<crossref>conf/iwslt/2012</crossref>
<url>db/conf/iwslt/iwslt2012.html#SaamMKHSKSSNTNW12</url>
</inproceedings>
</r>
<r><inproceedings key="conf/iwslt/HeckKSSSSKMNTNW12" mdate="2024-08-01">
<author pid="120/6962">Michael Heck</author>
<author pid="124/9006">Keigo Kubo</author>
<author pid="79/9103">Matthias Sperber</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="15/1807">Sebastian St&#252;ker</author>
<author pid="138/0300">Christian Saam</author>
<author pid="70/11337">Kevin Kilgour</author>
<author pid="60/7921">Christian Mohr</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<author pid="08/2456">Alex Waibel</author>
<title>The KIT-NAIST (contrastive) English ASR system for IWSLT 2012.</title>
<pages>91-95</pages>
<year>2012</year>
<booktitle>IWSLT</booktitle>
<ee type="oa">https://www.isca-archive.org/iwslt_2012/heck12_iwslt.html</ee>
<crossref>conf/iwslt/2012</crossref>
<url>db/conf/iwslt/iwslt2012.html#HeckKSSSSKMNTNW12</url>
</inproceedings>
</r>
<r><inproceedings key="conf/iwslt/KanoSTNTN12" mdate="2024-08-01">
<author pid="140/2697">Takatomo Kano</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="124/9221">Shinnosuke Takamichi</author>
<author pid="03/8155">Graham Neubig</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>A method for translation of paralinguistic information.</title>
<pages>158-163</pages>
<year>2012</year>
<booktitle>IWSLT</booktitle>
<ee type="oa">https://www.isca-archive.org/iwslt_2012/kano12_iwslt.html</ee>
<crossref>conf/iwslt/2012</crossref>
<url>db/conf/iwslt/iwslt2012.html#KanoSTNTN12</url>
</inproceedings>
</r>
<r><inproceedings key="conf/asru/IshiiTSSN11" mdate="2021-04-09">
<author pid="98/11427">Shunta Ishii</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="71/3717">Sakriani Sakti</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Blind noise suppression for Non-Audible Murmur recognition with stereo signal processing.</title>
<pages>494-499</pages>
<year>2011</year>
<booktitle>ASRU</booktitle>
<ee>https://doi.org/10.1109/ASRU.2011.6163981</ee>
<crossref>conf/asru/2011</crossref>
<url>db/conf/asru/asru2011.html#IshiiTSSN11</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/DoiNTSS11" mdate="2021-04-09">
<author pid="68/4056">Hironori Doi</author>
<author pid="31/8013">Keigo Nakamura</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>An evaluation of alaryngeal speech enhancement methods based on voice conversion techniques.</title>
<pages>5136-5139</pages>
<year>2011</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2011.5947513</ee>
<crossref>conf/icassp/2011</crossref>
<url>db/conf/icassp/icassp2011.html#DoiNTSS11</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/BabaniTSS11" mdate="2021-04-09">
<author pid="58/9879">Denis Babani</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Acoustic model training for non-audible murmur recognition using transformed normal speech data.</title>
<pages>5224-5227</pages>
<year>2011</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2011.5947535</ee>
<crossref>conf/icassp/2011</crossref>
<url>db/conf/icassp/icassp2011.html#BabaniTSS11</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/HattoriTKSS11" mdate="2023-06-23">
<author pid="30/10648">Nobuhiko Hattori</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="32/4341">Hisashi Kawai</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Speaker-Adaptive Speech Synthesis Based on Eigenvoice Conversion and Language-Dependent Prosodic Conversion in Speech-to-Speech Translation.</title>
<pages>2769-2772</pages>
<year>2011</year>
<booktitle>INTERSPEECH</booktitle>
<crossref>conf/interspeech/2011</crossref>
<url>db/conf/interspeech/interspeech2011.html#HattoriTKSS11</url>
<ee type="oa">https://doi.org/10.21437/Interspeech.2011-693</ee>
</inproceedings>
</r>
<r><article key="journals/ieicet/OhtaniTSS10" mdate="2021-04-09">
<author pid="34/8763">Yamato Ohtani</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Adaptive Training for Voice Conversion Based on Eigenvoices.</title>
<pages>1589-1598</pages>
<year>2010</year>
<volume>93-D</volume>
<journal>IEICE Trans. Inf. Syst.</journal>
<number>6</number>
<ee>https://doi.org/10.1587/transinf.E93.D.1589</ee>
<ee>http://search.ieice.org/bin/summary.php?id=e93-d_6_1589</ee>
<url>db/journals/ieicet/ieicet93d.html#OhtaniTSS10</url>
</article>
</r>
<r><article key="journals/ieicet/NakamuraTSS10" mdate="2021-04-09">
<author pid="31/8013">Keigo Nakamura</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Evaluation of Extremely Small Sound Source Signals Used in Speaking-Aid System with Statistical Voice Conversion.</title>
<pages>1909-1917</pages>
<year>2010</year>
<volume>93-D</volume>
<journal>IEICE Trans. Inf. Syst.</journal>
<number>7</number>
<ee>https://doi.org/10.1587/transinf.E93.D.1909</ee>
<ee>http://search.ieice.org/bin/summary.php?id=e93-d_7_1909</ee>
<url>db/journals/ieicet/ieicet93d.html#NakamuraTSS10</url>
</article>
</r>
<r><article key="journals/ieicet/DoiNTSS10" mdate="2021-04-09">
<author pid="68/4056">Hironori Doi</author>
<author pid="31/8013">Keigo Nakamura</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Esophageal Speech Enhancement Based on Statistical Voice Conversion with Gaussian Mixture Models.</title>
<pages>2472-2482</pages>
<year>2010</year>
<volume>93-D</volume>
<journal>IEICE Trans. Inf. Syst.</journal>
<number>9</number>
<ee>https://doi.org/10.1587/transinf.E93.D.2472</ee>
<ee>http://search.ieice.org/bin/summary.php?id=e93-d_9_2472</ee>
<url>db/journals/ieicet/ieicet93d.html#DoiNTSS10</url>
</article>
</r>
<r><article key="journals/ieicet/OhtaniTSS10a" mdate="2021-04-09">
<author pid="34/8763">Yamato Ohtani</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Improvements of the One-to-Many Eigenvoice Conversion System.</title>
<pages>2491-2499</pages>
<year>2010</year>
<volume>93-D</volume>
<journal>IEICE Trans. Inf. Syst.</journal>
<number>9</number>
<ee>https://doi.org/10.1587/transinf.E93.D.2491</ee>
<ee>http://search.ieice.org/bin/summary.php?id=e93-d_9_2491</ee>
<url>db/journals/ieicet/ieicet93d.html#OhtaniTSS10a</url>
</article>
</r>
<r><article key="journals/speech/HiraharaOSTNNS10" mdate="2022-10-02">
<author pid="24/4107">Tatsuya Hirahara</author>
<author orcid="0000-0001-8962-9304" pid="97/2873">Makoto Otani</author>
<author pid="39/8014">Shota Shimizu</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="31/8013">Keigo Nakamura</author>
<author pid="75/5531">Yoshitaka Nakajima</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Silent-speech enhancement using body-conducted vocal-tract resonance signals.</title>
<pages>301-313</pages>
<year>2010</year>
<volume>52</volume>
<journal>Speech Commun.</journal>
<number>4</number>
<ee>https://doi.org/10.1016/j.specom.2009.12.001</ee>
<url>db/journals/speech/speech52.html#HiraharaOSTNNS10</url>
</article>
</r>
<r><article key="journals/speech/TranBLT10" mdate="2025-11-15">
<author orcid="0009-0002-9023-6772" pid="69/8013">Viet-Anh Tran</author>
<author orcid="0000-0002-6053-0818" pid="48/1036">G&#233;rard Bailly</author>
<author pid="89/3110">H&#233;l&#232;ne Loevenbruck</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<title>Improvement to a NAM-captured whisper-to-speech system.</title>
<pages>314-326</pages>
<year>2010</year>
<volume>52</volume>
<journal>Speech Commun.</journal>
<number>4</number>
<ee>https://doi.org/10.1016/j.specom.2009.11.005</ee>
<url>db/journals/speech/speech52.html#TranBLT10</url>
</article>
</r>
<r><article key="journals/taslp/StylianouTWKR10" mdate="2024-07-25">
<author pid="14/1209">Yannis Stylianou</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author orcid="0000-0002-3947-2123" pid="20/924">Chung-Hsien Wu 0001</author>
<author pid="19/4966">Alexander Kain</author>
<author pid="11/6417">Olivier Rosec</author>
<title>Introduction to the Special Section on Voice Transformation.</title>
<pages>909-911</pages>
<year>2010</year>
<volume>18</volume>
<journal>IEEE Trans. Speech Audio Process.</journal>
<number>5</number>
<ee>https://doi.org/10.1109/TASL.2010.2051826</ee>
<url>db/journals/taslp/taslp18.html#StylianouTWKR10</url>
</article>
</r>
<r><inproceedings key="conf/blizzard/ShigaTSNKTT010" mdate="2024-09-20">
<author pid="02/2978">Yoshinori Shiga</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/448">Shinsuke Sakai</author>
<author pid="37/2394">Jinfu Ni</author>
<author pid="32/4341">Hisashi Kawai</author>
<author pid="25/369">Keiichi Tokuda</author>
<author pid="20/3844">Minoru Tsuzaki</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>NICT Blizzard Challenge 2010 Entry.</title>
<year>2010</year>
<booktitle>Blizzard Challenge</booktitle>
<ee type="oa">https://doi.org/10.21437/Blizzard.2010-15</ee>
<crossref>conf/blizzard/2010</crossref>
<url>db/conf/blizzard/blizzard2010.html#ShigaTSNKTT010</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/DoiNTSS10" mdate="2021-04-09">
<author pid="68/4056">Hironori Doi</author>
<author pid="31/8013">Keigo Nakamura</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Statistical approach to enhancing esophageal speech based on Gaussian mixture models.</title>
<pages>4250-4253</pages>
<year>2010</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2010.5495676</ee>
<crossref>conf/icassp/2010</crossref>
<url>db/conf/icassp/icassp2010.html#DoiNTSS10</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/OhtaniTSS10" mdate="2021-04-09">
<author pid="34/8763">Yamato Ohtani</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Non-parallel training for many-to-many eigenvoice conversion.</title>
<pages>4822-4825</pages>
<year>2010</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2010.5495139</ee>
<crossref>conf/icassp/2010</crossref>
<url>db/conf/icassp/icassp2010.html#OhtaniTSS10</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/ShigaTSK10" mdate="2023-06-23">
<author pid="02/2978">Yoshinori Shiga</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/448">Shinsuke Sakai</author>
<author pid="32/4341">Hisashi Kawai</author>
<title>Improved training of excitation for HMM-based parametric speech synthesis.</title>
<pages>809-812</pages>
<year>2010</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2010-179</ee>
<crossref>conf/interspeech/2010</crossref>
<url>db/conf/interspeech/interspeech2010.html#ShigaTSK10</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/NakamuraTSS10" mdate="2023-06-23">
<author pid="31/8013">Keigo Nakamura</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>The use of air-pressure sensor in electrolaryngeal speech enhancement based on statistical voice conversion.</title>
<pages>1628-1631</pages>
<year>2010</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2010-471</ee>
<crossref>conf/interspeech/2010</crossref>
<url>db/conf/interspeech/interspeech2010.html#NakamuraTSS10</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/OhtaTOSS10" mdate="2023-06-23">
<author pid="06/9233">Kumi Ohta</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="34/8763">Yamato Ohtani</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Adaptive voice-quality control based on one-to-many eigenvoice conversion.</title>
<pages>2158-2161</pages>
<year>2010</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2010-595</ee>
<crossref>conf/interspeech/2010</crossref>
<url>db/conf/interspeech/interspeech2010.html#OhtaTOSS10</url>
</inproceedings>
</r>
<r><inproceedings key="conf/ssw/HayashidaTOSS10" mdate="2024-07-31">
<author pid="142/7270">Chie Hayashida</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="34/8763">Yamato Ohtani</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Linear transformation approaches to many-to-one voice conversion.</title>
<pages>74-79</pages>
<year>2010</year>
<booktitle>SSW</booktitle>
<ee type="oa">https://www.isca-archive.org/ssw_2010/hayashida10_ssw.html</ee>
<crossref>conf/ssw/2010</crossref>
<url>db/conf/ssw/ssw2010.html#HayashidaTOSS10</url>
</inproceedings>
</r>
<r><article key="journals/speech/GomezTSS09" mdate="2021-04-09">
<author pid="44/2122">Randy Gomez</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Techniques in rapid unsupervised speaker adaptation based on HMM-Sufficient Statistics.</title>
<pages>42-57</pages>
<year>2009</year>
<volume>51</volume>
<journal>Speech Commun.</journal>
<number>1</number>
<ee>https://doi.org/10.1016/j.specom.2008.05.014</ee>
<url>db/journals/speech/speech51.html#GomezTSS09</url>
</article>
</r>
<r><article key="journals/taslp/YamagishiNZLTTKR09" mdate="2025-03-03">
<author pid="87/3979">Junichi Yamagishi</author>
<author orcid="0000-0002-2278-0429" pid="83/3783">Takashi Nose</author>
<author orcid="0000-0002-8959-5471" pid="42/7014">Heiga Zen</author>
<author pid="70/5210">Zhen-Hua Ling</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="25/369">Keiichi Tokuda</author>
<author pid="68/2005">Simon King 0001</author>
<author pid="33/3792">Steve Renals</author>
<title>Robust Speaker-Adaptive HMM-Based Text-to-Speech Synthesis.</title>
<pages>1208-1230</pages>
<year>2009</year>
<volume>17</volume>
<journal>IEEE Trans. Speech Audio Process.</journal>
<number>6</number>
<ee>https://doi.org/10.1109/TASL.2009.2016394</ee>
<url>db/journals/taslp/taslp17.html#YamagishiNZLTTKR09</url>
</article>
</r>
<r><inproceedings key="conf/blizzard/MaiaTSSNKTT009" mdate="2024-09-20">
<author pid="15/7825">Ranniery Maia</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="57/448">Shinsuke Sakai</author>
<author pid="02/2978">Yoshinori Shiga</author>
<author pid="37/2394">Jinfu Ni</author>
<author pid="32/4341">Hisashi Kawai</author>
<author pid="25/369">Keiichi Tokuda</author>
<author pid="20/3844">Minoru Tsuzaki</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>The NICT Entry for the Blizzard Challenge 2009: an Enhanced HMM-based Speech Synthesis System with Trajectory Training considering Global Variance and State-Dependent Mixed Excitation.</title>
<year>2009</year>
<booktitle>Blizzard Challenge</booktitle>
<ee type="oa">https://doi.org/10.21437/Blizzard.2009-11</ee>
<crossref>conf/blizzard/2009</crossref>
<url>db/conf/blizzard/blizzard2009.html#MaiaTSSNKTT009</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/TodaNSS09" mdate="2023-03-23">
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="31/8013">Keigo Nakamura</author>
<author pid="45/8053">Hidehiko Sekimoto</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Voice conversion for various types of body transmitted speech.</title>
<pages>3601-3604</pages>
<year>2009</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2009.4960405</ee>
<ee>https://doi.ieeecomputersociety.org/10.1109/ICASSP.2009.4960405</ee>
<crossref>conf/icassp/2009</crossref>
<url>db/conf/icassp/icassp2009.html#TodaNSS09</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/YuTGKMTY09" mdate="2023-03-23">
<author pid="197/1322-4">Kai Yu 0004</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="27/7520">Milica Gasic</author>
<author pid="80/6099">Simon Keizer</author>
<author pid="m/FrancoisMairesse">Fran&#231;ois Mairesse</author>
<author pid="44/2850">Blaise Thomson</author>
<author pid="11/9311">Steve J. Young</author>
<title>Probablistic modelling of F0 in unvoiced regions in HMM based speech synthesis.</title>
<pages>3773-3776</pages>
<year>2009</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2009.4960448</ee>
<ee>https://doi.ieeecomputersociety.org/10.1109/ICASSP.2009.4960448</ee>
<crossref>conf/icassp/2009</crossref>
<url>db/conf/icassp/icassp2009.html#YuTGKMTY09</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/MiyamotoNTSS09" mdate="2023-03-23">
<author pid="24/2856">Daisuke Miyamoto</author>
<author pid="31/8013">Keigo Nakamura</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Acoustic compensation methods for body transmitted speech conversion.</title>
<pages>3901-3904</pages>
<year>2009</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2009.4960480</ee>
<ee>https://doi.ieeecomputersociety.org/10.1109/ICASSP.2009.4960480</ee>
<crossref>conf/icassp/2009</crossref>
<url>db/conf/icassp/icassp2009.html#MiyamotoNTSS09</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/TodaY09" mdate="2023-03-23">
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="11/9311">Steve J. Young</author>
<title>Trajectory training considering global variance for HMM-based speech synthesis.</title>
<pages>4025-4028</pages>
<year>2009</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2009.4960511</ee>
<ee>https://doi.ieeecomputersociety.org/10.1109/ICASSP.2009.4960511</ee>
<crossref>conf/icassp/2009</crossref>
<url>db/conf/icassp/icassp2009.html#TodaY09</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/TodaNNKNS09" mdate="2023-06-23">
<author pid="85/741">Tomoki Toda</author>
<author pid="31/8013">Keigo Nakamura</author>
<author pid="09/1646">Takayuki Nagai</author>
<author pid="25/9192">Tomomi Kaino</author>
<author pid="75/5531">Yoshitaka Nakajima</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Technologies for processing body-conducted speech detected with non-audible murmur microphone.</title>
<pages>632-635</pages>
<year>2009</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2009-224</ee>
<crossref>conf/interspeech/2009</crossref>
<url>db/conf/interspeech/interspeech2009.html#TodaNNKNS09</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/TranBLT09" mdate="2025-11-15">
<author orcid="0009-0002-9023-6772" pid="69/8013">Viet-Anh Tran</author>
<author pid="48/1036">G&#233;rard Bailly</author>
<author pid="89/3110">H&#233;l&#232;ne Loevenbruck</author>
<author pid="85/741">Tomoki Toda</author>
<title>Multimodal HMM-based NAM-to-speech conversion.</title>
<pages>656-659</pages>
<year>2009</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2009-230</ee>
<crossref>conf/interspeech/2009</crossref>
<url>db/conf/interspeech/interspeech2009.html#TranBLT09</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/NakamuraTSS09" mdate="2023-06-23">
<author pid="31/8013">Keigo Nakamura</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Electrolaryngeal speech enhancement based on statistical voice conversion.</title>
<pages>1431-1434</pages>
<year>2009</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2009-439</ee>
<crossref>conf/interspeech/2009</crossref>
<url>db/conf/interspeech/interspeech2009.html#NakamuraTSS09</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/OhtaniTSS09" mdate="2023-06-23">
<author pid="34/8763">Yamato Ohtani</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Many-to-many eigenvoice conversion with reference voice.</title>
<pages>1623-1626</pages>
<year>2009</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2009-485</ee>
<crossref>conf/interspeech/2009</crossref>
<url>db/conf/interspeech/interspeech2009.html#OhtaniTSS09</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/CharlierOTMD09" mdate="2023-06-23">
<author pid="73/9237">Malorie Charlier</author>
<author pid="34/8763">Yamato Ohtani</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="19/7269">Alexis Moinet</author>
<author pid="79/2569">Thierry Dutoit</author>
<title>Cross-language voice conversion based on eigenvoices.</title>
<pages>1635-1638</pages>
<year>2009</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2009-488</ee>
<crossref>conf/interspeech/2009</crossref>
<url>db/conf/interspeech/interspeech2009.html#CharlierOTMD09</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/MaiaTTSN09" mdate="2023-06-23">
<author pid="15/7825">Ranniery Maia</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="25/369">Keiichi Tokuda</author>
<author pid="57/448">Shinsuke Sakai</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>A decision tree-based clustering approach to state definition in an excitation modeling framework for HMM-based speech synthesis.</title>
<pages>1783-1786</pages>
<year>2009</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2009-149</ee>
<crossref>conf/interspeech/2009</crossref>
<url>db/conf/interspeech/interspeech2009.html#MaiaTTSN09</url>
</inproceedings>
</r>
<r><article key="journals/ieicet/CincarekTSS08" mdate="2021-04-09">
<author pid="05/3057">Tobias Cincarek</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Cost Reduction of Acoustic Modeling for Real-Environment Applications Using Unsupervised and Selective Training.</title>
<pages>499-507</pages>
<year>2008</year>
<volume>91-D</volume>
<journal>IEICE Trans. Inf. Syst.</journal>
<number>3</number>
<ee>https://doi.org/10.1093/ietisy/e91-d.3.499</ee>
<url>db/journals/ieicet/ieicet91d.html#CincarekTSS08</url>
</article>
</r>
<r><article key="journals/ieicet/NaginoSTSS08" mdate="2021-04-09">
<author pid="18/910">Goshu Nagino</author>
<author pid="23/6099">Makoto Shozakai</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Building an Effective Speech Corpus by Utilizing Statistical Multidimensional Scaling Method.</title>
<pages>607-614</pages>
<year>2008</year>
<volume>91-D</volume>
<journal>IEICE Trans. Inf. Syst.</journal>
<number>3</number>
<ee>https://doi.org/10.1093/ietisy/e91-d.3.607</ee>
<url>db/journals/ieicet/ieicet91d.html#NaginoSTSS08</url>
</article>
</r>
<r><article key="journals/ieicet/ZenTT08" mdate="2024-10-06">
<author orcid="0000-0002-8959-5471" pid="42/7014">Heiga Zen</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="25/369">Keiichi Tokuda</author>
<title>The Nitech-NAIST HMM-Based Speech Synthesis System for the Blizzard Challenge 2006.</title>
<pages>1764-1773</pages>
<year>2008</year>
<volume>91-D</volume>
<journal>IEICE Trans. Inf. Syst.</journal>
<number>6</number>
<ee>https://doi.org/10.1093/ietisy/e91-d.6.1764</ee>
<url>db/journals/ieicet/ieicet91d.html#ZenTT08</url>
</article>
</r>
<r><article key="journals/speech/TodaBT08" mdate="2021-04-09">
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="b/AlanWBlack">Alan W. Black</author>
<author pid="25/369">Keiichi Tokuda</author>
<title>Statistical mapping between articulatory movements and acoustic spectrum using a Gaussian mixture model.</title>
<pages>215-227</pages>
<year>2008</year>
<volume>50</volume>
<journal>Speech Commun.</journal>
<number>3</number>
<ee>https://doi.org/10.1016/j.specom.2007.09.001</ee>
<url>db/journals/speech/speech50.html#TodaBT08</url>
</article>
</r>
<r><inproceedings key="conf/blizzard/MaiaNSTTS008" mdate="2024-09-19">
<author pid="15/7825">Ranniery Maia</author>
<author pid="37/2394">Jinfu Ni</author>
<author pid="57/448">Shinsuke Sakai</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="25/369">Keiichi Tokuda</author>
<author pid="08/1223">Tohru Shimizu</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>The NICT/ATR speech synthesis system for the Blizzard Challenge 2008.</title>
<year>2008</year>
<booktitle>Blizzard Challenge</booktitle>
<ee type="oa">https://doi.org/10.21437/Blizzard.2008-3</ee>
<crossref>conf/blizzard/2008</crossref>
<url>db/conf/blizzard/blizzard2008.html#MaiaNSTTS008</url>
</inproceedings>
</r>
<r><inproceedings key="conf/blizzard/YamagishiZWTT08" mdate="2024-09-19">
<author pid="87/3979">Junichi Yamagishi</author>
<author pid="42/7014">Heiga Zen</author>
<author pid="91/6000">Yi-Jian Wu</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="25/369">Keiichi Tokuda</author>
<title>The HTS-2008 System: Yet Another Evaluation of the Speaker-Adaptive HMM-based Speech Synthesis System in The 2008 Blizzard Challenge.</title>
<year>2008</year>
<booktitle>Blizzard Challenge</booktitle>
<ee type="oa">https://doi.org/10.21437/Blizzard.2008-7</ee>
<crossref>conf/blizzard/2008</crossref>
<url>db/conf/blizzard/blizzard2008.html#YamagishiZWTT08</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/TodaT08" mdate="2021-04-09">
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="25/369">Keiichi Tokuda</author>
<title>Statistical approach to vocal tract transfer function estimation based on factor analyzed trajectory HMM.</title>
<pages>3925-3928</pages>
<year>2008</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2008.4518512</ee>
<crossref>conf/icassp/2008</crossref>
<url>db/conf/icassp/icassp2008.html#TodaT08</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/YamagishiNZTT08" mdate="2025-03-03">
<author pid="87/3979">Junichi Yamagishi</author>
<author orcid="0000-0002-2278-0429" pid="83/3783">Takashi Nose</author>
<author orcid="0000-0002-8959-5471" pid="42/7014">Heiga Zen</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="25/369">Keiichi Tokuda</author>
<title>Performance evaluation of the speaker-independent HMM-based speech synthesis system &#34;HTS 2007&#34; for the Blizzard Challenge 2007.</title>
<pages>3957-3960</pages>
<year>2008</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2008.4518520</ee>
<crossref>conf/icassp/2008</crossref>
<url>db/conf/icassp/icassp2008.html#YamagishiNZTT08</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/MaiaTTSN08" mdate="2021-04-09">
<author pid="15/7825">Ranniery Maia</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="25/369">Keiichi Tokuda</author>
<author pid="04/7823">Shinichi Sakai</author>
<author pid="79/7819">Shun Nakamura</author>
<title>On the state definition for a trainable excitation model in HMM-based speech synthesis.</title>
<pages>3965-3968</pages>
<year>2008</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2008.4518522</ee>
<crossref>conf/icassp/2008</crossref>
<url>db/conf/icassp/icassp2008.html#MaiaTTSN08</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/YutaniUNTT08" mdate="2023-06-23">
<author pid="25/8056">Kaori Yutani</author>
<author pid="91/8053">Yosuke Uto</author>
<author pid="90/302">Yoshihiko Nankaku</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="25/369">Keiichi Tokuda</author>
<title>Simultaneous conversion of duration and spectrum based on statistical models including time-sequence matching.</title>
<pages>1072-1075</pages>
<year>2008</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2008-331</ee>
<crossref>conf/interspeech/2008</crossref>
<url>db/conf/interspeech/interspeech2008.html#YutaniUNTT08</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/MuramatsuOTSS08" mdate="2023-06-23">
<author pid="40/9234">Takashi Muramatsu</author>
<author pid="34/8763">Yamato Ohtani</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Low-delay voice conversion based on maximum likelihood estimation of spectral parameter trajectory.</title>
<pages>1076-1079</pages>
<year>2008</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2008-332</ee>
<crossref>conf/interspeech/2008</crossref>
<url>db/conf/interspeech/interspeech2008.html#MuramatsuOTSS08</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/OhtaniTSS08" mdate="2023-06-23">
<author pid="34/8763">Yamato Ohtani</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>An improved one-to-many eigenvoice conversion system.</title>
<pages>1080-1083</pages>
<year>2008</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2008-333</ee>
<crossref>conf/interspeech/2008</crossref>
<url>db/conf/interspeech/interspeech2008.html#OhtaniTSS08</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/TaniTOSS08" mdate="2023-06-23">
<author pid="56/9234">Daisuke Tani</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="34/8763">Yamato Ohtani</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Maximum a posteriori adaptation for many-to-one eigenvoice conversion.</title>
<pages>1461-1463</pages>
<year>2008</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2008-421</ee>
<crossref>conf/interspeech/2008</crossref>
<url>db/conf/interspeech/interspeech2008.html#TaniTOSS08</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/NakamuraTNSS08" mdate="2023-06-23">
<author pid="31/8013">Keigo Nakamura</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="75/5531">Yoshitaka Nakajima</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Evaluation of speaking-aid system with voice conversion for laryngectomees toward its use in practical environments.</title>
<pages>2209-2212</pages>
<year>2008</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2008-577</ee>
<crossref>conf/interspeech/2008</crossref>
<url>db/conf/interspeech/interspeech2008.html#NakamuraTNSS08</url>
</inproceedings>
</r>
<r><inproceedings key="conf/iscslp/OuraNTTMSN08" mdate="2024-09-18">
<author pid="22/7077">Keiichiro Oura</author>
<author pid="90/302">Yoshihiko Nankaku</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="25/369">Keiichi Tokuda</author>
<author pid="15/7825">Ranniery Maia</author>
<author pid="57/448">Shinsuke Sakai</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Simultaneous Acoustic, Prosodic, and Phrasing Model Training for TTs Conversion Systems.</title>
<pages>1-4</pages>
<year>2008</year>
<booktitle>ISCSLP</booktitle>
<ee>https://doi.org/10.1109/CHINSL.2008.ECP.12</ee>
<ee type="oa">https://www.isca-archive.org/iscslp_2008/oura08_iscslp.html</ee>
<crossref>conf/iscslp/2008</crossref>
<url>db/conf/iscslp/iscslp2008.html#OuraNTTMSN08</url>
</inproceedings>
</r>
<r><article key="journals/ieicet/ZenTNT07" mdate="2024-10-06">
<author orcid="0000-0002-8959-5471" pid="42/7014">Heiga Zen</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="80/1615">Masaru Nakamura</author>
<author pid="25/369">Keiichi Tokuda</author>
<title>Details of the Nitech HMM-Based Speech Synthesis System for the Blizzard Challenge 2005.</title>
<pages>325-333</pages>
<year>2007</year>
<volume>90-D</volume>
<journal>IEICE Trans. Inf. Syst.</journal>
<number>1</number>
<ee>https://doi.org/10.1093/ietisy/e90-1.1.325</ee>
<url>db/journals/ieicet/ieicet90d.html#ZenTNT07</url>
</article>
</r>
<r><article key="journals/ieicet/GomezTSS07" mdate="2021-04-09">
<author pid="44/2122">Randy Gomez</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Reducing Computation Time of the Rapid Unsupervised Speaker Adaptation Based on HMM-Sufficient Statistics.</title>
<pages>554-561</pages>
<year>2007</year>
<volume>90-D</volume>
<journal>IEICE Trans. Inf. Syst.</journal>
<number>2</number>
<ee>https://doi.org/10.1093/ietisy/e90-d.2.554</ee>
<url>db/journals/ieicet/ieicet90d.html#GomezTSS07</url>
</article>
</r>
<r><article key="journals/ieicet/TodaT07" mdate="2021-04-09">
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="25/369">Keiichi Tokuda</author>
<title>A Speech Parameter Generation Algorithm Considering Global Variance for HMM-Based Speech Synthesis.</title>
<pages>816-824</pages>
<year>2007</year>
<volume>90-D</volume>
<journal>IEICE Trans. Inf. Syst.</journal>
<number>5</number>
<ee>https://doi.org/10.1093/ietisy/e90-d.5.816</ee>
<url>db/journals/ieicet/ieicet90d.html#TodaT07</url>
</article>
</r>
<r><article key="journals/taslp/TodaBT07" mdate="2021-04-09">
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="b/AlanWBlack">Alan W. Black</author>
<author pid="25/369">Keiichi Tokuda</author>
<title>Voice Conversion Based on Maximum-Likelihood Estimation of Spectral Parameter Trajectory.</title>
<pages>2222-2235</pages>
<year>2007</year>
<volume>15</volume>
<journal>IEEE Trans. Speech Audio Process.</journal>
<number>8</number>
<ee>https://doi.org/10.1109/TASL.2007.907344</ee>
<url>db/journals/taslp/taslp15.html#TodaBT07</url>
</article>
</r>
<r><inproceedings key="conf/blizzard/NiHKTTTSM007" mdate="2024-09-19">
<author pid="37/2394">Jinfu Ni</author>
<author pid="01/2055">Toshio Hirai</author>
<author pid="32/4341">Hisashi Kawai</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="25/369">Keiichi Tokuda</author>
<author pid="20/3844">Minoru Tsuzaki</author>
<author pid="57/448">Shinsuke Sakai</author>
<author pid="15/7825">Ranniery Maia</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>ATRECSS - ATR English speech corpus for speech synthesis.</title>
<year>2007</year>
<booktitle>Blizzard Challenge</booktitle>
<ee type="oa">https://www.isca-archive.org/blizzard_2007/ni07_blizzard.html</ee>
<crossref>conf/blizzard/2007</crossref>
<url>db/conf/blizzard/blizzard2007.html#NiHKTTTSM007</url>
</inproceedings>
</r>
<r><inproceedings key="conf/blizzard/YamagishiZTT07" mdate="2024-09-19">
<author pid="87/3979">Junichi Yamagishi</author>
<author pid="42/7014">Heiga Zen</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="25/369">Keiichi Tokuda</author>
<title>Speaker-independent HMM-based speech synthesis system - HTS-2007 system for the Blizzard Challenge 2007.</title>
<year>2007</year>
<booktitle>Blizzard Challenge</booktitle>
<ee type="oa">https://www.isca-archive.org/blizzard_2007/yamagishi07_blizzard.html</ee>
<crossref>conf/blizzard/2007</crossref>
<url>db/conf/blizzard/blizzard2007.html#YamagishiZTT07</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/TodaOS07" mdate="2021-04-09">
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="34/8763">Yamato Ohtani</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>One-to-Many and Many-to-One Voice Conversion Based on Eigenvoices.</title>
<pages>1249-1252</pages>
<year>2007</year>
<booktitle>ICASSP (4)</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2007.367303</ee>
<crossref>conf/icassp/2007</crossref>
<url>db/conf/icassp/icassp2007.html#TodaOS07</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/GomezTSS07" mdate="2023-06-23">
<author pid="44/2122">Randy Gomez</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Rapid unsupervised speaker adaptation using single utterance based on MLLR and speaker selection.</title>
<pages>262-265</pages>
<year>2007</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2007-117</ee>
<crossref>conf/interspeech/2007</crossref>
<url>db/conf/interspeech/interspeech2007.html#GomezTSS07</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/CincarekSTSS07" mdate="2023-06-23">
<author pid="05/3057">Tobias Cincarek</author>
<author pid="04/9250">Izumi Shindo</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Development of preschool children subsystem for ASR and q&#38;a in a real-environment speech-oriented guidance task.</title>
<pages>1469-1472</pages>
<year>2007</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2007-426</ee>
<crossref>conf/interspeech/2007</crossref>
<url>db/conf/interspeech/interspeech2007.html#CincarekSTSS07</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/MaiaTZNT07" mdate="2023-06-23">
<author pid="15/7825">Ranniery Maia</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="42/7014">Heiga Zen</author>
<author pid="90/302">Yoshihiko Nankaku</author>
<author pid="25/369">Keiichi Tokuda</author>
<title>A trainable excitation model for HMM-based speech synthesis.</title>
<pages>1909-1912</pages>
<year>2007</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2007-530</ee>
<crossref>conf/interspeech/2007</crossref>
<url>db/conf/interspeech/interspeech2007.html#MaiaTZNT07</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/OhtaniTSS07" mdate="2023-06-23">
<author pid="34/8763">Yamato Ohtani</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Speaker adaptive training for one-to-many eigenvoice conversion based on Gaussian mixture model.</title>
<pages>1981-1984</pages>
<year>2007</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2007-554</ee>
<crossref>conf/interspeech/2007</crossref>
<url>db/conf/interspeech/interspeech2007.html#OhtaniTSS07</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/NakamuraTSS07" mdate="2023-06-23">
<author pid="31/8013">Keigo Nakamura</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Impact of various small sound source signals on voice conversion accuracy in speech communication aid for laryngectomees.</title>
<pages>2517-2520</pages>
<year>2007</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2007-669</ee>
<crossref>conf/interspeech/2007</crossref>
<url>db/conf/interspeech/interspeech2007.html#NakamuraTSS07</url>
</inproceedings>
</r>
<r><inproceedings key="conf/ssw/SakaiNMTTTKN07" mdate="2024-07-31">
<author pid="57/448">Shinsuke Sakai</author>
<author pid="37/2394">Jinfu Ni</author>
<author pid="15/7825">Ranniery Maia</author>
<author pid="25/369">Keiichi Tokuda</author>
<author pid="20/3844">Minoru Tsuzaki</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="32/4341">Hisashi Kawai</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Communicative speech synthesis with XIMERA: a first step.</title>
<pages>28-33</pages>
<year>2007</year>
<booktitle>SSW</booktitle>
<ee type="oa">https://www.isca-archive.org/ssw_2007/sakai07_ssw.html</ee>
<crossref>conf/ssw/2007</crossref>
<url>db/conf/ssw/ssw2007.html#SakaiNMTTTKN07</url>
</inproceedings>
</r>
<r><inproceedings key="conf/ssw/OhtaOTSS07" mdate="2024-07-31">
<author pid="06/9233">Kumi Ohta</author>
<author pid="34/8763">Yamato Ohtani</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Regression approaches to voice quality controll based on one-to-many eigenvoice conversion.</title>
<pages>101-106</pages>
<year>2007</year>
<booktitle>SSW</booktitle>
<ee type="oa">https://www.isca-archive.org/ssw_2007/ohta07_ssw.html</ee>
<crossref>conf/ssw/2007</crossref>
<url>db/conf/ssw/ssw2007.html#OhtaOTSS07</url>
</inproceedings>
</r>
<r><inproceedings key="conf/ssw/TaniOTSS07" mdate="2024-07-31">
<author pid="56/9234">Daisuke Tani</author>
<author pid="34/8763">Yamato Ohtani</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>An evaluation of many-to-one voice conversion algorithms with pre-stored speaker data sets.</title>
<pages>107-112</pages>
<year>2007</year>
<booktitle>SSW</booktitle>
<ee type="oa">https://www.isca-archive.org/ssw_2007/tani07_ssw.html</ee>
<crossref>conf/ssw/2007</crossref>
<url>db/conf/ssw/ssw2007.html#TaniOTSS07</url>
</inproceedings>
</r>
<r><inproceedings key="conf/ssw/YamagishiKRKZTT07" mdate="2024-07-31">
<author pid="87/3979">Junichi Yamagishi</author>
<author pid="88/1648">Takao Kobayashi</author>
<author pid="33/3792">Steve Renals</author>
<author pid="68/2005">Simon King 0001</author>
<author pid="42/7014">Heiga Zen</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="25/369">Keiichi Tokuda</author>
<title>Improved average-voice-based speech synthesis using gender-mixed modeling and a parameter generation algorithm considering GV.</title>
<pages>125-130</pages>
<year>2007</year>
<booktitle>SSW</booktitle>
<ee type="oa">https://www.isca-archive.org/ssw_2007/yamagishi07_ssw.html</ee>
<crossref>conf/ssw/2007</crossref>
<url>db/conf/ssw/ssw2007.html#YamagishiKRKZTT07</url>
</inproceedings>
</r>
<r><inproceedings key="conf/ssw/MaiaTZNT07" mdate="2024-07-31">
<author pid="15/7825">Ranniery Maia</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="42/7014">Heiga Zen</author>
<author pid="90/302">Yoshihiko Nankaku</author>
<author pid="25/369">Keiichi Tokuda</author>
<title>An excitation model for HMM-based speech synthesis based on residual modeling.</title>
<pages>131-136</pages>
<year>2007</year>
<booktitle>SSW</booktitle>
<ee type="oa">https://www.isca-archive.org/ssw_2007/maia07_ssw.html</ee>
<crossref>conf/ssw/2007</crossref>
<url>db/conf/ssw/ssw2007.html#MaiaTZNT07</url>
</inproceedings>
</r>
<r><inproceedings key="conf/ssw/NankakuNTT07" mdate="2024-07-31">
<author pid="90/302">Yoshihiko Nankaku</author>
<author pid="35/7736">Kenichi Nakamura</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="25/369">Keiichi Tokuda</author>
<title>Spectral conversion based on statistical models including time-sequence matching.</title>
<pages>333-338</pages>
<year>2007</year>
<booktitle>SSW</booktitle>
<ee type="oa">https://www.isca-archive.org/ssw_2007/nankaku07_ssw.html</ee>
<crossref>conf/ssw/2007</crossref>
<url>db/conf/ssw/ssw2007.html#NankakuNTT07</url>
</inproceedings>
</r>
<r><article key="journals/ieicet/CincarekTSS06" mdate="2021-04-09">
<author pid="05/3057">Tobias Cincarek</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Utterance-Based Selective Training for the Automatic Creation of Task-Dependent Acoustic Models.</title>
<pages>962-969</pages>
<year>2006</year>
<volume>89-D</volume>
<journal>IEICE Trans. Inf. Syst.</journal>
<number>3</number>
<ee>https://doi.org/10.1093/ietisy/e89-d.3.962</ee>
<url>db/journals/ieicet/ieicet89d.html#CincarekTSS06</url>
</article>
</r>
<r><article key="journals/ieicet/GomezLTSS06" mdate="2021-04-09">
<author pid="44/2122">Randy Gomez</author>
<author pid="85/6462">Akinobu Lee</author>
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Improving Rapid Unsupervised Speaker Adaptation Based on HMM-Sufficient Statistics in Noisy Environments Using Multi-Template Models.</title>
<pages>998-1005</pages>
<year>2006</year>
<volume>89-D</volume>
<journal>IEICE Trans. Inf. Syst.</journal>
<number>3</number>
<ee>https://doi.org/10.1093/ietisy/e89-d.3.998</ee>
<url>db/journals/ieicet/ieicet89d.html#GomezLTSS06</url>
</article>
</r>
<r><article key="journals/speech/TodaKTS06" mdate="2021-04-09">
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="32/4341">Hisashi Kawai</author>
<author pid="20/3844">Minoru Tsuzaki</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>An evaluation of cost functions sensitively capturing local degradation of naturalness for segment selection in concatenative speech synthesis.</title>
<pages>45-56</pages>
<year>2006</year>
<volume>48</volume>
<journal>Speech Commun.</journal>
<number>1</number>
<ee>https://doi.org/10.1016/j.specom.2005.05.011</ee>
<url>db/journals/speech/speech48.html#TodaKTS06</url>
</article>
</r>
<r><inproceedings key="conf/blizzard/TodaKHNNYTT006" mdate="2024-09-18">
<author pid="85/741">Tomoki Toda</author>
<author pid="32/4341">Hisashi Kawai</author>
<author pid="01/2055">Toshio Hirai</author>
<author pid="37/2394">Jinfu Ni</author>
<author pid="95/7824">Nobuyuki Nishizawa</author>
<author pid="87/3979">Junichi Yamagishi</author>
<author pid="20/3844">Minoru Tsuzaki</author>
<author pid="25/369">Keiichi Tokuda</author>
<author pid="57/1548-1">Satoshi Nakamura 0001</author>
<title>Developing a Test Bed of English Text-to-Speech System XIMERA for the Blizzard Challenge 2006.</title>
<year>2006</year>
<booktitle>Blizzard Challenge</booktitle>
<ee type="oa">https://doi.org/10.21437/Blizzard.2006-7</ee>
<crossref>conf/blizzard/2006</crossref>
<url>db/conf/blizzard/blizzard2006.html#TodaKHNNYTT006</url>
</inproceedings>
</r>
<r><inproceedings key="conf/blizzard/ZenTT06" mdate="2024-09-18">
<author pid="42/7014">Heiga Zen</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="25/369">Keiichi Tokuda</author>
<title>The Nitech-NAIST HMM-based speech synthesis system for the Blizzard Challenge 2006.</title>
<year>2006</year>
<booktitle>Blizzard Challenge</booktitle>
<ee type="oa">https://doi.org/10.21437/Blizzard.2006-3</ee>
<crossref>conf/blizzard/2006</crossref>
<url>db/conf/blizzard/blizzard2006.html#ZenTT06</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/NakamuraTNT06" mdate="2020-06-22">
<author pid="35/7736">Kenichi Nakamura</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="90/302">Yoshihiko Nankaku</author>
<author pid="25/369">Keiichi Tokuda</author>
<title>On the Use of Phonetic Information for Mapping from Articulatory Movements to Vocal Tract Spectrum.</title>
<pages>93-96</pages>
<year>2006</year>
<booktitle>ICASSP (1)</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2006.1659965</ee>
<crossref>conf/icassp/2006</crossref>
<url>db/conf/icassp/icassp2006.html#NakamuraTNT06</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/GomezTSS06" mdate="2020-06-22">
<author pid="44/2122">Randy Gomez</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Improving Rapid Unsupervised Speaker Adaptation Based On Hmm Sufficient Statistics.</title>
<pages>1001-1004</pages>
<year>2006</year>
<booktitle>ICASSP (1)</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2006.1660192</ee>
<crossref>conf/icassp/2006</crossref>
<url>db/conf/icassp/icassp2006.html#GomezTSS06</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/CincarekTSS06" mdate="2023-06-22">
<author pid="05/3057">Tobias Cincarek</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Acoustic modeling for spoken dialogue systems based on unsupervised utterance-based selective training.</title>
<year>2006</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2006-478</ee>
<crossref>conf/interspeech/2006</crossref>
<url>db/conf/interspeech/interspeech2006.html#CincarekTSS06</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/NakagiriTKS06" mdate="2023-06-22">
<author pid="54/9262">Mikihiro Nakagiri</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="92/6978">Hideki Kashioka</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Improving body transmitted unvoiced speech with statistical voice conversion.</title>
<year>2006</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2006-583</ee>
<crossref>conf/interspeech/2006</crossref>
<url>db/conf/interspeech/interspeech2006.html#NakagiriTKS06</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/NakamuraTSS06" mdate="2023-06-22">
<author pid="31/8013">Keigo Nakamura</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Speaking aid system for total laryngectomees using voice conversion of body transmitted artificial speech.</title>
<year>2006</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2006-419</ee>
<crossref>conf/interspeech/2006</crossref>
<url>db/conf/interspeech/interspeech2006.html#NakamuraTSS06</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/OhtaniTSS06" mdate="2023-06-22">
<author pid="34/8763">Yamato Ohtani</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Maximum likelihood voice conversion based on GMM with STRAIGHT mixed excitation.</title>
<year>2006</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2006-582</ee>
<crossref>conf/interspeech/2006</crossref>
<url>db/conf/interspeech/interspeech2006.html#OhtaniTSS06</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/TodaOS06" mdate="2023-06-22">
<author pid="85/741">Tomoki Toda</author>
<author pid="34/8763">Yamato Ohtani</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Eigenvoice conversion based on Gaussian mixture model.</title>
<year>2006</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2006-613</ee>
<crossref>conf/interspeech/2006</crossref>
<url>db/conf/interspeech/interspeech2006.html#TodaOS06</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/UtoNTLT06" mdate="2023-06-22">
<author pid="91/8053">Yosuke Uto</author>
<author pid="90/302">Yoshihiko Nankaku</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="85/6462">Akinobu Lee</author>
<author pid="25/369">Keiichi Tokuda</author>
<title>Voice conversion based on mixtures of factor analyzers.</title>
<year>2006</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2006-585</ee>
<crossref>conf/interspeech/2006</crossref>
<url>db/conf/interspeech/interspeech2006.html#UtoNTLT06</url>
</inproceedings>
</r>
<r><article key="journals/ieicet/AdachiTKSS05" mdate="2020-04-11">
<author pid="43/4541">Kazuki Adachi</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="02/6430">Hiromichi Kawanami</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Designing Target Cost Function Based on Prosody of Speech Database.</title>
<pages>519-524</pages>
<ee>http://search.ieice.org/bin/summary.php?id=e88-d_3_519&#38;category=D&#38;year=2005&#38;lang=E&#38;abst=</ee>
<year>2005</year>
<volume>88-D</volume>
<journal>IEICE Trans. Inf. Syst.</journal>
<number>3</number>
<url>db/journals/ieicet/ieicet88d.html#AdachiTKSS05</url>
</article>
</r>
<r><inproceedings key="conf/icassp/TodaBT05" mdate="2021-04-09">
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="b/AlanWBlack">Alan W. Black</author>
<author pid="25/369">Keiichi Tokuda</author>
<title>Spectral Conversion Based on Maximum Likelihood Estimation Considering Global Variance of Converted Parameter.</title>
<pages>9-12</pages>
<year>2005</year>
<booktitle>ICASSP (1)</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2005.1415037</ee>
<crossref>conf/icassp/2005</crossref>
<url>db/conf/icassp/icassp2005.html#TodaBT05</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/ZenT05" mdate="2023-06-22">
<author pid="42/7014">Heiga Zen</author>
<author pid="85/741">Tomoki Toda</author>
<title>An overview of nitech HMM-based speech synthesis system for blizzard challenge 2005.</title>
<pages>93-96</pages>
<year>2005</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2005-76</ee>
<crossref>conf/interspeech/2005</crossref>
<url>db/conf/interspeech/interspeech2005.html#ZenT05</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/TodaS05" mdate="2023-06-22">
<author pid="85/741">Tomoki Toda</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>NAM-to-speech conversion with Gaussian mixture models.</title>
<pages>1957-1960</pages>
<year>2005</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2005-611</ee>
<crossref>conf/interspeech/2005</crossref>
<url>db/conf/interspeech/interspeech2005.html#TodaS05</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/TodaT05" mdate="2023-06-22">
<author pid="85/741">Tomoki Toda</author>
<author pid="25/369">Keiichi Tokuda</author>
<title>Speech parameter generation algorithm considering global variance for HMM-based speech synthesis.</title>
<pages>2801-2804</pages>
<year>2005</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2005-617</ee>
<crossref>conf/interspeech/2005</crossref>
<url>db/conf/interspeech/interspeech2005.html#TodaT05</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/TodaKT04" mdate="2020-06-22">
<author pid="85/741">Tomoki Toda</author>
<author pid="32/4341">Hisashi Kawai</author>
<author pid="20/3844">Minoru Tsuzaki</author>
<title>Optimizing sub-cost functions for segment selection based on perceptual evaluations in concatenative speech synthesis.</title>
<pages>657-660</pages>
<year>2004</year>
<booktitle>ICASSP (1)</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2004.1326071</ee>
<crossref>conf/icassp/2004</crossref>
<url>db/conf/icassp/icassp2004.html#TodaKT04</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/KawaiT04" mdate="2020-06-22">
<author pid="32/4341">Hisashi Kawai</author>
<author pid="85/741">Tomoki Toda</author>
<title>An evaluation of automatic phone segmentation for concatenative speech synthesis.</title>
<pages>677-680</pages>
<year>2004</year>
<booktitle>ICASSP (1)</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2004.1326076</ee>
<crossref>conf/icassp/2004</crossref>
<url>db/conf/icassp/icassp2004.html#KawaiT04</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/TodaBT04" mdate="2023-06-22">
<author pid="85/741">Tomoki Toda</author>
<author pid="b/AlanWBlack">Alan W. Black</author>
<author pid="25/369">Keiichi Tokuda</author>
<title>Acoustic-to-articulatory inversion mapping with Gaussian mixture model.</title>
<year>2004</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2004-410</ee>
<crossref>conf/interspeech/2004</crossref>
<url>db/conf/interspeech/interspeech2004.html#TodaBT04</url>
<pages>1129-1132</pages>
</inproceedings>
</r>
<r><inproceedings key="conf/lrec/AdachiTKSS04" mdate="2019-08-19">
<author pid="43/4541">Kazuki Adachi</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="02/6430">Hiromichi Kawanami</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Perceptual Evaluation of Quality Deterioration Owing to Prosody Modification.</title>
<year>2004</year>
<booktitle>LREC</booktitle>
<ee type="oa">http://www.lrec-conf.org/proceedings/lrec2004/summaries/681.htm</ee>
<crossref>conf/lrec/2004</crossref>
<url>db/conf/lrec/lrec2004.html#AdachiTKSS04</url>
</inproceedings>
</r>
<r><inproceedings key="conf/ssw/TodaBT04" mdate="2024-07-31">
<author pid="85/741">Tomoki Toda</author>
<author pid="b/AlanWBlack">Alan W. Black</author>
<author pid="25/369">Keiichi Tokuda</author>
<title>Mapping from articulatory movements to vocal tract spectrum with Gaussian mixture model for articulatory speech synthesis.</title>
<pages>31-36</pages>
<year>2004</year>
<booktitle>SSW</booktitle>
<ee type="oa">https://www.isca-archive.org/ssw_2004/toda04_ssw.html</ee>
<crossref>conf/ssw/2004</crossref>
<url>db/conf/ssw/ssw2004.html#TodaBT04</url>
</inproceedings>
</r>
<r><inproceedings key="conf/ssw/KawaiTNTT04" mdate="2024-07-31">
<author pid="32/4341">Hisashi Kawai</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="37/2394">Jinfu Ni</author>
<author pid="20/3844">Minoru Tsuzaki</author>
<author pid="25/369">Keiichi Tokuda</author>
<title>XIMERA: a new TTS from ATR based on corpus-based technologies.</title>
<pages>179-184</pages>
<year>2004</year>
<booktitle>SSW</booktitle>
<ee type="oa">https://www.isca-archive.org/ssw_2004/kawai04_ssw.html</ee>
<crossref>conf/ssw/2004</crossref>
<url>db/conf/ssw/ssw2004.html#KawaiTNTT04</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/TodaKTS03" mdate="2020-06-22">
<author pid="85/741">Tomoki Toda</author>
<author pid="32/4341">Hisashi Kawai</author>
<author pid="20/3844">Minoru Tsuzaki</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Segment selection considering local degradation of naturalness in concatenative speech synthesis.</title>
<pages>696-699</pages>
<year>2003</year>
<booktitle>ICASSP (1)</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2003.1198876</ee>
<crossref>conf/icassp/2003</crossref>
<url>db/conf/icassp/icassp2003.html#TodaKTS03</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/TodaKT03" mdate="2023-06-22">
<author pid="85/741">Tomoki Toda</author>
<author pid="32/4341">Hisashi Kawai</author>
<author pid="20/3844">Minoru Tsuzaki</author>
<title>Optimizing integrated cost function for segment selection in concatenative speech synthesis based on perceptual evaluations.</title>
<year>2003</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Eurospeech.2003-123</ee>
<crossref>conf/interspeech/2003</crossref>
<url>db/conf/interspeech/interspeech2003.html#TodaKT03</url>
<pages>297-300</pages>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/ShiraishiTKSS03" mdate="2023-06-22">
<author pid="20/9286">Tatsuya Shiraishi</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="02/6430">Hiromichi Kawanami</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Simple designing methods of corpus-based visual speech synthesis.</title>
<year>2003</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Eurospeech.2003-627</ee>
<crossref>conf/interspeech/2003</crossref>
<url>db/conf/interspeech/interspeech2003.html#ShiraishiTKSS03</url>
<pages>2241-2244</pages>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/KawanamiITSS03" mdate="2023-06-22">
<author pid="02/6430">Hiromichi Kawanami</author>
<author pid="63/9283">Yohei Iwami</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>GMM-based voice conversion applied to emotional speech synthesis.</title>
<year>2003</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Eurospeech.2003-661</ee>
<crossref>conf/interspeech/2003</crossref>
<url>db/conf/interspeech/interspeech2003.html#KawanamiITSS03</url>
<pages>2401-2404</pages>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/TodaKTS02" mdate="2021-04-09">
<author orcid="0000-0001-8146-1279" pid="85/741">Tomoki Toda</author>
<author pid="32/4341">Hisashi Kawai</author>
<author pid="20/3844">Minoru Tsuzaki</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Unit selection algorithm for Japanese speech synthesis based on both phoneme unit and diphone unit.</title>
<pages>465-468</pages>
<year>2002</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2002.5743755</ee>
<crossref>conf/icassp/2002</crossref>
<url>db/conf/icassp/icassp2002.html#TodaKTS02</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/MashimoTKKSC02" mdate="2023-06-22">
<author pid="68/9241">Mikiko Mashimo</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="02/6430">Hiromichi Kawanami</author>
<author pid="92/6978">Hideki Kashioka</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<author pid="42/87">Nick Campbell 0001</author>
<title>Evaluation of cross-language voice conversion using bilingual and non-bilingual databases.</title>
<year>2002</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/ICSLP.2002-138</ee>
<crossref>conf/interspeech/2002</crossref>
<url>db/conf/interspeech/interspeech2002.html#MashimoTKKSC02</url>
<pages>293-296</pages>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/KawanamiMTS02" mdate="2023-06-22">
<author pid="02/6430">Hiromichi Kawanami</author>
<author pid="11/1981">Tsuyoshi Masuda</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Designing Japanese speech database covering wide range in prosody for hybrid speech synthesizer.</title>
<year>2002</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/ICSLP.2002-111</ee>
<crossref>conf/interspeech/2002</crossref>
<url>db/conf/interspeech/interspeech2002.html#KawanamiMTS02</url>
<pages>2425-2428</pages>
</inproceedings>
</r>
<r><inproceedings key="conf/lrec/KawanamiMTS02" mdate="2019-08-19">
<author pid="02/6430">Hiromichi Kawanami</author>
<author pid="11/1981">Tsuyoshi Masuda</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Designing speech database with prosodic variety for expressive TTS system.</title>
<year>2002</year>
<booktitle>LREC</booktitle>
<ee type="oa">http://www.lrec-conf.org/proceedings/lrec2002/sumarios/337.htm</ee>
<crossref>conf/lrec/2002</crossref>
<url>db/conf/lrec/lrec2002.html#KawanamiMTS02</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/TodaSS01" mdate="2023-03-23">
<author pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Voice conversion algorithm based on Gaussian mixture model with dynamic frequency warping of STRAIGHT spectrum.</title>
<pages>841-844</pages>
<year>2001</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2001.941046</ee>
<ee>https://doi.ieeecomputersociety.org/10.1109/ICASSP.2001.941046</ee>
<crossref>conf/icassp/2001</crossref>
<url>db/conf/icassp/icassp2001.html#TodaSS01</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/TodaSS01" mdate="2023-06-22">
<author pid="85/741">Tomoki Toda</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>High quality voice conversion based on Gaussian mixture model with dynamic frequency warping.</title>
<pages>349-352</pages>
<year>2001</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Eurospeech.2001-108</ee>
<crossref>conf/interspeech/2001</crossref>
<url>db/conf/interspeech/interspeech2001.html#TodaSS01</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/MashimoTSC01" mdate="2023-06-22">
<author pid="68/9241">Mikiko Mashimo</author>
<author pid="85/741">Tomoki Toda</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<author pid="42/87">Nick Campbell 0001</author>
<title>Evaluation of cross-language voice conversion based on GMM and straight.</title>
<pages>361-364</pages>
<year>2001</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Eurospeech.2001-111</ee>
<crossref>conf/interspeech/2001</crossref>
<url>db/conf/interspeech/interspeech2001.html#MashimoTSC01</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/TodaLSS00" mdate="2023-06-22">
<author pid="85/741">Tomoki Toda</author>
<author pid="59/9295">Jinlin Lu</author>
<author pid="87/629">Hiroshi Saruwatari</author>
<author pid="50/2173">Kiyohiro Shikano</author>
<title>Straight-based voice conversion algorithm based on Gaussian mixture model.</title>
<pages>279-282</pages>
<year>2000</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/ICSLP.2000-532</ee>
<crossref>conf/interspeech/2000</crossref>
<url>db/conf/interspeech/interspeech2000.html#TodaLSS00</url>
</inproceedings>
</r>
<coauthors n="388" nc="1">
<co c="0"><na f="a/Abur:Defne" pid="329/8124">Defne Abur</na></co>
<co c="0"><na f="a/Adachi:Fumihiro" pid="84/4862">Fumihiro Adachi</na></co>
<co c="0"><na f="a/Adachi:Kazuki" pid="43/4541">Kazuki Adachi</na></co>
<co c="0"><na f="a/Adachi:Yusuke" pid="168/2958">Yusuke Adachi</na></co>
<co c="0"><na f="a/Adriani:Mirna" pid="15/6057">Mirna Adriani</na></co>
<co c="0"><na f="a/Ahmadi:Farzaneh" pid="52/9232">Farzaneh Ahmadi</na></co>
<co c="0"><na f="a/Akabe:Koichi" pid="151/8457">Koichi Akabe</na></co>
<co c="0"><na f="a/Akagi:Masato" pid="18/2260">Masato Akagi</na></co>
<co c="0"><na f="a/Alku:Paavo" pid="99/2726">Paavo Alku</na></co>
<co c="0"><na f="a/Ando:Atsushi" pid="173/6654">Atsushi Ando</na></co>
<co c="0"><na f="a/Aono:Yushi" pid="175/7791">Yushi Aono</na></co>
<co c="0"><na f="a/Arthur:Philip" pid="150/5396">Philip Arthur</na></co>
<co c="0"><na f="a/Ashihara:Takanori" pid="260/4522">Takanori Ashihara</na></co>
<co c="0"><na f="a/Astudillo:Ram=oacute=n_Fernandez" pid="56/7987">Ram&#243;n Fernandez Astudillo</na></co>
<co c="0"><na f="b/Babani:Denis" pid="58/9879">Denis Babani</na></co>
<co c="0"><na f="b/Bailly:G=eacute=rard" pid="48/1036">G&#233;rard Bailly</na></co>
<co c="0"><na f="b/Banno:Hideki" pid="95/6268">Hideki Banno</na></co>
<co c="0"><na f="b/Becker:Markus" pid="63/4934">Markus Becker</na></co>
<co c="0"><na f="b/Black:Alan_W=" pid="b/AlanWBlack">Alan W. Black</na></co>
<co c="0"><na f="b/Bonastre:Jean=Fran=ccedil=ois" pid="72/3130">Jean-Fran&#231;ois Bonastre</na></co>
<co c="0"><na f="c/Campbell_0001:Nick" pid="42/87">Nick Campbell 0001</na></co>
<co c="0"><na f="c/Charlier:Malorie" pid="73/9237">Malorie Charlier</na></co>
<co c="0"><na f="c/Chen_0002:Kuan=Yu" pid="35/6313-2">Kuan-Yu Chen 0002</na></co>
<co c="0" n="2"><na f="c/Chen:Linghui" pid="85/9874">Linghui Chen</na><na>Ling-Hui Chen</na></co>
<co c="0"><na f="c/Chen:Liping" pid="88/1450">Liping Chen</na></co>
<co c="-1"><na f="c/Chen:Shaowen" pid="342/6470">Shaowen Chen</na></co>
<co c="0"><na f="c/Chen:Xiaohong" pid="02/1438">Xiaohong Chen</na></co>
<co c="0" n="2"><na f="c/Chen_0006:Yuwen" pid="49/8346-6">Yuwen Chen 0006</na><na>Yu-Wen Chen 0006</na></co>
<co c="0"><na f="c/Chiang:Hsin=Tien" pid="241/2983">Hsin-Tien Chiang</na></co>
<co c="0"><na f="c/Choi:Yeonjong" pid="323/5236">Yeonjong Choi</na></co>
<co c="0"><na f="c/Cincarek:Tobias" pid="05/3057">Tobias Cincarek</na></co>
<co c="0"><na f="c/Clark:Rob" pid="23/6618">Rob Clark</na></co>
<co c="0"><na f="c/Cooper:Erica" pid="03/7184">Erica Cooper</na></co>
<co c="0"><na f="d/Das:Rohan_Kumar" pid="158/4101">Rohan Kumar Das</na></co>
<co c="0"><na f="d/Deguchi:Daisuke" pid="00/2175">Daisuke Deguchi</na></co>
<co c="0"><na f="d/Delgado:H=eacute=ctor" pid="21/10462">H&#233;ctor Delgado</na></co>
<co c="0"><na f="d/Demiroglu:Cenk" pid="21/2706">Cenk Demiroglu</na></co>
<co c="0"><na f="d/Do:Quoc_Truong" pid="166/6489">Quoc Truong Do</na></co>
<co c="0"><na f="d/Doi:Hironori" pid="68/4056">Hironori Doi</na></co>
<co c="0"><na f="d/Duan:Zhiyao" pid="04/6716">Zhiyao Duan</na></co>
<co c="0"><na f="d/Duh:Kevin" pid="58/3217">Kevin Duh</na></co>
<co c="0"><na f="d/Dutoit:Thierry" pid="79/2569">Thierry Dutoit</na></co>
<co c="0"><na f="e/Endo:Hajime" pid="226/1934">Hajime Endo</na></co>
<co c="0"><na f="e/Eshghi:Mohammad" pid="42/3681">Mohammad Eshghi</na></co>
<co c="0"><na f="e/Evans:Nicholas_W=_D=" pid="84/366">Nicholas W. D. Evans</na></co>
<co c="0"><na f="f/Fan_0002:Fan" pid="20/4226-2">Fan Fan 0002</na></co>
<co c="0"><na f="f/Feng:Jingyi" pid="174/4737">Jingyi Feng</na></co>
<co c="0"><na f="f/Fu:Szu=Wei" pid="160/0591">Szu-Wei Fu</na></co>
<co c="0"><na f="f/Fudaba:Hiroyuki" pid="16/1027">Hiroyuki Fudaba</na></co>
<co c="0"><na f="f/Fujimura:Takuya" pid="216/3624">Takuya Fujimura</na></co>
<co c="0"><na f="f/Fujita:Tomoki" pid="12/9858">Tomoki Fujita</na></co>
<co c="0"><na f="f/Fujita:Yusuke" pid="50/6352">Yusuke Fujita</na></co>
<co c="0"><na f="g/Gao_0040:Yuan" pid="76/2452-40">Yuan Gao 0040</na></co>
<co c="0"><na f="g/Gasic:Milica" pid="27/7520">Milica Gasic</na></co>
<co c="0"><na f="g/Gomez:Randy" pid="44/2122">Randy Gomez</na></co>
<co c="0"><na f="g/Goto:Masataka" pid="85/1753">Masataka Goto</na></co>
<co c="0"><na f="g/Govender:Avashna" pid="173/6404">Avashna Govender</na></co>
<co c="0"><na f="g/Guo:Jing" pid="95/5907">Jing Guo</na></co>
<co c="0"><na f="h/Halpern:Bence_Mark" pid="271/4266">Bence Mark Halpern</na></co>
<co c="0"><na f="h/Han:Jionghao" pid="378/4662">Jionghao Han</na></co>
<co c="0"><na f="h/Hashizume:Yuka" pid="334/0168">Yuka Hashizume</na></co>
<co c="0"><na f="h/Hata:Hideaki" pid="57/4288">Hideaki Hata</na></co>
<co c="0"><na f="h/Hatakoshi:Yuto" pid="04/11293">Yuto Hatakoshi</na></co>
<co c="0"><na f="h/Hattori:Kimihiro" pid="429/5419">Kimihiro Hattori</na></co>
<co c="0"><na f="h/Hattori:Nobuhiko" pid="30/10648">Nobuhiko Hattori</na></co>
<co c="0"><na f="h/Hayamizu:Satoru" pid="24/64">Satoru Hayamizu</na></co>
<co c="0"><na f="h/Hayashi:Tomoki" pid="82/8616">Tomoki Hayashi</na></co>
<co c="0"><na f="h/Hayashida:Chie" pid="142/7270">Chie Hayashida</na></co>
<co c="0"><na f="h/He:Jiajun" pid="205/5074">Jiajun He</na></co>
<co c="0"><na f="h/Heck:Michael" pid="120/6962">Michael Heck</na></co>
<co c="0"><na f="h/Henderson:Fergus" pid="h/FergusHenderson">Fergus Henderson</na></co>
<co c="0"><na f="h/Hikosaka:Shu" pid="277/3528">Shu Hikosaka</na></co>
<co c="0"><na f="h/Hirahara:Tatsuya" pid="24/4107">Tatsuya Hirahara</na></co>
<co c="0"><na f="h/Hirai:Toshio" pid="01/2055">Toshio Hirai</na></co>
<co c="0"><na f="h/Hiraoka:Takuya" pid="31/977">Takuya Hiraoka</na></co>
<co c="0"><na f="h/Hojo:Nobukatsu" pid="158/4110">Nobukatsu Hojo</na></co>
<co c="0"><na f="h/Hori:Takaaki" pid="46/3941">Takaaki Hori</na></co>
<co c="0"><na f="h/Horio:Kento" pid="226/1799">Kento Horio</na></co>
<co c="0"><na f="h/Hsu:Wei=Ning" pid="160/9923">Wei-Ning Hsu</na></co>
<co c="0"><na f="h/Hu:Cheng=Hung" pid="156/0822">Cheng-Hung Hu</na></co>
<co c="0"><na f="h/Hu:Desheng" pid="02/4151">Desheng Hu</na></co>
<co c="0"><na f="h/Hu:Xinhui" pid="04/1555">Xinhui Hu</na></co>
<co c="0"><na f="h/Hu:Yih=Chun" pid="51/6805">Yih-Chun Hu</na></co>
<co c="0"><na f="h/Huang:Wen=Chin" pid="225/7821">Wen-Chin Huang</na></co>
<co c="0"><na f="h/Hwang:Hsin=Te" pid="85/2787">Hsin-Te Hwang</na></co>
<co c="0"><na f="i/Ijima:Yusuke" pid="67/8052">Yusuke Ijima</na></co>
<co c="0"><na f="i/Ikeda:Yukichi" pid="213/0653">Yukichi Ikeda</na></co>
<co c="0"><na f="i/Ilham:Faiz" pid="175/8915">Faiz Ilham</na></co>
<co c="0"><na f="i/Imamura:Takehiro" pid="398/1356">Takehiro Imamura</na></co>
<co c="0"><na f="i/Imoto:Keisuke" pid="140/2806">Keisuke Imoto</na></co>
<co c="0"><na f="i/Inoue:Katsuki" pid="214/2329">Katsuki Inoue</na></co>
<co c="0"><na f="i/Inukai:Tatsuo" pid="178/0956">Tatsuo Inukai</na></co>
<co c="0"><na f="i/Irino:Toshio" pid="72/3903">Toshio Irino</na></co>
<co c="0"><na f="i/Ishii:Shunta" pid="98/11427">Shunta Ishii</na></co>
<co c="0"><na f="i/Isotani:Ryosuke" pid="36/6317">Ryosuke Isotani</na></co>
<co c="0"><na f="i/Ito:Ryuya" pid="220/9537">Ryuya Ito</na></co>
<co c="0"><na f="i/Itoi:Miyuki" pid="140/1975">Miyuki Itoi</na></co>
<co c="0"><na f="i/Iwahashi:Naoto" pid="22/2156">Naoto Iwahashi</na></co>
<co c="0"><na f="i/Iwami:Yohei" pid="63/9283">Yohei Iwami</na></co>
<co c="0"><na f="i/Iwasaka:Hidemi" pid="160/4287">Hidemi Iwasaka</na></co>
<co c="0"><na f="j/Jang:Jyh=Shing_Roger" pid="81/829">Jyh-Shing Roger Jang</na></co>
<co c="0"><na f="j/Jia:Ye" pid="217/2520">Ye Jia</na></co>
<co c="0"><na f="j/Jiang_0006:Yuan" pid="02/393-6">Yuan Jiang 0006</na></co>
<co c="0"><na f="j/Jinbo:Nozomi" pid="158/4125">Nozomi Jinbo</na></co>
<co c="0"><na f="j/Juvela:Lauri" pid="158/4213">Lauri Juvela</na></co>
<co c="0"><na f="k/Kain:Alexander" pid="19/4966">Alexander Kain</na></co>
<co c="0"><na f="k/Kaino:Tomomi" pid="25/9192">Tomomi Kaino</na></co>
<co c="0"><na f="k/Kameoka:Hirokazu" pid="97/941">Hirokazu Kameoka</na></co>
<co c="0"><na f="k/Kamiyama:Hosana" pid="212/6389">Hosana Kamiyama</na></co>
<co c="0"><na f="k/Kaneda:Takashi" pid="218/9267">Takashi Kaneda</na></co>
<co c="0"><na f="k/Kaneko:Masataka" pid="95/1951">Masataka Kaneko</na></co>
<co c="0"><na f="k/Kaneko:Takuhiro" pid="119/1623">Takuhiro Kaneko</na></co>
<co c="0"><na f="k/Kano:Takatomo" pid="140/2697">Takatomo Kano</na></co>
<co c="0"><na f="k/Kashioka:Hideki" pid="92/6978">Hideki Kashioka</na></co>
<co c="0"><na f="k/Kawahara:Hideki" pid="84/3249">Hideki Kawahara</na></co>
<co c="0"><na f="k/Kawai:Hisashi" pid="32/4341">Hisashi Kawai</na></co>
<co c="0"><na f="k/Kawamura:Masaya" pid="192/7186">Masaya Kawamura</na></co>
<co c="0"><na f="k/Kawanami:Hiromichi" pid="02/6430">Hiromichi Kawanami</na></co>
<co c="0"><na f="k/Keizer:Simon" pid="80/6099">Simon Keizer</na></co>
<co c="0"><na f="k/Khodabakhsh_0001:Ali" pid="149/9670">Ali Khodabakhsh 0001</na></co>
<co c="0"><na f="k/Kilgour:Kevin" pid="70/11337">Kevin Kilgour</na></co>
<co c="0"><na f="k/Kim:Sehun" pid="47/4994">Sehun Kim</na></co>
<co c="0"><na f="k/King_0001:Simon" pid="68/2005">Simon King 0001</na></co>
<co c="0"><na f="k/Kinnunen:Tomi" pid="94/4754">Tomi Kinnunen</na></co>
<co c="0"><na f="k/Kiso:Tetsuo" pid="97/9767">Tetsuo Kiso</na></co>
<co c="0"><na f="k/Kitaoka:Norihide" pid="07/6964">Norihide Kitaoka</na></co>
<co c="0"><na f="k/Kizuki:Hideaki" pid="175/8818">Hideaki Kizuki</na></co>
<co c="0"><na f="k/Kobashikawa:Satoshi" pid="09/3769">Satoshi Kobashikawa</na></co>
<co c="0"><na f="k/Kobayashi:Kazuhiro" pid="70/8239">Kazuhiro Kobayashi</na></co>
<co c="0"><na f="k/Kobayashi:Takao" pid="88/1648">Takao Kobayashi</na></co>
<co c="0"><na f="k/Komatsu:Tatsuya" pid="136/5113">Tatsuya Komatsu</na></co>
<co c="0"><na f="k/Kondo:Reishi" pid="180/2712">Reishi Kondo</na></co>
<co c="0"><na f="k/Koriyama:Tomoki" pid="44/9232">Tomoki Koriyama</na></co>
<co c="0"><na f="k/Koto:Fajri" pid="160/0019">Fajri Koto</na></co>
<co c="0"><na f="k/Ku:Pin=Jui" pid="289/7336">Pin-Jui Ku</na></co>
<co c="0"><na f="k/Kubo:Kazutaka" pid="214/2335">Kazutaka Kubo</na></co>
<co c="0"><na f="k/Kubo:Keigo" pid="124/9006">Keigo Kubo</na></co>
<co c="0"><na f="k/Kurita:Yusuke" pid="131/8830">Yusuke Kurita</na></co>
<co c="0"><na f="k/Kuroyanagi:Ibuki" pid="294/8499">Ibuki Kuroyanagi</na></co>
<co c="0"><na f="l/Lee:Akinobu" pid="85/6462">Akinobu Lee</na></co>
<co c="0"><na f="l/Lee:Hung=Shin" pid="13/8052">Hung-Shin Lee</na></co>
<co c="0" n="2"><na f="l/Lee:Hung=yi" pid="81/8056">Hung-yi Lee</na><na>Hung-Yi Lee</na></co>
<co c="0" n="2"><na f="l/Lee:Kong=Aik" pid="35/4621">Kong-Aik Lee</na><na>Kong Aik Lee</na></co>
<co c="0"><na f="l/Leon:Phillip_L=_De" pid="15/4307">Phillip L. De Leon</na></co>
<co c="0"><na f="l/Li:Fengji" pid="372/9374">Fengji Li</na></co>
<co c="0"><na f="l/Li_0001:Haizhou" pid="36/4118">Haizhou Li 0001</na></co>
<co c="0"><na f="l/Li:Junjie" pid="83/5144">Junjie Li</na></co>
<co c="0"><na f="l/Li_0063:Li" pid="53/2189-63">Li Li 0063</na></co>
<co c="0"><na f="l/Li_0001:Xingfeng" pid="63/9121-1">Xingfeng Li 0001</na></co>
<co c="0"><na f="l/Li:Xinjian" pid="69/8079">Xinjian Li</na></co>
<co c="0"><na f="l/Li:Yongwei" pid="172/4495">Yongwei Li</na></co>
<co c="0" n="2"><na f="l/Ling:Zhen=Hua" pid="70/5210">Zhen-Hua Ling</na><na>Zhenhua Ling</na></co>
<co c="0"><na f="l/Liou:Yi=Syuan" pid="301/7716">Yi-Syuan Liou</na></co>
<co c="0"><na f="l/Liu:Cheng" pid="15/2288">Cheng Liu</na></co>
<co c="0"><na f="l/Liu:Ching=Feng" pid="196/4607">Ching-Feng Liu</na></co>
<co c="0"><na f="l/Liu:Li=Juan" pid="93/6344">Li-Juan Liu</na></co>
<co c="0"><na f="l/Liu:Songxiang" pid="226/2031">Songxiang Liu</na></co>
<co c="0"><na f="l/Liu:Tao" pid="43/656">Tao Liu</na></co>
<co c="0"><na f="l/Liu:Zeyan" pid="284/4048">Zeyan Liu</na></co>
<co c="0"><na f="l/Livescu:Karen" pid="51/2464">Karen Livescu</na></co>
<co c="0"><na f="l/Lo:Chen=Chou" pid="234/6665">Chen-Chou Lo</na></co>
<co c="0"><na f="l/Loevenbruck:H=eacute=l=egrave=ne" pid="89/3110">H&#233;l&#232;ne Loevenbruck</na></co>
<co c="0"><na f="l/Lorenzo=Trueba:Jaime" pid="124/9038">Jaime Lorenzo-Trueba</na></co>
<co c="0"><na f="l/Lu:Jinlin" pid="59/9295">Jinlin Lu</na></co>
<co c="0"><na f="l/Luan:Shuming" pid="331/6661">Shuming Luan</na></co>
<co c="0"><na f="l/Lubis:Nurul" pid="160/8114">Nurul Lubis</na></co>
<co c="0"><na f="l/Luo:Shang=Bao" pid="234/6347">Shang-Bao Luo</na></co>
<co c="0"><na f="m/Ma:Ding" pid="30/6718">Ding Ma</na></co>
<co c="0"><na f="m/Maguer:S=eacute=bastien_Le" pid="47/10649">S&#233;bastien Le Maguer</na></co>
<co c="0"><na f="m/Maia:Ranniery" pid="15/7825">Ranniery Maia</na></co>
<co c="0"><na f="m/Mairesse:Fran=ccedil=ois" pid="m/FrancoisMairesse">Fran&#231;ois Mairesse</na></co>
<co c="0"><na f="m/Maki:Hayato" pid="166/6698">Hayato Maki</na></co>
<co c="0"><na f="m/Makino:Shoji" pid="31/6801">Shoji Makino</na></co>
<co c="0"><na f="m/Mashimo:Mikiko" pid="68/9241">Mikiko Mashimo</na></co>
<co c="0"><na f="m/Masuda:Tsuyoshi" pid="11/1981">Tsuyoshi Masuda</na></co>
<co c="0"><na f="m/Masumura:Ryo" pid="08/10650">Ryo Masumura</na></co>
<co c="0"><na f="m/Matrouf:Driss" pid="53/1360">Driss Matrouf</na></co>
<co c="0"><na f="m/Matsubara:Keisuke" pid="157/7171">Keisuke Matsubara</na></co>
<co c="0"><na f="m/Matsumiya:Sho" pid="146/3942">Sho Matsumiya</na></co>
<co c="0"><na f="m/Matsumoto_0001:Yuji" pid="11/4619">Yuji Matsumoto 0001</na></co>
<co c="0"><na f="m/Matsunaga:Noriyuki" pid="265/6508">Noriyuki Matsunaga</na></co>
<co c="0"><na f="m/Mi:Jinyi" pid="387/9787">Jinyi Mi</na></co>
<co c="0"><na f="m/Mieno:Takashi" pid="173/6560">Takashi Mieno</na></co>
<co c="0"><na f="m/Miura:Akiva" pid="166/1748">Akiva Miura</na></co>
<co c="0"><na f="m/Miyaji:Hikari" pid="429/5349">Hikari Miyaji</na></co>
<co c="0"><na f="m/Miyamoto:Daisuke" pid="24/2856">Daisuke Miyamoto</na></co>
<co c="0"><na f="m/Miyashita:Atsushi" pid="357/2546">Atsushi Miyashita</na></co>
<co c="0"><na f="m/Miyazaki:Koichi" pid="27/9637">Koichi Miyazaki</na></co>
<co c="0"><na f="m/Miyazaki:Ryoichi" pid="53/10650">Ryoichi Miyazaki</na></co>
<co c="0"><na f="m/Mizukami:Masahiro" pid="169/3173">Masahiro Mizukami</na></co>
<co c="0"><na f="m/Mohr:Christian" pid="60/7921">Christian Mohr</na></co>
<co c="0"><na f="m/Moinet:Alexis" pid="19/7269">Alexis Moinet</na></co>
<co c="0"><na f="m/Moriguchi:Takuto" pid="140/2835">Takuto Moriguchi</na></co>
<co c="0"><na f="m/Morikawa:Kazuho" pid="214/2225">Kazuho Morikawa</na></co>
<co c="0"><na f="m/Morise:Masanori" pid="75/6649">Masanori Morise</na></co>
<co c="0"><na f="m/Moriya:Takafumi" pid="175/8811">Takafumi Moriya</na></co>
<co c="0"><na f="m/Muramatsu:Takashi" pid="40/9234">Takashi Muramatsu</na></co>
<co c="0"><na f="m/Murata:Masato" pid="332/1313">Masato Murata</na></co>
<co c="0"><na f="m/Mushika:Koji" pid="252/4963">Koji Mushika</na></co>
<co c="0"><na f="n/Nagai:Takayuki" pid="09/1646">Takayuki Nagai</na></co>
<co c="0"><na f="n/Nagino:Goshu" pid="18/910">Goshu Nagino</na></co>
<co c="0"><na f="n/Nakagiri:Mikihiro" pid="54/9262">Mikihiro Nakagiri</na></co>
<co c="0"><na f="n/Nakajima:Yoshitaka" pid="75/5531">Yoshitaka Nakajima</na></co>
<co c="0"><na f="n/Nakamura:Keigo" pid="31/8013">Keigo Nakamura</na></co>
<co c="0"><na f="n/Nakamura:Kenichi" pid="35/7736">Kenichi Nakamura</na></co>
<co c="0"><na f="n/Nakamura:Masaru" pid="80/1615">Masaru Nakamura</na></co>
<co c="0"><na f="n/Nakamura_0001:Satoshi" pid="57/1548-1">Satoshi Nakamura 0001</na></co>
<co c="0"><na f="n/Nakamura:Shun" pid="79/7819">Shun Nakamura</na></co>
<co c="0"><na f="n/Nakamura:Tomoaki" pid="63/4486">Tomoaki Nakamura</na></co>
<co c="0"><na f="n/Nakano:Tomoyasu" pid="49/5220">Tomoyasu Nakano</na></co>
<co c="0"><na f="n/Nakata:Yuuto" pid="429/5477">Yuuto Nakata</na></co>
<co c="0"><na f="n/Nakatani:Hikaru" pid="278/9159">Hikaru Nakatani</na></co>
<co c="0"><na f="n/Nankaku:Yoshihiko" pid="90/302">Yoshihiko Nankaku</na></co>
<co c="0"><na f="n/Nautsch:Andreas" pid="49/10647">Andreas Nautsch</na></co>
<co c="0"><na f="n/Negoro:Hideki" pid="160/4279">Hideki Negoro</na></co>
<co c="0"><na f="n/Neubig:Graham" pid="03/8155">Graham Neubig</na></co>
<co c="0"><na f="n/Nguyen:The_Tung" pid="03/7981">The Tung Nguyen</na></co>
<co c="0"><na f="n/Ni:Jinfu" pid="37/2394">Jinfu Ni</na></co>
<co c="0"><na f="n/Nio:Lasguido" pid="146/8261">Lasguido Nio</na></co>
<co c="0"><na f="n/Nishida:Masafumi" pid="28/2662">Masafumi Nishida</na></co>
<co c="0"><na f="n/Nishizawa:Kaito" pid="409/3900">Kaito Nishizawa</na></co>
<co c="0"><na f="n/Nishizawa:Nobuyuki" pid="95/7824">Nobuyuki Nishizawa</na></co>
<co c="0"><na f="n/Nisimura:Ryuichi" pid="75/2915">Ryuichi Nisimura</na></co>
<co c="0"><na f="n/Niu:Haijun" pid="36/2169">Haijun Niu</na></co>
<co c="0"><na f="n/Niwa:Kiseki" pid="429/5304">Kiseki Niwa</na></co>
<co c="0"><na f="n/Nomura:Toshio" pid="43/4382">Toshio Nomura</na></co>
<co c="0"><na f="n/Nose:Takashi" pid="83/3783">Takashi Nose</na></co>
<co c="0"><na f="o/Oda:Yusuke" pid="148/4505">Yusuke Oda</na></co>
<co c="0"><na f="o/Odagaki:Yu" pid="159/9930">Yu Odagaki</na></co>
<co c="0"><na f="o/Ogita:Kenichi" pid="422/2585">Kenichi Ogita</na></co>
<co c="0"><na f="o/Ogura:Tadashi" pid="243/3880">Tadashi Ogura</na></co>
<co c="0" n="2"><na f="o/Ohgushi:Masaya" pid="140/2857">Masaya Ohgushi</na><na>Masaya Ogushi</na></co>
<co c="0"><na f="o/Ohira:Shigeki" pid="36/2324">Shigeki Ohira</na></co>
<co c="0"><na f="o/Ohta:Kumi" pid="06/9233">Kumi Ohta</na></co>
<co c="0"><na f="o/Ohtani:Yamato" pid="34/8763">Yamato Ohtani</na></co>
<co c="0"><na f="o/Okada:Hiroyuki" pid="24/2307">Hiroyuki Okada</na></co>
<co c="0"><na f="o/Okamoto:Kosuke" pid="220/9460">Kosuke Okamoto</na></co>
<co c="0"><na f="o/Okamoto:Takuma" pid="132/9091">Takuma Okamoto</na></co>
<co c="0"><na f="o/Omori:Takashi" pid="21/4085">Takashi Omori</na></co>
<co c="0"><na f="o/Onuma:Kai" pid="252/5717">Kai Onuma</na></co>
<co c="0"><na f="o/Oshima:Yuji" pid="130/8033">Yuji Oshima</na></co>
<co c="0"><na f="o/Otani:Makoto" pid="97/2873">Makoto Otani</na></co>
<co c="0"><na f="o/Oura:Keiichiro" pid="22/7077">Keiichiro Oura</na></co>
<co c="0"><na f="p/Peng:Yu=Huai" pid="214/2303">Yu-Huai Peng</na></co>
<co c="0"><na f="p/Purwarianti:Ayu" pid="15/2527">Ayu Purwarianti</na></co>
<co c="0"><na f="q/Qian:Zhaopeng" pid="193/6214">Zhaopeng Qian</na></co>
<co c="0"><na f="q/Qin:Yong" pid="20/4298">Yong Qin</na></co>
<co c="0"><na f="r/Rebernik:Teja" pid="317/5018">Teja Rebernik</na></co>
<co c="0"><na f="r/Renals:Steve" pid="33/3792">Steve Renals</na></co>
<co c="0"><na f="r/Ronanki:Srikanth" pid="180/2841">Srikanth Ronanki</na></co>
<co c="0"><na f="r/Rosec:Olivier" pid="11/6417">Olivier Rosec</na></co>
<co c="0"><na f="r/Roux:Jonathan_Le" pid="36/4575">Jonathan Le Roux</na></co>
<co c="0"><na f="s/Saam:Christian" pid="138/0300">Christian Saam</na></co>
<co c="0"><na f="s/Sahidullah:Md=" pid="93/9671">Md. Sahidullah</na></co>
<co c="0"><na f="s/Saito:Daisuke" pid="17/7825">Daisuke Saito</na></co>
<co c="0"><na f="s/Sakai:Shinichi" pid="04/7823">Shinichi Sakai</na></co>
<co c="0"><na f="s/Sakai:Shinsuke" pid="57/448">Shinsuke Sakai</na></co>
<co c="0"><na f="s/Sakakibara:Ken=Ichi" pid="11/9231">Ken-Ichi Sakakibara</na></co>
<co c="0"><na f="s/Sakti:Sakriani" pid="71/3717">Sakriani Sakti</na></co>
<co c="0"><na f="s/Sano:Motoaki" pid="140/2888">Motoaki Sano</na></co>
<co c="0"><na f="s/Saruwatari:Hiroshi" pid="87/629">Hiroshi Saruwatari</na></co>
<co c="0"><na f="s/Sasakura:Takafumi" pid="159/9937">Takafumi Sasakura</na></co>
<co c="0"><na f="s/Sato_0002:Hiroshi" pid="55/6900-2">Hiroshi Sato 0002</na></co>
<co c="0"><na f="s/Sawada:Keito" pid="429/5300">Keito Sawada</na></co>
<co c="0"><na f="s/Sawada:Naoki" pid="70/2093">Naoki Sawada</na></co>
<co c="0"><na f="s/Scharenborg:Odette" pid="87/1365">Odette Scharenborg</na></co>
<co c="0"><na f="s/Seiya:Shunya" pid="220/9281">Shunya Seiya</na></co>
<co c="0"><na f="s/Seki:Shogo" pid="194/1307">Shogo Seki</na></co>
<co c="0"><na f="s/Sekimoto:Hidehiko" pid="45/8053">Hidehiko Sekimoto</na></co>
<co c="0"><na f="s/Shen:Fei" pid="99/6100">Fei Shen</na></co>
<co c="0"><na f="s/Shi:Jiatong" pid="229/3529">Jiatong Shi</na></co>
<co c="0"><na f="s/Shi:Xiaohan" pid="89/3530">Xiaohan Shi</na></co>
<co c="0"><na f="s/Shiga:Yoshinori" pid="02/2978">Yoshinori Shiga</na></co>
<co c="0"><na f="s/Shikano:Kiyohiro" pid="50/2173">Kiyohiro Shikano</na></co>
<co c="0"><na f="s/Shimizu:Hiroaki" pid="44/5142">Hiroaki Shimizu</na></co>
<co c="0"><na f="s/Shimizu:Shota" pid="39/8014">Shota Shimizu</na></co>
<co c="0"><na f="s/Shimizu:Sota" pid="14/5586">Sota Shimizu</na></co>
<co c="0"><na f="s/Shimizu:Tohru" pid="08/1223">Tohru Shimizu</na></co>
<co c="0"><na f="s/Shindo:Hiroyuki" pid="97/8317">Hiroyuki Shindo</na></co>
<co c="0"><na f="s/Shindo:Izumi" pid="04/9250">Izumi Shindo</na></co>
<co c="0"><na f="s/Shiraishi:Tatsuya" pid="20/9286">Tatsuya Shiraishi</na></co>
<co c="0"><na f="s/Shozakai:Makoto" pid="23/6099">Makoto Shozakai</na></co>
<co c="0" n="2"><na f="s/Son:R=_J=_J=_H=_van" pid="34/4715">R. J. J. H. van Son</na><na>Rob J. J. H. van Son</na></co>
<co c="0"><na f="s/Sperber:Matthias" pid="79/9103">Matthias Sperber</na></co>
<co c="0"><na f="s/Steiner:Ingmar" pid="64/9231">Ingmar Steiner</na></co>
<co c="0"><na f="s/Stewart:Bryan" pid="124/9083">Bryan Stewart</na></co>
<co c="0"><na f="s/St=uuml=ker:Sebastian" pid="15/1807">Sebastian St&#252;ker</na></co>
<co c="0"><na f="s/Stylianou:Yannis" pid="14/1209">Yannis Stylianou</na></co>
<co c="0"><na f="s/Sugiura:Komei" pid="77/2654">Komei Sugiura</na></co>
<co c="0"><na f="s/Sugiyama:Kyoshiro" pid="154/5283">Kyoshiro Sugiyama</na></co>
<co c="0"><na f="t/Tachibana:Kentaro" pid="35/8056">Kentaro Tachibana</na></co>
<co c="0"><na f="t/Taga:Haruka" pid="306/9710">Haruka Taga</na></co>
<co c="0"><na f="t/Tajiri:Yusuke" pid="173/6532">Yusuke Tajiri</na></co>
<co c="0"><na f="t/Takada:Moe" pid="237/0010">Moe Takada</na></co>
<co c="0"><na f="t/Takamichi:Shinnosuke" pid="124/9221">Shinnosuke Takamichi</na></co>
<co c="0"><na f="t/Takashima:Ryoichi" pid="22/8758">Ryoichi Takashima</na></co>
<co c="0"><na f="t/Takeda:Kazuya" pid="38/3483">Kazuya Takeda</na></co>
<co c="0"><na f="t/Takiguchi:Tetsuya" pid="79/4485">Tetsuya Takiguchi</na></co>
<co c="0"><na f="t/Tamamori:Akira" pid="01/8760">Akira Tamamori</na></co>
<co c="0"><na f="t/Tamura:Satoshi" pid="32/4043">Satoshi Tamura</na></co>
<co c="0"><na f="t/Tan_0003:Xu" pid="96/10484-3">Xu Tan 0003</na></co>
<co c="0"><na f="t/Tan:Zheng=Hua" pid="39/4898">Zheng-Hua Tan</na></co>
<co c="0"><na f="t/Tanaka:Hiroki" pid="66/2476">Hiroki Tanaka</na></co>
<co c="0"><na f="t/Tanaka:Kou" pid="140/2727">Kou Tanaka</na></co>
<co c="0"><na f="t/Tang:Shaoqi" pid="429/5417">Shaoqi Tang</na></co>
<co c="0"><na f="t/Tang:Yuxun" pid="358/9271">Yuxun Tang</na></co>
<co c="0"><na f="t/Tani:Daisuke" pid="56/9234">Daisuke Tani</na></co>
<co c="0"><na f="t/Tanikawa:Ukyo" pid="200/0164">Ukyo Tanikawa</na></co>
<co c="0"><na f="t/Terashima:Ryo" pid="67/8451">Ryo Terashima</na></co>
<co c="0"><na f="t/Thomson:Blaise" pid="44/2850">Blaise Thomson</na></co>
<co c="0"><na f="t/Tian:Jingguang" pid="313/9922">Jingguang Tian</na></co>
<co c="0"><na f="t/Tian:Xiaohai" pid="153/0728">Xiaohai Tian</na></co>
<co c="0" n="2"><na f="t/Tienkamp:Thomas" pid="317/5237">Thomas Tienkamp</na><na>Thomas B. Tienkamp</na></co>
<co c="0"><na f="t/Tjandra:Andros" pid="166/6471">Andros Tjandra</na></co>
<co c="0"><na f="t/Tobing:Patrick_Lumban" pid="158/4210">Patrick Lumban Tobing</na></co>
<co c="0"><na f="t/Todisco:Massimiliano" pid="89/6736">Massimiliano Todisco</na></co>
<co c="0"><na f="t/Tokuda:Keiichi" pid="25/369">Keiichi Tokuda</na></co>
<co c="0"><na f="t/Toshniwal:Shubham" pid="160/4302">Shubham Toshniwal</na></co>
<co c="0"><na f="t/Tran:Viet=Anh" pid="69/8013">Viet-Anh Tran</na></co>
<co c="0"><na f="t/Tsai:Shu=Wei" pid="301/7776">Shu-Wei Tsai</na></co>
<co c="0"><na f="t/Tsao_0001:Yu" pid="66/7146-1">Yu Tsao 0001</na></co>
<co c="0"><na f="t/Tsunomori:Yuiko" pid="195/4988">Yuiko Tsunomori</na></co>
<co c="0"><na f="t/Tsuruta:Sakura" pid="160/0028">Sakura Tsuruta</na></co>
<co c="0"><na f="t/Tsuzaki:Minoru" pid="20/3844">Minoru Tsuzaki</na></co>
<co c="0"><na f="u/Unoki:Masashi" pid="73/1580">Masashi Unoki</na></co>
<co c="0"><na f="u/Uto:Yosuke" pid="91/8053">Yosuke Uto</na></co>
<co c="0"><na f="v/Vestman:Ville" pid="201/7499">Ville Vestman</na></co>
<co c="0"><na f="v/Vijayan:Karthika" pid="148/9726">Karthika Vijayan</na></co>
<co c="0"><na f="v/Villavicencio:Fernando" pid="64/2286">Fernando Villavicencio</na></co>
<co c="0"><na f="v/Violeta:Lester_Phillip" pid="304/2864">Lester Phillip Violeta</na></co>
<co c="0"><na f="v/Visscher:Sebastiaan_A=_H=_J=_de" pid="329/8363">Sebastiaan A. H. J. de Visscher</na></co>
<co c="0"><na f="v/Vu:Hoa_Trong" pid="117/4989">Hoa Trong Vu</na></co>
<co c="0"><na f="w/Waibel:Alex" pid="08/2456">Alex Waibel</na></co>
<co c="0"><na f="w/Wakabayashi:Yukoh" pid="158/4215">Yukoh Wakabayashi</na></co>
<co c="0"><na f="w/Wang:Hsin=Min" pid="28/5019">Hsin-Min Wang</na></co>
<co c="0"><na f="w/Wang:Hui" pid="39/721">Hui Wang</na></co>
<co c="0"><na f="w/Wang:Jiachen" pid="145/6290">Jiachen Wang</na></co>
<co c="0"><na f="w/Wang_0111:Li" pid="58/6810-111">Li Wang 0111</na></co>
<co c="0"><na f="w/Wang:Quan" pid="86/5728">Quan Wang</na></co>
<co c="0"><na f="w/Wang_0169:Rui" pid="06/2293-169">Rui Wang 0169</na></co>
<co c="0"><na f="w/Wang_0037:Xin" pid="10/5630-37">Xin Wang 0037</na></co>
<co c="0"><na f="w/Watanabe_0001:Shinji" pid="39/3245-1">Shinji Watanabe 0001</na></co>
<co c="0"><na f="w/Wester:Mirjam" pid="30/1879">Mirjam Wester</na></co>
<co c="0"><na f="w/Wieling_0001:Martijn" pid="35/2985">Martijn Wieling 0001</na></co>
<co c="0"><na f="w/Wilkinghoff:Kevin" pid="207/9559">Kevin Wilkinghoff</na></co>
<co c="0"><na f="w/Witjes:Max_J=_H=" pid="227/2349">Max J. H. Witjes</na></co>
<co c="0"><na f="w/Wu:Chia=Hua" pid="160/2718">Chia-Hua Wu</na></co>
<co c="0"><na f="w/Wu_0001:Chung=Hsien" pid="20/924">Chung-Hsien Wu 0001</na></co>
<co c="0"><na f="w/Wu:Yi=Chiao" pid="188/5943">Yi-Chiao Wu</na></co>
<co c="0"><na f="w/Wu:Yi=Jian" pid="91/6000">Yi-Jian Wu</na></co>
<co c="0"><na f="w/Wu_0001:Zhizheng" pid="29/8054-1">Zhizheng Wu 0001</na></co>
<co c="0"><na f="x/Xie:Chao" pid="04/5576">Chao Xie</na></co>
<co c="0"><na f="x/Xu:Shengyuan" pid="386/5432">Shengyuan Xu</na></co>
<co c="0"><na f="x/Xu:Xinkang" pid="277/3578">Xinkang Xu</na></co>
<co c="0"><na f="y/Yamagishi:Junichi" pid="87/3979">Junichi Yamagishi</na></co>
<co c="0"><na f="y/Yamamoto:Kenzo" pid="120/6772">Kenzo Yamamoto</na></co>
<co c="0"><na f="y/Yamamoto:Ryuichi" pid="81/3672">Ryuichi Yamamoto</na></co>
<co c="0"><na f="y/Yamane:Soichi" pid="180/2666">Soichi Yamane</na></co>
<co c="0"><na f="y/Yamashita:Haruki" pid="368/3515">Haruki Yamashita</na></co>
<co c="0"><na f="y/Yamauchi:Yuki" pid="139/5461">Yuki Yamauchi</na></co>
<co c="0"><na f="y/Yang:Shu=Wen" pid="246/0774">Shu-Wen Yang</na></co>
<co c="0"><na f="y/Yang:Zekun" pid="176/5718">Zekun Yang</na></co>
<co c="0"><na f="y/Yasuda:Yusuke" pid="228/9342">Yusuke Yasuda</na></co>
<co c="0"><na f="y/Yasuhara:Kazuki" pid="265/6506">Kazuki Yasuhara</na></co>
<co c="0"><na f="y/Yen:Ming=Chi" pid="155/7903">Ming-Chi Yen</na></co>
<co c="0"><na f="y/Yoneyama:Reo" pid="290/1755">Reo Yoneyama</na></co>
<co c="-1"><na f="y/Yoon:Dohyun" pid="429/5371">Dohyun Yoon</na></co>
<co c="0"><na f="y/Yoshida:Riki" pid="159/9947">Riki Yoshida</na></co>
<co c="0"><na f="y/Yoshimoto:Akifumi" pid="162/7852">Akifumi Yoshimoto</na></co>
<co c="0"><na f="y/Yoshimura:Takenori" pid="173/6383">Takenori Yoshimura</na></co>
<co c="0"><na f="y/Yoshino:Koichiro" pid="59/8883">Koichiro Yoshino</na></co>
<co c="0"><na f="y/Yoshioka:Daiki" pid="200/0472">Daiki Yoshioka</na></co>
<co c="0"><na f="y/Young:Steve_J=" pid="11/9311">Steve J. Young</na></co>
<co c="0"><na f="y/Yu:Cheng" pid="43/3340">Cheng Yu</na></co>
<co c="0"><na f="y/Yu_0004:Kai" pid="197/1322-4">Kai Yu 0004</na></co>
<co c="0"><na f="y/Yutani:Kaori" pid="25/8056">Kaori Yutani</na></co>
<co c="0"><na f="z/Zang:Yongyi" pid="348/9712">Yongyi Zang</na></co>
<co c="0"><na f="z/Zen:Heiga" pid="42/7014">Heiga Zen</na></co>
<co c="0"><na f="z/Zezario:Ryandhimas_E=" pid="199/7652">Ryandhimas E. Zezario</na></co>
<co c="0"><na f="z/Zhang:Jing=Xuan" pid="223/5831">Jing-Xuan Zhang</na></co>
<co c="0"><na f="z/Zhang:Shaochuan" pid="313/2304">Shaochuan Zhang</na></co>
<co c="0"><na f="z/Zhang:Xueyao" pid="237/9484">Xueyao Zhang</na></co>
<co c="0"><na f="z/Zhang_0001:You" pid="26/3166-1">You Zhang 0001</na></co>
<co c="0"><na f="z/Zhang_0033:Yu" pid="50/671-33">Yu Zhang 0033</na></co>
<co c="0" n="2"><na f="z/Zhao:Wen=Xiao" pid="77/8083">Wen-Xiao Zhao</na><na>Wenxiao Zhao</na></co>
<co c="0"><na f="z/Zhao_0006:Yi" pid="51/4138-6">Yi Zhao 0006</na></co>
<co c="0"><na f="z/Zhou:Jie" pid="00/5012">Jie Zhou</na></co>
<co c="0"><na f="z/Zhou_0024:Xiao" pid="267/2864-24">Xiao Zhou 0024</na></co>
</coauthors>
</dblpperson>

