|
#author("2026-05-25T15:34:01+00:00","default:spadmin","spadmin") [[ホーム/発表論文/2025]] #author("2026-05-25T15:34:51+00:00","default:spadmin","spadmin") * 発表論文 - 2025 [#j580b137] * 発表論文 - 2026 [#j580b137] #contents ** 論文誌 [#i5f82a69] + Yoshihiko Nankaku, Takato Fujimoto, Takenori Yoshimura, Shinji Takaki, Kei Hashimoto, Keiichiro Oura, and Keiichi Tokuda, "Deep Hidden Semi-Markov Model-Based Speech Synthesis," in IEEE Access, vol. 14, pp. 58495-58514, 2026. (Full paper peer reviewed) [[link>https://ieeexplore.ieee.org/document/11481063]] ** 研究会 [#r5474f12] + %%%田牧宏都%%%, 橋本佳, 南角吉彦, 徳田恵一,``評価情報を付与した音声・声質説明文ペアデータを用いた自然言語声質制御音声合成の検討,'' 音声言語情報処理研究会, vol. 2026-SLP-159, no. 46, pp. 1-6, 沖縄, 日本, 2026年3月. &publication(2026/20260303_TReport_SLP_Tamaki_Hiroto_paper.pdf, paper); &publication(2026/20260303_TReport_SLP_Tamaki_Hiroto_slide.pptx, slide); + %%%恩田将人%%%, 橋本佳, 南角吉彦, 徳田恵一,``ニューラルコーデック言語モデルに基づく音声プロンプトを利用したzero-shot声質変換,'' 音声言語情報処理研究会, vol. 2026-SLP-159, no. 45, pp. 1-7, 沖縄, 日本, 2026年3月. &publication(2026/20260303_TReport_SLP_Masato_Onda_paper.pdf, paper); &publication(2026/20260303_TReport_SLP_Masato_Onda_slide.pptx, slide); ** 全国大会 [#f48328c8] + %%%高瀬昴%%%, 西原美玖, 法野行哉, 橋本佳, 南角吉彦, 徳田恵一,``ソロ歌声データを用いたユニゾン歌声を生成可能なニューラルボコーダの学習,'' 日本音響学会2026年春季研究発表会, pp. 965-968, 東京, 日本, 2026年3月. &publication(2026/20260318_DConference_ASJS_Subaru_Takase_paper.pdf, paper); &publication(2026/20260318_DConference_ASJS_Subaru_Takase_slide.pptx, slide); &publication(2026/20260318_DConference_ASJS_Subaru_Takase_abst.pdf, abst); + %%%菊池遥斗%%%, 能勢隆, 林崎由, 小林清流, 橋本佳, 伊藤彰則,``アニメ調キャラクター顔画像に適するクロスモーダル音声合成のための再現性・親和性改善の検討,'' 日本音響学会2026年春季研究発表会, pp. 1109-1112, 東京, 日本, 2026年3月. &publication(2026/20260317_DConference_ASJS_Haruto_Kikuchi_paper.pdf, paper); + %%%石井信吾%%%, 橋本佳, 南角吉彦, 徳田恵一,``共役事後分布に基づいた深層隠れセミマルコフモデルに基づく音声合成,'' 日本音響学会2026年春季研究発表会, pp. 915-918, 東京, 日本, 2026年3月. &publication(2026/20260317_DConference_ASJS_Shingo_Ishii_paper.pdf, paper); &publication(2026/20260317_DConference_ASJS_Shingo_Ishii_slide.pptx, slide); &publication(2026/20260317_DConference_ASJS_Shingo_Ishii_abst.pdf, abst); + %%%淺野秀輝%%%, 橋本佳, 南角吉彦, 徳田恵一,``深層隠れセミマルコフモデルに基づいたEnd-to-Endテキスト音声合成,'' 日本音響学会2026年春季研究発表会, pp. 911-914, 東京, 日本, 2026年3月. &publication(2026/20260317_DConference_ASJS_Hideki_Asano_paper.pdf, paper); &publication(2026/20260317_DConference_ASJS_Hideki_Asano_slide.pptx, slide); &publication(2026/20260317_DConference_ASJS_Hideki_Asano_abst.pdf, abst); ** 学位論文 [#qdc8f7d1] + %%%安永真那斗%%%, ``顔画像を用いたクロスモーダル音声合成のための拡散モデルに基づく話者埋め込み推定,'' 卒業論文, 名古屋工業大学, 2026年2月. //%%%Manato Yasunaga%%%, ``Diffusion model-based speaker embedding estimation for cross-modal speech synthesis using facial images,'' Bachelor thesis, Nagoya Institute of Technology, February, 2026. &publication(2026/20260217_Thesis_Bachelor_Manato_Yasunaga_paper.pdf, paper); &publication(2026/20260217_Thesis_Bachelor_Manato_Yasunaga_slide.pptx, slide); &publication(2026/20260217_Thesis_Bachelor_Manato_Yasunaga_abst.pdf, abst); + %%%服部直起%%%, ``リアルタイム声質変換のための自己発声骨導音のアクティブキャンセリングを導入した聴覚フィードバック制御,'' 卒業論文, 名古屋工業大学, 2026年2月. //%%%Naoki Hattori%%%, ``Auditory feedback control incorporating active canceling of self-voiced bone-conducted sounds for real-time voice conversion,'' Bachelor thesis, Nagoya Institute of Technology, February, 2026. &publication(2026/20260217_Thesis_Bachelor_Naoki_Hattori_paper.pdf, paper); &publication(2026/20260217_Thesis_Bachelor_Naoki_Hattori_slide.pptx, slide); &publication(2026/20260217_Thesis_Bachelor_Naoki_Hattori_abst.pdf, abst); // ikeda-kun + %%%池田陸人%%%, ``触覚・音声のクロスモーダルインターフェースのための感情・話者認識,'' 卒業論文, 名古屋工業大学, 2026年2月. //%%%Rikuto Ikeda%%%, ``Emotion and speaker recognition for a cross-modal haptic audio interface,'' Bachelor thesis,Nagoya Institute of Technology, Febuary, 2026. &publication(2026/20260217_Thesis_Bachelor_Rikuto_Ikeda_paper.pdf, paper); &publication(2026/20260217_Thesis_Bachelor_Rikuto_Ikeda_slide.pptx, slide); &publication(2026/20260217_Thesis_Bachelor_Rikuto_Ikeda_abst.pdf, abst); + %%%淺野秀輝%%%, ``深層隠れセミマルコフモデルと波形生成モデルのEnd-to-End学習によるテキスト音声合成,'' 卒業論文, 名古屋工業大学, 2026年2月. //%%%Hideki Asano%%%, ``Text-to-speech synthesis via end-to-end learning with deep hidden semi-Markov models and waveform generative models,'' Bachelor thesis, Nagoya Institute of Technology, February, 2026. &publication(2026/20260217_Thesis_Bachelor_Hideki_Asano_paper.pdf, paper); &publication(2026/20260217_Thesis_Bachelor_Hideki_Asano_slide.pptx, slide); &publication(2026/20260217_Thesis_Bachelor_Hideki_Asano_abst.pdf, abst); // ishii-kun + %%%石井信吾%%%, ``共役事後分布を用いた深層隠れセミマルコフモデルに基づく音声合成,'' 卒業論文, 名古屋工業大学, 2026年2月. //%%%Shingo Ishii%%%, ``Speech synthesis based on deep hidden semi-Markov models using conjugate posterior distributions,'' Bachelor thesis, Nagoya Institute of Technology, February, 2026. &publication(2026/20260217_Thesis_Bachelor_Shingo_Ishii_paper.pdf, paper); &publication(2026/20260217_Thesis_Bachelor_Shingo_Ishii_slide.pptx, slide); &publication(2026/20260217_Thesis_Bachelor_Shingo_Ishii_abst.pdf, abst); + %%%宮下翔%%%, ``周期信号の位相情報を用いたフレーム駆動型ニューラルボコーダの不特定話者化,'' 卒業論文, 名古屋工業大学, 2026年2月. //%%%Sho Miyashita%%%, ``Speaker-independent frame-level neural vocoder utilizing phase information of periodic signals,'' Bachelor thesis, Nagoya Institute of Technology, February, 2026. &publication(2026/20260217_Thesis_Bachelor_Sho_Miyashita_paper.pdf, paper); &publication(2026/20260217_Thesis_Bachelor_Sho_Miyashita_slide.pptx, slide); &publication(2026/20260217_Thesis_Bachelor_Sho_Miyashita_abst.pdf, abst); // takase-kun + %%%高瀬昴%%%, ``疑似ユニゾン歌声データセットの構築とユニゾン歌声を生成可能なニューラルボコーダの学習,'' 卒業論文, 名古屋工業大学, 2026年2月. //%%%Subaru Takase%%%, ``Construction of pseudo unison singing voice dataset and training of neural vocoder for unison singing generation,'' Bachelor thesis, Nagoya Institute of Technology, February, 2026. &publication(2026/20260217_Thesis_Bachelor_Subaru_Takase_paper.pdf, paper); &publication(2026/20260217_Thesis_Bachelor_Subaru_Takase_slide.pptx, slide); &publication(2026/20260217_Thesis_Bachelor_Subaru_Takase_abst.pdf, abst); + %%%杉浦篤志%%%, ``自己教師あり学習モデルを用いたWhisper-to-Normal音声変換におけるモデル学習法の検討,'' 卒業論文, 名古屋工業大学, 2026年2月. //%%%Atsushi Sugiura%%%, ``A study on training methods for whisper-to-normal speech conversion using self-supervised learning models,'' Bachelor thesis, Nagoya Institute of Technology, February, 2026. &publication(2026/20260217_Thesis_Bachelor_Atsushi_Sugiura_paper.pdf, paper); &publication(2026/20260217_Thesis_Bachelor_Atsushi_Sugiura_slide.pptx, slide); &publication(2026/20260217_Thesis_Bachelor_Atsushi_Sugiura_abst.pdf, abst); + %%%篠原海斗%%%, ``離散トークンに基づくzero-shot音声合成のための声質変換と話速変更による学習データ拡張法,'' 卒業論文, 名古屋工業大学, 2026年2月. //%%%Kaito Shinohara%%%, ``A data augmentation method for zero-shot speech synthesis using voice conversion and speaking rate modification based on discrete tokens,'' Bachelor thesis, Nagoya Institute of Technology, February, 2026. &publication(2026/20260217_Thesis_Bachelor_Kaito_Shinohara_paper.pdf, paper); &publication(2026/20260217_Thesis_Bachelor_Kaito_Shinohara_slide.pptx, slide); &publication(2026/20260217_Thesis_Bachelor_Kaito_Shinohara_abst.pdf, abst); + %%%木村駿%%%, ``話者照合における発話単位アテンション機構に基づく話者埋め込み抽出法の検討,'' 卒業論文, 名古屋工業大学, 2026年2月. //%%%Shun Kimura%%%, ``A study on speaker embedding extraction methods based on utterance-level attention mechanisms for speaker verification,'' Bachelor thesis, Nagoya Institute of Technology, February, 2026. &publication(2026/20260217_Thesis_Bachelor_Shun_Kimura_paper.pdf, paper); &publication(2026/20260217_Thesis_Bachelor_Shun_Kimura_slide.pptx, slide); &publication(2026/20260217_Thesis_Bachelor_Shun_Kimura_abst.pdf, abst); + %%%岩佐慧晶%%%, ``ソースフィルタ型ニューラルボコーダにおける破裂音モジュールの導入,'' 卒業論文, 名古屋工業大学, 2026年2月. //%%%Keisho Iwasa%%%, ``Introducing a plosive module in a source-filter neural vocoder,'' Bachelor thesis, Nagoya Institute of Technology, February, 2026. &publication(2026/20260217_Thesis_Bachelor_Keisho_Iwasa_paper.pdf, paper); &publication(2026/20260217_Thesis_Bachelor_Keisho_Iwasa_slide.pptx, slide); &publication(2026/20260217_Thesis_Bachelor_Keisho_Iwasa_abst.pdf, abst); // yasuda-sempai + %%%安田周平%%%, ``音素識別損失を導入した深層隠れセミマルコフモデルに基づく音声合成,'' 修士論文, 名古屋工業大学, 2026年2月. //%%%Shuhei Yasuda%%%, ``Speech synthesis based on deep hidden semi-Markov models incorporating phoneme classification loss,'' Master thesis, Nagoya Institute of Technology, February, 2026. &publication(2026/20260210_Thesis_Master_Shuhei_Yasuda_paper.pdf, paper); &publication(2026/20260210_Thesis_Master_Shuhei_Yasuda_slide.pptx, slide); &publication(2026/20260210_Thesis_Master_Shuhei_Yasuda_abst.pdf, abst); // miyake-sempai + %%%三宅恭平%%%, ``複数時間単位の変動を考慮した深層隠れセミマルコフモデルに基づく音声合成,'' 修士論文, 名古屋工業大学, 2026年2月. //%%%Kyohei Miyake%%%, ``Speech synthesis based on deep hidden semi-Markov models incorporating latent variables at multiple temporal resolutions,'' Master thesis, Nagoya Institute of Technology, February, 2026. &publication(2026/20260210_Thesis_Master_Kyohei_Miyake_paper.pdf, paper); &publication(2026/20260210_Thesis_Master_Kyohei_Miyake_slide.pptx, slide); &publication(2026/20260210_Thesis_Master_Kyohei_Miyake_abst.pdf, abst); // hirachi-sempai + %%%平地巧%%%, ``階層型変分オートエンコーダに基づくニューラルボコーダに対する敵対的学習の導入,'' 修士論文, 名古屋工業大学, 2026年2月. //%%%Takumi Hirachi%%%, ``Applying adversarial learning to hierarchical variational autoencoder-based neural vocoders,'' Master thesis, Nagoya Institute of Technology, February, 2026. &publication(2026/20260210_Thesis_Master_Takumi_Hirachi_paper.pdf, paper); &publication(2026/20260210_Thesis_Master_Takumi_Hirachi_slide.pptx, slide); &publication(2026/20260210_Thesis_Master_Takumi_Hirachi_abst.pdf, abst); // nishida-sempai + %%%西田開登%%%, ``End-to-End音声合成のための自己回帰構造を導入した階層化生成モデル,'' 修士論文, 名古屋工業大学, 2026年2月. //%%%Kaito Nishida%%%, ``Hierarchical generative models with autoregressive structures for end-to-end speech synthesis,'' Master thesis, Nagoya Institute of Technology, February, 2026. &publication(2026/20260210_Thesis_Master_Kaito_Nishida_paper.pdf, paper); &publication(2026/20260210_Thesis_Master_Kaito_Nishida_slide.pptx, slide); &publication(2026/20260210_Thesis_Master_Kaito_Nishida_abst.pdf, abst); // ida-sempai + %%%井田侑作%%%, ``ゼロショット音声生成タスクのための話者表現の自己教師有り学習,'' 修士論文, 名古屋工業大学, 2026年2月. //%%%Yusaku Ida%%%, ``Self-supervised learning of speaker representation for zero-shot speech generation tasks,'' Master thesis, Nagoya Institute of Technology, February, 2026. &publication(2026/20260210_Thesis_Master_Yusaku_Ida_paper.pdf, paper); &publication(2026/20260210_Thesis_Master_Yusaku_Ida_slide.pptx, slide); &publication(2026/20260210_Thesis_Master_Yusaku_Ida_abst.pdf, abst); // iida-sempai + %%%飯田諒%%%, ``双方向の自己回帰構造を組み込んだ深層隠れセミマルコフモデルに基づく音声合成,'' 修士論文, 名古屋工業大学, 2026年2月. //%%%Ryo Iida%%%, ``Speech synthesis based on deep hidden semi-Markov models incorporating bidirectional autoregressive structures,'' Master thesis, Nagoya Institute of Technology, February, 2026. &publication(2026/20260210_Thesis_Master_Ryo_Iida_paper.pdf, paper); &publication(2026/20260210_Thesis_Master_Ryo_Iida_slide.pptx, slide); &publication(2026/20260210_Thesis_Master_Ryo_Iida_abst.pdf, abst); + %%%山田美晴%%%, ``対照学習による顔画像と音声のマルチモーダルモデルを用いた顔画像を入力とするゼロショット音声合成,'' 修士論文, 名古屋工業大学, 2026年2月. //%%%Miharu Yamada%%%, ``Zero-shot speech synthesis with face image inputs based on multimodal model of face images and speech trained with contrastive learning,'' Master thesis, Nagoya Institute of Technology, February, 2026. &publication(2026/20260210_Thesis_Master_Miharu_Yamada_paper.pdf, paper); &publication(2026/20260210_Thesis_Master_Miharu_Yamada_slide.pptx, slide); &publication(2026/20260210_Thesis_Master_Miharu_Yamada_abst.pdf, abst); // hodotsuka-sempai + %%%程塚海月%%%, ``話者照合における対照損失を導入した教師なし自己蒸留の検討,'' 修士論文, 名古屋工業大学, 2026年2月. //%%%Mizuki Hodotsuka%%%, ``An investigation of unsupervised self-distillation incorporating contrastive loss in speaker verification,'' Master thesis, Nagoya Institute of Technology, February, 2026. &publication(2026/20260210_Thesis_Master_Mizuki_Hodotsuka_paper.pdf, paper); &publication(2026/20260210_Thesis_Master_Mizuki_Hodotsuka_slide.pptx, slide); &publication(2026/20260210_Thesis_Master_Mizuki_Hodotsuka_abst.pdf, abst); // tamaki-sempai + %%%田牧宏都%%%, ``自然言語による声質制御可能な音声合成のための声質説明文データの作成・評価システム,'' 修士論文, 名古屋工業大学, 2026年2月. //%%%Hiroto Tamaki%%%, ``A system for creation and evaluation of voice-characteristic description data for natural-language-controllable speech synthesis,'' Master thesis, Nagoya Institute of Technology, February, 2026. &publication(2026/20260210_Thesis_Master_Hiroto_Tamaki_paper.pdf, paper); &publication(2026/20260210_Thesis_Master_Hiroto_Tamaki_slide.pptx, slide); &publication(2026/20260210_Thesis_Master_Hiroto_Tamaki_abst.pdf, abst); // onda-sempai + %%%恩田将人%%%, ``ニューラルコーデック言語モデルに基づく音声プロンプトを用いたzero-shot声質変換,'' 修士論文, 名古屋工業大学, 2026年2月. //%%%Masato Onda%%%, ``Zero-shot voice conversion using speech prompts based on neural codec language models,'' Master thesis, Nagoya Institute of Technology, February, 2026. &publication(2026/20260210_Thesis_Master_Masato_Onda_paper.pdf, paper); &publication(2026/20260210_Thesis_Master_Masato_Onda_slide.pptx, slide); &publication(2026/20260210_Thesis_Master_Masato_Onda_abst.pdf, abst); // wen-sempai + %%%文晨航%%%, ``テキスト音声合成のための深層ガウス混合モデルに基づく特徴表現抽出,'' 修士論文, 名古屋工業大学, 2026年2月. //%%%Chenhang Wen%%%, ``Feature representation extraction based on deep gaussian mixture models for text-to-speech synthesis,'' Master thesis, Nagoya Institute of Technology, February, 2026. &publication(2026/20260210_Thesis_Master_Chenhang_Wen_paper.pdf, paper); &publication(2026/20260210_Thesis_Master_Chenhang_Wen_slide.pptx, slide); &publication(2026/20260210_Thesis_Master_Chenhang_Wen_abst.pdf, abst); // kim-sempai + %%%金賢優%%%, ``ニューラル歌声合成における楽譜特徴表現と音響モデルの文脈学習能力との関係に関する調査,'' 修士論文, 名古屋工業大学, 2026年2月. //%%%Hyunwoo Kim%%%, ``An investigation of the relationship between score feature representations and context modeling capability of acoustic models for neural singing voice synthesis,'' Master thesis, Nagoya Institute of Technology, February, 2026. &publication(2026/20260210_Thesis_Master_Hyunwoo_Kim_paper.pdf, paper); &publication(2026/20260210_Thesis_Master_Hyunwoo_Kim_slide.pptx, slide); &publication(2026/20260210_Thesis_Master_Hyunwoo_Kim_abst.pdf, abst); // hamada-sempai + %%%濱田脩斗%%%, ``ニューラルコーデック言語モデルを用いた方言音声合成に関する検討,'' 修士論文, 名古屋工業大学, 2026年2月. //%%%Shuto Hamada%%%, ``An investigation of japanese dialect speech synthesis with neural codec language models,'' Master thesis, Nagoya Institute of Technology, February, 2026. &publication(2026/20260210_Thesis_Master_Shuto_Hamada_paper.pdf, paper); &publication(2026/20260210_Thesis_Master_Shuto_Hamada_slide.pptx, slide); &publication(2026/20260210_Thesis_Master_Shuto_Hamada_abst.pdf, abst); // aohara-sempai + %%%青原光%%%, ``自己教師あり学習に基づく軽量な音声内容特徴抽出モデルの設計と評価,'' 修士論文, 名古屋工業大学, 2026年2月. //%%%Hikaru Aohara%%%, ``Design and evaluation of a lightweight speech content feature extraction model based on self-supervised learning,'' Master thesis, Nagoya Institute of Technology, February, 2026. &publication(2026/20260210_Thesis_Master_Hikaru_Aohara_paper.pdf, paper); &publication(2026/20260210_Thesis_Master_Hikaru_Aohara_slide.pptx, slide); &publication(2026/20260210_Thesis_Master_Hikaru_Aohara_abst.pdf, abst); //** 講演 [#f58d14fc] //** プレプリント [#t650b610] //** 著書 [#t58b17b3] //** その他 [#u12e6a23] **過去の発表論文 [#ifa8da9a] #ls2(ホーム/発表論文/,reverse);