<?xml version='1.0' encoding='utf-8'?>
<urlset xmlns="http://www.sitemaps.org/schemas/sitemap/0.9" xmlns:news="http://www.google.com/schemas/sitemap-news/0.9">
  <url>
    <loc>https://speech-expert.ru/news/hybrid-real-complex-phase-aware-speech-enhancement</loc>
    <news:news>
      <news:publication>
        <news:name>Speech Expert</news:name>
        <news:language>ru</news:language>
      </news:publication>
      <news:publication_date>2026-08-14T00:00:00Z</news:publication_date>
      <news:title>Hybrid Real- and Complex-Valued Neural Network Architecture for Speech Enhancement</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://speech-expert.ru/news/rt-semamba-realtime-speech-enhancement-distillation</loc>
    <news:news>
      <news:publication>
        <news:name>Speech Expert</news:name>
        <news:language>ru</news:language>
      </news:publication>
      <news:publication_date>2026-08-14T00:00:00Z</news:publication_date>
      <news:title>RT-SEMamba: Real-Time Speech Enhancement Mamba via Progressive Knowledge Distillation</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://speech-expert.ru/news/six-codec-latent-generative-speech-enhancement-paradigms</loc>
    <news:news>
      <news:publication>
        <news:name>Speech Expert</news:name>
        <news:language>ru</news:language>
      </news:publication>
      <news:publication_date>2026-08-14T00:00:00Z</news:publication_date>
      <news:title>Rethinking Language Model-Based Generative Speech Enhancement in the Latent Space of a Neural Audio Codec</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://speech-expert.ru/news/slt-2026-smartglasses-egocentric-speech-benchmark</loc>
    <news:news>
      <news:publication>
        <news:name>Speech Expert</news:name>
        <news:language>ru</news:language>
      </news:publication>
      <news:publication_date>2026-08-14T00:00:00Z</news:publication_date>
      <news:title>The SLT 2026 SmartGlasses Challenge: Benchmarking Egocentric Multi-Talker Speech Recognition and Understanding with Audio-Language Models</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://speech-expert.ru/news/on-policy-self-distillation-multi-dialect-asr</loc>
    <news:news>
      <news:publication>
        <news:name>Speech Expert</news:name>
        <news:language>ru</news:language>
      </news:publication>
      <news:publication_date>2026-08-14T00:00:00Z</news:publication_date>
      <news:title>On-Policy Self-Distillation for Multi-Dialect ASR: Mastering Dialects, Retaining Mandarin</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://speech-expert.ru/news/midashenglm-gen-unified-audio-scenes</loc>
    <news:news>
      <news:publication>
        <news:name>Speech Expert</news:name>
        <news:language>ru</news:language>
      </news:publication>
      <news:publication_date>2026-08-14T00:00:00Z</news:publication_date>
      <news:title>MiDashengLM-Gen: Unified Audio Scene Generation via LLM-Driven Autoregressive Flow Matching</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://speech-expert.ru/news/uniswap-streaming-audio-visual-identity</loc>
    <news:news>
      <news:publication>
        <news:name>Speech Expert</news:name>
        <news:language>ru</news:language>
      </news:publication>
      <news:publication_date>2026-08-14T00:00:00Z</news:publication_date>
      <news:title>UniSwap: Streaming Audio-Visual Identity Swapping for Talking Videos</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://speech-expert.ru/news/phoenix-tts-joint-tokenizer-flow-matching</loc>
    <news:news>
      <news:publication>
        <news:name>Speech Expert</news:name>
        <news:language>ru</news:language>
      </news:publication>
      <news:publication_date>2026-08-14T00:00:00Z</news:publication_date>
      <news:title>Phoenix TTS: High-Fidelity Synthesis and Voice Conversion via Flow-Matching-Driven Speech Tokenization</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://speech-expert.ru/news/confucius4-tts-transcript-free-cross-lingual</loc>
    <news:news>
      <news:publication>
        <news:name>Speech Expert</news:name>
        <news:language>ru</news:language>
      </news:publication>
      <news:publication_date>2026-08-14T00:00:00Z</news:publication_date>
      <news:title>Confucius4-TTS: Transcript-Free Cross-Lingual Zero-Shot TTS with a Learnable Speaker Encoder</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://speech-expert.ru/news/easper-accessible-asr-language-documentation</loc>
    <news:news>
      <news:publication>
        <news:name>Speech Expert</news:name>
        <news:language>ru</news:language>
      </news:publication>
      <news:publication_date>2026-08-14T00:00:00Z</news:publication_date>
      <news:title>Easper: An Accessible ASR Pipeline for Language Documentation</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://speech-expert.ru/news/relative-transfer-matrix-multisource-microphones</loc>
    <news:news>
      <news:publication>
        <news:name>Speech Expert</news:name>
        <news:language>ru</news:language>
      </news:publication>
      <news:publication_date>2026-08-14T00:00:00Z</news:publication_date>
      <news:title>Deep Learning Based Relative Transfer Matrix Estimation for Multiple Sources and Multiple Microphones</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://speech-expert.ru/news/luna-tts-diffusion-family-realtime</loc>
    <news:news>
      <news:publication>
        <news:name>Speech Expert</news:name>
        <news:language>ru</news:language>
      </news:publication>
      <news:publication_date>2026-08-14T00:00:00Z</news:publication_date>
      <news:title>Luna-TTS Family Technical Report</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://speech-expert.ru/news/cookvoice-controllable-speech-singing-generation</loc>
    <news:news>
      <news:publication>
        <news:name>Speech Expert</news:name>
        <news:language>ru</news:language>
      </news:publication>
      <news:publication_date>2026-08-14T00:00:00Z</news:publication_date>
      <news:title>CookVoice: Unified Framework for Style Controllable Multi-Modal Human Voice Generation</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://speech-expert.ru/news/infant-audio-whisper-structured-speaker-conditioning</loc>
    <news:news>
      <news:publication>
        <news:name>Speech Expert</news:name>
        <news:language>ru</news:language>
      </news:publication>
      <news:publication_date>2026-08-14T00:00:00Z</news:publication_date>
      <news:title>Robust Multi-Tier Infant-Centered Audio Understanding with Whisper via Structured Speaker Conditioning</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://speech-expert.ru/news/donorrank-language-selection-low-resource-asr</loc>
    <news:news>
      <news:publication>
        <news:name>Speech Expert</news:name>
        <news:language>ru</news:language>
      </news:publication>
      <news:publication_date>2026-08-14T00:00:00Z</news:publication_date>
      <news:title>DonorRank: Donor Language Selection for Low-Resource Cross-Lingual Speech Recognition</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://speech-expert.ru/news/qwen-musicavqa-7b-temporal-audio-tokens</loc>
    <news:news>
      <news:publication>
        <news:name>Speech Expert</news:name>
        <news:language>ru</news:language>
      </news:publication>
      <news:publication_date>2026-08-14T00:00:00Z</news:publication_date>
      <news:title>Qwen-MusicAVQA-7B: A Multimodal Model for Music Audio-Visual QA</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://speech-expert.ru/news/relfx-relative-audio-effects-representation</loc>
    <news:news>
      <news:publication>
        <news:name>Speech Expert</news:name>
        <news:language>ru</news:language>
      </news:publication>
      <news:publication_date>2026-08-13T00:00:00Z</news:publication_date>
      <news:title>RelFx: относительные представления аудиоэффектов без dry reference</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://speech-expert.ru/news/bioacoustic-denoising-time-frequency-ridge-tracking</loc>
    <news:news>
      <news:publication>
        <news:name>Speech Expert</news:name>
        <news:language>ru</news:language>
      </news:publication>
      <news:publication_date>2026-08-13T00:00:00Z</news:publication_date>
      <news:title>Bioacoustic denoising: частотно-временное восстановление ультразвуковых вокализаций</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://speech-expert.ru/news/dino-a-audio-self-distillation</loc>
    <news:news>
      <news:publication>
        <news:name>Speech Expert</news:name>
        <news:language>ru</news:language>
      </news:publication>
      <news:publication_date>2026-08-13T00:00:00Z</news:publication_date>
      <news:title>DINO-A: отрицательный результат для самодистилляции аудио</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://speech-expert.ru/news/garhwali-asr-multiseed-evaluation</loc>
    <news:news>
      <news:publication>
        <news:name>Speech Expert</news:name>
        <news:language>ru</news:language>
      </news:publication>
      <news:publication_date>2026-08-13T00:00:00Z</news:publication_date>
      <news:title>Garhwali ASR: почему пять seed важнее новой функции потерь</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://speech-expert.ru/news/voxsumm-speech-summarization-translation</loc>
    <news:news>
      <news:publication>
        <news:name>Speech Expert</news:name>
        <news:language>ru</news:language>
      </news:publication>
      <news:publication_date>2026-08-13T00:00:00Z</news:publication_date>
      <news:title>VoxSumm: длинная устная новость, перевод и краткое изложение</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://speech-expert.ru/news/never-stop-speaking-dos-speech-language-models</loc>
    <news:news>
      <news:publication>
        <news:name>Speech Expert</news:name>
        <news:language>ru</news:language>
      </news:publication>
      <news:publication_date>2026-08-13T00:00:00Z</news:publication_date>
      <news:title>Never Stop Speaking: DoS-риск для речевых языковых моделей</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://speech-expert.ru/news/duplexworld-voice-agent-benchmark</loc>
    <news:news>
      <news:publication>
        <news:name>Speech Expert</news:name>
        <news:language>ru</news:language>
      </news:publication>
      <news:publication_date>2026-08-13T00:00:00Z</news:publication_date>
      <news:title>DuplexWorld: голосовые агенты звучат хорошо, но задачи решают не всегда</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://speech-expert.ru/news/asr-roundtrip-masked-tts-reading-errors</loc>
    <news:news>
      <news:publication>
        <news:name>Speech Expert</news:name>
        <news:language>ru</news:language>
      </news:publication>
      <news:publication_date>2026-08-13T00:00:00Z</news:publication_date>
      <news:title>ASR-roundtrip: как распознавание скрывает ошибки китайского TTS</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://speech-expert.ru/news/x2-turn-streaming-asr-turn-state</loc>
    <news:news>
      <news:publication>
        <news:name>Speech Expert</news:name>
        <news:language>ru</news:language>
      </news:publication>
      <news:publication_date>2026-08-13T00:00:00Z</news:publication_date>
      <news:title>X2-Turn: потоковое распознавание речи вместе с состоянием реплики</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://speech-expert.ru/news/whisper-aware-llm-whispered-speech-recognition</loc>
    <news:news>
      <news:publication>
        <news:name>Speech Expert</news:name>
        <news:language>ru</news:language>
      </news:publication>
      <news:publication_date>2026-08-13T00:00:00Z</news:publication_date>
      <news:title>Whisper-Aware LLM: распознавание шепотной речи без галлюцинаций</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://speech-expert.ru/news/age-aware-child-speech-phoneme-recognition</loc>
    <news:news>
      <news:publication>
        <news:name>Speech Expert</news:name>
        <news:language>ru</news:language>
      </news:publication>
      <news:publication_date>2026-08-13T00:00:00Z</news:publication_date>
      <news:title>Age-aware phoneme recognition: детская речь на устройстве</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://speech-expert.ru/news/mymediwhisper-burmese-medical-asr</loc>
    <news:news>
      <news:publication>
        <news:name>Speech Expert</news:name>
        <news:language>ru</news:language>
      </news:publication>
      <news:publication_date>2026-08-13T00:00:00Z</news:publication_date>
      <news:title>myMediWhisper: бирманская медицинская речь и адаптация Whisper</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://speech-expert.ru/news/voice-anonymization-worst-case-privacy-disclosure</loc>
    <news:news>
      <news:publication>
        <news:name>Speech Expert</news:name>
        <news:language>ru</news:language>
      </news:publication>
      <news:publication_date>2026-08-13T00:00:00Z</news:publication_date>
      <news:title>Worst-case privacy disclosure: почему EER мало для анонимизации голоса</news:title>
    </news:news>
  </url>
  <url>
    <loc>https://speech-expert.ru/news/bitse-binaural-target-speaker-extraction</loc>
    <news:news>
      <news:publication>
        <news:name>Speech Expert</news:name>
        <news:language>ru</news:language>
      </news:publication>
      <news:publication_date>2026-08-13T00:00:00Z</news:publication_date>
      <news:title>BiTSE: бинауральное выделение целевого диктора для AR-очков</news:title>
    </news:news>
  </url>
</urlset>