@misc{indiciae8dbf33217479, title = {SiamCTC: Learning Speech Representations through Monotonic Temporal Alignment}, author = {SooHwan Eom and Mark Hasegawa-Johnson and Chang D. Yoo}, year = {2026}, doi = {10.21437/interspeech.2025-2746}, url = {https://arxiv.org/abs/2606.02220}, note = {Source identifier: 2606.02220} }