@misc{indiciae957f62a24a23, title = {DiffSoundStream: Efficient Speech Tokenization via Diffusion Decoding}, author = {Yang Yang and Yunpeng Li and George Sung and Shao-Fu Shih and Craig Dooley and Alessio Centazzo and Ramanan Rajeswaran}, year = {2025}, url = {https://arxiv.org/abs/2506.22362}, note = {Source identifier: 2506.22362} }