@misc{indiciae8316661e970c, title = {Speech Synthesis From Continuous Features Using Per-Token Latent Diffusion}, author = {Arnon Turetzky and Avihu Dekel and Nimrod Shabtay and Slava Shechtman and David Haws and Hagai Aronowitz and Ron Hoory and Yossi Adi}, year = {2025}, url = {https://arxiv.org/abs/2410.16048}, note = {Source identifier: 2410.16048} }