@misc{indiciae2285e93c5037, title = {SVLA: A Unified Speech-Vision-Language Assistant with Multimodal Reasoning and Speech Generation}, author = {Ngoc Dung Huynh and Mohamed Reda Bouadjenek and Imran Razzak and Hakim Hacid and Sunil Aryal}, year = {2025}, url = {https://arxiv.org/abs/2503.24164}, note = {Source identifier: 2503.24164} }