@misc{indiciaef5d1bb5c87bf, title = {R-AVST: Empowering Video-LLMs with Fine-Grained Spatio-Temporal Reasoning in Complex Audio-Visual Scenarios}, author = {Lu Zhu and Tiantian Geng and Yangye Chen and Teng Wang and Ping Lu and Feng Zheng}, year = {2025}, url = {https://arxiv.org/abs/2511.16901}, note = {Source identifier: 2511.16901} }