@misc{indiciaee4c3a2267f65, title = {Can Large Audio Language Models Understand Audio Well? Speech, Scene and Events Understanding Benchmark for LALMs}, author = {Han Yin and Jung-Woo Choi}, year = {2026}, url = {https://arxiv.org/abs/2509.13148}, note = {Source identifier: 2509.13148} }