@misc{indiciae63496fda5ac6, title = {AV-Reasoner: Improving and Benchmarking Clue-Grounded Audio-Visual Counting for MLLMs}, author = {Lidong Lu and Guo Chen and Zhiqi Li and Yicheng Liu and Tong Lu}, year = {2025}, url = {https://arxiv.org/abs/2506.05328}, note = {Source identifier: 2506.05328} }