@misc{indiciaef6442ba31e27, title = {Pinpointing Trigger Moment for Grounded Video QA: Enhancing Spatio-temporal Grounding in Multimodal Large Language Models}, author = {Jinhwan Seo and Yoonki Cho and Junhyug Noh and Sung-eui Yoon}, year = {2025}, url = {https://arxiv.org/abs/2511.02182}, note = {Source identifier: 2511.02182} }