@misc{indiciaecd57bb206ef6, title = {Universal Video Temporal Grounding with Generative Multi-modal Large Language Models}, author = {Zeqian Li and Shangzhe Di and Zhonghua Zhai and Weilin Huang and Yanfeng Wang and Weidi Xie}, year = {2025}, url = {https://arxiv.org/abs/2506.18883}, note = {Source identifier: 2506.18883} }