@misc{indiciae7a713e4abc3b, title = {Factorized Learning for Temporally Grounded Video-Language Models}, author = {Wenzheng Zeng and Difei Gao and Mike Zheng Shou and Hwee Tou Ng}, year = {2025}, url = {https://arxiv.org/abs/2512.24097}, note = {Source identifier: 2512.24097} }