@misc{indiciaef0f95383ea29, title = {ReVisionLLM: Recursive Vision-Language Model for Temporal Grounding in Hour-Long Videos}, author = {Tanveer Hannan and Md Mohaiminul Islam and Jindong Gu and Thomas Seidl and Gedas Bertasius}, year = {2024}, url = {https://arxiv.org/abs/2411.14901}, note = {Source identifier: 2411.14901} }