@misc{indiciae39291f08f9a6, title = {Planner-Refiner: Dynamic Space-Time Refinement for Vision-Language Alignment in Videos}, author = {Tuyen Tran and Thao Minh Le and Quang-Hung Le and Truyen Tran}, year = {2025}, url = {https://arxiv.org/abs/2508.07330}, note = {Source identifier: 2508.07330} }