@misc{indiciaebe30f1c6dde5, title = {Enhancing Localized Reasoning for Long Video Understanding via Efficient Segment-to-Video Supervision}, author = {Beibei Zhang and Chao Xu and Jun Lan and Zongyi Li and Lai Wei and Huijia Zhu and Tongwei Ren}, year = {2026}, url = {https://arxiv.org/abs/2608.20814}, note = {Source identifier: 2608.20814} }