@misc{indiciae2d374f859ae4, title = {CineCap: Structured Reasoning with Spatio-Temporal Anchors for Cinematographic Video Captioning}, author = {Xinyu Mao and Yuhui Zeng and Xiaokun Liu and Wenyu Qin and Meng Wang and Xin Tao and Pengfei Wan and Xiaohan Xing and Max Meng}, year = {2026}, url = {https://arxiv.org/abs/2606.24636}, note = {Source identifier: 2606.24636} }