@misc{indiciae838ea61c4c72, title = {LVD-2M: A Long-take Video Dataset with Temporally Dense Captions}, author = {Tianwei Xiong and Yuqing Wang and Daquan Zhou and Zhijie Lin and Jiashi Feng and Xihui Liu}, year = {2024}, url = {https://arxiv.org/abs/2410.10816}, note = {Source identifier: 2410.10816} }