@misc{indiciae299c9863c70b, title = {More than a Moment: Towards Coherent Sequences of Audio Descriptions}, author = {Eshika Khandelwal and Junyu Xie and Tengda Han and Max Bain and Arsha Nagrani and Andrew Zisserman and Gül Varol and Makarand Tapaswi}, year = {2025}, url = {https://arxiv.org/abs/2510.25440}, note = {Source identifier: 2510.25440} }