@misc{indiciae56553693fcf6, title = {Shot-by-Shot: Film-Grammar-Aware Training-Free Audio Description Generation}, author = {Junyu Xie and Tengda Han and Max Bain and Arsha Nagrani and Eshika Khandelwal and Gül Varol and Weidi Xie and Andrew Zisserman}, year = {2025}, url = {https://arxiv.org/abs/2504.01020}, note = {Source identifier: 2504.01020} }