@misc{indiciae4bfc8a8f1ada, title = {What You Say Is What You Show: Visual Narration Detection in Instructional Videos}, author = {Kumar Ashutosh and Rohit Girdhar and Lorenzo Torresani and Kristen Grauman}, year = {2023}, url = {https://arxiv.org/abs/2301.02307}, note = {Source identifier: 2301.02307} }