@misc{indiciae1613774e1132, title = {Implicit and Explicit Commonsense for Multi-sentence Video Captioning}, author = {Shih-Han Chou and James J. Little and Leonid Sigal}, year = {2024}, url = {https://arxiv.org/abs/2303.07545}, note = {Source identifier: 2303.07545} }