@misc{indiciaeb5c5113961c9, title = {Watch, Listen and Tell: Multi-modal Weakly Supervised Dense Event Captioning}, author = {Tanzila Rahman and Bicheng Xu and Leonid Sigal}, year = {2019}, url = {https://arxiv.org/abs/1909.09944}, note = {Source identifier: 1909.09944} }