@misc{indiciaebf76ecbe0006, title = {Multimodal Pretraining for Dense Video Captioning}, author = {Gabriel Huang and Bo Pang and Zhenhai Zhu and Clara Rivera and Radu Soricut}, year = {2020}, url = {https://arxiv.org/abs/2011.11760}, note = {Source identifier: 2011.11760} }