@misc{indiciae65c838935186, title = {Clue: Cross-modal Coherence Modeling for Caption Generation}, author = {Malihe Alikhani and Piyush Sharma and Shengjie Li and Radu Soricut and Matthew Stone}, year = {2020}, doi = {10.18653/v1/2020.acl-main.583}, url = {https://arxiv.org/abs/2005.00908}, note = {Source identifier: 2005.00908} }