@misc{indiciae13c4391e9f7f, title = {Visual Clues: Bridging Vision and Language Foundations for Image Paragraph Captioning}, author = {Yujia Xie and Luowei Zhou and Xiyang Dai and Lu Yuan and Nguyen Bach and Ce Liu and Michael Zeng}, year = {2022}, url = {https://arxiv.org/abs/2206.01843}, note = {Source identifier: 2206.01843} }