@misc{indiciae5d5515d19312, title = {Improving Vision-and-Language Navigation with Image-Text Pairs from the Web}, author = {Arjun Majumdar and Ayush Shrivastava and Stefan Lee and Peter Anderson and Devi Parikh and Dhruv Batra}, year = {2020}, url = {https://arxiv.org/abs/2004.14973}, note = {Source identifier: 2004.14973} }