@misc{indiciae6be2a9ff7a15, title = {Audio Captioning using Pre-Trained Large-Scale Language Model Guided by Audio-based Similar Caption Retrieval}, author = {Yuma Koizumi and Yasunori Ohishi and Daisuke Niizumi and Daiki Takeuchi and Masahiro Yasuda}, year = {2020}, url = {https://arxiv.org/abs/2012.07331}, note = {Source identifier: 2012.07331} }