@misc{indiciae048c3ff01a20, title = {Joint Speech Recognition and Audio Captioning}, author = {Chaitanya Narisetty and Emiru Tsunoo and Xuankai Chang and Yosuke Kashiwagi and Michael Hentschel and Shinji Watanabe}, year = {2022}, url = {https://arxiv.org/abs/2202.01405}, note = {Source identifier: 2202.01405} }