@misc{indiciae06d8023d9b29, title = {Unifying Event Detection and Captioning as Sequence Generation via Pre-Training}, author = {Qi Zhang and Yuqing Song and Qin Jin}, year = {2022}, url = {https://arxiv.org/abs/2207.08625}, note = {Source identifier: 2207.08625} }