@misc{indiciae099c8be3f989, title = {VSTAR: A Video-grounded Dialogue Dataset for Situated Semantic Understanding with Scene and Topic Transitions}, author = {Yuxuan Wang and Zilong Zheng and Xueliang Zhao and Jinpeng Li and Yueqian Wang and Dongyan Zhao}, year = {2023}, url = {https://arxiv.org/abs/2305.18756}, note = {Source identifier: 2305.18756} }