@misc{indiciae4f7204d120d6, title = {PolyVivid: Vivid Multi-Subject Video Generation with Cross-Modal Interaction and Enhancement}, author = {Teng Hu and Zhentao Yu and Zhengguang Zhou and Jiangning Zhang and Yuan Zhou and Qinglin Lu and Ran Yi}, year = {2025}, url = {https://arxiv.org/abs/2506.07848}, note = {Source identifier: 2506.07848} }