@misc{indiciae73331ebfe4b9, title = {Inter-3D VQA: A Roadside Multimodal Benchmark for 3D Spatiotemporally Grounded Visual Question Answering}, author = {Shaozu Ding and Linan Song and Dajiang Suo}, year = {2026}, url = {https://arxiv.org/abs/2608.28762}, note = {Source identifier: 2608.28762} }