<?xml version="1.0"?>
<dblpperson name="Qin Jin" pid="47/2670" n="343">
<person key="homepages/47/2670" mdate="2025-01-10">
<author pid="47/2670">Qin Jin</author>
<url>https://orcid.org/0000-0001-6486-6020</url>
<url>https://www.wikidata.org/entity/Q130830978</url>
</person>
<r><inproceedings key="conf/aaai/ChenXMLYZYWJ26" mdate="2026-04-08">
<author pid="82/353-8">Liangyu Chen 0008</author>
<author pid="177/6703-3">Yichen Xu 0003</author>
<author orcid="0009-0006-7605-5923" pid="253/1100">Jianzhe Ma</author>
<author pid="35/9071-3">Yuqi Liu 0003</author>
<author pid="413/7448">Donglu Yang</author>
<author pid="50/6759">Liang Zhang</author>
<author pid="339/2864">Zihao Yue</author>
<author pid="203/1536-1">Wenxuan Wang 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>ChartEditor: A Reinforcement Learning Framework for Robust Chart Editing.</title>
<pages>20199-20207</pages>
<year>2026</year>
<booktitle>AAAI</booktitle>
<ee type="oa">https://doi.org/10.1609/aaai.v40i24.39107</ee>
<crossref>conf/aaai/2026</crossref>
<url>db/conf/aaai/aaai2026.html#ChenXMLYZYWJ26</url>
</inproceedings>
</r>
<r><inproceedings key="conf/aaai/HuangDWWLGWJ26" mdate="2026-03-23">
<author pid="305/3439">Zhaopei Huang</author>
<author pid="157/2769">Qifeng Dai</author>
<author pid="262/6203">Guozheng Wu</author>
<author pid="153/4201">Xiaopeng Wu</author>
<author pid="228/7857">Xubin Li</author>
<author pid="135/4944">Tiezheng Ge</author>
<author pid="203/1536-1">Wenxuan Wang 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>Mem-PAL: Towards Memory-based Personalized Dialogue Assistants for Long-term User-Agent Interaction.</title>
<pages>31229-31237</pages>
<year>2026</year>
<booktitle>AAAI</booktitle>
<ee type="oa">https://doi.org/10.1609/aaai.v40i37.40385</ee>
<crossref>conf/aaai/2026</crossref>
<url>db/conf/aaai/aaai2026.html#HuangDWWLGWJ26</url>
</inproceedings>
</r>
<r><inproceedings key="conf/acl/WangYWJ26" mdate="2026-07-30">
<author pid="79/10743">Ziheng Wang</author>
<author pid="339/2864">Zihao Yue</author>
<author pid="203/1536-1">Wenxuan Wang 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>Exploring Attention Attractors in Large Language Models.</title>
<pages>1148-1160</pages>
<year>2026</year>
<booktitle>ACL (1)</booktitle>
<ee type="oa">https://doi.org/10.18653/v1/2026.acl-long.51</ee>
<ee type="oa">https://aclanthology.org/2026.acl-long.51/</ee>
<crossref>conf/acl/2026-1</crossref>
<url>db/conf/acl/acl2026-1.html#WangYWJ26</url>
</inproceedings>
</r>
<r><inproceedings key="conf/acl/LiHWJ26" mdate="2026-07-30">
<author pid="11/2724">Zhuoqun Li</author>
<author pid="305/3439">Zhaopei Huang</author>
<author pid="203/1536-1">Wenxuan Wang 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>LongMP-Bench: A Benchmark for Multimodal Persona Understanding in Long-Term Dialogues.</title>
<pages>23132-23160</pages>
<year>2026</year>
<booktitle>ACL (Findings)</booktitle>
<ee type="oa">https://doi.org/10.18653/v1/2026.findings-acl.1159</ee>
<ee type="oa">https://aclanthology.org/2026.findings-acl.1159/</ee>
<crossref>conf/acl/2026f</crossref>
<url>db/conf/acl/acl2026f.html#LiHWJ26</url>
</inproceedings>
</r>
<r><inproceedings key="conf/acl/XuMWCCWJ26" mdate="2026-08-03">
<author pid="177/6703-3">Yichen Xu 0003</author>
<author orcid="0009-0006-7605-5923" pid="253/1100">Jianzhe Ma</author>
<author pid="210/1379">Chuhan Wang</author>
<author pid="363/6984">Zhonghao Cao</author>
<author pid="82/353-8">Liangyu Chen 0008</author>
<author pid="203/1536-1">Wenxuan Wang 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>A Survey of Large Models in Sports.</title>
<pages>37154-37189</pages>
<year>2026</year>
<booktitle>ACL (Findings)</booktitle>
<ee type="oa">https://doi.org/10.18653/v1/2026.findings-acl.1851</ee>
<ee type="oa">https://aclanthology.org/2026.findings-acl.1851/</ee>
<crossref>conf/acl/2026f</crossref>
<url>db/conf/acl/acl2026f.html#XuMWCCWJ26</url>
</inproceedings>
</r>
<r><inproceedings key="conf/acl/WangYJ26" mdate="2026-07-30">
<author pid="37/5333">Xueyan Wang</author>
<author pid="266/2264">Dingyi Yang</author>
<author pid="47/2670">Qin Jin</author>
<title>HowToNarrate: A General-Domain Benchmark for Synchronized Video Narration with External Knowledge.</title>
<pages>39110-39131</pages>
<year>2026</year>
<booktitle>ACL (1)</booktitle>
<ee type="oa">https://doi.org/10.18653/v1/2026.acl-long.1815</ee>
<ee type="oa">https://aclanthology.org/2026.acl-long.1815/</ee>
<crossref>conf/acl/2026-1</crossref>
<url>db/conf/acl/acl2026-1.html#WangYJ26</url>
</inproceedings>
</r>
<r><inproceedings key="conf/acl/MaWJ26" mdate="2026-08-03">
<author orcid="0009-0006-7605-5923" pid="253/1100">Jianzhe Ma</author>
<author pid="203/1536-1">Wenxuan Wang 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>A Survey of Deep Learning for Geometry Problem Solving.</title>
<pages>39416-39448</pages>
<year>2026</year>
<booktitle>ACL (1)</booktitle>
<ee type="oa">https://doi.org/10.18653/v1/2026.acl-long.1829</ee>
<ee type="oa">https://aclanthology.org/2026.acl-long.1829/</ee>
<crossref>conf/acl/2026-1</crossref>
<url>db/conf/acl/acl2026-1.html#MaWJ26</url>
</inproceedings>
</r>
<r><inproceedings key="conf/acl/YangWJ26" mdate="2026-07-30">
<author pid="266/2264">Dingyi Yang</author>
<author pid="141/0952">Mingshuo Wang</author>
<author pid="47/2670">Qin Jin</author>
<title>ChangJuan: A Comprehensive Benchmark for Book-Length Chinese Story Evaluation.</title>
<pages>41116-41134</pages>
<year>2026</year>
<booktitle>ACL (Findings)</booktitle>
<ee type="oa">https://doi.org/10.18653/v1/2026.findings-acl.2044</ee>
<ee type="oa">https://aclanthology.org/2026.findings-acl.2044/</ee>
<crossref>conf/acl/2026f</crossref>
<url>db/conf/acl/acl2026f.html#YangWJ26</url>
</inproceedings>
</r>
<r><inproceedings key="conf/acl/XuCZYMWJ26" mdate="2026-08-03">
<author pid="177/6703-3">Yichen Xu 0003</author>
<author pid="82/353-8">Liangyu Chen 0008</author>
<author pid="50/6759">Liang Zhang</author>
<author pid="339/2864">Zihao Yue</author>
<author orcid="0009-0006-7605-5923" pid="253/1100">Jianzhe Ma</author>
<author pid="203/1536-1">Wenxuan Wang 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>POLYCHARTQA: Benchmarking Large Vision-Language Models with Multilingual Chart Question Answering.</title>
<pages>44154-44186</pages>
<year>2026</year>
<booktitle>ACL (1)</booktitle>
<ee type="oa">https://doi.org/10.18653/v1/2026.acl-long.2043</ee>
<ee type="oa">https://aclanthology.org/2026.acl-long.2043/</ee>
<crossref>conf/acl/2026-1</crossref>
<url>db/conf/acl/acl2026-1.html#XuCZYMWJ26</url>
</inproceedings>
</r>
<r><article publtype="informal" key="journals/corr/abs-2602-09722" mdate="2026-03-25">
<author pid="44/6292">Ye Wang</author>
<author pid="251/3691">Sipeng Zheng</author>
<author pid="14/3727-11">Hao Luo 0011</author>
<author pid="73/10693-2">Wanpeng Zhang 0002</author>
<author pid="254/2084">Haoqi Yuan</author>
<author pid="358/4076">Chaoyi Xu</author>
<author pid="426/9329">Haiweng Xu</author>
<author pid="340/4016">Yicheng Feng</author>
<author pid="135/9009">Mingyang Yu</author>
<author pid="56/6876">Zhiyu Kang</author>
<author pid="99/965-2">Zongqing Lu 0002</author>
<author pid="47/2670">Qin Jin</author>
<title>Rethinking Visual-Language-Action Model Scaling: Alignment, Mixture, and Regularization.</title>
<year>2026</year>
<month>February</month>
<volume>abs/2602.09722</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2602.09722</ee>
<url>db/journals/corr/corr2602.html#abs-2602-09722</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2602-11464" mdate="2026-03-29">
<author pid="15/4777">Tao Zhang</author>
<author pid="01/9629">Song Xia</author>
<author pid="44/6292">Ye Wang</author>
<author pid="47/2670">Qin Jin</author>
<title>EasyMimic: A Low-Cost Framework for Robot Imitation Learning from Human Videos.</title>
<year>2026</year>
<month>February</month>
<volume>abs/2602.11464</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2602.11464</ee>
<url>db/journals/corr/corr2602.html#abs-2602-11464</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2604-12320" mdate="2026-07-10">
<author orcid="0009-0006-7605-5923" pid="253/1100">Jianzhe Ma</author>
<author pid="363/6984">Zhonghao Cao</author>
<author pid="435/3086">Shangkui Chen</author>
<author pid="177/6703-3">Yichen Xu 0003</author>
<author pid="203/1536-1">Wenxuan Wang 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>EgoEsportsQA: An Egocentric Video Benchmark for Perception and Reasoning in Esports.</title>
<year>2026</year>
<month>April</month>
<volume>abs/2604.12320</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2604.12320</ee>
<url>db/journals/corr/corr2604.html#abs-2604-12320</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2604-15127" mdate="2026-05-28">
<author pid="385/1380-1">Huanran Hu 0001</author>
<author pid="20/7764">Zihui Ren</author>
<author pid="266/2264">Dingyi Yang</author>
<author pid="82/353-8">Liangyu Chen 0008</author>
<author pid="378/3056">Qixiang Gao</author>
<author pid="135/4944">Tiezheng Ge</author>
<author pid="47/2670">Qin Jin</author>
<title>MCSC-Bench: Multimodal Context-to-Script Creation for Realistic Video Production.</title>
<year>2026</year>
<month>April</month>
<volume>abs/2604.15127</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2604.15127</ee>
<url>db/journals/corr/corr2604.html#abs-2604-15127</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2604-15736" mdate="2026-07-10">
<author pid="177/6703-3">Yichen Xu 0003</author>
<author pid="174/0438">Yuanhang Liu</author>
<author pid="210/1379">Chuhan Wang</author>
<author pid="216/4838">Zihan Zhao</author>
<author pid="435/6686">Jinghan Luo</author>
<author orcid="0009-0006-7605-5923" pid="253/1100">Jianzhe Ma</author>
<author pid="203/1536-1">Wenxuan Wang 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>RefereeBench: Are Video MLLMs Ready to be Multi-Sport Referees.</title>
<year>2026</year>
<month>April</month>
<volume>abs/2604.15736</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2604.15736</ee>
<url>db/journals/corr/corr2604.html#abs-2604-15736</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2604-18356" mdate="2026-05-19">
<author pid="305/3439">Zhaopei Huang</author>
<author pid="149/1280">Yanfeng Jia</author>
<author pid="117/4068">Jiayi Zhao</author>
<author pid="66/10435">Xinjie Zhang</author>
<author pid="203/1536-1">Wenxuan Wang 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>ComPASS: Towards Personalized Agentic Social Support via Tool-Augmented Companionship.</title>
<year>2026</year>
<month>April</month>
<volume>abs/2604.18356</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2604.18356</ee>
<url>db/journals/corr/corr2604.html#abs-2604-18356</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2605-16381" mdate="2026-06-24">
<author pid="54/2788">Ao Li</author>
<author pid="319/7252-1">Zihan Xiao 0001</author>
<author pid="339/2864">Zihao Yue</author>
<author pid="293/8958">Boshen Xu</author>
<author pid="262/6579">Linli Yao</author>
<author pid="157/2088">Jiaze Li</author>
<author pid="202/1707">Pei Fu</author>
<author pid="289/4197">Jianzhong Ju</author>
<author pid="61/3233-1">Jian Luan 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>StreamPro: From Reactive Perception to Proactive Decision-Making in Streaming Video.</title>
<year>2026</year>
<month>May</month>
<volume>abs/2605.16381</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2605.16381</ee>
<url>db/journals/corr/corr2605.html#abs-2605-16381</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2605-28073" mdate="2026-06-14">
<author pid="260/7814">Hanwen Cui</author>
<author pid="334/2644">Yuting Mei</author>
<author pid="322/6277">Yuhang Fu</author>
<author pid="266/2264">Dingyi Yang</author>
<author pid="47/2670">Qin Jin</author>
<title>StoryLens: Preference-Aligned Story Rewriting via Context-Aware Narrative Enrichment.</title>
<year>2026</year>
<month>May</month>
<volume>abs/2605.28073</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2605.28073</ee>
<url>db/journals/corr/corr2605.html#abs-2605-28073</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2606-17520" mdate="2026-07-08">
<author pid="10/239">Jiawei Zhang</author>
<author pid="121/6500">Yiming Yan</author>
<author pid="10/3072">Chao Liang</author>
<author pid="74/6847">Nuo Xu</author>
<author pid="440/1367">Seson Sun</author>
<author pid="154/7819">Qichen Zhang</author>
<author pid="205/8666">Yuhao Xu</author>
<author pid="331/1645">Yantai Yang</author>
<author pid="430/3481">Yingqiao Wang</author>
<author pid="47/2670">Qin Jin</author>
<author pid="49/4941">Zhipeng Zhang</author>
<title>GASE: Gaussian Splatting-Based Automated System for Reconstructing Embodied-Simulation Environments.</title>
<year>2026</year>
<month>June</month>
<volume>abs/2606.17520</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2606.17520</ee>
<url>db/journals/corr/corr2606.html#abs-2606-17520</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2607-06986" mdate="2026-08-05">
<author pid="50/9959">Wenhao Feng</author>
<author pid="358/9271">Yuxun Tang</author>
<author pid="229/3529">Jiatong Shi</author>
<author pid="47/2670">Qin Jin</author>
<title>MMGenre: Benchmarking Singing Voice Synthesis across Multiple Musical Genres.</title>
<year>2026</year>
<month>July</month>
<volume>abs/2607.06986</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2607.06986</ee>
<url>db/journals/corr/corr2607.html#abs-2607-06986</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2607-13056" mdate="2026-08-07">
<author pid="52/5716">Chang Liu</author>
<author pid="10/239">Jiawei Zhang</author>
<author pid="15/4777">Tao Zhang</author>
<author pid="44/6292">Ye Wang</author>
<author pid="76/7825">Hongyu Zhou</author>
<author pid="47/2670">Qin Jin</author>
<title>HRIBench: Benchmarking Interaction-Centric Human-Robot Collaboration.</title>
<year>2026</year>
<month>July</month>
<volume>abs/2607.13056</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2607.13056</ee>
<url>db/journals/corr/corr2607.html#abs-2607-13056</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article key="journals/imst/ChenWJ25" mdate="2025-09-11">
<author orcid="0009-0008-8719-8828" pid="27/7017">Peng Chen</author>
<author pid="240/2025">Huihui Wang</author>
<author pid="47/2670">Qin Jin</author>
<title>MPKU-Net: A U-Shaped Medical Image Segmentation Network Based on MLP and KAN.</title>
<year>2025</year>
<volume>35</volume>
<journal>Int. J. Imaging Syst. Technol.</journal>
<number>3</number>
<ee>https://doi.org/10.1002/ima.70105</ee>
<url>db/journals/imst/imst35.html#ChenWJ25</url>
<stream>streams/journals/imst</stream>
</article>
</r>
<r><article key="journals/tomccap/ZhangZSZJ25" mdate="2025-09-11">
<author orcid="0009-0001-7672-6778" pid="66/10435">Xinjie Zhang</author>
<author orcid="0000-0001-9794-9578" pid="305/3353">Tenggan Zhang</author>
<author orcid="0009-0004-1035-407X" pid="02/2264">Lei Sun</author>
<author orcid="0000-0002-1497-8865" pid="121/8902">Jinming Zhao</author>
<author orcid="0000-0001-6486-6020" pid="47/2670">Qin Jin</author>
<title>Exploring Interpretability in Deep Learning for Affective Computing: A Comprehensive Review.</title>
<pages>194:1-194:28</pages>
<year>2025</year>
<month>July</month>
<volume>21</volume>
<journal>ACM Trans. Multim. Comput. Commun. Appl.</journal>
<number>7</number>
<ee>https://doi.org/10.1145/3723005</ee>
<url>db/journals/tomccap/tomccap21.html#ZhangZSZJ25</url>
<stream>streams/journals/tomccap</stream>
</article>
</r>
<r><inproceedings key="conf/3dim/XuZJ25" mdate="2025-09-08">
<author pid="293/8958">Boshen Xu</author>
<author pid="251/3691">Sipeng Zheng</author>
<author pid="47/2670">Qin Jin</author>
<title>SPAFormer: Sequential 3D Part Assembly with Transformers.</title>
<pages>1317-1327</pages>
<year>2025</year>
<booktitle>3DV</booktitle>
<ee>https://doi.org/10.1109/3DV66043.2025.00125</ee>
<crossref>conf/3dim/2025</crossref>
<url>db/conf/3dim/3dv2025.html#XuZJ25</url>
<stream>streams/conf/3dim</stream>
</inproceedings>
</r>
<r><inproceedings key="conf/acl/Hu0ZY0ZJ0025" mdate="2026-06-10">
<author pid="249/1182">Anwen Hu</author>
<author pid="80/1339-1">Haiyang Xu 0001</author>
<author pid="50/6759">Liang Zhang</author>
<author pid="304/1336">Jiabo Ye</author>
<author pid="51/5332-8">Ming Yan 0008</author>
<author pid="86/1953-11">Ji Zhang 0011</author>
<author pid="47/2670">Qin Jin</author>
<author pid="h/FeiHuang-2">Fei Huang 0002</author>
<author pid="84/2644-1">Jingren Zhou 0001</author>
<title>mPLUG-DocOwl2: High-resolution Compressing for OCR-free Multi-page Document Understanding.</title>
<pages>5817-5834</pages>
<year>2025</year>
<booktitle>ACL (1)</booktitle>
<ee type="oa">https://doi.org/10.18653/v1/2025.acl-long.291</ee>
<ee type="oa">https://aclanthology.org/2025.acl-long.291/</ee>
<crossref>conf/acl/2025-1</crossref>
<url>db/conf/acl/acl2025-1.html#Hu0ZY0ZJ0025</url>
</inproceedings>
</r>
<r><inproceedings key="conf/acl/YangJ25" mdate="2026-06-10">
<author pid="266/2264">Dingyi Yang</author>
<author pid="47/2670">Qin Jin</author>
<title>What Matters in Evaluating Book-Length Stories? A Systematic Study of Long Story Evaluation.</title>
<pages>16375-16398</pages>
<year>2025</year>
<booktitle>ACL (1)</booktitle>
<ee type="oa">https://doi.org/10.18653/v1/2025.acl-long.799</ee>
<ee type="oa">https://aclanthology.org/2025.acl-long.799/</ee>
<crossref>conf/acl/2025-1</crossref>
<url>db/conf/acl/acl2025-1.html#YangJ25</url>
</inproceedings>
</r>
<r><inproceedings key="conf/acl/YueZWJ25" mdate="2026-06-10">
<author pid="339/2864">Zihao Yue</author>
<author pid="355/8240">Yepeng Zhang</author>
<author pid="79/10743">Ziheng Wang</author>
<author pid="47/2670">Qin Jin</author>
<title>Movie101v2: Improved Movie Narration Benchmark.</title>
<pages>17081-17095</pages>
<year>2025</year>
<booktitle>ACL (1)</booktitle>
<ee type="oa">https://doi.org/10.18653/v1/2025.acl-long.836</ee>
<ee type="oa">https://aclanthology.org/2025.acl-long.836/</ee>
<crossref>conf/acl/2025-1</crossref>
<url>db/conf/acl/acl2025-1.html#YueZWJ25</url>
</inproceedings>
</r>
<r><inproceedings key="conf/acl/ZhangWJ25" mdate="2026-06-10">
<author pid="66/10435">Xinjie Zhang</author>
<author pid="203/1536-1">Wenxuan Wang 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>IntentionESC: An Intention-Centered Framework for Enhancing Emotional Support in Dialogue Systems.</title>
<pages>26494-26516</pages>
<year>2025</year>
<booktitle>ACL (Findings)</booktitle>
<ee type="oa">https://doi.org/10.18653/v1/2025.findings-acl.1358</ee>
<ee type="oa">https://aclanthology.org/2025.findings-acl.1358/</ee>
<crossref>conf/acl/2025f</crossref>
<url>db/conf/acl/acl2025f.html#ZhangWJ25</url>
</inproceedings>
</r>
<r><inproceedings key="conf/emnlp/DuLSWZGZJ25" mdate="2026-02-04">
<author pid="51/3199-11">Yang Du 0011</author>
<author pid="223/9450">Zhuoran Lin</author>
<author pid="222/2920">Kaiqiang Song</author>
<author pid="15/2887">Biao Wang</author>
<author pid="50/9014">Zhicheng Zheng</author>
<author pid="135/4944">Tiezheng Ge</author>
<author pid="33/1610-7">Bo Zheng 0007</author>
<author pid="47/2670">Qin Jin</author>
<title>VC4VG: Optimizing Video Captions for Text-to-Video Generation.</title>
<pages>1124-1138</pages>
<year>2025</year>
<booktitle>EMNLP</booktitle>
<ee type="oa">https://doi.org/10.18653/v1/2025.emnlp-main.59</ee>
<crossref>conf/emnlp/2025</crossref>
<url>db/conf/emnlp/emnlp2025.html#DuLSWZGZJ25</url>
<stream>streams/conf/emnlp</stream>
</inproceedings>
</r>
<r><inproceedings key="conf/iccv/CaoZWXWJLL25" mdate="2026-05-13">
<author pid="17/1169">Bin Cao</author>
<author pid="251/3691">Sipeng Zheng</author>
<author pid="44/6292">Ye Wang</author>
<author pid="352/3374">Lujie Xia</author>
<author pid="378/1061">Qianshan Wei</author>
<author pid="47/2670">Qin Jin</author>
<author pid="72/2590">Jing Liu</author>
<author pid="99/965-2">Zongqing Lu 0002</author>
<title>MotionCtrl: A Real-Time Controllable Vision-Language-Motion Model.</title>
<year>2025</year>
<booktitle>ICCV</booktitle>
<pages>12253-12262</pages>
<crossref>conf/iccv/2025</crossref>
<ee>https://doi.org/10.1109/ICCV51701.2025.01139</ee>
<url>db/conf/iccv/iccv2025.html#CaoZWXWJLL25</url>
<stream>streams/conf/iccv</stream>
</inproceedings>
</r>
<r><inproceedings key="conf/iclr/XuWDSZJ25" mdate="2026-01-31">
<author pid="293/8958">Boshen Xu</author>
<author pid="79/10743">Ziheng Wang</author>
<author pid="51/3199-11">Yang Du 0011</author>
<author pid="378/1643">Zhinan Song</author>
<author pid="251/3691">Sipeng Zheng</author>
<author pid="47/2670">Qin Jin</author>
<title>Do Egocentric Video-Language Models Truly Understand Hand-Object Interactions?</title>
<year>2025</year>
<booktitle>ICLR</booktitle>
<ee type="oa">https://openreview.net/forum?id=M8gXSFGkn2</ee>
<crossref>conf/iclr/2025</crossref>
<url>db/conf/iclr/iclr2025.html#XuWDSZJ25</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icml/WangZCWZJ025" mdate="2026-02-04">
<author pid="44/6292">Ye Wang</author>
<author pid="251/3691">Sipeng Zheng</author>
<author pid="17/1169">Bin Cao</author>
<author pid="378/1061">Qianshan Wei</author>
<author pid="385/3798">Weishuai Zeng</author>
<author pid="47/2670">Qin Jin</author>
<author pid="99/965-2">Zongqing Lu 0002</author>
<title>Scaling Large Motion Models with Million-Level Human Motions.</title>
<year>2025</year>
<booktitle>ICML</booktitle>
<ee type="oa">https://proceedings.mlr.press/v267/wang25fb.html</ee>
<ee type="oa">https://openreview.net/forum?id=TO6jrwuxi4</ee>
<crossref>conf/icml/2025</crossref>
<url>db/conf/icml/icml2025.html#WangZCWZJ025</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/YangZY0X0J25" mdate="2026-04-07">
<author orcid="0009-0003-7763-0094" pid="413/7448">Donglu Yang</author>
<author orcid="0000-0002-6187-3628" pid="50/6759">Liang Zhang</author>
<author orcid="0000-0002-3470-5442" pid="339/2864">Zihao Yue</author>
<author orcid="0000-0003-2099-4025" pid="82/353-8">Liangyu Chen 0008</author>
<author orcid="0009-0005-8113-6464" pid="177/6703-3">Yichen Xu 0003</author>
<author orcid="0000-0002-9803-8204" pid="203/1536-1">Wenxuan Wang 0001</author>
<author orcid="0000-0001-6486-6020" pid="47/2670">Qin Jin</author>
<title>ChartM<sup>3</sup>: Benchmarking Chart Editing with Multimodal Instructions.</title>
<pages>5001-5009</pages>
<year>2025</year>
<booktitle>ACM Multimedia</booktitle>
<ee>https://doi.org/10.1145/3746027.3755714</ee>
<crossref>conf/mm/2025</crossref>
<url>db/conf/mm/mm2025.html#YangZY0X0J25</url>
</inproceedings>
</r>
<r><inproceedings key="conf/nips/WangWXDLXYJZYFH25" mdate="2026-06-28">
<author pid="44/6292">Ye Wang</author>
<author pid="79/10743">Ziheng Wang</author>
<author pid="293/8958">Boshen Xu</author>
<author pid="51/3199-11">Yang Du 0011</author>
<author pid="251/8268">Kejun Lin</author>
<author pid="319/7252-1">Zihan Xiao 0001</author>
<author pid="339/2864">Zihao Yue</author>
<author pid="289/4197">Jianzhong Ju</author>
<author pid="50/6759">Liang Zhang</author>
<author pid="266/2264">Dingyi Yang</author>
<author pid="439/1035">Xiangnan Fang</author>
<author pid="237/8845">Zewen He</author>
<author pid="152/8206">Zhenbo Luo</author>
<author pid="203/1536-1">Wenxuan Wang 0001</author>
<author pid="268/5425">Junqi Lin</author>
<author pid="61/3233-1">Jian Luan 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>Time-R1: Post-Training Large Vision Language Model for Temporal Video Grounding.</title>
<year>2025</year>
<booktitle>NeurIPS</booktitle>
<ee type="oa">http://papers.nips.cc/paper_files/paper/2025/hash/7801b29c93b599b8d0c44138596bdeed-Abstract-Conference.html</ee>
<crossref>conf/nips/2025</crossref>
<url>db/conf/nips/neurips2025.html#WangWXDLXYJZYFH25</url>
</inproceedings>
</r>
<r><inproceedings key="conf/nips/WuMYLLRWZWJH25" mdate="2026-07-29">
<author pid="301/4852-1">Yuning Wu 0001</author>
<author pid="354/8774">Jiahao Mei</author>
<author pid="51/5332-8">Ming Yan 0008</author>
<author pid="52/9457-3">Chenliang Li 0003</author>
<author pid="285/4166">Shaopeng Lai</author>
<author pid="401/8284">Yuran Ren</author>
<author pid="121/0900">Zijia Wang</author>
<author pid="86/1953-11">Ji Zhang 0011</author>
<author pid="82/2416">Mengyue Wu</author>
<author pid="47/2670">Qin Jin</author>
<author pid="h/FeiHuang-2">Fei Huang 0002</author>
<title>WritingBench: A Comprehensive Benchmark for Generative Writing.</title>
<year>2025</year>
<booktitle>NeurIPS</booktitle>
<ee type="oa">http://papers.nips.cc/paper_files/paper/2025/hash/4aedf0cba303537fcb6cf948bb41b2df-Abstract-Datasets_and_Benchmarks_Track.html</ee>
<crossref>conf/nips/2025</crossref>
<url>db/conf/nips/neurips2025.html#WuMYLLRWZWJH25</url>
</inproceedings>
</r>
<r><inproceedings key="conf/nips/XuMLZJ25" mdate="2026-06-23">
<author pid="293/8958">Boshen Xu</author>
<author pid="334/2644">Yuting Mei</author>
<author pid="402/1495">Xinbi Liu</author>
<author pid="251/3691">Sipeng Zheng</author>
<author pid="47/2670">Qin Jin</author>
<title>EgoDTM: Towards 3D-Aware Egocentric Video-Language Pretraining.</title>
<year>2025</year>
<booktitle>NeurIPS</booktitle>
<ee type="oa">http://papers.nips.cc/paper_files/paper/2025/hash/b1a81aa84530d9fa99e420b2305cc11a-Abstract-Conference.html</ee>
<crossref>conf/nips/2025</crossref>
<url>db/conf/nips/neurips2025.html#XuMLZJ25</url>
</inproceedings>
</r>
<r><article publtype="informal" key="journals/corr/abs-2503-05244" mdate="2025-10-08">
<author pid="301/4852-1">Yuning Wu 0001</author>
<author pid="354/8774">Jiahao Mei</author>
<author pid="51/5332-8">Ming Yan 0008</author>
<author pid="52/9457-3">Chenliang Li 0003</author>
<author pid="285/4166">Shaopeng Lai</author>
<author pid="401/8284">Yuran Ren</author>
<author pid="121/0900">Zijia Wang</author>
<author pid="86/1953-11">Ji Zhang 0011</author>
<author pid="82/2416">Mengyue Wu</author>
<author pid="47/2670">Qin Jin</author>
<author pid="h/FeiHuang-2">Fei Huang 0002</author>
<title>WritingBench: A Comprehensive Benchmark for Generative Writing.</title>
<year>2025</year>
<month>March</month>
<volume>abs/2503.05244</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2503.05244</ee>
<url>db/journals/corr/corr2503.html#abs-2503-05244</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2503-13377" mdate="2026-06-24">
<author pid="44/6292">Ye Wang</author>
<author pid="293/8958">Boshen Xu</author>
<author pid="339/2864">Zihao Yue</author>
<author pid="319/7252-1">Zihan Xiao 0001</author>
<author pid="79/10743">Ziheng Wang</author>
<author pid="50/6759">Liang Zhang</author>
<author pid="266/2264">Dingyi Yang</author>
<author pid="203/1536-1">Wenxuan Wang 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>TimeZero: Temporal Video Grounding with Reasoning-Guided LVLM.</title>
<year>2025</year>
<month>March</month>
<volume>abs/2503.13377</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2503.13377</ee>
<url>db/journals/corr/corr2503.html#abs-2503-13377</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2503-15470" mdate="2025-04-14">
<author pid="293/8958">Boshen Xu</author>
<author pid="334/2644">Yuting Mei</author>
<author pid="402/1495">Xinbi Liu</author>
<author pid="251/3691">Sipeng Zheng</author>
<author pid="47/2670">Qin Jin</author>
<title>EgoDTM: Towards 3D-Aware Egocentric Video-Language Pretraining.</title>
<year>2025</year>
<month>March</month>
<volume>abs/2503.15470</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2503.15470</ee>
<url>db/journals/corr/corr2503.html#abs-2503-15470</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2505-19125" mdate="2026-01-31">
<author pid="35/9071-3">Yuqi Liu 0003</author>
<author pid="47/2670">Qin Jin</author>
<author pid="366/0733">Tianyuan Qu</author>
<author pid="13/5407">Xuan Liu</author>
<author pid="51/3199-11">Yang Du 0011</author>
<author pid="28/4556-1">Bei Yu 0001</author>
<author pid="31/5649">Jiaya Jia</author>
<title>RTime-QA: A Benchmark for Atomic Temporal Event Understanding in Large Multi-modal Models.</title>
<year>2025</year>
<month>May</month>
<volume>abs/2505.19125</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2505.19125</ee>
<url>db/journals/corr/corr2505.html#abs-2505-19125</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2506-05947" mdate="2025-10-01">
<author pid="66/10435">Xinjie Zhang</author>
<author pid="203/1536-1">Wenxuan Wang 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>IntentionESC: An Intention-Centered Framework for Enhancing Emotional Support in Dialogue Systems.</title>
<year>2025</year>
<month>June</month>
<volume>abs/2506.05947</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2506.05947</ee>
<url>db/journals/corr/corr2506.html#abs-2506-05947</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2507-11936" mdate="2026-02-01">
<author orcid="0009-0006-7605-5923" pid="253/1100">Jianzhe Ma</author>
<author pid="203/1536-1">Wenxuan Wang 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>A Survey of Deep Learning for Geometry Problem Solving.</title>
<year>2025</year>
<month>July</month>
<volume>abs/2507.11936</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2507.11936</ee>
<url>db/journals/corr/corr2507.html#abs-2507-11936</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2507-11939" mdate="2026-04-08">
<author pid="177/6703-3">Yichen Xu 0003</author>
<author pid="82/353-8">Liangyu Chen 0008</author>
<author pid="50/6759">Liang Zhang</author>
<author pid="203/1536-1">Wenxuan Wang 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>POLYCHARTQA: Benchmarking Large Vision-Language Models with Multilingual Chart Question Answering.</title>
<year>2025</year>
<month>July</month>
<volume>abs/2507.11939</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2507.11939</ee>
<url>db/journals/corr/corr2507.html#abs-2507-11939</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2507-15597" mdate="2026-01-04">
<author pid="14/3727-11">Hao Luo 0011</author>
<author pid="340/4016">Yicheng Feng</author>
<author orcid="0000-0001-5351-3449" pid="73/10693-2">Wanpeng Zhang 0002</author>
<author pid="251/3691">Sipeng Zheng</author>
<author pid="44/6292">Ye Wang</author>
<author pid="254/2084">Haoqi Yuan</author>
<author pid="191/0307">Jiazheng Liu</author>
<author pid="358/4076">Chaoyi Xu</author>
<author pid="47/2670">Qin Jin</author>
<author pid="99/965-2">Zongqing Lu 0002</author>
<title>Being-H0: Vision-Language-Action Pretraining from Large-Scale Human Videos.</title>
<year>2025</year>
<month>July</month>
<volume>abs/2507.15597</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2507.15597</ee>
<url>db/journals/corr/corr2507.html#abs-2507-15597</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2507-21167" mdate="2026-04-08">
<author pid="413/7448">Donglu Yang</author>
<author pid="50/6759">Liang Zhang</author>
<author pid="339/2864">Zihao Yue</author>
<author pid="82/353-8">Liangyu Chen 0008</author>
<author pid="177/6703-3">Yichen Xu 0003</author>
<author pid="203/1536-1">Wenxuan Wang 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>ChartM<sup>3</sup>: Benchmarking Chart Editing with Multimodal Instructions.</title>
<year>2025</year>
<month>July</month>
<volume>abs/2507.21167</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2507.21167</ee>
<url>db/journals/corr/corr2507.html#abs-2507-21167</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2508-07863" mdate="2025-09-13">
<author pid="17/1169">Bin Cao</author>
<author pid="251/3691">Sipeng Zheng</author>
<author pid="44/6292">Ye Wang</author>
<author pid="352/3374">Lujie Xia</author>
<author pid="378/1061">Qianshan Wei</author>
<author pid="47/2670">Qin Jin</author>
<author pid="72/2590">Jing Liu</author>
<author pid="99/965-2">Zongqing Lu 0002</author>
<title>Being-M0.5: A Real-Time Controllable Vision-Language-Motion Model.</title>
<year>2025</year>
<month>August</month>
<volume>abs/2508.07863</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2508.07863</ee>
<url>db/journals/corr/corr2508.html#abs-2508-07863</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2510-01812" mdate="2025-11-08">
<author pid="358/9271">Yuxun Tang</author>
<author pid="43/6387">Lan Liu</author>
<author pid="50/9959">Wenhao Feng</author>
<author pid="122/3760">Yiwen Zhao</author>
<author pid="378/4662">Jionghao Han</author>
<author pid="09/9833">Yifeng Yu</author>
<author pid="229/3529">Jiatong Shi</author>
<author pid="47/2670">Qin Jin</author>
<title>SingMOS-Pro: An Comprehensive Benchmark for Singing Quality Assessment.</title>
<year>2025</year>
<month>October</month>
<volume>abs/2510.01812</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2510.01812</ee>
<url>db/journals/corr/corr2510.html#abs-2510-01812</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2510-20286" mdate="2025-11-28">
<author pid="82/353-8">Liangyu Chen 0008</author>
<author pid="295/8180">Hanzhang Zhou</author>
<author pid="238/3506">Chenglin Cai</author>
<author pid="69/10223">Jianan Zhang</author>
<author pid="223/6918">Panrong Tong</author>
<author pid="209/9882">Quyu Kong</author>
<author pid="98/5660">Xu Zhang</author>
<author pid="10/2639">Chen Liu</author>
<author pid="35/9071-3">Yuqi Liu 0003</author>
<author pid="203/1536-1">Wenxuan Wang 0001</author>
<author pid="33/4822-39">Yue Wang 0039</author>
<author pid="47/2670">Qin Jin</author>
<author pid="324/3832">Steven Hoi</author>
<title>UI-Ins: Enhancing GUI Grounding with Multi-Perspective Instruction-as-Reasoning.</title>
<year>2025</year>
<month>October</month>
<volume>abs/2510.20286</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2510.20286</ee>
<url>db/journals/corr/corr2510.html#abs-2510-20286</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2510-24134" mdate="2026-01-31">
<author pid="51/3199-11">Yang Du 0011</author>
<author pid="223/9450">Zhuoran Lin</author>
<author pid="222/2920">Kaiqiang Song</author>
<author pid="15/2887">Biao Wang</author>
<author pid="50/9014">Zhicheng Zheng</author>
<author pid="135/4944">Tiezheng Ge</author>
<author pid="33/1610-7">Bo Zheng 0007</author>
<author pid="47/2670">Qin Jin</author>
<title>VC4VG: Optimizing Video Captions for Text-to-Video Generation.</title>
<year>2025</year>
<month>October</month>
<volume>abs/2510.24134</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2510.24134</ee>
<url>db/journals/corr/corr2510.html#abs-2510-24134</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2511-13410" mdate="2026-01-15">
<author pid="305/3439">Zhaopei Huang</author>
<author pid="157/2769">Qifeng Dai</author>
<author pid="262/6203">Guozheng Wu</author>
<author pid="153/4201">Xiaopeng Wu</author>
<author pid="139/2261">Kehan Chen</author>
<author pid="50/790">Chuan Yu</author>
<author pid="228/7857">Xubin Li</author>
<author pid="135/4944">Tiezheng Ge</author>
<author pid="203/1536-1">Wenxuan Wang 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>Mem-PAL: Towards Memory-based Personalized Dialogue Assistants for Long-term User-Agent Interaction.</title>
<year>2025</year>
<month>November</month>
<volume>abs/2511.13410</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2511.13410</ee>
<url>db/journals/corr/corr2511.html#abs-2511-13410</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2511-15266" mdate="2026-04-08">
<author pid="82/353-8">Liangyu Chen 0008</author>
<author pid="177/6703-3">Yichen Xu 0003</author>
<author orcid="0009-0006-7605-5923" pid="253/1100">Jianzhe Ma</author>
<author pid="35/9071-3">Yuqi Liu 0003</author>
<author pid="413/7448">Donglu Yang</author>
<author pid="50/6759">Liang Zhang</author>
<author pid="203/1536-1">Wenxuan Wang 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>ChartEditor: A Reinforcement Learning Framework for Robust Chart Editing.</title>
<year>2025</year>
<month>November</month>
<volume>abs/2511.15266</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2511.15266</ee>
<url>db/journals/corr/corr2511.html#abs-2511-15266</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2511-16595" mdate="2026-06-24">
<author pid="293/8958">Boshen Xu</author>
<author pid="319/7252-1">Zihan Xiao 0001</author>
<author pid="157/2088">Jiaze Li</author>
<author pid="289/4197">Jianzhong Ju</author>
<author pid="152/8206">Zhenbo Luo</author>
<author pid="61/3233-1">Jian Luan 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>TimeViper: A Hybrid Mamba-Transformer Vision-Language Model for Efficient Long Video Understanding.</title>
<year>2025</year>
<month>November</month>
<volume>abs/2511.16595</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2511.16595</ee>
<url>db/journals/corr/corr2511.html#abs-2511-16595</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2512-01715" mdate="2026-01-24">
<author pid="73/10693-2">Wanpeng Zhang 0002</author>
<author pid="44/6292">Ye Wang</author>
<author pid="14/3727-11">Hao Luo 0011</author>
<author pid="254/2084">Haoqi Yuan</author>
<author pid="340/4016">Yicheng Feng</author>
<author pid="251/3691">Sipeng Zheng</author>
<author pid="47/2670">Qin Jin</author>
<author pid="99/965-2">Zongqing Lu 0002</author>
<title>DiG-Flow: Discrepancy-Guided Flow Matching for Robust VLA Models.</title>
<year>2025</year>
<month>December</month>
<volume>abs/2512.01715</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2512.01715</ee>
<url>db/journals/corr/corr2512.html#abs-2512-01715</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2512-05991" mdate="2026-01-23">
<author pid="52/5716">Chang Liu</author>
<author pid="402/6306">Tianjiao Jing</author>
<author pid="189/3741">Chengcheng Ma</author>
<author pid="399/8327">XuanQi Zhou</author>
<author pid="419/4007">Zhengxuan Lian</author>
<author pid="47/2670">Qin Jin</author>
<author pid="178/5601">Hongliang Yuan</author>
<author pid="58/8294">Shi-Sheng Huang</author>
<title>EmoDiffTalk:Emotion-aware Diffusion for Editable 3D Gaussian Talking Head.</title>
<year>2025</year>
<month>December</month>
<volume>abs/2512.05991</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2512.05991</ee>
<url>db/journals/corr/corr2512.html#abs-2512-05991</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2512-12839" mdate="2026-01-23">
<author pid="266/2264">Dingyi Yang</author>
<author pid="47/2670">Qin Jin</author>
<title>What Matters in Evaluating Book-Length Stories? A Systematic Study of Long Story Evaluation.</title>
<year>2025</year>
<month>December</month>
<volume>abs/2512.12839</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2512.12839</ee>
<url>db/journals/corr/corr2512.html#abs-2512-12839</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><phdthesis key="phd/basesearch/Jin24a" mdate="2024-10-15">
<author pid="47/2670">Qin Jin</author>
<title>Robust Speaker Recognition</title>
<school>Karlsruhe University, Germany</school>
<year>2024</year>
<note type="source">base-search.net (ftubkarlsruhe:oai:EVASTAR-Karlsruhe.de:1000166651)</note>
<ee>https://publikationen.bibliothek.kit.edu/1000166651</ee>
<ee>https://www.base-search.net/Record/e8c6928e829b7f3cdd8dcd18927845cfbfbd0817070a91e3db4ff222071cc07c</ee>
</phdthesis>
</r>
<r><article key="journals/tmm/ZengHPJ24" mdate="2026-07-14">
<author orcid="0000-0003-1908-1157" pid="240/5421">Yawen Zeng</author>
<author orcid="0000-0002-0654-6026" pid="68/690-5">Ning Han 0005</author>
<author orcid="0009-0008-4555-9146" pid="294/1953">Keyu Pan</author>
<author orcid="0000-0001-6486-6020" pid="47/2670">Qin Jin</author>
<title>Temporally Language Grounding With Multi-Modal Multi-Prompt Tuning.</title>
<pages>3366-3377</pages>
<year>2024</year>
<volume>26</volume>
<journal>IEEE Trans. Multim.</journal>
<ee>https://doi.org/10.1109/TMM.2023.3310282</ee>
<url>db/journals/tmm/tmm26.html#ZengHPJ24</url>
</article>
</r>
<r><inproceedings key="conf/acl/ZhangJHZW24" mdate="2025-01-19">
<author pid="50/6759">Liang Zhang</author>
<author pid="47/2670">Qin Jin</author>
<author pid="248/7736">Haoyang Huang</author>
<author pid="02/621">Dongdong Zhang</author>
<author pid="72/5870">Furu Wei</author>
<title>Respond in my Language: Mitigating Language Inconsistency in Response Generation based on Large Language Models.</title>
<pages>4177-4192</pages>
<year>2024</year>
<booktitle>ACL (1)</booktitle>
<ee type="oa">https://doi.org/10.18653/v1/2024.acl-long.229</ee>
<ee type="oa">https://aclanthology.org/2024.acl-long.229</ee>
<ee>https://www.wikidata.org/entity/Q131457856</ee>
<crossref>conf/acl/2024-1</crossref>
<url>db/conf/acl/acl2024-1.html#ZhangJHZW24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/acl/YangZWWG0J24" mdate="2025-01-19">
<author pid="266/2264">Dingyi Yang</author>
<author pid="335/2023">Chunru Zhan</author>
<author pid="79/10743">Ziheng Wang</author>
<author pid="15/2887">Biao Wang</author>
<author pid="135/4944">Tiezheng Ge</author>
<author pid="33/1610-7">Bo Zheng 0007</author>
<author pid="47/2670">Qin Jin</author>
<title>Synchronized Video Storytelling: Generating Video Narrations with Structured Storyline.</title>
<pages>9479-9493</pages>
<year>2024</year>
<booktitle>ACL (1)</booktitle>
<ee type="oa">https://doi.org/10.18653/v1/2024.acl-long.513</ee>
<ee type="oa">https://aclanthology.org/2024.acl-long.513</ee>
<ee>https://www.wikidata.org/entity/Q131458223</ee>
<crossref>conf/acl/2024-1</crossref>
<url>db/conf/acl/acl2024-1.html#YangZWWG0J24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/acl/YueZJ24" mdate="2024-09-24">
<author pid="339/2864">Zihao Yue</author>
<author pid="50/6759">Liang Zhang</author>
<author pid="47/2670">Qin Jin</author>
<title>Less is More: Mitigating Multimodal Hallucination from an EOS Decision Perspective.</title>
<pages>11766-11781</pages>
<year>2024</year>
<booktitle>ACL (1)</booktitle>
<ee type="oa">https://doi.org/10.18653/v1/2024.acl-long.633</ee>
<ee type="oa">https://aclanthology.org/2024.acl-long.633</ee>
<crossref>conf/acl/2024-1</crossref>
<url>db/conf/acl/acl2024-1.html#YueZJ24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/acl/ZhangZZZJ24" mdate="2026-04-07">
<author pid="305/3353">Tenggan Zhang</author>
<author orcid="0009-0001-7672-6778" pid="66/10435">Xinjie Zhang</author>
<author pid="121/8902">Jinming Zhao</author>
<author pid="54/40">Li Zhou</author>
<author pid="47/2670">Qin Jin</author>
<title>ESCoT: Towards Interpretable Emotional Support Dialogue Systems.</title>
<pages>13395-13412</pages>
<year>2024</year>
<booktitle>ACL (1)</booktitle>
<ee type="oa">https://doi.org/10.18653/v1/2024.acl-long.723</ee>
<ee type="oa">https://aclanthology.org/2024.acl-long.723</ee>
<ee>https://www.wikidata.org/entity/Q131458540</ee>
<crossref>conf/acl/2024-1</crossref>
<url>db/conf/acl/acl2024-1.html#ZhangZZZJ24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/eccv/ChenYXJ24" mdate="2025-07-08">
<author pid="82/353-8">Liangyu Chen 0008</author>
<author pid="339/2864">Zihao Yue</author>
<author pid="293/8958">Boshen Xu</author>
<author pid="47/2670">Qin Jin</author>
<title>Unveiling Visual Biases in Audio-Visual Localization Benchmarks.</title>
<pages>227-237</pages>
<year>2024</year>
<booktitle>ECCV Workshops (19)</booktitle>
<ee>https://doi.org/10.1007/978-3-031-93806-1_17</ee>
<crossref>conf/eccv/2024-w19</crossref>
<url>db/conf/eccv/eccv2024-w19.html#ChenYXJ24</url>
<stream>streams/conf/eccv</stream>
</inproceedings>
</r>
<r><inproceedings key="conf/emnlp/ZhangHXYXJZ024" mdate="2026-04-08">
<author pid="50/6759">Liang Zhang</author>
<author pid="249/1182">Anwen Hu</author>
<author pid="80/1339-1">Haiyang Xu 0001</author>
<author pid="51/5332-8">Ming Yan 0008</author>
<author pid="177/6703-3">Yichen Xu 0003</author>
<author pid="47/2670">Qin Jin</author>
<author pid="86/1953-11">Ji Zhang 0011</author>
<author pid="h/FeiHuang-2">Fei Huang 0002</author>
<title>TinyChart: Efficient Chart Understanding with Program-of-Thoughts Learning and Visual Token Merging.</title>
<pages>1882-1898</pages>
<year>2024</year>
<booktitle>EMNLP</booktitle>
<ee type="oa">https://doi.org/10.18653/v1/2024.emnlp-main.112</ee>
<ee type="oa">https://aclanthology.org/2024.emnlp-main.112</ee>
<crossref>conf/emnlp/2024</crossref>
<url>db/conf/emnlp/emnlp2024.html#ZhangHXYXJZ024</url>
</inproceedings>
</r>
<r><inproceedings key="conf/emnlp/HuXYYZZZJHZ24" mdate="2025-06-25">
<author pid="249/1182">Anwen Hu</author>
<author pid="80/1339-1">Haiyang Xu 0001</author>
<author pid="304/1336">Jiabo Ye</author>
<author pid="51/5332-8">Ming Yan 0008</author>
<author pid="50/6759">Liang Zhang</author>
<author pid="36/2259-71">Bo Zhang 0071</author>
<author pid="86/1953-11">Ji Zhang 0011</author>
<author pid="47/2670">Qin Jin</author>
<author pid="h/FeiHuang-2">Fei Huang 0002</author>
<author pid="84/2644-1">Jingren Zhou 0001</author>
<title>mPLUG-DocOwl 1.5: Unified Structure Learning for OCR-free Document Understanding.</title>
<pages>3096-3120</pages>
<year>2024</year>
<booktitle>EMNLP (Findings)</booktitle>
<ee type="oa">https://doi.org/10.18653/v1/2024.findings-emnlp.175</ee>
<ee type="oa">https://aclanthology.org/2024.findings-emnlp.175</ee>
<crossref>conf/emnlp/2024f</crossref>
<url>db/conf/emnlp/emnlp2024f.html#HuXYYZZZJHZ24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/emnlp/SunZJ24" mdate="2025-06-13">
<author pid="02/2264">Lei Sun</author>
<author pid="121/8902">Jinming Zhao</author>
<author pid="47/2670">Qin Jin</author>
<title>Revealing Personality Traits: A New Benchmark Dataset for Explainable Personality Recognition on Dialogues.</title>
<pages>19988-20002</pages>
<year>2024</year>
<booktitle>EMNLP</booktitle>
<ee type="oa">https://doi.org/10.18653/v1/2024.emnlp-main.1115</ee>
<ee type="oa">https://aclanthology.org/2024.emnlp-main.1115</ee>
<crossref>conf/emnlp/2024</crossref>
<url>db/conf/emnlp/emnlp2024.html#SunZJ24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icmcs/ZhangHZJ24" mdate="2025-12-07">
<author pid="192/1145">Fengyuan Zhang</author>
<author orcid="0000-0001-6325-791X" pid="305/3439">Zhaopei Huang</author>
<author pid="66/10435">Xinjie Zhang</author>
<author pid="47/2670">Qin Jin</author>
<title>Adaptive Temporal Motion Guided Graph Convolution Network for Micro-expression Recognition.</title>
<pages>1-6</pages>
<year>2024</year>
<booktitle>ICME</booktitle>
<ee>https://doi.org/10.1109/ICME57554.2024.10688180</ee>
<crossref>conf/icmcs/2024</crossref>
<url>db/conf/icmcs/icme2024.html#ZhangHZJ24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/ijcai/HuangZJ24" mdate="2024-10-18">
<author pid="305/3439">Zhaopei Huang</author>
<author pid="121/8902">Jinming Zhao</author>
<author pid="47/2670">Qin Jin</author>
<title>ECR-Chain: Advancing Generative Language Models to Better Emotion-Cause Reasoners through Reasoning Chains.</title>
<pages>6288-6296</pages>
<year>2024</year>
<booktitle>IJCAI</booktitle>
<ee type="oa">https://www.ijcai.org/proceedings/2024/695</ee>
<crossref>conf/ijcai/2024</crossref>
<url>db/conf/ijcai/ijcai2024.html#HuangZJ24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/ChangSTWTW0A0J24" mdate="2026-02-07">
<author pid="194/1149">Xuankai Chang</author>
<author pid="229/3529">Jiatong Shi</author>
<author pid="249/2901">Jinchuan Tian</author>
<author pid="301/4852-1">Yuning Wu 0001</author>
<author pid="358/9271">Yuxun Tang</author>
<author pid="72/10107-8">Yihan Wu 0008</author>
<author orcid="0000-0002-5970-8631" pid="39/3245-1">Shinji Watanabe 0001</author>
<author pid="171/0957">Yossi Adi</author>
<author pid="86/11429-1">Xie Chen 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>The Interspeech 2024 Challenge on Speech Processing Using Discrete Units.</title>
<year>2024</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2024-1878</ee>
<crossref>conf/interspeech/2024</crossref>
<url>db/conf/interspeech/interspeech2024.html#ChangSTWTW0A0J24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/ShiLBZWTYJ024" mdate="2025-12-25">
<author pid="229/3529">Jiatong Shi</author>
<author pid="349/5036">Yueqian Lin</author>
<author pid="221/6115">Xinyi Bai</author>
<author pid="164/6584">Keyi Zhang</author>
<author pid="301/4852-1">Yuning Wu 0001</author>
<author pid="358/9271">Yuxun Tang</author>
<author pid="09/9833">Yifeng Yu</author>
<author pid="47/2670">Qin Jin</author>
<author orcid="0000-0002-5970-8631" pid="39/3245-1">Shinji Watanabe 0001</author>
<title>Singing Voice Data Scaling-up: An Introduction to ACE-Opencpop and ACE-KiSing.</title>
<year>2024</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2024-33</ee>
<crossref>conf/interspeech/2024</crossref>
<url>db/conf/interspeech/interspeech2024.html#ShiLBZWTYJ024</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/TangWSJ24" mdate="2025-07-16">
<author pid="358/9271">Yuxun Tang</author>
<author pid="301/4852-1">Yuning Wu 0001</author>
<author pid="229/3529">Jiatong Shi</author>
<author pid="47/2670">Qin Jin</author>
<title>SingOMD: Singing Oriented Multi-resolution Discrete Representation Construction from Speech Models.</title>
<year>2024</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2024-2291</ee>
<crossref>conf/interspeech/2024</crossref>
<url>db/conf/interspeech/interspeech2024.html#TangWSJ24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/WuZSTYJ24" mdate="2025-07-16">
<author pid="301/4852-1">Yuning Wu 0001</author>
<author pid="65/6731">Chunlei Zhang</author>
<author pid="229/3529">Jiatong Shi</author>
<author pid="358/9271">Yuxun Tang</author>
<author pid="72/8479">Shan Yang</author>
<author pid="47/2670">Qin Jin</author>
<title>TokSing: Singing Voice Synthesis based on Discrete Tokens.</title>
<year>2024</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2024-2360</ee>
<crossref>conf/interspeech/2024</crossref>
<url>db/conf/interspeech/interspeech2024.html#WuZSTYJ24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/iscslp/WuYSQJ24" mdate="2025-07-16">
<author pid="301/4852-1">Yuning Wu 0001</author>
<author pid="09/9833">Yifeng Yu</author>
<author pid="229/3529">Jiatong Shi</author>
<author pid="05/5638">Tao Qian</author>
<author pid="47/2670">Qin Jin</author>
<title>A Systematic Exploration of Joint-Training for Singing Voice Synthesis.</title>
<pages>289-293</pages>
<year>2024</year>
<booktitle>ISCSLP</booktitle>
<ee>https://doi.org/10.1109/ISCSLP63861.2024.10799952</ee>
<crossref>conf/iscslp/2024</crossref>
<url>db/conf/iscslp/iscslp2024.html#WuYSQJ24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/iscslp/TangSWJ24" mdate="2025-07-16">
<author pid="358/9271">Yuxun Tang</author>
<author pid="229/3529">Jiatong Shi</author>
<author pid="301/4852-1">Yuning Wu 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>An Exploration on Singing MOS Prediction.</title>
<pages>651-655</pages>
<year>2024</year>
<booktitle>ISCSLP</booktitle>
<ee>https://doi.org/10.1109/ISCSLP63861.2024.10800519</ee>
<crossref>conf/iscslp/2024</crossref>
<url>db/conf/iscslp/iscslp2024.html#TangSWJ24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mir/MeiYJ24" mdate="2024-06-18">
<author orcid="0009-0009-6036-8687" pid="334/2644">Yuting Mei</author>
<author orcid="0000-0002-9809-8864" pid="262/6579">Linli Yao</author>
<author orcid="0000-0001-6486-6020" pid="47/2670">Qin Jin</author>
<title>UBiSS: A Unified Framework for Bimodal Semantic Summarization of Videos.</title>
<pages>1034-1042</pages>
<year>2024</year>
<booktitle>ICMR</booktitle>
<ee>https://doi.org/10.1145/3652583.3658038</ee>
<crossref>conf/mir/2024</crossref>
<url>db/conf/mir/icmr2024.html#MeiYJ24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/YaoZWHGJ0J24" mdate="2025-02-14">
<author orcid="0000-0002-9809-8864" pid="262/6579">Linli Yao</author>
<author orcid="0009-0009-4483-262X" pid="304/8398">Yuanmeng Zhang</author>
<author orcid="0009-0000-9366-994X" pid="79/10743">Ziheng Wang</author>
<author orcid="0009-0008-7065-4671" pid="319/3945">Xinglin Hou</author>
<author orcid="0000-0003-1381-2692" pid="135/4944">Tiezheng Ge</author>
<author orcid="0000-0003-1665-3025" pid="99/7950-1">Yuning Jiang 0001</author>
<author orcid="0000-0001-8241-9320" pid="37/1971-1">Xu Sun 0001</author>
<author orcid="0000-0001-6486-6020" pid="47/2670">Qin Jin</author>
<title>Edit As You Wish: Video Caption Editing with Multi-grained User Control.</title>
<pages>1924-1933</pages>
<year>2024</year>
<booktitle>ACM Multimedia</booktitle>
<ee>https://doi.org/10.1145/3664647.3680724</ee>
<ee>https://www.wikidata.org/entity/Q131165487</ee>
<crossref>conf/mm/2024</crossref>
<url>db/conf/mm/mm2024.html#YaoZWHGJ0J24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/Du0J24" mdate="2026-01-30">
<author orcid="0009-0007-8636-7258" pid="51/3199-11">Yang Du 0011</author>
<author orcid="0000-0002-2119-0881" pid="35/9071-3">Yuqi Liu 0003</author>
<author orcid="0000-0001-6486-6020" pid="47/2670">Qin Jin</author>
<title>Reversed in Time: A Novel Temporal-Emphasized Benchmark for Cross-Modal Video-Text Retrieval.</title>
<pages>5260-5269</pages>
<year>2024</year>
<booktitle>ACM Multimedia</booktitle>
<ee>https://doi.org/10.1145/3664647.3680731</ee>
<ee>https://www.wikidata.org/entity/Q131165512</ee>
<crossref>conf/mm/2024</crossref>
<url>db/conf/mm/mm2024.html#Du0J24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/WuSYTQLHB0J24" mdate="2025-07-16">
<author orcid="0009-0007-9859-1425" pid="301/4852-1">Yuning Wu 0001</author>
<author orcid="0000-0002-9050-8304" pid="229/3529">Jiatong Shi</author>
<author orcid="0009-0007-8286-5778" pid="09/9833">Yifeng Yu</author>
<author orcid="0009-0002-4538-7440" pid="358/9271">Yuxun Tang</author>
<author orcid="0009-0009-1479-9179" pid="05/5638">Tao Qian</author>
<author orcid="0000-0003-1473-8981" pid="349/5036">Yueqian Lin</author>
<author orcid="0009-0004-0604-4992" pid="378/4662">Jionghao Han</author>
<author orcid="0009-0005-5479-9692" pid="221/6115">Xinyi Bai</author>
<author orcid="0000-0002-5970-8631" pid="39/3245-1">Shinji Watanabe 0001</author>
<author orcid="0000-0001-6486-6020" pid="47/2670">Qin Jin</author>
<title>Muskits-ESPnet: A Comprehensive Toolkit for Singing Voice Synthesis in New Paradigm.</title>
<pages>11279-11281</pages>
<year>2024</year>
<booktitle>ACM Multimedia</booktitle>
<ee>https://doi.org/10.1145/3664647.3685000</ee>
<ee>https://www.wikidata.org/entity/Q131177205</ee>
<crossref>conf/mm/2024</crossref>
<url>db/conf/mm/mm2024.html#WuSYTQLHB0J24</url>
</inproceedings>
</r>
<r><inproceedings key="conf/slt/ShiTWJYMCWTBAZDSWLRJSW24" mdate="2026-02-05">
<author pid="229/3529">Jiatong Shi</author>
<author pid="249/2901">Jinchuan Tian</author>
<author pid="72/10107-8">Yihan Wu 0008</author>
<author pid="212/6435">Jee-Weon Jung</author>
<author pid="329/5838">Jia Qi Yip</author>
<author pid="226/5828">Yoshiki Masuyama</author>
<author pid="82/2443">William Chen</author>
<author pid="301/4852-1">Yuning Wu 0001</author>
<author pid="358/9271">Yuxun Tang</author>
<author pid="253/3004">Massa Baali</author>
<author pid="358/5007">Dareen Alharthi</author>
<author pid="68/3245">Dong Zhang</author>
<author pid="388/3709">Ruifan Deng</author>
<author pid="332/5917">Tejes Srivastava</author>
<author pid="95/10219">Haibin Wu</author>
<author pid="227/2380">Alexander H. Liu</author>
<author pid="60/3996">Bhiksha Raj</author>
<author pid="47/2670">Qin Jin</author>
<author pid="s/RuihuaSong">Ruihua Song</author>
<author orcid="0000-0002-5970-8631" pid="39/3245-1">Shinji Watanabe 0001</author>
<title>ESPnet-Codec: Comprehensive Training and Evaluation of Neural Codecs For Audio, Music, and Speech.</title>
<pages>562-569</pages>
<year>2024</year>
<booktitle>SLT</booktitle>
<ee>https://doi.org/10.1109/SLT61566.2024.10832289</ee>
<crossref>conf/slt/2024</crossref>
<url>db/conf/slt/slt2024.html#ShiTWJYMCWTBAZDSWLRJSW24</url>
<stream>streams/conf/slt</stream>
</inproceedings>
</r>
<r><inproceedings key="conf/trec/WangD0J24" mdate="2026-01-31">
<author pid="37/5333">Xueyan Wang</author>
<author pid="51/3199-11">Yang Du 0011</author>
<author pid="35/9071-3">Yuqi Liu 0003</author>
<author pid="47/2670">Qin Jin</author>
<title>RUC_AIM3 at TRECVID 2024: Ad-hoc Video Search.</title>
<year>2024</year>
<booktitle>TREC</booktitle>
<ee type="oa">https://trec.nist.gov/pubs/trec33/papers/ruc_aim3.avs.pdf</ee>
<crossref>conf/trec/2024</crossref>
<url>db/conf/trec/trec2024.html#WangD0J24</url>
</inproceedings>
</r>
<r><proceedings key="conf/iscslp/2024" mdate="2025-11-13">
<editor pid="07/8638">Yanmin Qian</editor>
<editor pid="47/2670">Qin Jin</editor>
<editor pid="44/5256">Zhijian Ou</editor>
<editor pid="70/5210">Zhenhua Ling</editor>
<editor pid="24/968-1">Zhiyong Wu 0001</editor>
<editor pid="51/4056-1">Ya Li 0001</editor>
<editor pid="70/1741-1">Lei Xie 0001</editor>
<editor pid="46/2916-1">Jianhua Tao 0001</editor>
<title>14th IEEE International Symposium on Chinese Spoken Language Processing, ISCSLP 2024, Beijing, China, November 7-10, 2024</title>
<booktitle>ISCSLP</booktitle>
<publisher>IEEE</publisher>
<year>2024</year>
<isbn>979-8-3315-1682-6</isbn>
<ee>https://doi.org/10.1109/ISCSLP63861.2024</ee>
<url>db/conf/iscslp/iscslp2024.html</url>
</proceedings>
</r>
<r><article publtype="informal" key="journals/corr/abs-2401-17619" mdate="2025-07-16">
<author pid="229/3529">Jiatong Shi</author>
<author pid="349/5036">Yueqian Lin</author>
<author pid="221/6115">Xinyi Bai</author>
<author pid="164/6584">Keyi Zhang</author>
<author pid="301/4852-1">Yuning Wu 0001</author>
<author pid="358/9271">Yuxun Tang</author>
<author pid="09/9833">Yifeng Yu</author>
<author pid="47/2670">Qin Jin</author>
<author orcid="0000-0002-5970-8631" pid="39/3245-1">Shinji Watanabe 0001</author>
<title>Singing Voice Data Scaling-up: An Introduction to ACE-Opencpop and KiSing-v2.</title>
<year>2024</year>
<volume>abs/2401.17619</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2401.17619</ee>
<url>db/journals/corr/corr2401.html#abs-2401-17619</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2402-14545" mdate="2024-03-22">
<author pid="339/2864">Zihao Yue</author>
<author pid="50/6759">Liang Zhang</author>
<author pid="47/2670">Qin Jin</author>
<title>Less is More: Mitigating Multimodal Hallucination from an EOS Decision Perspective.</title>
<year>2024</year>
<volume>abs/2402.14545</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2402.14545</ee>
<url>db/journals/corr/corr2402.html#abs-2402-14545</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2403-05856" mdate="2024-04-04">
<author pid="293/8958">Boshen Xu</author>
<author pid="251/3691">Sipeng Zheng</author>
<author pid="47/2670">Qin Jin</author>
<title>POV: Prompt-Oriented View-Agnostic Learning for Egocentric Hand-Object Interaction in the Multi-View World.</title>
<year>2024</year>
<volume>abs/2403.05856</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2403.05856</ee>
<url>db/journals/corr/corr2403.html#abs-2403-05856</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2403-05874" mdate="2024-04-04">
<author pid="293/8958">Boshen Xu</author>
<author pid="251/3691">Sipeng Zheng</author>
<author pid="47/2670">Qin Jin</author>
<title>SPAFormer: Sequential 3D Part Assembly with Transformers.</title>
<year>2024</year>
<volume>abs/2403.05874</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2403.05874</ee>
<url>db/journals/corr/corr2403.html#abs-2403-05874</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2403-12895" mdate="2025-06-25">
<author pid="249/1182">Anwen Hu</author>
<author pid="80/1339-1">Haiyang Xu 0001</author>
<author pid="304/1336">Jiabo Ye</author>
<author pid="51/5332-8">Ming Yan 0008</author>
<author pid="50/6759">Liang Zhang</author>
<author pid="36/2259-71">Bo Zhang 0071</author>
<author pid="l/ChenLi1">Chen Li 0001</author>
<author pid="86/1953-11">Ji Zhang 0011</author>
<author pid="47/2670">Qin Jin</author>
<author pid="h/FeiHuang-2">Fei Huang 0002</author>
<author pid="84/2644-1">Jingren Zhou 0001</author>
<title>mPLUG-DocOwl 1.5: Unified Structure Learning for OCR-free Document Understanding.</title>
<year>2024</year>
<volume>abs/2403.12895</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2403.12895</ee>
<url>db/journals/corr/corr2403.html#abs-2403-12895</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2404-13370" mdate="2024-05-25">
<author pid="339/2864">Zihao Yue</author>
<author pid="355/8240">Yepeng Zhang</author>
<author pid="79/10743">Ziheng Wang</author>
<author pid="47/2670">Qin Jin</author>
<title>Movie101v2: Improved Movie Narration Benchmark.</title>
<year>2024</year>
<volume>abs/2404.13370</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2404.13370</ee>
<url>db/journals/corr/corr2404.html#abs-2404-13370</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2404-14705" mdate="2024-05-25">
<author pid="376/0782">Qingrong He</author>
<author pid="251/8268">Kejun Lin</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="249/1182">Anwen Hu</author>
<author pid="47/2670">Qin Jin</author>
<title>Think-Program-reCtify: 3D Situated Reasoning with Large Language Models.</title>
<year>2024</year>
<volume>abs/2404.14705</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2404.14705</ee>
<url>db/journals/corr/corr2404.html#abs-2404-14705</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2404-16635" mdate="2026-04-08">
<author pid="50/6759">Liang Zhang</author>
<author pid="249/1182">Anwen Hu</author>
<author pid="80/1339-1">Haiyang Xu 0001</author>
<author pid="51/5332-8">Ming Yan 0008</author>
<author pid="177/6703-3">Yichen Xu 0003</author>
<author pid="47/2670">Qin Jin</author>
<author pid="86/1953-11">Ji Zhang 0011</author>
<author pid="h/FeiHuang-2">Fei Huang 0002</author>
<title>TinyChart: Efficient Chart Understanding with Visual Token Merging and Program-of-Thoughts Learning.</title>
<year>2024</year>
<volume>abs/2404.16635</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2404.16635</ee>
<url>db/journals/corr/corr2404.html#abs-2404-16635</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2405-10860" mdate="2024-06-12">
<author pid="305/3439">Zhaopei Huang</author>
<author pid="121/8902">Jinming Zhao</author>
<author pid="47/2670">Qin Jin</author>
<title>ECR-Chain: Advancing Generative Language Models to Better Emotion-Cause Reasoners through Reasoning Chains.</title>
<year>2024</year>
<volume>abs/2405.10860</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2405.10860</ee>
<url>db/journals/corr/corr2405.html#abs-2405-10860</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2405-14040" mdate="2024-06-19">
<author pid="266/2264">Dingyi Yang</author>
<author pid="335/2023">Chunru Zhan</author>
<author pid="79/10743">Ziheng Wang</author>
<author pid="15/2887">Biao Wang</author>
<author pid="135/4944">Tiezheng Ge</author>
<author pid="33/1610-7">Bo Zheng 0007</author>
<author pid="47/2670">Qin Jin</author>
<title>Synchronized Video Storytelling: Generating Video Narrations with Structured Storyline.</title>
<year>2024</year>
<volume>abs/2405.14040</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2405.14040</ee>
<url>db/journals/corr/corr2405.html#abs-2405-14040</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2405-17719" mdate="2026-01-31">
<author pid="293/8958">Boshen Xu</author>
<author pid="79/10743">Ziheng Wang</author>
<author pid="51/3199-11">Yang Du 0011</author>
<author pid="378/1643">Zhinan Song</author>
<author pid="251/3691">Sipeng Zheng</author>
<author pid="47/2670">Qin Jin</author>
<title>EgoNCE++: Do Egocentric Video-Language Models Really Understand Hand-Object Interactions?</title>
<year>2024</year>
<volume>abs/2405.17719</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2405.17719</ee>
<url>db/journals/corr/corr2405.html#abs-2405-17719</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2406-07725" mdate="2026-02-07">
<author pid="194/1149">Xuankai Chang</author>
<author pid="229/3529">Jiatong Shi</author>
<author pid="249/2901">Jinchuan Tian</author>
<author pid="301/4852-1">Yuning Wu 0001</author>
<author pid="358/9271">Yuxun Tang</author>
<author pid="72/10107-8">Yihan Wu 0008</author>
<author orcid="0000-0002-5970-8631" pid="39/3245-1">Shinji Watanabe 0001</author>
<author pid="171/0957">Yossi Adi</author>
<author pid="86/11429-1">Xie Chen 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>The Interspeech 2024 Challenge on Speech Processing Using Discrete Units.</title>
<year>2024</year>
<volume>abs/2406.07725</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2406.07725</ee>
<url>db/journals/corr/corr2406.html#abs-2406-07725</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2406-08416" mdate="2025-07-16">
<author pid="301/4852-1">Yuning Wu 0001</author>
<author pid="65/6731">Chunlei Zhang</author>
<author pid="229/3529">Jiatong Shi</author>
<author pid="358/9271">Yuxun Tang</author>
<author pid="72/8479">Shan Yang</author>
<author pid="47/2670">Qin Jin</author>
<title>TokSing: Singing Voice Synthesis based on Discrete Tokens.</title>
<year>2024</year>
<volume>abs/2406.08416</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2406.08416</ee>
<url>db/journals/corr/corr2406.html#abs-2406-08416</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2406-08905" mdate="2025-07-16">
<author pid="358/9271">Yuxun Tang</author>
<author pid="301/4852-1">Yuning Wu 0001</author>
<author pid="229/3529">Jiatong Shi</author>
<author pid="47/2670">Qin Jin</author>
<title>SingOMD: Singing Oriented Multi-resolution Discrete Representation Construction from Speech Models.</title>
<year>2024</year>
<volume>abs/2406.08905</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2406.08905</ee>
<url>db/journals/corr/corr2406.html#abs-2406-08905</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2406-08997" mdate="2024-07-09">
<author pid="192/1145">Fengyuan Zhang</author>
<author pid="305/3439">Zhaopei Huang</author>
<author pid="66/10435">Xinjie Zhang</author>
<author pid="47/2670">Qin Jin</author>
<title>Adaptive Temporal Motion Guided Graph Convolution Network for Micro-expression Recognition.</title>
<year>2024</year>
<volume>abs/2406.08997</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2406.08997</ee>
<url>db/journals/corr/corr2406.html#abs-2406-08997</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2406-10911" mdate="2025-07-16">
<author pid="358/9271">Yuxun Tang</author>
<author pid="229/3529">Jiatong Shi</author>
<author pid="301/4852-1">Yuning Wu 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>SingMOS: An extensive Open-Source Singing Voice Dataset for MOS Prediction.</title>
<year>2024</year>
<volume>abs/2406.10911</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2406.10911</ee>
<url>db/journals/corr/corr2406.html#abs-2406-10911</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2406-10960" mdate="2024-07-18">
<author pid="305/3353">Tenggan Zhang</author>
<author pid="66/10435">Xinjie Zhang</author>
<author pid="121/8902">Jinming Zhao</author>
<author pid="54/40">Li Zhou</author>
<author pid="47/2670">Qin Jin</author>
<title>ESCoT: Towards Interpretable Emotional Support Dialogue Systems.</title>
<year>2024</year>
<volume>abs/2406.10960</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2406.10960</ee>
<url>db/journals/corr/corr2406.html#abs-2406-10960</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2406-16301" mdate="2024-07-16">
<author pid="334/2644">Yuting Mei</author>
<author pid="262/6579">Linli Yao</author>
<author pid="47/2670">Qin Jin</author>
<title>UBiSS: A Unified Framework for Bimodal Semantic Summarization of Videos.</title>
<year>2024</year>
<volume>abs/2406.16301</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2406.16301</ee>
<url>db/journals/corr/corr2406.html#abs-2406-16301</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2406-16578" mdate="2024-07-16">
<author pid="44/6292">Ye Wang</author>
<author pid="334/2644">Yuting Mei</author>
<author pid="251/3691">Sipeng Zheng</author>
<author pid="47/2670">Qin Jin</author>
<title>QuadrupedGPT: Towards a Versatile Quadruped Agent in Open-ended Worlds.</title>
<year>2024</year>
<volume>abs/2406.16578</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2406.16578</ee>
<url>db/journals/corr/corr2406.html#abs-2406-16578</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2408-14622" mdate="2024-09-28">
<author pid="266/2264">Dingyi Yang</author>
<author pid="47/2670">Qin Jin</author>
<title>What Makes a Good Story and How Can We Measure It? A Comprehensive Survey of Story Evaluation.</title>
<year>2024</year>
<volume>abs/2408.14622</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2408.14622</ee>
<url>db/journals/corr/corr2408.html#abs-2408-14622</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2409-03420" mdate="2025-06-25">
<author pid="249/1182">Anwen Hu</author>
<author pid="80/1339-1">Haiyang Xu 0001</author>
<author pid="50/6759">Liang Zhang</author>
<author pid="304/1336">Jiabo Ye</author>
<author pid="51/5332-8">Ming Yan 0008</author>
<author pid="86/1953-11">Ji Zhang 0011</author>
<author pid="47/2670">Qin Jin</author>
<author pid="h/FeiHuang-2">Fei Huang 0002</author>
<author pid="84/2644-1">Jingren Zhou 0001</author>
<title>mPLUG-DocOwl2: High-resolution Compressing for OCR-free Multi-page Document Understanding.</title>
<year>2024</year>
<volume>abs/2409.03420</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2409.03420</ee>
<url>db/journals/corr/corr2409.html#abs-2409-03420</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2409-06709" mdate="2025-07-08">
<author pid="82/353-8">Liangyu Chen 0008</author>
<author pid="339/2864">Zihao Yue</author>
<author pid="293/8958">Boshen Xu</author>
<author pid="47/2670">Qin Jin</author>
<title>Unveiling Visual Biases in Audio-Visual Localization Benchmarks.</title>
<year>2024</year>
<volume>abs/2409.06709</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2409.06709</ee>
<url>db/journals/corr/corr2409.html#abs-2409-06709</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2409-07226" mdate="2025-12-25">
<author pid="301/4852-1">Yuning Wu 0001</author>
<author pid="229/3529">Jiatong Shi</author>
<author pid="09/9833">Yifeng Yu</author>
<author pid="358/9271">Yuxun Tang</author>
<author pid="05/5638">Tao Qian</author>
<author pid="349/5036">Yueqian Lin</author>
<author pid="378/4662">Jionghao Han</author>
<author pid="221/6115">Xinyi Bai</author>
<author orcid="0000-0002-5970-8631" pid="39/3245-1">Shinji Watanabe 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>Muskits-ESPnet: A Comprehensive Toolkit for Singing Voice Synthesis in New Paradigm.</title>
<year>2024</year>
<volume>abs/2409.07226</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2409.07226</ee>
<url>db/journals/corr/corr2409.html#abs-2409-07226</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2409-15897" mdate="2026-02-05">
<author pid="229/3529">Jiatong Shi</author>
<author pid="249/2901">Jinchuan Tian</author>
<author pid="72/10107-8">Yihan Wu 0008</author>
<author pid="212/6435">Jee-weon Jung</author>
<author pid="329/5838">Jia Qi Yip</author>
<author pid="226/5828">Yoshiki Masuyama</author>
<author pid="82/2443">William Chen</author>
<author pid="301/4852-1">Yuning Wu 0001</author>
<author pid="358/9271">Yuxun Tang</author>
<author pid="253/3004">Massa Baali</author>
<author pid="358/5007">Dareen Alharthi</author>
<author pid="68/3245">Dong Zhang</author>
<author pid="388/3709">Ruifan Deng</author>
<author pid="332/5917">Tejes Srivastava</author>
<author pid="95/10219">Haibin Wu</author>
<author pid="227/2380">Alexander H. Liu</author>
<author pid="60/3996">Bhiksha Raj</author>
<author pid="47/2670">Qin Jin</author>
<author pid="s/RuihuaSong">Ruihua Song</author>
<author orcid="0000-0002-5970-8631" pid="39/3245-1">Shinji Watanabe 0001</author>
<title>ESPnet-Codec: Comprehensive Training and Evaluation of Neural Codecs for Audio, Music, and Speech.</title>
<year>2024</year>
<volume>abs/2409.15897</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2409.15897</ee>
<url>db/journals/corr/corr2409.html#abs-2409-15897</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2409-19723" mdate="2024-10-18">
<author pid="02/2264">Lei Sun</author>
<author pid="121/8902">Jinming Zhao</author>
<author pid="47/2670">Qin Jin</author>
<title>Revealing Personality Traits: A New Benchmark Dataset for Explainable Personality Recognition on Dialogues.</title>
<year>2024</year>
<volume>abs/2409.19723</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2409.19723</ee>
<url>db/journals/corr/corr2409.html#abs-2409-19723</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2410-03311" mdate="2025-04-14">
<author pid="44/6292">Ye Wang</author>
<author pid="251/3691">Sipeng Zheng</author>
<author pid="17/1169">Bin Cao</author>
<author pid="378/1061">Qianshan Wei</author>
<author pid="47/2670">Qin Jin</author>
<author pid="99/965-2">Zongqing Lu 0002</author>
<title>Quo Vadis, Motion Generation? From Large Language Models to Large Motion Models.</title>
<year>2024</year>
<volume>abs/2410.03311</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2410.03311</ee>
<url>db/journals/corr/corr2410.html#abs-2410-03311</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2412-19178" mdate="2026-01-31">
<author pid="51/3199-11">Yang Du 0011</author>
<author pid="35/9071-3">Yuqi Liu 0003</author>
<author pid="47/2670">Qin Jin</author>
<title>Reversed in Time: A Novel Temporal-Emphasized Benchmark for Cross-Modal Video-Text Retrieval.</title>
<year>2024</year>
<volume>abs/2412.19178</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2412.19178</ee>
<url>db/journals/corr/corr2412.html#abs-2412-19178</url>
<stream>streams/journals/corr</stream>
</article>
</r>
<r><article key="journals/ijautcomp/ZhangRHJ23" mdate="2023-04-29">
<author orcid="0000-0002-6187-3628" pid="50/6759">Liang Zhang</author>
<author pid="262/6356">Ludan Ruan</author>
<author pid="249/1182">Anwen Hu</author>
<author orcid="0000-0001-6486-6020" pid="47/2670">Qin Jin</author>
<title>Multimodal Pretraining from Monolingual to Multilingual.</title>
<pages>220-232</pages>
<year>2023</year>
<month>April</month>
<volume>20</volume>
<journal>Mach. Intell. Res.</journal>
<number>2</number>
<ee>https://doi.org/10.1007/s11633-022-1414-4</ee>
<url>db/journals/ijautcomp/ijautcomp20.html#ZhangRHJ23</url>
</article>
</r>
<r><article key="journals/staeors/ZhangLJMYHHHCL23" mdate="2023-01-31">
<author pid="02/6428-12">Yun Zhang 0012</author>
<author pid="41/4012">Qi Lu</author>
<author pid="47/2670">Qin Jin</author>
<author pid="142/6162">Wanting Meng</author>
<author orcid="0000-0001-9967-7756" pid="253/4903">Shuhu Yang</author>
<author pid="90/839">Shen Huang</author>
<author pid="47/468">Yanling Han</author>
<author orcid="0000-0003-0045-1066" pid="142/6255">Zhonghua Hong</author>
<author pid="176/4824">Zhansheng Chen</author>
<author pid="77/6507">Weiliang Liu</author>
<title>Global Sea Surface Height Measurement From CYGNSS Based on Machine Learning.</title>
<pages>841-852</pages>
<year>2023</year>
<volume>16</volume>
<journal>IEEE J. Sel. Top. Appl. Earth Obs. Remote. Sens.</journal>
<ee type="oa">https://doi.org/10.1109/JSTARS.2022.3231916</ee>
<url>db/journals/staeors/staeors16.html#ZhangLJMYHHHCL23</url>
</article>
</r>
<r><inproceedings key="conf/aaai/LiuXXJ23" mdate="2024-10-25">
<author pid="35/9071-3">Yuqi Liu 0003</author>
<author pid="229/5603">Luhui Xu</author>
<author pid="48/8617">Pengfei Xiong</author>
<author pid="47/2670">Qin Jin</author>
<title>Token Mixing: Parameter-Efficient Transfer Learning from Image-Language to Video-Language.</title>
<pages>1781-1789</pages>
<year>2023</year>
<booktitle>AAAI</booktitle>
<ee type="oa">https://doi.org/10.1609/aaai.v37i2.25267</ee>
<crossref>conf/aaai/2023</crossref>
<url>db/conf/aaai/aaai2023.html#LiuXXJ23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/aaai/ZengJBL23" mdate="2023-09-04">
<author pid="240/5421">Yawen Zeng</author>
<author pid="47/2670">Qin Jin</author>
<author pid="38/8548">Tengfei Bao</author>
<author pid="20/179">Wenfeng Li</author>
<title>Multi-Modal Knowledge Hypergraph for Diverse Image Retrieval.</title>
<pages>3376-3383</pages>
<year>2023</year>
<booktitle>AAAI</booktitle>
<ee type="oa">https://doi.org/10.1609/aaai.v37i3.25445</ee>
<crossref>conf/aaai/2023</crossref>
<url>db/conf/aaai/aaai2023.html#ZengJBL23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/aaai/RuanH0ZZJ23" mdate="2023-09-04">
<author pid="262/6356">Ludan Ruan</author>
<author pid="249/1182">Anwen Hu</author>
<author pid="222/3101">Yuqing Song 0003</author>
<author pid="50/6759">Liang Zhang</author>
<author pid="251/3691">Sipeng Zheng</author>
<author pid="47/2670">Qin Jin</author>
<title>Accommodating Audio Modality in CLIP for Multimodal Processing.</title>
<pages>9641-9649</pages>
<year>2023</year>
<booktitle>AAAI</booktitle>
<ee type="oa">https://doi.org/10.1609/aaai.v37i8.26153</ee>
<crossref>conf/aaai/2023</crossref>
<url>db/conf/aaai/aaai2023.html#RuanH0ZZJ23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/aaai/ZhangHZHJ23" mdate="2023-09-04">
<author pid="50/6759">Liang Zhang</author>
<author pid="249/1182">Anwen Hu</author>
<author pid="05/3499">Jing Zhang</author>
<author pid="35/7808">Shuo Hu</author>
<author pid="47/2670">Qin Jin</author>
<title>MPMQA: Multimodal Question Answering on Product Manuals.</title>
<pages>13958-13966</pages>
<year>2023</year>
<booktitle>AAAI</booktitle>
<ee type="oa">https://doi.org/10.1609/aaai.v37i11.26634</ee>
<crossref>conf/aaai/2023</crossref>
<url>db/conf/aaai/aaai2023.html#ZhangHZHJ23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/acl/QianLSWGYJ23" mdate="2025-07-16">
<author pid="05/5638">Tao Qian</author>
<author pid="282/9088">Fan Lou</author>
<author pid="229/3529">Jiatong Shi</author>
<author pid="301/4852-1">Yuning Wu 0001</author>
<author pid="07/272">Shuai Guo</author>
<author pid="18/1022-6">Xiang Yin 0006</author>
<author pid="47/2670">Qin Jin</author>
<title>UniLG: A Unified Structure-aware Framework for Lyrics Generation.</title>
<pages>983-1001</pages>
<year>2023</year>
<booktitle>ACL (1)</booktitle>
<ee type="oa">https://doi.org/10.18653/v1/2023.acl-long.56</ee>
<ee type="oa">https://aclanthology.org/2023.acl-long.56</ee>
<ee>https://www.wikidata.org/entity/Q131458009</ee>
<crossref>conf/acl/2023-1</crossref>
<url>db/conf/acl/acl2023-1.html#QianLSWGYJ23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/acl/HuCZJ23" mdate="2025-01-19">
<author pid="249/1182">Anwen Hu</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="50/6759">Liang Zhang</author>
<author pid="47/2670">Qin Jin</author>
<title>InfoMetIC: An Informative Metric for Reference-free Image Caption Evaluation.</title>
<pages>3171-3185</pages>
<year>2023</year>
<booktitle>ACL (1)</booktitle>
<ee type="oa">https://doi.org/10.18653/v1/2023.acl-long.178</ee>
<ee type="oa">https://aclanthology.org/2023.acl-long.178</ee>
<ee>https://www.wikidata.org/entity/Q131458131</ee>
<crossref>conf/acl/2023-1</crossref>
<url>db/conf/acl/acl2023-1.html#HuCZJ23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/acl/YueZHZWJ23" mdate="2023-08-10">
<author pid="339/2864">Zihao Yue</author>
<author pid="52/323">Qi Zhang</author>
<author pid="249/1182">Anwen Hu</author>
<author pid="50/6759">Liang Zhang</author>
<author pid="79/10743">Ziheng Wang</author>
<author pid="47/2670">Qin Jin</author>
<title>Movie101: A New Movie Understanding Benchmark.</title>
<pages>4669-4684</pages>
<year>2023</year>
<booktitle>ACL (1)</booktitle>
<ee type="oa">https://doi.org/10.18653/v1/2023.acl-long.257</ee>
<ee type="oa">https://aclanthology.org/2023.acl-long.257</ee>
<crossref>conf/acl/2023-1</crossref>
<url>db/conf/acl/acl2023-1.html#YueZHZWJ23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/acl/YangJ23" mdate="2023-08-10">
<author pid="266/2264">Dingyi Yang</author>
<author pid="47/2670">Qin Jin</author>
<title>Attractive Storyteller: Stylized Visual Storytelling with Unpaired Text.</title>
<pages>11053-11066</pages>
<year>2023</year>
<booktitle>ACL (1)</booktitle>
<ee type="oa">https://doi.org/10.18653/v1/2023.acl-long.619</ee>
<ee type="oa">https://aclanthology.org/2023.acl-long.619</ee>
<crossref>conf/acl/2023-1</crossref>
<url>db/conf/acl/acl2023-1.html#YangJ23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/cvpr/RuanMYH0FYJG23" mdate="2025-03-03">
<author pid="262/6356">Ludan Ruan</author>
<author pid="324/2590">Yiyang Ma</author>
<author pid="86/4843-5">Huan Yang 0005</author>
<author orcid="0000-0003-1419-059X" pid="270/6402">Huiguo He</author>
<author pid="39/3711-1">Bei Liu 0001</author>
<author pid="83/8692">Jianlong Fu</author>
<author pid="131/4855">Nicholas Jing Yuan</author>
<author pid="47/2670">Qin Jin</author>
<author pid="44/3946">Baining Guo</author>
<title>MM-Diffusion: Learning Multi-Modal Diffusion Models for Joint Audio and Video Generation.</title>
<pages>10219-10228</pages>
<year>2023</year>
<booktitle>CVPR</booktitle>
<ee>https://doi.org/10.1109/CVPR52729.2023.00985</ee>
<crossref>conf/cvpr/2023</crossref>
<url>db/conf/cvpr/cvpr2023.html#RuanMYH0FYJG23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/cvpr/ZhengXJ23" mdate="2023-11-12">
<author pid="251/3691">Sipeng Zheng</author>
<author orcid="0009-0000-1896-9600" pid="293/8958">Boshen Xu</author>
<author pid="47/2670">Qin Jin</author>
<title>Open-Category Human-Object Interaction Pre-training via Language Modeling Framework.</title>
<pages>19392-19402</pages>
<year>2023</year>
<booktitle>CVPR</booktitle>
<ee>https://doi.org/10.1109/CVPR52729.2023.01858</ee>
<crossref>conf/cvpr/2023</crossref>
<url>db/conf/cvpr/cvpr2023.html#ZhengXJ23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/emnlp/YeHXYYXLT0ZJHLH23" mdate="2025-10-08">
<author pid="304/1336">Jiabo Ye</author>
<author pid="249/1182">Anwen Hu</author>
<author pid="80/1339-1">Haiyang Xu 0001</author>
<author pid="254/3247">Qinghao Ye</author>
<author pid="51/5332-8">Ming Yan 0008</author>
<author pid="205/7621">Guohai Xu</author>
<author pid="52/9457-3">Chenliang Li 0003</author>
<author pid="93/1076">Junfeng Tian</author>
<author pid="05/2084-1">Qi Qian 0001</author>
<author pid="86/1953-11">Ji Zhang 0011</author>
<author pid="47/2670">Qin Jin</author>
<author pid="42/963-1">Liang He 0001</author>
<author pid="50/3323-1">Xin Lin 0001</author>
<author pid="h/FeiHuang-2">Fei Huang 0002</author>
<title>UReader: Universal OCR-free Visually-situated Language Understanding with Multimodal Large Language Model.</title>
<pages>2841-2858</pages>
<year>2023</year>
<booktitle>EMNLP (Findings)</booktitle>
<ee type="oa">https://doi.org/10.18653/v1/2023.findings-emnlp.187</ee>
<ee type="oa">https://aclanthology.org/2023.findings-emnlp.187</ee>
<crossref>conf/emnlp/2023f</crossref>
<url>db/conf/emnlp/emnlp2023f.html#YeHXYYXLT0ZJHLH23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/WuSQGJ23" mdate="2025-07-16">
<author pid="301/4852-1">Yuning Wu 0001</author>
<author pid="229/3529">Jiatong Shi</author>
<author pid="05/5638">Tao Qian</author>
<author pid="226/5309">Dongji Gao</author>
<author pid="47/2670">Qin Jin</author>
<title>Phoneix: Acoustic Feature Processing Strategy for Enhanced Singing Pronunciation With Phoneme Distribution Predictor.</title>
<pages>1-5</pages>
<year>2023</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP49357.2023.10097204</ee>
<crossref>conf/icassp/2023</crossref>
<url>db/conf/icassp/icassp2023.html#WuSQGJ23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/iccv/HuCZJ23" mdate="2024-01-19">
<author pid="249/1182">Anwen Hu</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="50/6759">Liang Zhang</author>
<author pid="47/2670">Qin Jin</author>
<title>Explore and Tell: Embodied Visual Captioning in 3D Environments.</title>
<pages>2482-2491</pages>
<year>2023</year>
<booktitle>ICCV</booktitle>
<ee>https://doi.org/10.1109/ICCV51070.2023.00235</ee>
<crossref>conf/iccv/2023</crossref>
<url>db/conf/iccv/iccv2023.html#HuCZJ23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icmcs/ChenDCJ23" mdate="2023-09-05">
<author pid="274/6597">Jieting Chen</author>
<author pid="346/0048">Junkai Ding</author>
<author pid="04/1620">Wenping Chen</author>
<author pid="47/2670">Qin Jin</author>
<title>Knowledge Enhanced Model for Live Video Comment Generation.</title>
<pages>2267-2272</pages>
<year>2023</year>
<booktitle>ICME</booktitle>
<ee>https://doi.org/10.1109/ICME55011.2023.00387</ee>
<crossref>conf/icmcs/2023</crossref>
<url>db/conf/icmcs/icme2023.html#ChenDCJ23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/LinRX0WX0SZJ023" mdate="2025-01-19">
<author orcid="0009-0000-4742-4247" pid="337/9434">Hongpeng Lin</author>
<author orcid="0009-0009-1039-4940" pid="262/6356">Ludan Ruan</author>
<author orcid="0009-0000-1597-9512" pid="337/9800">Wenke Xia</author>
<author orcid="0000-0002-2974-9184" pid="85/670-2">Peiyu Liu 0002</author>
<author orcid="0009-0008-1875-961X" pid="287/4833">Jingyuan Wen</author>
<author orcid="0009-0007-8315-6842" pid="04/5753">Yixin Xu</author>
<author orcid="0000-0002-7118-6733" pid="49/8496-1">Di Hu 0001</author>
<author orcid="0000-0002-2163-7401" pid="s/RuihuaSong">Ruihua Song</author>
<author orcid="0000-0002-8333-6196" pid="52/8700">Wayne Xin Zhao</author>
<author orcid="0000-0001-6486-6020" pid="47/2670">Qin Jin</author>
<author orcid="0000-0003-0280-7724" pid="53/5234">Zhiwu Lu 0001</author>
<title>TikTalk: A Video-Based Dialogue Dataset for Multi-Modal Chitchat in Real World.</title>
<pages>1303-1313</pages>
<year>2023</year>
<booktitle>ACM Multimedia</booktitle>
<ee>https://doi.org/10.1145/3581783.3612425</ee>
<ee>https://www.wikidata.org/entity/Q130940950</ee>
<crossref>conf/mm/2023</crossref>
<url>db/conf/mm/mm2023.html#LinRX0WX0SZJ023</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/XuZJ23" mdate="2023-11-09">
<author orcid="0009-0000-1896-9600" pid="293/8958">Boshen Xu</author>
<author orcid="0000-0001-5331-6314" pid="251/3691">Sipeng Zheng</author>
<author orcid="0000-0001-6486-6020" pid="47/2670">Qin Jin</author>
<title>POV: Prompt-Oriented View-Agnostic Learning for Egocentric Hand-Object Interaction in the Multi-view World.</title>
<pages>2807-2816</pages>
<year>2023</year>
<booktitle>ACM Multimedia</booktitle>
<ee>https://doi.org/10.1145/3581783.3612484</ee>
<crossref>conf/mm/2023</crossref>
<url>db/conf/mm/mm2023.html#XuZJ23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/YangCHGJJ23" mdate="2025-02-14">
<author orcid="0009-0006-8924-5259" pid="266/2264">Dingyi Yang</author>
<author orcid="0009-0001-7682-1775" pid="28/3046-5">Hongyu Chen 0005</author>
<author orcid="0000-0002-5116-3454" pid="319/3945">Xinglin Hou</author>
<author orcid="0000-0003-1381-2692" pid="135/4944">Tiezheng Ge</author>
<author orcid="0000-0003-1665-3025" pid="99/7950-1">Yuning Jiang 0001</author>
<author orcid="0000-0001-6486-6020" pid="47/2670">Qin Jin</author>
<title>Visual Captioning at Will: Describing Images and Videos Guided by a Few Stylized Sentences.</title>
<pages>5705-5715</pages>
<year>2023</year>
<booktitle>ACM Multimedia</booktitle>
<ee>https://doi.org/10.1145/3581783.3612263</ee>
<crossref>conf/mm/2023</crossref>
<url>db/conf/mm/mm2023.html#YangCHGJJ23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/0003ZLYMJ23" mdate="2025-11-25">
<author orcid="0000-0002-6873-7530" pid="69/10440-3">Yuchen Liu 0003</author>
<author orcid="0009-0008-9646-9802" pid="168/0332">Haoyu Zhang</author>
<author orcid="0009-0008-7218-3877" pid="134/5661-3">Shichao Liu 0003</author>
<author orcid="0000-0003-0472-2783" pid="18/1022-6">Xiang Yin 0006</author>
<author orcid="0000-0001-5508-1328" pid="18/10648-1">Zejun Ma 0001</author>
<author orcid="0000-0001-6486-6020" pid="47/2670">Qin Jin</author>
<title>Emotionally Situated Text-to-Speech Synthesis in User-Agent Conversation.</title>
<pages>5966-5974</pages>
<year>2023</year>
<booktitle>ACM Multimedia</booktitle>
<ee>https://doi.org/10.1145/3581783.3613823</ee>
<ee>https://www.wikidata.org/entity/Q131011811</ee>
<crossref>conf/mm/2023</crossref>
<url>db/conf/mm/mm2023.html#0003ZLYMJ23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/nips/YueHZJ23" mdate="2024-03-01">
<author pid="339/2864">Zihao Yue</author>
<author pid="249/1182">Anwen Hu</author>
<author pid="50/6759">Liang Zhang</author>
<author pid="47/2670">Qin Jin</author>
<title>Learning Descriptive Image Captioning via Semipermeable Maximum Likelihood Estimation.</title>
<year>2023</year>
<booktitle>NeurIPS</booktitle>
<ee type="oa">http://papers.nips.cc/paper_files/paper/2023/hash/fa1cfe4e956d85e016b1f8f49b189a0b-Abstract-Conference.html</ee>
<crossref>conf/nips/2023</crossref>
<url>db/conf/nips/neurips2023.html#YueHZJ23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/nlpcc/HuangZJ23" mdate="2023-10-11">
<author pid="305/3439">Zhaopei Huang</author>
<author pid="121/8902">Jinming Zhao</author>
<author pid="47/2670">Qin Jin</author>
<title>Two-Stage Adaptation for Cross-Corpus Multimodal Emotion Recognition.</title>
<pages>431-443</pages>
<year>2023</year>
<booktitle>NLPCC (2)</booktitle>
<ee>https://doi.org/10.1007/978-3-031-44696-2_34</ee>
<crossref>conf/nlpcc/2023-2</crossref>
<url>db/conf/nlpcc/nlpcc2023-2.html#HuangZJ23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/sigir/ChenYJ23" mdate="2023-07-21">
<author orcid="0000-0001-9371-5256" pid="197/5539">Weijing Chen</author>
<author orcid="0000-0002-9809-8864" pid="262/6579">Linli Yao</author>
<author orcid="0000-0001-6486-6020" pid="47/2670">Qin Jin</author>
<title>Rethinking Benchmarks for Cross-modal Image-text Retrieval.</title>
<pages>1241-1251</pages>
<year>2023</year>
<booktitle>SIGIR</booktitle>
<ee>https://doi.org/10.1145/3539618.3591758</ee>
<crossref>conf/sigir/2023</crossref>
<url>db/conf/sigir/sigir2023.html#ChenYJ23</url>
</inproceedings>
</r>
<r><inproceedings key="conf/www/YaoCJ23" mdate="2025-01-19">
<author orcid="0000-0002-9809-8864" pid="262/6579">Linli Yao</author>
<author orcid="0000-0001-9371-5256" pid="197/5539">Weijing Chen</author>
<author orcid="0000-0001-6486-6020" pid="47/2670">Qin Jin</author>
<title>CapEnrich: Enriching Caption Semantics for Web Images via Cross-modal Pre-trained Knowledge.</title>
<pages>2392-2401</pages>
<year>2023</year>
<booktitle>WWW</booktitle>
<ee>https://doi.org/10.1145/3543507.3583232</ee>
<ee>https://www.wikidata.org/entity/Q130830981</ee>
<crossref>conf/www/2023</crossref>
<url>db/conf/www/www2023.html#YaoCJ23</url>
</inproceedings>
</r>
<r><article publtype="informal" key="journals/corr/abs-2301-05880" mdate="2023-01-19">
<author pid="337/9434">Hongpeng Lin</author>
<author pid="262/6356">Ludan Ruan</author>
<author pid="337/9800">Wenke Xia</author>
<author pid="85/670-2">Peiyu Liu 0002</author>
<author pid="287/4833">Jingyuan Wen</author>
<author pid="04/5753">Yixin Xu</author>
<author pid="49/8496-1">Di Hu 0001</author>
<author pid="s/RuihuaSong">Ruihua Song</author>
<author pid="52/8700">Wayne Xin Zhao</author>
<author pid="47/2670">Qin Jin</author>
<author pid="53/5234">Zhiwu Lu 0001</author>
<title>TikTalk: A Multi-Modal Dialogue Dataset for Real-World Chitchat.</title>
<year>2023</year>
<volume>abs/2301.05880</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2301.05880</ee>
<url>db/journals/corr/corr2301.html#abs-2301-05880</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2303-06591" mdate="2023-03-16">
<author pid="262/6356">Ludan Ruan</author>
<author pid="249/1182">Anwen Hu</author>
<author pid="222/3101">Yuqing Song 0003</author>
<author pid="50/6759">Liang Zhang</author>
<author pid="251/3691">Sipeng Zheng</author>
<author pid="47/2670">Qin Jin</author>
<title>Accommodating Audio Modality in CLIP for Multimodal Processing.</title>
<year>2023</year>
<volume>abs/2303.06591</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2303.06591</ee>
<url>db/journals/corr/corr2303.html#abs-2303-06591</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2303-08607" mdate="2025-07-16">
<author pid="301/4852-1">Yuning Wu 0001</author>
<author pid="229/3529">Jiatong Shi</author>
<author pid="05/5638">Tao Qian</author>
<author pid="226/5309">Dongji Gao</author>
<author pid="47/2670">Qin Jin</author>
<title>PHONEix: Acoustic Feature Processing Strategy for Enhanced Singing Pronunciation with Phoneme Distribution Predictor.</title>
<year>2023</year>
<volume>abs/2303.08607</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2303.08607</ee>
<url>db/journals/corr/corr2303.html#abs-2303-08607</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2304-09660" mdate="2023-04-24">
<author pid="50/6759">Liang Zhang</author>
<author pid="249/1182">Anwen Hu</author>
<author pid="05/3499">Jing Zhang</author>
<author pid="35/7808">Shuo Hu</author>
<author pid="47/2670">Qin Jin</author>
<title>MPMQA: Multimodal Question Answering on Product Manuals.</title>
<year>2023</year>
<volume>abs/2304.09660</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2304.09660</ee>
<url>db/journals/corr/corr2304.html#abs-2304-09660</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2304-10824" mdate="2023-05-02">
<author pid="197/5539">Weijing Chen</author>
<author pid="262/6579">Linli Yao</author>
<author pid="47/2670">Qin Jin</author>
<title>Rethinking Benchmarks for Cross-modal Image-text Retrieval.</title>
<year>2023</year>
<volume>abs/2304.10824</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2304.10824</ee>
<url>db/journals/corr/corr2304.html#abs-2304-10824</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2304-14657" mdate="2023-05-04">
<author pid="274/6597">Jieting Chen</author>
<author pid="346/0048">Junkai Ding</author>
<author pid="04/1620">Wenping Chen</author>
<author pid="47/2670">Qin Jin</author>
<title>Knowledge Enhanced Model for Live Video Comment Generation.</title>
<year>2023</year>
<volume>abs/2304.14657</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2304.14657</ee>
<url>db/journals/corr/corr2304.html#abs-2304-14657</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2305-06002" mdate="2023-05-16">
<author pid="249/1182">Anwen Hu</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="50/6759">Liang Zhang</author>
<author pid="47/2670">Qin Jin</author>
<title>InfoMetIC: An Informative Metric for Reference-free Image Caption Evaluation.</title>
<year>2023</year>
<volume>abs/2305.06002</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2305.06002</ee>
<url>db/journals/corr/corr2305.html#abs-2305-06002</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2305-08389" mdate="2025-02-14">
<author pid="262/6579">Linli Yao</author>
<author pid="304/8398">Yuanmeng Zhang</author>
<author pid="79/10743">Ziheng Wang</author>
<author pid="319/3945">Xinglin Hou</author>
<author pid="135/4944">Tiezheng Ge</author>
<author pid="99/7950-1">Yuning Jiang 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>Edit As You Wish: Video Description Editing with Multi-grained Commands.</title>
<year>2023</year>
<volume>abs/2305.08389</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2305.08389</ee>
<url>db/journals/corr/corr2305.html#abs-2305-08389</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2305-12140" mdate="2023-05-26">
<author pid="339/2864">Zihao Yue</author>
<author pid="52/323">Qi Zhang</author>
<author pid="249/1182">Anwen Hu</author>
<author pid="50/6759">Liang Zhang</author>
<author pid="79/10743">Ziheng Wang</author>
<author pid="47/2670">Qin Jin</author>
<title>Movie101: A New Movie Understanding Benchmark.</title>
<year>2023</year>
<volume>abs/2305.12140</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2305.12140</ee>
<url>db/journals/corr/corr2305.html#abs-2305-12140</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2306-13460" mdate="2023-06-27">
<author pid="339/2864">Zihao Yue</author>
<author pid="249/1182">Anwen Hu</author>
<author pid="50/6759">Liang Zhang</author>
<author pid="47/2670">Qin Jin</author>
<title>Learning Descriptive Image Captioning via Semipermeable Maximum Likelihood Estimation.</title>
<year>2023</year>
<volume>abs/2306.13460</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2306.13460</ee>
<url>db/journals/corr/corr2306.html#abs-2306-13460</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2307-10567" mdate="2023-07-26">
<author pid="52/323">Qi Zhang</author>
<author pid="251/3691">Sipeng Zheng</author>
<author pid="47/2670">Qin Jin</author>
<title>No-frills Temporal Video Grounding: Multi-Scale Neighboring Attention and Zoom-in Boundary Detection.</title>
<year>2023</year>
<volume>abs/2307.10567</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2307.10567</ee>
<url>db/journals/corr/corr2307.html#abs-2307-10567</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2307-16399" mdate="2025-02-14">
<author pid="266/2264">Dingyi Yang</author>
<author pid="28/3046-5">Hongyu Chen 0005</author>
<author pid="319/3945">Xinglin Hou</author>
<author pid="135/4944">Tiezheng Ge</author>
<author pid="99/7950-1">Yuning Jiang 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>Visual Captioning at Will: Describing Images and Videos Guided by a Few Stylized Sentences.</title>
<year>2023</year>
<volume>abs/2307.16399</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2307.16399</ee>
<url>db/journals/corr/corr2307.html#abs-2307-16399</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2308-02867" mdate="2025-07-16">
<author pid="301/4852-1">Yuning Wu 0001</author>
<author pid="09/9833">Yifeng Yu</author>
<author pid="229/3529">Jiatong Shi</author>
<author pid="05/5638">Tao Qian</author>
<author pid="47/2670">Qin Jin</author>
<title>A Systematic Exploration of Joint-training for Singing Voice Synthesis.</title>
<year>2023</year>
<volume>abs/2308.02867</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2308.02867</ee>
<url>db/journals/corr/corr2308.html#abs-2308-02867</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2308-10447" mdate="2023-08-30">
<author pid="249/1182">Anwen Hu</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="50/6759">Liang Zhang</author>
<author pid="47/2670">Qin Jin</author>
<title>Explore and Tell: Embodied Visual Captioning in 3D Environments.</title>
<year>2023</year>
<volume>abs/2308.10447</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2308.10447</ee>
<url>db/journals/corr/corr2308.html#abs-2308-10447</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2310-05126" mdate="2025-10-08">
<author pid="304/1336">Jiabo Ye</author>
<author pid="249/1182">Anwen Hu</author>
<author pid="80/1339-1">Haiyang Xu 0001</author>
<author pid="254/3247">Qinghao Ye</author>
<author pid="51/5332-8">Ming Yan 0008</author>
<author pid="205/7621">Guohai Xu</author>
<author pid="52/9457-3">Chenliang Li 0003</author>
<author pid="93/1076">Junfeng Tian</author>
<author pid="05/2084-1">Qi Qian 0001</author>
<author pid="86/1953-11">Ji Zhang 0011</author>
<author pid="47/2670">Qin Jin</author>
<author pid="42/963-1">Liang He 0001</author>
<author pid="50/3323-1">Xin Alex Lin</author>
<author pid="h/FeiHuang-2">Fei Huang 0002</author>
<title>UReader: Universal OCR-free Visually-situated Language Understanding with Multimodal Large Language Model.</title>
<year>2023</year>
<volume>abs/2310.05126</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2310.05126</ee>
<url>db/journals/corr/corr2310.html#abs-2310-05126</url>
</article>
</r>
<r><article key="journals/aiopen/RuanJ22" mdate="2023-02-10">
<author pid="262/6356">Ludan Ruan</author>
<author orcid="0000-0001-6486-6020" pid="47/2670">Qin Jin</author>
<title>Survey: Transformer based video-language pre-training.</title>
<pages>1-13</pages>
<year>2022</year>
<month>January</month>
<volume>3</volume>
<journal>AI Open</journal>
<ee type="oa">https://doi.org/10.1016/j.aiopen.2022.01.001</ee>
<url>db/journals/aiopen/aiopen3.html#RuanJ22</url>
</article>
</r>
<r><article key="journals/tmm/SongCJLXH22" mdate="2025-06-11">
<author orcid="0000-0002-0246-4330" pid="222/3101">Yuqing Song 0003</author>
<author orcid="0000-0002-7313-9703" pid="153/0734">Shizhe Chen</author>
<author orcid="0000-0001-6486-6020" pid="47/2670">Qin Jin</author>
<author pid="05/6715">Wei Luo</author>
<author pid="33/3881">Jun Xie</author>
<author orcid="0000-0002-3709-5053" pid="h/FeiHuang-2">Fei Huang 0002</author>
<title>Enhancing Neural Machine Translation With Dual-Side Multimodal Awareness.</title>
<pages>3013-3024</pages>
<year>2022</year>
<volume>24</volume>
<journal>IEEE Trans. Multim.</journal>
<ee>https://doi.org/10.1109/TMM.2021.3092187</ee>
<url>db/journals/tmm/tmm24.html#SongCJLXH22</url>
</article>
</r>
<r><inproceedings key="conf/aaai/YaoWJ22" mdate="2023-09-30">
<author orcid="0000-0002-9809-8864" pid="262/6579">Linli Yao</author>
<author pid="227/7810">Weiying Wang</author>
<author pid="47/2670">Qin Jin</author>
<title>Image Difference Captioning with Pre-training and Contrastive Learning.</title>
<pages>3108-3116</pages>
<year>2022</year>
<booktitle>AAAI</booktitle>
<ee type="oa">https://doi.org/10.1609/aaai.v36i3.20218</ee>
<crossref>conf/aaai/2022</crossref>
<url>db/conf/aaai/aaai2022.html#YaoWJ22</url>
</inproceedings>
</r>
<r><inproceedings key="conf/acl/ZhaoZ0LJW022" mdate="2025-01-19">
<author pid="121/8902">Jinming Zhao</author>
<author pid="305/3353">Tenggan Zhang</author>
<author pid="159/1747-3">Jingwen Hu 0003</author>
<author pid="69/10440-3">Yuchen Liu 0003</author>
<author pid="47/2670">Qin Jin</author>
<author pid="23/8015">Xinchao Wang</author>
<author pid="36/4118">Haizhou Li 0001</author>
<title>M3ED: Multi-modal Multi-scene Multi-label Emotional Dialogue Database.</title>
<pages>5699-5710</pages>
<year>2022</year>
<booktitle>ACL (1)</booktitle>
<ee type="oa">https://doi.org/10.18653/v1/2022.acl-long.391</ee>
<ee type="oa">https://aclanthology.org/2022.acl-long.391</ee>
<ee>https://www.wikidata.org/entity/Q131458744</ee>
<crossref>conf/acl/2022-1</crossref>
<url>db/conf/acl/acl2022-1.html#ZhaoZ0LJW022</url>
</inproceedings>
</r>
<r><inproceedings key="conf/coling/LiuZ0LJ22" mdate="2022-11-22">
<author pid="69/10440-3">Yuchen Liu 0003</author>
<author pid="121/8902">Jinming Zhao</author>
<author pid="159/1747-3">Jingwen Hu 0003</author>
<author pid="231/3910">Ruichen Li</author>
<author pid="47/2670">Qin Jin</author>
<title>DialogueEIN: Emotion Interaction Network for Dialogue Affective Analysis.</title>
<pages>684-693</pages>
<year>2022</year>
<booktitle>COLING</booktitle>
<ee type="oa">https://aclanthology.org/2022.coling-1.57</ee>
<crossref>conf/coling/2022</crossref>
<url>db/conf/coling/coling2022.html#LiuZ0LJ22</url>
</inproceedings>
</r>
<r><inproceedings key="conf/cvpr/MengLLHJZLJ22" mdate="2025-12-07">
<author pid="276/9497">Liyu Meng</author>
<author pid="69/10440-3">Yuchen Liu 0003</author>
<author pid="48/1674">Xiaolong Liu</author>
<author orcid="0000-0001-6325-791X" pid="305/3439">Zhaopei Huang</author>
<author pid="276/9880">Wenqiang Jiang</author>
<author pid="305/3353">Tenggan Zhang</author>
<author pid="227/5169">Chuanhe Liu</author>
<author pid="47/2670">Qin Jin</author>
<title>Valence and Arousal Estimation based on Multimodal Temporal-Aware Features for Videos in the Wild.</title>
<pages>2344-2351</pages>
<year>2022</year>
<booktitle>CVPR Workshops</booktitle>
<ee>https://doi.org/10.1109/CVPRW56347.2022.00261</ee>
<crossref>conf/cvpr/2022w</crossref>
<url>db/conf/cvpr/cvpr2022w.html#MengLLHJZLJ22</url>
</inproceedings>
</r>
<r><inproceedings key="conf/cvpr/ZhengCJ22" mdate="2022-10-05">
<author pid="251/3691">Sipeng Zheng</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<title>VRDFormer: End-to-End Video Visual Relation Detection with Transformers.</title>
<pages>18814-18824</pages>
<year>2022</year>
<booktitle>CVPR</booktitle>
<ee>https://doi.org/10.1109/CVPR52688.2022.01827</ee>
<crossref>conf/cvpr/2022</crossref>
<url>db/conf/cvpr/cvpr2022.html#ZhengCJ22</url>
</inproceedings>
</r>
<r><inproceedings key="conf/eccv/ZhangLLLMSJZZJ22" mdate="2023-02-21">
<author pid="305/3353">Tenggan Zhang</author>
<author pid="227/5169">Chuanhe Liu</author>
<author pid="48/1674">Xiaolong Liu</author>
<author pid="69/10440-3">Yuchen Liu 0003</author>
<author pid="276/9497">Liyu Meng</author>
<author pid="02/2264">Lei Sun</author>
<author pid="276/9880">Wenqiang Jiang</author>
<author pid="192/1145">Fengyuan Zhang</author>
<author pid="121/8902">Jinming Zhao</author>
<author pid="47/2670">Qin Jin</author>
<title>Multi-Task Learning Framework for Emotion Recognition In-the-Wild.</title>
<pages>143-156</pages>
<year>2022</year>
<booktitle>ECCV Workshops (6)</booktitle>
<ee>https://doi.org/10.1007/978-3-031-25075-0_11</ee>
<crossref>conf/eccv/2022-w6</crossref>
<url>db/conf/eccv/eccv2022-w6.html#ZhangLLLMSJZZJ22</url>
</inproceedings>
</r>
<r><inproceedings key="conf/eccv/ZhengCJ22" mdate="2022-11-10">
<author pid="251/3691">Sipeng Zheng</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<title>Few-Shot Action Recognition with Hierarchical Matching and Contrastive Learning.</title>
<pages>297-313</pages>
<year>2022</year>
<booktitle>ECCV (4)</booktitle>
<ee>https://doi.org/10.1007/978-3-031-19772-7_18</ee>
<crossref>conf/eccv/2022-4</crossref>
<url>db/conf/eccv/eccv2022-4.html#ZhengCJ22</url>
</inproceedings>
</r>
<r><inproceedings key="conf/eccv/LiuXXCJ22" mdate="2024-10-25">
<author orcid="0000-0002-2119-0881" pid="35/9071-3">Yuqi Liu 0003</author>
<author pid="48/8617">Pengfei Xiong</author>
<author pid="229/5603">Luhui Xu</author>
<author orcid="0000-0002-9342-1837" pid="316/8117">Shengming Cao</author>
<author orcid="0000-0001-6486-6020" pid="47/2670">Qin Jin</author>
<title>TS2-Net: Token Shift and Selection Transformer for Text-Video Retrieval.</title>
<pages>319-335</pages>
<year>2022</year>
<booktitle>ECCV (14)</booktitle>
<ee>https://doi.org/10.1007/978-3-031-19781-9_19</ee>
<crossref>conf/eccv/2022-14</crossref>
<url>db/conf/eccv/eccv2022-14.html#LiuXXCJ22</url>
</inproceedings>
</r>
<r><inproceedings key="conf/eccv/Zhang0J22" mdate="2022-11-13">
<author orcid="0000-0001-5902-1955" pid="52/323">Qi Zhang</author>
<author pid="222/3101">Yuqing Song 0003</author>
<author orcid="0000-0001-6486-6020" pid="47/2670">Qin Jin</author>
<title>Unifying Event Detection and Captioning as Sequence Generation via Pre-training.</title>
<pages>363-379</pages>
<year>2022</year>
<booktitle>ECCV (36)</booktitle>
<ee>https://doi.org/10.1007/978-3-031-20059-5_21</ee>
<crossref>conf/eccv/2022-36</crossref>
<url>db/conf/eccv/eccv2022-36.html#Zhang0J22</url>
</inproceedings>
</r>
<r><inproceedings key="conf/emnlp/ZhangYHWJ22" mdate="2023-08-10">
<author pid="52/323">Qi Zhang</author>
<author pid="339/2864">Zihao Yue</author>
<author pid="249/1182">Anwen Hu</author>
<author pid="79/10743">Ziheng Wang</author>
<author pid="47/2670">Qin Jin</author>
<title>MovieUN: A Dataset for Movie Understanding and Narrating.</title>
<pages>1873-1885</pages>
<year>2022</year>
<booktitle>EMNLP (Findings)</booktitle>
<ee type="oa">https://doi.org/10.18653/v1/2022.findings-emnlp.135</ee>
<ee type="oa">https://aclanthology.org/2022.findings-emnlp.135</ee>
<crossref>conf/emnlp/2022f</crossref>
<url>db/conf/emnlp/emnlp2022f.html#ZhangYHWJ22</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/ZhaoLJWL22" mdate="2023-06-26">
<author pid="121/8902">Jinming Zhao</author>
<author pid="231/3910">Ruichen Li</author>
<author pid="47/2670">Qin Jin</author>
<author orcid="0000-0003-0057-1404" pid="23/8015">Xinchao Wang</author>
<author pid="36/4118">Haizhou Li 0001</author>
<title>Memobert: Pre-Training Model with Prompt-Based Learning for Multimodal Emotion Recognition.</title>
<pages>4703-4707</pages>
<year>2022</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP43922.2022.9746910</ee>
<crossref>conf/icassp/2022</crossref>
<url>db/conf/icassp/icassp2022.html#ZhaoLJWL22</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/QianSGWJ22" mdate="2022-06-07">
<author pid="05/5638">Tao Qian</author>
<author pid="229/3529">Jiatong Shi</author>
<author pid="07/272">Shuai Guo</author>
<author pid="44/3072">Peter Wu</author>
<author pid="47/2670">Qin Jin</author>
<title>Training Strategies for Automatic Song Writing: A Unified Framework Perspective.</title>
<pages>4738-4742</pages>
<year>2022</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP43922.2022.9746818</ee>
<crossref>conf/icassp/2022</crossref>
<url>db/conf/icassp/icassp2022.html#QianSGWJ22</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/GuoSQ0J22" mdate="2023-06-21">
<author pid="07/272">Shuai Guo</author>
<author pid="229/3529">Jiatong Shi</author>
<author pid="05/5638">Tao Qian</author>
<author orcid="0000-0002-5970-8631" pid="39/3245-1">Shinji Watanabe 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>SingAug: Data Augmentation for Singing Voice Synthesis with Cycle-consistent Training Strategy.</title>
<pages>4272-4276</pages>
<year>2022</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2022-978</ee>
<crossref>conf/interspeech/2022</crossref>
<url>db/conf/interspeech/interspeech2022.html#GuoSQ0J22</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/ShiGQHWXCLW0J22" mdate="2025-07-16">
<author pid="229/3529">Jiatong Shi</author>
<author pid="07/272">Shuai Guo</author>
<author pid="05/5638">Tao Qian</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="301/4852-1">Yuning Wu 0001</author>
<author pid="257/8301">Fangzheng Xu</author>
<author pid="194/1149">Xuankai Chang</author>
<author pid="320/0415">Huazhe Li</author>
<author pid="44/3072">Peter Wu</author>
<author orcid="0000-0002-5970-8631" pid="39/3245-1">Shinji Watanabe 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>Muskits: an End-to-end Music Processing Toolkit for Singing Voice Synthesis.</title>
<pages>4277-4281</pages>
<year>2022</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2022-10039</ee>
<crossref>conf/interspeech/2022</crossref>
<url>db/conf/interspeech/interspeech2022.html#ShiGQHWXCLW0J22</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/Alameda-PinedaJ22" mdate="2022-10-14">
<author pid="22/10486">Xavier Alameda-Pineda</author>
<author pid="47/2670">Qin Jin</author>
<author pid="o/VincentOria">Vincent Oria</author>
<author pid="81/7871">Laura Toni</author>
<title>M4MM '22: 1st International Workshop on Methodologies for Multimedia.</title>
<pages>7394-7396</pages>
<year>2022</year>
<booktitle>ACM Multimedia</booktitle>
<ee>https://doi.org/10.1145/3503161.3554769</ee>
<crossref>conf/mm/2022</crossref>
<url>db/conf/mm/mm2022.html#Alameda-PinedaJ22</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/0001JLTL22" mdate="2025-01-19">
<author pid="60/7642">Si Liu 0001</author>
<author pid="47/2670">Qin Jin</author>
<author pid="29/8842">Luoqi Liu</author>
<author pid="278/2959">Zongheng Tang</author>
<author pid="331/1495">Linli Lin</author>
<title>PIC'22: 4th Person in Context Workshop.</title>
<pages>7418-7419</pages>
<year>2022</year>
<booktitle>ACM Multimedia</booktitle>
<ee>https://doi.org/10.1145/3503161.3554766</ee>
<ee>https://www.wikidata.org/entity/Q131013585</ee>
<crossref>conf/mm/2022</crossref>
<url>db/conf/mm/mm2022.html#0001JLTL22</url>
</inproceedings>
</r>
<r><inproceedings key="conf/nips/ZhangHJ22" mdate="2024-01-08">
<author pid="50/6759">Liang Zhang</author>
<author pid="249/1182">Anwen Hu</author>
<author pid="47/2670">Qin Jin</author>
<title>Multi-Lingual Acquisition on Multimodal Pre-training for Cross-modal Retrieval.</title>
<year>2022</year>
<crossref>conf/nips/2022</crossref>
<booktitle>NeurIPS</booktitle>
<ee type="oa">http://papers.nips.cc/paper_files/paper/2022/hash/bfadef437ed27372648714c930c3a77a-Abstract-Conference.html</ee>
<url>db/conf/nips/neurips2022.html#ZhangHJ22</url>
</inproceedings>
</r>
<r><inproceedings key="conf/sigir/Zhao0J22" mdate="2022-07-08">
<author pid="222/2669">Yida Zhao</author>
<author pid="222/3101">Yuqing Song 0003</author>
<author pid="47/2670">Qin Jin</author>
<title>Progressive Learning for Image Retrieval with Hybrid-Modality Queries.</title>
<pages>1012-1021</pages>
<year>2022</year>
<booktitle>SIGIR</booktitle>
<ee>https://doi.org/10.1145/3477495.3532047</ee>
<crossref>conf/sigir/2022</crossref>
<url>db/conf/sigir/sigir2022.html#Zhao0J22</url>
</inproceedings>
</r>
<r><proceedings key="conf/mm/2022" mdate="2022-10-13">
<editor pid="08/1790">Jo&#227;o Magalh&#227;es</editor>
<editor pid="b/AlbertoDelBimbo">Alberto Del Bimbo</editor>
<editor pid="50/290">Shin'ichi Satoh 0001</editor>
<editor pid="20/3519">Nicu Sebe</editor>
<editor pid="22/10486">Xavier Alameda-Pineda</editor>
<editor pid="47/2670">Qin Jin</editor>
<editor pid="o/VincentOria">Vincent Oria</editor>
<editor pid="81/7871">Laura Toni</editor>
<title>MM '22: The 30th ACM International Conference on Multimedia, Lisboa, Portugal, October 10 - 14, 2022</title>
<publisher>ACM</publisher>
<booktitle>ACM Multimedia</booktitle>
<year>2022</year>
<isbn>978-1-4503-9203-7</isbn>
<ee>https://doi.org/10.1145/3503161</ee>
<url>db/conf/mm/mm2022.html</url>
</proceedings>
</r>
<r><article publtype="informal" key="journals/corr/abs-2202-04298" mdate="2022-02-18">
<author pid="262/6579">Linli Yao</author>
<author pid="227/7810">Weiying Wang</author>
<author pid="47/2670">Qin Jin</author>
<title>Image Difference Captioning with Pre-training and Contrastive Learning.</title>
<year>2022</year>
<volume>abs/2202.04298</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2202.04298</ee>
<url>db/journals/corr/corr2202.html#abs-2202-04298</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2203-13032" mdate="2022-11-22">
<author pid="276/9497">Liyu Meng</author>
<author pid="69/10440-3">Yuchen Liu 0003</author>
<author pid="48/1674">Xiaolong Liu</author>
<author pid="305/3439">Zhaopei Huang</author>
<author pid="15/5135">Yuan Cheng</author>
<author pid="93/6765">Meng Wang</author>
<author pid="227/5169">Chuanhe Liu</author>
<author pid="47/2670">Qin Jin</author>
<title>Multi-modal Emotion Estimation for in-the-wild Videos.</title>
<year>2022</year>
<volume>abs/2203.13032</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2203.13032</ee>
<url>db/journals/corr/corr2203.html#abs-2203-13032</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2203-17001" mdate="2023-03-21">
<author pid="07/272">Shuai Guo</author>
<author pid="229/3529">Jiatong Shi</author>
<author pid="05/5638">Tao Qian</author>
<author orcid="0000-0002-5970-8631" pid="39/3245-1">Shinji Watanabe 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>SingAug: Data Augmentation for Singing Voice Synthesis with Cycle-consistent Training Strategy.</title>
<year>2022</year>
<volume>abs/2203.17001</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2203.17001</ee>
<url>db/journals/corr/corr2203.html#abs-2203-17001</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2204-11212" mdate="2022-04-28">
<author pid="222/2669">Yida Zhao</author>
<author pid="222/3101">Yuqing Song 0003</author>
<author pid="47/2670">Qin Jin</author>
<title>Progressive Learning for Image Retrieval with Hybrid-Modality Queries.</title>
<year>2022</year>
<volume>abs/2204.11212</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2204.11212</ee>
<url>db/journals/corr/corr2204.html#abs-2204-11212</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2205-04029" mdate="2026-02-03">
<author pid="229/3529">Jiatong Shi</author>
<author pid="07/272">Shuai Guo</author>
<author pid="05/5638">Tao Qian</author>
<author pid="272/8774">Nan Huo</author>
<author pid="82/8616">Tomoki Hayashi</author>
<author pid="301/4852-1">Yuning Wu 0001</author>
<author pid="190/4519">Frank Xu 0001</author>
<author pid="194/1149">Xuankai Chang</author>
<author pid="320/0415">Huazhe Li</author>
<author pid="44/3072">Peter Wu</author>
<author orcid="0000-0002-5970-8631" pid="39/3245-1">Shinji Watanabe 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>Muskits: an End-to-End Music Processing Toolkit for Singing Voice Synthesis.</title>
<year>2022</year>
<volume>abs/2205.04029</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2205.04029</ee>
<url>db/journals/corr/corr2205.html#abs-2205-04029</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2205-10237" mdate="2023-06-26">
<author pid="121/8902">Jinming Zhao</author>
<author pid="305/3353">Tenggan Zhang</author>
<author pid="159/1747-3">Jingwen Hu 0003</author>
<author pid="69/10440-3">Yuchen Liu 0003</author>
<author pid="47/2670">Qin Jin</author>
<author orcid="0000-0003-0057-1404" pid="23/8015">Xinchao Wang</author>
<author pid="36/4118">Haizhou Li 0001</author>
<title>M3ED: Multi-modal Multi-scene Multi-label Emotional Dialogue Database.</title>
<year>2022</year>
<volume>abs/2205.10237</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2205.10237</ee>
<url>db/journals/corr/corr2205.html#abs-2205-10237</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2206-11091" mdate="2022-06-27">
<author pid="50/6759">Liang Zhang</author>
<author pid="249/1182">Anwen Hu</author>
<author pid="47/2670">Qin Jin</author>
<title>Generalizing Multimodal Pre-training into Multilingual via Language Acquisition.</title>
<year>2022</year>
<volume>abs/2206.11091</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2206.11091</ee>
<url>db/journals/corr/corr2206.html#abs-2206-11091</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2207-07852" mdate="2024-10-25">
<author pid="35/9071-3">Yuqi Liu 0003</author>
<author pid="48/8617">Pengfei Xiong</author>
<author pid="229/5603">Luhui Xu</author>
<author pid="316/8117">Shengming Cao</author>
<author pid="47/2670">Qin Jin</author>
<title>TS2-Net: Token Shift and Selection Transformer for Text-Video Retrieval.</title>
<year>2022</year>
<volume>abs/2207.07852</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2207.07852</ee>
<url>db/journals/corr/corr2207.html#abs-2207-07852</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2207-08625" mdate="2022-07-19">
<author pid="52/323">Qi Zhang</author>
<author pid="222/3101">Yuqing Song 0003</author>
<author pid="47/2670">Qin Jin</author>
<title>Unifying Event Detection and Captioning as Sequence Generation via Pre-Training.</title>
<year>2022</year>
<volume>abs/2207.08625</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2207.08625</ee>
<url>db/journals/corr/corr2207.html#abs-2207-08625</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2208-05375" mdate="2022-09-26">
<author pid="251/3691">Sipeng Zheng</author>
<author pid="52/323">Qi Zhang</author>
<author pid="39/3711-1">Bei Liu 0001</author>
<author pid="47/2670">Qin Jin</author>
<author pid="83/8692">Jianlong Fu</author>
<title>Exploring Anchor-based Detection for Ego4D Natural Language Query.</title>
<year>2022</year>
<volume>abs/2208.05375</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2208.05375</ee>
<url>db/journals/corr/corr2208.html#abs-2208-05375</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2211-09371" mdate="2022-11-23">
<author pid="262/6579">Linli Yao</author>
<author pid="197/5539">Weijing Chen</author>
<author pid="47/2670">Qin Jin</author>
<title>CapEnrich: Enriching Caption Semantics for Web Images via Cross-modal Pre-trained Knowledge.</title>
<year>2022</year>
<volume>abs/2211.09371</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2211.09371</ee>
<url>db/journals/corr/corr2211.html#abs-2211-09371</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2212-09478" mdate="2025-03-03">
<author pid="262/6356">Ludan Ruan</author>
<author pid="324/2590">Yiyang Ma</author>
<author pid="86/4843-5">Huan Yang 0005</author>
<author orcid="0000-0003-1419-059X" pid="270/6402">Huiguo He</author>
<author pid="39/3711-1">Bei Liu 0001</author>
<author pid="83/8692">Jianlong Fu</author>
<author pid="131/4855">Nicholas Jing Yuan</author>
<author pid="47/2670">Qin Jin</author>
<author pid="44/3946">Baining Guo</author>
<title>MM-Diffusion: Learning Multi-Modal Diffusion Models for Joint Audio and Video Generation.</title>
<year>2022</year>
<volume>abs/2212.09478</volume>
<journal>CoRR</journal>
<ee type="oa">https://doi.org/10.48550/arXiv.2212.09478</ee>
<url>db/journals/corr/corr2212.html#abs-2212-09478</url>
</article>
</r>
<r><article key="journals/aiopen/HanZDGLHQYZZHHJ21" mdate="2026-04-07">
<author orcid="0000-0002-4726-7621" pid="19/3011-7">Xu Han 0007</author>
<author pid="23/10446">Zhengyan Zhang</author>
<author pid="04/4910-2">Ning Ding 0002</author>
<author pid="248/2313">Yuxian Gu</author>
<author orcid="0000-0002-9226-4569" pid="82/1364-36">Xiao Liu 0036</author>
<author pid="219/6931">Yuqi Huo</author>
<author pid="152/1733">Jiezhong Qiu</author>
<author pid="25/4120-13">Yuan Yao 0013</author>
<author pid="187/6243">Ao Zhang</author>
<author pid="50/6759">Liang Zhang</author>
<author pid="10/8243">Wentao Han</author>
<author pid="47/6668">Minlie Huang</author>
<author pid="47/2670">Qin Jin</author>
<author pid="00/6040">Yanyan Lan</author>
<author pid="51/3710-5">Yang Liu 0005</author>
<author pid="53/3245-1">Zhiyuan Liu 0001</author>
<author pid="53/5234">Zhiwu Lu 0001</author>
<author pid="69/1395">Xipeng Qiu</author>
<author pid="s/RuihuaSong">Ruihua Song</author>
<author pid="t/JieTang">Jie Tang 0001</author>
<author pid="w/JRWen">Ji-Rong Wen</author>
<author pid="58/3397">Jinhui Yuan</author>
<author pid="52/8700">Wayne Xin Zhao</author>
<author pid="50/2644-1">Jun Zhu 0001</author>
<title>Pre-trained models: Past, present and future.</title>
<pages>225-250</pages>
<year>2021</year>
<volume>2</volume>
<journal>AI Open</journal>
<ee type="oa">https://doi.org/10.1016/j.aiopen.2021.08.002</ee>
<ee>https://www.wikidata.org/entity/Q116758143</ee>
<url>db/journals/aiopen/aiopen2.html#HanZDGLHQYZZHHJ21</url>
</article>
</r>
<r><inproceedings key="conf/acl/ZhaoLJ20" mdate="2025-01-19">
<author pid="121/8902">Jinming Zhao</author>
<author pid="231/3910">Ruichen Li</author>
<author pid="47/2670">Qin Jin</author>
<title>Missing Modality Imagination Network for Emotion Recognition with Uncertain Missing Modalities.</title>
<pages>2608-2618</pages>
<year>2021</year>
<booktitle>ACL/IJCNLP (1)</booktitle>
<ee type="oa">https://doi.org/10.18653/v1/2021.acl-long.203</ee>
<ee type="oa">https://aclanthology.org/2021.acl-long.203</ee>
<ee>https://www.wikidata.org/entity/Q131458738</ee>
<crossref>conf/acl/2021-1</crossref>
<url>db/conf/acl/acl2021-1.html#ZhaoLJ20</url>
</inproceedings>
</r>
<r><inproceedings key="conf/acl/HuLZJ20" mdate="2022-11-22">
<author pid="159/1747-3">Jingwen Hu 0003</author>
<author orcid="0000-0002-6873-7530" pid="69/10440-3">Yuchen Liu 0003</author>
<author pid="121/8902">Jinming Zhao</author>
<author pid="47/2670">Qin Jin</author>
<title>MMGCN: Multimodal Fusion via Deep Graph Convolution Network for Emotion Recognition in Conversation.</title>
<pages>5666-5675</pages>
<year>2021</year>
<booktitle>ACL/IJCNLP (1)</booktitle>
<ee type="oa">https://doi.org/10.18653/v1/2021.acl-long.440</ee>
<ee type="oa">https://aclanthology.org/2021.acl-long.440</ee>
<crossref>conf/acl/2021-1</crossref>
<url>db/conf/acl/acl2021-1.html#HuLZJ20</url>
</inproceedings>
</r>
<r><inproceedings key="conf/cvpr/0003CJ21" mdate="2022-07-18">
<author pid="222/3101">Yuqing Song 0003</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<title>Towards Diverse Paragraph Captioning for Untrimmed Videos.</title>
<pages>11245-11254</pages>
<year>2021</year>
<booktitle>CVPR</booktitle>
<ee type="oa">https://openaccess.thecvf.com/content/CVPR2021/html/Song_Towards_Diverse_Paragraph_Captioning_for_Untrimmed_Videos_CVPR_2021_paper.html</ee>
<ee>https://doi.org/10.1109/CVPR46437.2021.01109</ee>
<crossref>conf/cvpr/2021</crossref>
<url>db/conf/cvpr/cvpr2021.html#0003CJ21</url>
</inproceedings>
</r>
<r><inproceedings key="conf/emnlp/0001WZJ21" mdate="2026-02-05">
<author pid="99/6879-1">Jia Chen 0001</author>
<author orcid="0000-0001-7384-8836" pid="246/5764-2">Yike Wu 0002</author>
<author orcid="0000-0001-5068-025X" pid="73/1947">Shiwan Zhao</author>
<author pid="47/2670">Qin Jin</author>
<title>Language Resource Efficient Learning for Captioning.</title>
<pages>1887-1895</pages>
<year>2021</year>
<booktitle>EMNLP (Findings)</booktitle>
<ee type="oa">https://doi.org/10.18653/v1/2021.findings-emnlp.162</ee>
<ee type="oa">https://aclanthology.org/2021.findings-emnlp.162</ee>
<crossref>conf/emnlp/2021f</crossref>
<url>db/conf/emnlp/emnlp2021f.html#0001WZJ21</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/ShiGHZJ21" mdate="2021-07-08">
<author pid="229/3529">Jiatong Shi</author>
<author pid="07/272">Shuai Guo</author>
<author pid="272/8774">Nan Huo</author>
<author pid="191/5931">Yuekai Zhang</author>
<author pid="47/2670">Qin Jin</author>
<title>Sequence-To-Sequence Singing Voice Synthesis With Perceptual Entropy Loss.</title>
<pages>76-80</pages>
<year>2021</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP39728.2021.9414348</ee>
<crossref>conf/icassp/2021</crossref>
<url>db/conf/icassp/icassp2021.html#ShiGHZJ21</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/LiZJ21" mdate="2023-06-21">
<author pid="231/3910">Ruichen Li</author>
<author pid="121/8902">Jinming Zhao</author>
<author pid="47/2670">Qin Jin</author>
<title>Speech Emotion Recognition via Multi-Level Cross-Modal Distillation.</title>
<pages>4488-4492</pages>
<year>2021</year>
<booktitle>Interspeech</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2021-785</ee>
<crossref>conf/interspeech/2021</crossref>
<url>db/conf/interspeech/interspeech2021.html#LiZJ21</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mir/LiuFCJHR21" mdate="2021-12-06">
<author pid="39/3711-1">Bei Liu 0001</author>
<author pid="83/8692">Jianlong Fu</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<author pid="h/AlexanderGHauptmann">Alexander G. Hauptmann</author>
<author pid="r/YongRui">Yong Rui</author>
<title>MMPT'21: International Joint Workshop on Multi-Modal Pre-Training for Multimedia Understanding.</title>
<pages>694-695</pages>
<year>2021</year>
<booktitle>ICMR</booktitle>
<ee>https://doi.org/10.1145/3460426.3470947</ee>
<crossref>conf/mir/2021</crossref>
<url>db/conf/mir/icmr2021.html#LiuFCJHR21</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/ZhangHLZJ21" mdate="2022-10-02">
<author pid="305/3353">Tenggan Zhang</author>
<author orcid="0000-0001-6325-791X" pid="305/3439">Zhaopei Huang</author>
<author pid="231/3910">Ruichen Li</author>
<author pid="121/8902">Jinming Zhao</author>
<author pid="47/2670">Qin Jin</author>
<title>Multimodal Fusion Strategies for Physiological-emotion Analysis.</title>
<pages>43-50</pages>
<year>2021</year>
<booktitle>MuSe @ ACM Multimedia</booktitle>
<ee>https://doi.org/10.1145/3475957.3484452</ee>
<crossref>conf/mm/2021muse</crossref>
<url>db/conf/mm/muse2021.html#ZhangHLZJ21</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/0003CJLXH21" mdate="2025-06-11">
<author pid="222/3101">Yuqing Song 0003</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<author pid="05/6715">Wei Luo</author>
<author pid="33/3881">Jun Xie</author>
<author pid="h/FeiHuang-2">Fei Huang 0002</author>
<title>Product-oriented Machine Translation with Cross-modal Cross-lingual Pre-training.</title>
<pages>2843-2852</pages>
<year>2021</year>
<booktitle>ACM Multimedia</booktitle>
<ee>https://doi.org/10.1145/3474085.3475303</ee>
<crossref>conf/mm/2021</crossref>
<url>db/conf/mm/mm2021.html#0003CJLXH21</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/HuCJ21" mdate="2021-10-20">
<author pid="249/1182">Anwen Hu</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<title>Question-controlled Text-aware Image Captioning.</title>
<pages>3097-3105</pages>
<year>2021</year>
<booktitle>ACM Multimedia</booktitle>
<ee>https://doi.org/10.1145/3474085.3475452</ee>
<crossref>conf/mm/2021</crossref>
<url>db/conf/mm/mm2021.html#HuCJ21</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mmasia/RuanJ21" mdate="2022-01-12">
<author pid="262/6356">Ludan Ruan</author>
<author pid="47/2670">Qin Jin</author>
<title>Efficient Proposal Generation with U-shaped Network for Temporal Sentence Grounding.</title>
<pages>26:1-26:7</pages>
<year>2021</year>
<booktitle>MMAsia</booktitle>
<ee>https://doi.org/10.1145/3469877.3490606</ee>
<crossref>conf/mmasia/2021</crossref>
<url>db/conf/mmasia/mmasia2021.html#RuanJ21</url>
</inproceedings>
</r>
<r><proceedings key="conf/mir/2021mmpt" mdate="2021-10-06">
<editor pid="39/3711-1">Bei Liu 0001</editor>
<editor pid="83/8692">Jianlong Fu</editor>
<editor pid="153/0734">Shizhe Chen</editor>
<editor pid="47/2670">Qin Jin</editor>
<editor pid="h/AlexanderGHauptmann">Alexander G. Hauptmann</editor>
<editor pid="r/YongRui">Yong Rui</editor>
<title>MMPT@ICMR2021: Proceedings of the 2021 Workshop on Multi-Modal Pre-Training for Multimedia Understanding, Taipei, Taiwan, August 21, 2021</title>
<publisher>ACM</publisher>
<booktitle>MMPT@ICMR</booktitle>
<year>2021</year>
<isbn>978-1-4503-8530-5</isbn>
<ee>https://doi.org/10.1145/3463945</ee>
<url>db/conf/mir/mmpt2021.html</url>
</proceedings>
</r>
<r><article publtype="informal" key="journals/corr/abs-2103-06561" mdate="2026-01-08">
<author pid="219/6931">Yuqi Huo</author>
<author pid="223/5376">Manli Zhang</author>
<author pid="186/6864">Guangzhen Liu</author>
<author pid="240/2720">Haoyu Lu</author>
<author pid="132/7629-4">Yizhao Gao 0004</author>
<author pid="271/9521">Guoxing Yang</author>
<author pid="287/4833">Jingyuan Wen</author>
<author pid="55/826">Heng Zhang</author>
<author pid="287/5028">Baogui Xu</author>
<author pid="193/7989">Weihao Zheng</author>
<author pid="287/4970">Zongzheng Xi</author>
<author pid="287/5041">Yueqian Yang</author>
<author pid="249/1182">Anwen Hu</author>
<author pid="121/8902">Jinming Zhao</author>
<author pid="231/3910">Ruichen Li</author>
<author pid="222/2669">Yida Zhao</author>
<author pid="50/6759">Liang Zhang</author>
<author pid="222/3101">Yuqing Song 0003</author>
<author pid="54/1309">Xin Hong</author>
<author pid="260/2159">Wanqing Cui</author>
<author pid="228/0778">Dan Yang Hou</author>
<author pid="287/4789">Yingyan Li</author>
<author pid="28/6612-1">Junyi Li 0001</author>
<author pid="85/670-2">Peiyu Liu 0002</author>
<author pid="85/5448-1">Zheng Gong 0001</author>
<author pid="287/4999">Chuhao Jin</author>
<author pid="206/8045">Yuchong Sun</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="53/5234">Zhiwu Lu 0001</author>
<author pid="18/5740">Zhicheng Dou</author>
<author pid="47/2670">Qin Jin</author>
<author pid="00/6040">Yanyan Lan</author>
<author pid="52/8700">Wayne Xin Zhao</author>
<author pid="s/RuihuaSong">Ruihua Song</author>
<author pid="w/JRWen">Ji-Rong Wen</author>
<title>WenLan: Bridging Vision and Language by Large-Scale Multi-Modal Pre-Training.</title>
<year>2021</year>
<volume>abs/2103.06561</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2103.06561</ee>
<url>db/journals/corr/corr2103.html#abs-2103-06561</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2105-14477" mdate="2021-06-02">
<author pid="222/3101">Yuqing Song 0003</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<title>Towards Diverse Paragraph Captioning for Untrimmed Videos.</title>
<year>2021</year>
<volume>abs/2105.14477</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2105.14477</ee>
<url>db/journals/corr/corr2105.html#abs-2105-14477</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2106-06138" mdate="2021-06-15">
<author pid="262/6356">Ludan Ruan</author>
<author pid="274/6597">Jieting Chen</author>
<author pid="222/3101">Yuqing Song 0003</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<title>Team RUC_AIM3 Technical Report at ActivityNet 2021: Entities Object Localization.</title>
<year>2021</year>
<volume>abs/2106.06138</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2106.06138</ee>
<url>db/journals/corr/corr2106.html#abs-2106-06138</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2106-07139" mdate="2023-09-04">
<author pid="19/3011-7">Xu Han 0007</author>
<author pid="23/10446">Zhengyan Zhang</author>
<author pid="04/4910-2">Ning Ding 0002</author>
<author pid="248/2313">Yuxian Gu</author>
<author pid="82/1364-36">Xiao Liu 0036</author>
<author pid="219/6931">Yuqi Huo</author>
<author pid="152/1733">Jiezhong Qiu</author>
<author pid="50/6759">Liang Zhang</author>
<author pid="10/8243">Wentao Han</author>
<author pid="47/6668">Minlie Huang</author>
<author pid="47/2670">Qin Jin</author>
<author pid="00/6040">Yanyan Lan</author>
<author pid="51/3710-5">Yang Liu 0005</author>
<author pid="53/3245-1">Zhiyuan Liu 0001</author>
<author pid="53/5234">Zhiwu Lu 0001</author>
<author pid="69/1395">Xipeng Qiu</author>
<author pid="s/RuihuaSong">Ruihua Song</author>
<author pid="t/JieTang">Jie Tang 0001</author>
<author pid="w/JRWen">Ji-Rong Wen</author>
<author pid="58/3397">Jinhui Yuan</author>
<author pid="52/8700">Wayne Xin Zhao</author>
<author pid="50/2644-1">Jun Zhu 0001</author>
<title>Pre-Trained Models: Past, Present and Future.</title>
<year>2021</year>
<volume>abs/2106.07139</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2106.07139</ee>
<url>db/journals/corr/corr2106.html#abs-2106-07139</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2107-06779" mdate="2022-11-22">
<author pid="159/1747-3">Jingwen Hu 0003</author>
<author pid="69/10440-3">Yuchen Liu 0003</author>
<author pid="121/8902">Jinming Zhao</author>
<author pid="47/2670">Qin Jin</author>
<title>MMGCN: Multimodal Fusion via Deep Graph Convolution Network for Emotion Recognition in Conversation.</title>
<year>2021</year>
<volume>abs/2107.06779</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2107.06779</ee>
<url>db/journals/corr/corr2107.html#abs-2107-06779</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2108-02050" mdate="2021-08-05">
<author pid="249/1182">Anwen Hu</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<title>ICECAP: Information Concentrated Entity-aware Image Captioning.</title>
<year>2021</year>
<volume>abs/2108.02050</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2108.02050</ee>
<url>db/journals/corr/corr2108.html#abs-2108-02050</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2108-02059" mdate="2021-08-05">
<author pid="249/1182">Anwen Hu</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<title>Question-controlled Text-aware Image Captioning.</title>
<year>2021</year>
<volume>abs/2108.02059</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2108.02059</ee>
<url>db/journals/corr/corr2108.html#abs-2108-02059</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2108-11119" mdate="2025-06-11">
<author pid="222/3101">Yuqing Song 0003</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<author pid="05/6715">Wei Luo</author>
<author pid="33/3881">Jun Xie</author>
<author pid="h/FeiHuang-2">Fei Huang 0002</author>
<title>Product-oriented Machine Translation with Cross-modal Cross-lingual Pre-training.</title>
<year>2021</year>
<volume>abs/2108.11119</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2108.11119</ee>
<url>db/journals/corr/corr2108.html#abs-2108-11119</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2109-09920" mdate="2021-09-27">
<author pid="262/6356">Ludan Ruan</author>
<author pid="47/2670">Qin Jin</author>
<title>Survey: Transformer based Video-Language Pre-training.</title>
<year>2021</year>
<volume>abs/2109.09920</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2109.09920</ee>
<url>db/journals/corr/corr2109.html#abs-2109-09920</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2111-00865" mdate="2021-11-05">
<author pid="121/8902">Jinming Zhao</author>
<author pid="231/3910">Ruichen Li</author>
<author pid="47/2670">Qin Jin</author>
<author pid="23/8015">Xinchao Wang</author>
<author pid="36/4118">Haizhou Li 0001</author>
<title>MEmoBERT: Pre-training Model with Prompt-based Learning for Multimodal Emotion Recognition.</title>
<year>2021</year>
<volume>abs/2111.00865</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2111.00865</ee>
<url>db/journals/corr/corr2111.html#abs-2111-00865</url>
</article>
</r>
<r><inproceedings key="conf/cvpr/ChenJWW20" mdate="2021-08-30">
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<author pid="95/4442-15">Peng Wang 0015</author>
<author pid="96/3446-1">Qi Wu 0001</author>
<title>Say As You Wish: Fine-Grained Control of Image Caption Generation With Abstract Scene Graphs.</title>
<pages>9959-9968</pages>
<year>2020</year>
<booktitle>CVPR</booktitle>
<ee type="oa">https://openaccess.thecvf.com/content_CVPR_2020/html/Chen_Say_As_You_Wish_Fine-Grained_Control_of_Image_Caption_Generation_CVPR_2020_paper.html</ee>
<ee>https://doi.org/10.1109/CVPR42600.2020.00998</ee>
<crossref>conf/cvpr/2020</crossref>
<url>db/conf/cvpr/cvpr2020.html#ChenJWW20</url>
</inproceedings>
</r>
<r><inproceedings key="conf/cvpr/ChenZJW20" mdate="2021-08-30">
<author pid="153/0734">Shizhe Chen</author>
<author pid="222/2669">Yida Zhao</author>
<author pid="47/2670">Qin Jin</author>
<author pid="96/3446-1">Qi Wu 0001</author>
<title>Fine-Grained Video-Text Retrieval With Hierarchical Graph Reasoning.</title>
<pages>10635-10644</pages>
<year>2020</year>
<booktitle>CVPR</booktitle>
<ee type="oa">https://openaccess.thecvf.com/content_CVPR_2020/html/Chen_Fine-Grained_Video-Text_Retrieval_With_Hierarchical_Graph_Reasoning_CVPR_2020_paper.html</ee>
<ee>https://doi.org/10.1109/CVPR42600.2020.01065</ee>
<crossref>conf/cvpr/2020</crossref>
<url>db/conf/cvpr/cvpr2020.html#ChenZJW20</url>
</inproceedings>
</r>
<r><inproceedings key="conf/cvpr/ChenJ20" mdate="2021-08-30">
<author pid="99/6879-1">Jia Chen 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>Better Captioning With Sequence-Level Exploration.</title>
<pages>10887-10896</pages>
<year>2020</year>
<booktitle>CVPR</booktitle>
<ee type="oa">https://openaccess.thecvf.com/content_CVPR_2020/html/Chen_Better_Captioning_With_Sequence-Level_Exploration_CVPR_2020_paper.html</ee>
<ee>https://doi.org/10.1109/CVPR42600.2020.01090</ee>
<crossref>conf/cvpr/2020</crossref>
<url>db/conf/cvpr/cvpr2020.html#ChenJ20</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icmcs/ZhengCJ20" mdate="2020-08-18">
<author pid="251/3691">Sipeng Zheng</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<title>Skeleton-Based Interactive Graph Network For Human Object Interaction Detection.</title>
<pages>1-6</pages>
<year>2020</year>
<booktitle>ICME</booktitle>
<ee>https://doi.org/10.1109/ICME46284.2020.9102755</ee>
<crossref>conf/icmcs/2020</crossref>
<url>db/conf/icmcs/icme2020.html#ZhengCJ20</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/ShiHJ20" mdate="2021-01-29">
<author pid="229/3529">Jiatong Shi</author>
<author pid="272/8774">Nan Huo</author>
<author pid="47/2670">Qin Jin</author>
<title>Context-Aware Goodness of Pronunciation for Computer-Assisted Pronunciation Training.</title>
<pages>3057-3061</pages>
<year>2020</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2020-2953</ee>
<crossref>conf/interspeech/2020</crossref>
<url>db/conf/interspeech/interspeech2020.html#ShiHJ20</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/LiZHGJ20" mdate="2021-10-06">
<author pid="231/3910">Ruichen Li</author>
<author pid="121/8902">Jinming Zhao</author>
<author pid="159/1747-3">Jingwen Hu 0003</author>
<author pid="07/272">Shuai Guo</author>
<author pid="47/2670">Qin Jin</author>
<title>Multi-modal Fusion for Video Sentiment Analysis.</title>
<pages>19-25</pages>
<year>2020</year>
<booktitle>MuSe @ ACM Multimedia</booktitle>
<ee>https://doi.org/10.1145/3423327.3423671</ee>
<crossref>conf/mm/2020muse</crossref>
<url>db/conf/mm/muse2020.html#LiZHGJ20</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/WangCJ20" mdate="2020-10-15">
<author pid="227/7810">Weiying Wang</author>
<author pid="274/6597">Jieting Chen</author>
<author pid="47/2670">Qin Jin</author>
<title>VideoIC: A Video Interactive Comments Dataset and Multimodal Multitask Learning for Comments Generation.</title>
<pages>2599-2607</pages>
<year>2020</year>
<booktitle>ACM Multimedia</booktitle>
<ee>https://doi.org/10.1145/3394171.3413890</ee>
<crossref>conf/mm/2020</crossref>
<url>db/conf/mm/mm2020.html#WangCJ20</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/LiangLJ20" mdate="2020-10-15">
<author pid="243/6806">Jingjun Liang</author>
<author pid="231/3910">Ruichen Li</author>
<author pid="47/2670">Qin Jin</author>
<title>Semi-supervised Multi-modal Emotion Recognition with Cross-Modal Distribution Matching.</title>
<pages>2852-2861</pages>
<year>2020</year>
<booktitle>ACM Multimedia</booktitle>
<ee>https://doi.org/10.1145/3394171.3413579</ee>
<crossref>conf/mm/2020</crossref>
<url>db/conf/mm/mm2020.html#LiangLJ20</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/HuCJ20" mdate="2020-10-15">
<author pid="249/1182">Anwen Hu</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<title>ICECAP: Information Concentrated Entity-aware Image Captioning.</title>
<pages>4217-4225</pages>
<year>2020</year>
<booktitle>ACM Multimedia</booktitle>
<ee>https://doi.org/10.1145/3394171.3413576</ee>
<crossref>conf/mm/2020</crossref>
<url>db/conf/mm/mm2020.html#HuCJ20</url>
</inproceedings>
</r>
<r><inproceedings key="conf/trecvid/ZhaoSCJ20" mdate="2021-07-15">
<author pid="222/2669">Yida Zhao</author>
<author pid="222/3101">Yuqing Song 0003</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<title>RUC_AIM3 at TRECVID 2020: Ad-hoc Video Search &#38; Video to Text Description.</title>
<year>2020</year>
<booktitle>TRECVID</booktitle>
<ee type="oa">https://www-nlpir.nist.gov/projects/tvpubs/tv20.papers/ruc_aim3.pdf</ee>
<crossref>conf/trecvid/2020</crossref>
<url>db/conf/trecvid/trecvid2020.html#ZhaoSCJ20</url>
</inproceedings>
</r>
<r><article publtype="informal" key="journals/corr/abs-2003-00387" mdate="2020-11-12">
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<author pid="95/4442-15">Peng Wang 0015</author>
<author pid="96/3446-1">Qi Wu 0001</author>
<title>Say As You Wish: Fine-grained Control of Image Caption Generation with Abstract Scene Graphs.</title>
<year>2020</year>
<volume>abs/2003.00387</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2003.00387</ee>
<url>db/journals/corr/corr2003.html#abs-2003-00387</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2003-00392" mdate="2024-04-24">
<author pid="153/0734">Shizhe Chen</author>
<author pid="222/2669">Yida Zhao</author>
<author pid="47/2670">Qin Jin</author>
<author pid="96/3446-1">Qi Wu 0001</author>
<title>Fine-grained Video-Text Retrieval with Hierarchical Graph Reasoning.</title>
<year>2020</year>
<volume>abs/2003.00392</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2003.00392</ee>
<url>db/journals/corr/corr2003.html#abs-2003-00392</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2003-03749" mdate="2020-10-08">
<author pid="99/6879-1">Jia Chen 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>Better Captioning with Sequence-Level Exploration.</title>
<year>2020</year>
<volume>abs/2003.03749</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2003.03749</ee>
<url>db/journals/corr/corr2003.html#abs-2003-03749</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2004-05573" mdate="2020-04-14">
<author pid="153/0734">Shizhe Chen</author>
<author pid="227/7810">Weiying Wang</author>
<author pid="262/6356">Ludan Ruan</author>
<author pid="262/6579">Linli Yao</author>
<author pid="47/2670">Qin Jin</author>
<title>YouMakeup VQA Challenge: Towards Fine-grained Action Understanding in Domain-Specific Videos.</title>
<year>2020</year>
<volume>abs/2004.05573</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2004.05573</ee>
<url>db/journals/corr/corr2004.html#abs-2004-05573</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2006-07896" mdate="2020-06-17">
<author pid="222/3101">Yuqing Song 0003</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="222/2669">Yida Zhao</author>
<author pid="47/2670">Qin Jin</author>
<title>Team RUC_AIM3 Technical Report at Activitynet 2020 Task 2: Exploring Sequential Events Detection for Dense Video Captioning.</title>
<year>2020</year>
<volume>abs/2006.07896</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2006.07896</ee>
<url>db/journals/corr/corr2006.html#abs-2006-07896</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2008-00744" mdate="2020-10-20">
<author pid="188/5765">Samuel Albanie</author>
<author pid="51/3710-105">Yang Liu 0105</author>
<author pid="202/1922">Arsha Nagrani</author>
<author pid="202/1721">Antoine Miech</author>
<author pid="02/4957">Ernesto Coto</author>
<author pid="41/1854">Ivan Laptev</author>
<author pid="57/3775">Rahul Sukthankar</author>
<author pid="37/2516">Bernard Ghanem</author>
<author pid="z/AndrewZisserman">Andrew Zisserman</author>
<author pid="246/5141">Valentin Gabeur</author>
<author pid="01/6072-2">Chen Sun 0002</author>
<author pid="a/KarteekAlahari">Karteek Alahari</author>
<author pid="s/CordeliaSchmid">Cordelia Schmid</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="222/2669">Yida Zhao</author>
<author pid="47/2670">Qin Jin</author>
<author pid="251/8611">Kaixu Cui</author>
<author pid="93/4010">Hui Liu</author>
<author pid="82/4206">Chen Wang</author>
<author pid="150/4256">Yudong Jiang</author>
<author pid="271/8403">Xiaoshuai Hao</author>
<title>The End-of-End-to-End: A Video Understanding Pentathlon Challenge (2020).</title>
<year>2020</year>
<volume>abs/2008.00744</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2008.00744</ee>
<url>db/journals/corr/corr2008.html#abs-2008-00744</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2008-08647" mdate="2020-08-21">
<author pid="229/3529">Jiatong Shi</author>
<author pid="272/8774">Nan Huo</author>
<author pid="47/2670">Qin Jin</author>
<title>Context-aware Goodness of Pronunciation for Computer-Assisted Pronunciation Training.</title>
<year>2020</year>
<volume>abs/2008.08647</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2008.08647</ee>
<url>db/journals/corr/corr2008.html#abs-2008-08647</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2009-02598" mdate="2020-09-18">
<author pid="243/6806">Jingjun Liang</author>
<author pid="231/3910">Ruichen Li</author>
<author pid="47/2670">Qin Jin</author>
<title>Semi-supervised Multi-modal Emotion Recognition with Cross-Modal Distribution Matching.</title>
<year>2020</year>
<volume>abs/2009.02598</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2009.02598</ee>
<url>db/journals/corr/corr2009.html#abs-2009-02598</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-2010-12024" mdate="2020-10-27">
<author pid="229/3529">Jiatong Shi</author>
<author pid="07/272">Shuai Guo</author>
<author pid="272/8774">Nan Huo</author>
<author pid="191/5931">Yuekai Zhang</author>
<author pid="47/2670">Qin Jin</author>
<title>Sequence-to-sequence Singing Voice Synthesis with Perceptual Entropy Loss.</title>
<year>2020</year>
<volume>abs/2010.12024</volume>
<journal>CoRR</journal>
<ee type="oa">https://arxiv.org/abs/2010.12024</ee>
<url>db/journals/corr/corr2010.html#abs-2010-12024</url>
</article>
</r>
<r><article key="journals/tmm/ChenJCH19" mdate="2025-01-19">
<author orcid="0000-0002-7313-9703" pid="153/0734">Shizhe Chen</author>
<author orcid="0000-0001-6486-6020" pid="47/2670">Qin Jin</author>
<author pid="99/6879-1">Jia Chen 0001</author>
<author pid="h/AlexanderGHauptmann">Alexander G. Hauptmann</author>
<title>Generating Video Descriptions With Latent Topic Guidance.</title>
<pages>2407-2418</pages>
<year>2019</year>
<volume>21</volume>
<journal>IEEE Trans. Multim.</journal>
<number>9</number>
<ee>https://doi.org/10.1109/TMM.2019.2896515</ee>
<ee>https://www.wikidata.org/entity/Q128455538</ee>
<url>db/journals/tmm/tmm21.html#ChenJCH19</url>
</article>
</r>
<r><inproceedings key="conf/aaai/ChenJH19" mdate="2021-02-02">
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<author pid="h/AlexanderGHauptmann">Alexander G. Hauptmann</author>
<title>Unsupervised Bilingual Lexicon Induction from Mono-Lingual Multimodal Data.</title>
<pages>8207-8214</pages>
<year>2019</year>
<booktitle>AAAI</booktitle>
<ee type="oa">https://doi.org/10.1609/aaai.v33i01.33018207</ee>
<crossref>conf/aaai/2019</crossref>
<url>db/conf/aaai/aaai2019.html#ChenJH19</url>
</inproceedings>
</r>
<r><inproceedings key="conf/apsipa/LiangCJ19" mdate="2020-03-13">
<author pid="243/6806">Jingjun Liang</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<title>Semi-supervised Multimodal Emotion Recognition with Improved Wasserstein GANs.</title>
<pages>695-703</pages>
<year>2019</year>
<booktitle>APSIPA</booktitle>
<ee>https://doi.org/10.1109/APSIPAASC47483.2019.9023144</ee>
<crossref>conf/apsipa/2019</crossref>
<url>db/conf/apsipa/apsipa2019.html#LiangCJ19</url>
</inproceedings>
</r>
<r><inproceedings key="conf/emnlp/WangWCJ19" mdate="2019-12-12">
<author pid="227/7810">Weiying Wang</author>
<author pid="06/4175">Yongcheng Wang</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<title>YouMakeup: A Large-Scale Domain-Specific Multimodal Dataset for Fine-Grained Semantic Comprehension.</title>
<pages>5132-5142</pages>
<year>2019</year>
<booktitle>EMNLP/IJCNLP (1)</booktitle>
<ee type="oa">https://doi.org/10.18653/v1/D19-1517</ee>
<crossref>conf/emnlp/2019-1</crossref>
<url>db/conf/emnlp/emnlp2019-1.html#WangWCJ19</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/LiangCZJLL19" mdate="2024-10-06">
<author pid="243/6806">Jingjun Liang</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="121/8902">Jinming Zhao</author>
<author pid="47/2670">Qin Jin</author>
<author orcid="0000-0002-4213-2883" pid="83/1694">Haibo Liu</author>
<author pid="416/3281">Li Lu</author>
<title>Cross-culture Multimodal Emotion Recognition with Adversarial Learning.</title>
<pages>4000-4004</pages>
<year>2019</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2019.8683725</ee>
<crossref>conf/icassp/2019</crossref>
<url>db/conf/icassp/icassp2019.html#LiangCZJLL19</url>
</inproceedings>
</r>
<r><inproceedings key="conf/ijcai/ChenJF19" mdate="2019-08-20">
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<author pid="83/8692">Jianlong Fu</author>
<title>From Words to Sentences: A Progressive Learning Approach for Zero-resource Machine Translation with Visual Pivots.</title>
<year>2019</year>
<booktitle>IJCAI</booktitle>
<ee type="oa">https://doi.org/10.24963/ijcai.2019/685</ee>
<crossref>conf/ijcai/2019</crossref>
<url>db/conf/ijcai/ijcai2019.html#ChenJF19</url>
<pages>4932-4938</pages>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/ZhaoCLJ19" mdate="2021-01-29">
<author pid="121/8902">Jinming Zhao</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="243/6806">Jingjun Liang</author>
<author pid="47/2670">Qin Jin</author>
<title>Speech Emotion Recognition in Dyadic Dialogues with Attentive Interaction Modeling.</title>
<pages>1671-1675</pages>
<year>2019</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2019-2103</ee>
<crossref>conf/interspeech/2019</crossref>
<url>db/conf/interspeech/interspeech2019.html#ZhaoCLJ19</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mediaeval/WangYCJ19" mdate="2023-03-10">
<author pid="42/1503">Shuai Wang</author>
<author pid="262/6579">Linli Yao</author>
<author pid="274/6597">Jieting Chen</author>
<author pid="47/2670">Qin Jin</author>
<title>RUC at MediaEval 2019: Video Memorability Prediction Based on Visual Textual and Concept Related Features.</title>
<year>2019</year>
<booktitle>MediaEval</booktitle>
<ee type="oa">https://ceur-ws.org/Vol-2670/MediaEval_19_paper_33.pdf</ee>
<crossref>conf/mediaeval/2019</crossref>
<url>db/conf/mediaeval/mediaeval2019.html#WangYCJ19</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/ZhaoLLCJ19" mdate="2025-01-19">
<author pid="121/8902">Jinming Zhao</author>
<author pid="231/3910">Ruichen Li</author>
<author pid="243/6806">Jingjun Liang</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<title>Adversarial Domain Adaption for Multi-Cultural Dimensional Emotion Recognition in Dyadic Interactions.</title>
<pages>37-45</pages>
<year>2019</year>
<booktitle>AVEC@MM</booktitle>
<ee>https://doi.org/10.1145/3347320.3357692</ee>
<ee>https://www.wikidata.org/entity/Q130877208</ee>
<crossref>conf/mm/2019avec</crossref>
<url>db/conf/mm/avec2019.html#ZhaoLLCJ19</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/ZhengCJ19" mdate="2019-10-23">
<author pid="251/3691">Sipeng Zheng</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<title>Visual Relation Detection with Multi-Level Attention.</title>
<pages>121-129</pages>
<year>2019</year>
<booktitle>ACM Multimedia</booktitle>
<ee>https://doi.org/10.1145/3343031.3350962</ee>
<crossref>conf/mm/2019</crossref>
<url>db/conf/mm/mm2019.html#ZhengCJ19</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/0003CZJ19" mdate="2019-10-23">
<author pid="222/3101">Yuqing Song 0003</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="222/2669">Yida Zhao</author>
<author pid="47/2670">Qin Jin</author>
<title>Unpaired Cross-lingual Image Caption Generation with Self-Supervised Rewards.</title>
<pages>784-792</pages>
<year>2019</year>
<booktitle>ACM Multimedia</booktitle>
<ee>https://doi.org/10.1145/3343031.3350996</ee>
<crossref>conf/mm/2019</crossref>
<url>db/conf/mm/mm2019.html#0003CZJ19</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/ChenLFSJLQWZ19" mdate="2021-08-13">
<author pid="153/0734">Shizhe Chen</author>
<author pid="39/3711-1">Bei Liu 0001</author>
<author pid="83/8692">Jianlong Fu</author>
<author pid="s/RuihuaSong">Ruihua Song</author>
<author pid="47/2670">Qin Jin</author>
<author pid="39/7629">Pingping Lin</author>
<author pid="166/6091">Xiaoyu Qi</author>
<author pid="239/3587">Chunting Wang</author>
<author pid="98/2691">Jin Zhou</author>
<title>Neural Storyboard Artist: Visualizing Stories with Coherent Image Sequences.</title>
<pages>2236-2244</pages>
<year>2019</year>
<booktitle>ACM Multimedia</booktitle>
<ee>https://doi.org/10.1145/3343031.3350571</ee>
<crossref>conf/mm/2019</crossref>
<url>db/conf/mm/mm2019.html#ChenLFSJLQWZ19</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/ZhengCCJ19" mdate="2019-10-23">
<author pid="251/3691">Sipeng Zheng</author>
<author pid="84/7543">Xiangyu Chen</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<title>Relation Understanding in Videos.</title>
<pages>2662-2666</pages>
<year>2019</year>
<booktitle>ACM Multimedia</booktitle>
<ee>https://doi.org/10.1145/3343031.3356080</ee>
<crossref>conf/mm/2019</crossref>
<url>db/conf/mm/mm2019.html#ZhengCCJ19</url>
</inproceedings>
</r>
<r><inproceedings key="conf/trecvid/0003ZCJ19" mdate="2020-03-13">
<author pid="222/3101">Yuqing Song 0003</author>
<author pid="222/2669">Yida Zhao</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<title>RUC_AIM3 at TRECVID 2019: Video to Text.</title>
<year>2019</year>
<booktitle>TRECVID</booktitle>
<ee type="oa">https://www-nlpir.nist.gov/projects/tvpubs/tv19.papers/ruc_aim3.pdf</ee>
<crossref>conf/trecvid/2019</crossref>
<url>db/conf/trecvid/trecvid2019.html#0003ZCJ19</url>
</inproceedings>
</r>
<r><article publtype="informal" key="journals/corr/abs-1906-00378" mdate="2019-06-13">
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<author pid="h/AlexanderGHauptmann">Alexander G. Hauptmann</author>
<title>Unsupervised Bilingual Lexicon Induction from Mono-lingual Multimodal Data.</title>
<year>2019</year>
<volume>abs/1906.00378</volume>
<journal>CoRR</journal>
<ee type="oa">http://arxiv.org/abs/1906.00378</ee>
<url>db/journals/corr/corr1906.html#abs-1906-00378</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-1906-00872" mdate="2019-06-13">
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<author pid="83/8692">Jianlong Fu</author>
<title>From Words to Sentences: A Progressive Learning Approach for Zero-resource Machine Translation with Visual Pivots.</title>
<year>2019</year>
<volume>abs/1906.00872</volume>
<journal>CoRR</journal>
<ee type="oa">http://arxiv.org/abs/1906.00872</ee>
<url>db/journals/corr/corr1906.html#abs-1906-00872</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-1907-05092" mdate="2021-08-13">
<author pid="153/0734">Shizhe Chen</author>
<author pid="222/3101">Yuqing Song 0003</author>
<author pid="222/2669">Yida Zhao</author>
<author pid="47/2670">Qin Jin</author>
<author pid="207/1887">Zhaoyang Zeng</author>
<author pid="39/3711-1">Bei Liu 0001</author>
<author pid="83/8692">Jianlong Fu</author>
<author pid="h/AlexanderGHauptmann">Alexander G. Hauptmann</author>
<title>Activitynet 2019 Task 3: Exploring Contexts for Dense Captioning Events in Videos.</title>
<year>2019</year>
<volume>abs/1907.05092</volume>
<journal>CoRR</journal>
<ee type="oa">http://arxiv.org/abs/1907.05092</ee>
<url>db/journals/corr/corr1907.html#abs-1907-05092</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-1908-05407" mdate="2019-08-19">
<author pid="222/3101">Yuqing Song 0003</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="222/2669">Yida Zhao</author>
<author pid="47/2670">Qin Jin</author>
<title>Unpaired Cross-lingual Image Caption Generation with Self-Supervised Rewards.</title>
<year>2019</year>
<volume>abs/1908.05407</volume>
<journal>CoRR</journal>
<ee type="oa">http://arxiv.org/abs/1908.05407</ee>
<url>db/journals/corr/corr1908.html#abs-1908-05407</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-1910-06737" mdate="2024-07-17">
<author pid="153/0734">Shizhe Chen</author>
<author pid="222/2669">Yida Zhao</author>
<author pid="222/3101">Yuqing Song 0003</author>
<author pid="47/2670">Qin Jin</author>
<author pid="96/3446-1">Qi Wu 0001</author>
<title>Integrating Temporal and Spatial Attentions for VATEX Video Captioning Challenge 2019.</title>
<year>2019</year>
<volume>abs/1910.06737</volume>
<journal>CoRR</journal>
<ee type="oa">http://arxiv.org/abs/1910.06737</ee>
<url>db/journals/corr/corr1910.html#abs-1910-06737</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-1911-10460" mdate="2021-08-13">
<author pid="153/0734">Shizhe Chen</author>
<author pid="39/3711-1">Bei Liu 0001</author>
<author pid="83/8692">Jianlong Fu</author>
<author pid="s/RuihuaSong">Ruihua Song</author>
<author pid="47/2670">Qin Jin</author>
<author pid="39/7629">Pingping Lin</author>
<author pid="166/6091">Xiaoyu Qi</author>
<author pid="239/3587">Chunting Wang</author>
<author pid="98/2691">Jin Zhou</author>
<title>Neural Storyboard Artist: Visualizing Stories with Coherent Image Sequences.</title>
<year>2019</year>
<volume>abs/1911.10460</volume>
<journal>CoRR</journal>
<ee type="oa">http://arxiv.org/abs/1911.10460</ee>
<url>db/journals/corr/corr1911.html#abs-1911-10460</url>
</article>
</r>
<r><inproceedings key="conf/mediaeval/WangWCJ18" mdate="2023-03-10">
<author pid="42/1503">Shuai Wang</author>
<author pid="227/7810">Weiying Wang</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<title>RUC at MediaEval 2018: Visual and Textual Features Exploration for Predicting Media Memorability.</title>
<year>2018</year>
<booktitle>MediaEval</booktitle>
<ee type="oa">https://ceur-ws.org/Vol-2283/MediaEval_18_paper_51.pdf</ee>
<crossref>conf/mediaeval/2018</crossref>
<url>db/conf/mediaeval/mediaeval2018.html#WangWCJ18</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mir/ChenCJH18" mdate="2020-10-08">
<author pid="153/0734">Shizhe Chen</author>
<author pid="99/6879-1">Jia Chen 0001</author>
<author pid="47/2670">Qin Jin</author>
<author pid="h/AlexanderGHauptmann">Alexander G. Hauptmann</author>
<title>Class-aware Self-Attention for Audio Event Recognition.</title>
<pages>28-36</pages>
<year>2018</year>
<booktitle>ICMR</booktitle>
<ee>https://doi.org/10.1145/3206025.3206067</ee>
<crossref>conf/mir/2018</crossref>
<url>db/conf/mir/icmr2018.html#ChenCJH18</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/Jin18" mdate="2018-11-29">
<author pid="47/2670">Qin Jin</author>
<title>Session details: Deep-2 (Recognition).</title>
<year>2018</year>
<booktitle>ACM Multimedia</booktitle>
<ee>https://dl.acm.org/citation.cfm?id=3286931</ee>
<crossref>conf/mm/2018</crossref>
<url>db/conf/mm/mm2018.html#Jin18</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/ZhaoLCJ18" mdate="2025-01-19">
<author pid="121/8902">Jinming Zhao</author>
<author pid="231/3910">Ruichen Li</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<title>Multi-modal Multi-cultural Dimensional Continues Emotion Recognition in Dyadic Interactions.</title>
<pages>65-72</pages>
<year>2018</year>
<booktitle>AVEC@MM</booktitle>
<ee>https://doi.org/10.1145/3266302.3266313</ee>
<ee>https://www.wikidata.org/entity/Q130883968</ee>
<crossref>conf/mm/2018avec</crossref>
<url>db/conf/mm/avec2018.html#ZhaoLCJ18</url>
</inproceedings>
</r>
<r><inproceedings key="conf/pcm/LinJC0Z18" mdate="2018-09-18">
<author pid="07/1292">Xiaozhu Lin</author>
<author pid="47/2670">Qin Jin</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="222/3101">Yuqing Song 0003</author>
<author pid="222/2669">Yida Zhao</author>
<title>iMakeup: Makeup Instructional Video Dataset for Fine-Grained Dense Video Captioning.</title>
<pages>78-88</pages>
<year>2018</year>
<booktitle>PCM (3)</booktitle>
<ee>https://doi.org/10.1007/978-3-030-00764-5_8</ee>
<crossref>conf/pcm/2018-3</crossref>
<url>db/conf/pcm/pcm2018-3.html#LinJC0Z18</url>
</inproceedings>
</r>
<r><inproceedings key="conf/pcm/ZhaoCJ18" mdate="2018-09-19">
<author pid="121/8902">Jinming Zhao</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<title>Multimodal Dimensional and Continuous Emotion Recognition in Dyadic Video Interactions.</title>
<pages>301-312</pages>
<year>2018</year>
<booktitle>PCM (1)</booktitle>
<ee>https://doi.org/10.1007/978-3-030-00776-8_28</ee>
<crossref>conf/pcm/2018-1</crossref>
<url>db/conf/pcm/pcm2018-1.html#ZhaoCJ18</url>
</inproceedings>
</r>
<r><inproceedings key="conf/trecvid/ChenCJH00VCLHLK18" mdate="2024-08-13">
<author pid="99/6879-1">Jia Chen 0001</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<author pid="h/AlexanderGHauptmann">Alexander G. Hauptmann</author>
<author pid="154/3943-1">Po-Yao Huang 0001</author>
<author pid="62/10704-1">Junwei Liang 0001</author>
<author pid="177/2366">Vaibhav</author>
<author pid="116/8412">Xiaojun Chang</author>
<author pid="23/108-11">Jiang Liu 0011</author>
<author pid="76/10031">Ting-Yao Hu</author>
<author pid="184/5777">Wenhe Liu</author>
<author pid="52/7566">Wei Ke</author>
<author pid="208/4288">Wayner Barrios</author>
<author pid="19/8535">Haroon Idrees</author>
<author pid="207/7436">Donghyun Yoo</author>
<author pid="71/3516">Yaser Sheikh</author>
<author pid="62/5884">Ruslan Salakhutdinov</author>
<author pid="42/163">Kris Kitani</author>
<author pid="94/3756-7">Dong Huang 0007</author>
<title>Informedia @ TRECVID 2018: Ad-hoc Video Search, Video to Text Description, Activities in Extended video.</title>
<year>2018</year>
<booktitle>TRECVID</booktitle>
<ee type="oa">https://www-nlpir.nist.gov/projects/tvpubs/tv18.papers/inf.pdf</ee>
<crossref>conf/trecvid/2018</crossref>
<url>db/conf/trecvid/trecvid2018.html#ChenCJH00VCLHLK18</url>
</inproceedings>
</r>
<r><article publtype="informal" key="journals/corr/abs-1806-08854" mdate="2018-08-13">
<author pid="153/0734">Shizhe Chen</author>
<author pid="222/3101">Yuqing Song 0003</author>
<author pid="222/2669">Yida Zhao</author>
<author pid="222/3007">Jiarong Qiu</author>
<author pid="47/2670">Qin Jin</author>
<author pid="h/AlexanderGHauptmann">Alexander G. Hauptmann</author>
<title>RUC+CMU: System Report for Dense Captioning Events in Videos.</title>
<year>2018</year>
<volume>abs/1806.08854</volume>
<journal>CoRR</journal>
<ee type="oa">http://arxiv.org/abs/1806.08854</ee>
<url>db/journals/corr/corr1806.html#abs-1806-08854</url>
</article>
</r>
<r><article key="journals/ijids/ZhangJH17" mdate="2022-08-23">
<author pid="133/6346">Mengying Zhang</author>
<author pid="47/2670">Qin Jin</author>
<author pid="43/5900-4">Hongwei Liu 0004</author>
<title>Group division based on common weights in cross efficiency evaluation.</title>
<pages>209-223</pages>
<year>2017</year>
<volume>9</volume>
<journal>Int. J. Inf. Decis. Sci.</journal>
<number>3</number>
<ee>https://doi.org/10.1504/IJIDS.2017.10007803</ee>
<url>db/journals/ijids/ijids9.html#ZhangJH17</url>
</article>
</r>
<r><inproceedings key="conf/fgr/LiCJ17" mdate="2023-03-24">
<author pid="151/4579">Xinrui Li</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<title>Facial Action Units Detection with Multi-Features and -AUs Fusion.</title>
<pages>860-865</pages>
<year>2017</year>
<booktitle>FG</booktitle>
<ee>https://doi.org/10.1109/FG.2017.110</ee>
<ee>https://doi.ieeecomputersociety.org/10.1109/FG.2017.110</ee>
<crossref>conf/fgr/2017</crossref>
<url>db/conf/fgr/fg2017.html#LiCJ17</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icmi/WangWZCJZQ17" mdate="2025-10-02">
<author pid="42/1503">Shuai Wang</author>
<author pid="203/1536-1">Wenxuan Wang 0001</author>
<author pid="121/8902">Jinming Zhao</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<author pid="10/3523">Shilei Zhang</author>
<author pid="20/4298-1">Yong Qin 0001</author>
<title>Emotion recognition with multimodal features and temporal models.</title>
<pages>598-602</pages>
<year>2017</year>
<booktitle>ICMI</booktitle>
<ee>https://doi.org/10.1145/3136755.3143016</ee>
<ee>https://www.wikidata.org/entity/Q130988150</ee>
<crossref>conf/icmi/2017</crossref>
<url>db/conf/icmi/icmi2017.html#WangWZCJZQ17</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mediaeval/WangCZWJ17" mdate="2025-10-02">
<author pid="42/1503">Shuai Wang</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="121/8902">Jinming Zhao</author>
<author pid="203/1536-1">Wenxuan Wang 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>RUC at MediaEval 2017: Predicting Media Interestingness Task.</title>
<year>2017</year>
<booktitle>MediaEval</booktitle>
<ee type="oa">https://ceur-ws.org/Vol-1984/Mediaeval_2017_paper_30.pdf</ee>
<crossref>conf/mediaeval/2017</crossref>
<url>db/conf/mediaeval/mediaeval2017.html#WangCZWJ17</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mir/ChenCJ17" mdate="2020-10-08">
<author pid="153/0734">Shizhe Chen</author>
<author pid="99/6879-1">Jia Chen 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>Generating Video Descriptions with Topic Guidance.</title>
<pages>5-13</pages>
<year>2017</year>
<booktitle>ICMR</booktitle>
<ee>https://doi.org/10.1145/3078971.3079000</ee>
<crossref>conf/mir/2017</crossref>
<url>db/conf/mir/icmr2017.html#ChenCJ17</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/ChenJZW17" mdate="2018-11-06">
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<author pid="121/8902">Jinming Zhao</author>
<author pid="42/1503">Shuai Wang</author>
<title>Multimodal Multi-task Learning for Dimensional and Continuous Emotion Recognition.</title>
<pages>19-26</pages>
<year>2017</year>
<booktitle>AVEC@ACM Multimedia</booktitle>
<ee>https://doi.org/10.1145/3133944.3133949</ee>
<crossref>conf/mm/2017avec</crossref>
<url>db/conf/mm/avec2017.html#ChenJZW17</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/ChenCJH17" mdate="2020-10-08">
<author pid="153/0734">Shizhe Chen</author>
<author pid="99/6879-1">Jia Chen 0001</author>
<author pid="47/2670">Qin Jin</author>
<author pid="h/AlexanderGHauptmann">Alexander G. Hauptmann</author>
<title>Video Captioning with Guidance of Multimodal Latent Topics.</title>
<pages>1838-1846</pages>
<year>2017</year>
<booktitle>ACM Multimedia</booktitle>
<ee>https://doi.org/10.1145/3123266.3123420</ee>
<crossref>conf/mm/2017</crossref>
<url>db/conf/mm/mm2017.html#ChenCJH17</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/JinCCH17" mdate="2020-10-08">
<author pid="47/2670">Qin Jin</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="99/6879-1">Jia Chen 0001</author>
<author pid="h/AlexanderGHauptmann">Alexander G. Hauptmann</author>
<title>Knowing Yourself: Improving Video Caption via In-depth Recap.</title>
<pages>1906-1911</pages>
<year>2017</year>
<booktitle>ACM Multimedia</booktitle>
<ee>https://doi.org/10.1145/3123266.3127901</ee>
<crossref>conf/mm/2017</crossref>
<url>db/conf/mm/mm2017.html#JinCCH17</url>
</inproceedings>
</r>
<r><inproceedings key="conf/trecvid/Chen00CGJH17" mdate="2020-10-08">
<author pid="99/6879-1">Jia Chen 0001</author>
<author pid="62/10704-1">Junwei Liang 0001</author>
<author pid="23/108-11">Jiang Liu 0011</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="82/9201">Chenqiang Gao</author>
<author pid="47/2670">Qin Jin</author>
<author pid="h/AlexanderGHauptmann">Alexander G. Hauptmann</author>
<title>Informedia @ TRECVID 2017.</title>
<year>2017</year>
<booktitle>TRECVID</booktitle>
<ee type="oa">https://www-nlpir.nist.gov/projects/tvpubs/tv17.papers/inf.pdf</ee>
<crossref>conf/trecvid/2017</crossref>
<url>db/conf/trecvid/trecvid2017.html#Chen00CGJH17</url>
</inproceedings>
</r>
<r><article publtype="informal" key="journals/corr/abs-1708-09666" mdate="2020-10-08">
<author pid="153/0734">Shizhe Chen</author>
<author pid="99/6879-1">Jia Chen 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>Generating Video Descriptions with Topic Guidance.</title>
<year>2017</year>
<volume>abs/1708.09666</volume>
<journal>CoRR</journal>
<ee type="oa">http://arxiv.org/abs/1708.09666</ee>
<url>db/journals/corr/corr1708.html#abs-1708-09666</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-1708-09667" mdate="2020-10-08">
<author pid="153/0734">Shizhe Chen</author>
<author pid="99/6879-1">Jia Chen 0001</author>
<author pid="47/2670">Qin Jin</author>
<author pid="h/AlexanderGHauptmann">Alexander G. Hauptmann</author>
<title>Video Captioning with Guidance of Multimodal Latent Topics.</title>
<year>2017</year>
<volume>abs/1708.09667</volume>
<journal>CoRR</journal>
<ee type="oa">http://arxiv.org/abs/1708.09667</ee>
<url>db/journals/corr/corr1708.html#abs-1708-09667</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/abs-1709-02251" mdate="2018-08-13">
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<title>Multi-modal Conditional Attention Fusion for Dimensional Emotion Prediction.</title>
<year>2017</year>
<volume>abs/1709.02251</volume>
<journal>CoRR</journal>
<ee type="oa">http://arxiv.org/abs/1709.02251</ee>
<url>db/journals/corr/corr1709.html#abs-1709-02251</url>
</article>
</r>
<r><article key="journals/asc/YangWJX16" mdate="2019-09-16">
<author pid="36/4658-1">Gang Yang 0001</author>
<author pid="136/9323">Shaohui Wu</author>
<author pid="47/2670">Qin Jin</author>
<author orcid="0000-0001-6228-0317" pid="96/979">Jieping Xu</author>
<title>A hybrid approach based on stochastic competitive Hopfield neural network and efficient genetic algorithm for frequency assignment problem.</title>
<pages>104-116</pages>
<year>2016</year>
<volume>39</volume>
<journal>Appl. Soft Comput.</journal>
<ee>https://doi.org/10.1016/j.asoc.2015.10.056</ee>
<url>db/journals/asc/asc39.html#YangWJX16</url>
</article>
</r>
<r><article key="journals/ijisscm/ZhangJ16" mdate="2020-09-08">
<author pid="133/6346">Mengying Zhang</author>
<author pid="47/2670">Qin Jin</author>
<title>Coordinate the Express Delivery Supply Chain with Option Contracts.</title>
<pages>1-21</pages>
<year>2016</year>
<volume>9</volume>
<journal>Int. J. Inf. Syst. Supply Chain Manag.</journal>
<number>4</number>
<ee>https://doi.org/10.4018/IJISSCM.2016100101</ee>
<url>db/journals/ijisscm/ijisscm9.html#ZhangJ16</url>
</article>
</r>
<r><article key="journals/ijkbo/ZhangJH16" mdate="2022-08-23">
<author pid="133/6346">Mengying Zhang</author>
<author pid="47/2670">Qin Jin</author>
<author pid="43/5900-4">Hongwei Liu 0004</author>
<title>The Study of the Entrepreneurial Leadership Style of Real Estate Industry in China: Based on the Content Analysis of Microblog.</title>
<pages>45-57</pages>
<year>2016</year>
<volume>6</volume>
<journal>Int. J. Knowl. Based Organ.</journal>
<number>3</number>
<ee>https://doi.org/10.4018/IJKBO.2016070103</ee>
<url>db/journals/ijkbo/ijkbo6.html#ZhangJH16</url>
</article>
</r>
<r><article key="journals/tois/ChenJZBZSY16" mdate="2022-10-02">
<author pid="99/6879-1">Jia Chen 0001</author>
<author pid="47/2670">Qin Jin</author>
<author orcid="0000-0001-5068-025X" pid="73/1947">Shiwan Zhao</author>
<author pid="31/529">Shenghua Bao</author>
<author pid="89/5992-7">Li Zhang 0007</author>
<author pid="87/1363">Zhong Su</author>
<author pid="43/5685-1">Yong Yu 0001</author>
<title>Boosting Recommendation in Unexplored Categories by User Price Preference.</title>
<pages>12:1-12:27</pages>
<year>2016</year>
<volume>35</volume>
<journal>ACM Trans. Inf. Syst.</journal>
<number>2</number>
<ee>https://doi.org/10.1145/2978579</ee>
<url>db/journals/tois/tois35.html#ChenJZBZSY16</url>
</article>
</r>
<r><inproceedings key="conf/ccpr/MuCJ16" mdate="2017-05-24">
<author pid="187/9132">Guankun Mu</author>
<author pid="169/0393">Haibing Cao</author>
<author pid="47/2670">Qin Jin</author>
<title>Violent Scene Detection Using Convolutional Neural Networks and Deep Audio Features.</title>
<pages>451-463</pages>
<year>2016</year>
<booktitle>CCPR (2)</booktitle>
<ee>https://doi.org/10.1007/978-981-10-3005-5_37</ee>
<crossref>conf/ccpr/2016-2</crossref>
<url>db/conf/ccpr/ccpr2016-2.html#MuCJ16</url>
</inproceedings>
</r>
<r><inproceedings key="conf/ccpr/ChenDLLJLL16" mdate="2017-05-24">
<author pid="153/0734">Shizhe Chen</author>
<author pid="187/9165">Yujie Dian</author>
<author pid="151/4579">Xinrui Li</author>
<author pid="07/1292">Xiaozhu Lin</author>
<author pid="47/2670">Qin Jin</author>
<author pid="83/1694">Haibo Liu</author>
<author pid="416/3281">Li Lu</author>
<title>Emotion Recognition in Videos via Fusing Multimodal Features.</title>
<pages>632-644</pages>
<year>2016</year>
<booktitle>CCPR (2)</booktitle>
<ee>https://doi.org/10.1007/978-981-10-3005-5_52</ee>
<crossref>conf/ccpr/2016-2</crossref>
<url>db/conf/ccpr/ccpr2016-2.html#ChenDLLJLL16</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icmi/ChenLJZQ16" mdate="2025-01-19">
<author pid="153/0734">Shizhe Chen</author>
<author pid="151/4579">Xinrui Li</author>
<author pid="47/2670">Qin Jin</author>
<author pid="10/3523">Shilei Zhang</author>
<author pid="20/4298-1">Yong Qin 0001</author>
<title>Video emotion recognition in the wild based on fusion of multimodal features.</title>
<pages>494-500</pages>
<year>2016</year>
<booktitle>ICMI</booktitle>
<ee>https://doi.org/10.1145/2993148.2997629</ee>
<ee>https://www.wikidata.org/entity/Q130983941</ee>
<crossref>conf/icmi/2016</crossref>
<url>db/conf/icmi/icmi2016.html#ChenLJZQ16</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/JinLL16" mdate="2024-10-06">
<author pid="47/2670">Qin Jin</author>
<author orcid="0000-0003-2219-5569" pid="62/10704-1">Junwei Liang 0001</author>
<author pid="07/1292">Xiaozhu Lin</author>
<title>Generating Natural Video Descriptions via Multimodal Processing.</title>
<pages>570-574</pages>
<year>2016</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2016-380</ee>
<ee>https://www.wikidata.org/entity/Q121878676</ee>
<crossref>conf/interspeech/2016</crossref>
<url>db/conf/interspeech/interspeech2016.html#JinLL16</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mediaeval/ChenDJ16" mdate="2023-03-10">
<author pid="153/0734">Shizhe Chen</author>
<author pid="187/9165">Yujie Dian</author>
<author pid="47/2670">Qin Jin</author>
<title>RUC at MediaEval 2016: Predicting Media Interestingness Task.</title>
<year>2016</year>
<booktitle>MediaEval</booktitle>
<crossref>conf/mediaeval/2016</crossref>
<url>db/conf/mediaeval/mediaeval2016.html#ChenDJ16</url>
<ee type="oa">https://ceur-ws.org/Vol-1739/MediaEval_2016_paper_33.pdf</ee>
</inproceedings>
</r>
<r><inproceedings key="conf/mediaeval/ChenJ16" mdate="2023-03-10">
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<title>RUC at MediaEval 2016 Emotional Impact of Movies Task: Fusion of Multimodal Features.</title>
<year>2016</year>
<booktitle>MediaEval</booktitle>
<crossref>conf/mediaeval/2016</crossref>
<url>db/conf/mediaeval/mediaeval2016.html#ChenJ16</url>
<ee type="oa">https://ceur-ws.org/Vol-1739/MediaEval_2016_paper_37.pdf</ee>
</inproceedings>
</r>
<r><inproceedings key="conf/mir/JinL16" mdate="2024-10-06">
<author pid="47/2670">Qin Jin</author>
<author orcid="0000-0003-2219-5569" pid="62/10704-1">Junwei Liang 0001</author>
<title>Video Description Generation using Audio and Visual Cues.</title>
<pages>239-242</pages>
<year>2016</year>
<booktitle>ICMR</booktitle>
<ee>https://doi.org/10.1145/2911996.2912043</ee>
<ee>https://www.wikidata.org/entity/Q121878679</ee>
<crossref>conf/mir/2016</crossref>
<url>db/conf/mir/icmr2016.html#JinL16</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/ChenJ16" mdate="2018-11-06">
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<title>Multi-modal Conditional Attention Fusion for Dimensional Emotion Prediction.</title>
<pages>571-575</pages>
<year>2016</year>
<booktitle>ACM Multimedia</booktitle>
<ee>https://doi.org/10.1145/2964284.2967286</ee>
<crossref>conf/mm/2016</crossref>
<url>db/conf/mm/mm2016.html#ChenJ16</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/LiHJX16" mdate="2025-03-03">
<author orcid="0000-0002-0220-8310" pid="58/5856">Xirong Li 0001</author>
<author pid="167/4776">Yujia Huo</author>
<author pid="47/2670">Qin Jin</author>
<author orcid="0000-0001-6228-0317" pid="96/979">Jieping Xu</author>
<title>Detecting Violence in Video using Subclasses.</title>
<pages>586-590</pages>
<year>2016</year>
<booktitle>ACM Multimedia</booktitle>
<ee>https://doi.org/10.1145/2964284.2967289</ee>
<crossref>conf/mm/2016</crossref>
<url>db/conf/mm/mm2016.html#LiHJX16</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/XiongCJZ16" mdate="2023-03-20">
<author pid="188/1149-1">Yifan Xiong 0001</author>
<author pid="99/6879-1">Jia Chen 0001</author>
<author pid="47/2670">Qin Jin</author>
<author pid="94/3019">Chao Zhang</author>
<title>History Rhyme: Searching Historic Events by Multimedia Knowledge.</title>
<pages>749-751</pages>
<year>2016</year>
<booktitle>ACM Multimedia</booktitle>
<ee>https://doi.org/10.1145/2964284.2973832</ee>
<crossref>conf/mm/2016</crossref>
<url>db/conf/mm/mm2016.html#XiongCJZ16</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/ChenJX16" mdate="2023-03-20">
<author pid="99/6879-1">Jia Chen 0001</author>
<author pid="47/2670">Qin Jin</author>
<author pid="188/1149-1">Yifan Xiong 0001</author>
<title>Semantic Image Profiling for Historic Events: Linking Images to Phrases.</title>
<pages>1028-1037</pages>
<year>2016</year>
<booktitle>ACM Multimedia</booktitle>
<ee>https://doi.org/10.1145/2964284.2964306</ee>
<crossref>conf/mm/2016</crossref>
<url>db/conf/mm/mm2016.html#ChenJX16</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/JinCCXH16" mdate="2023-03-20">
<author pid="47/2670">Qin Jin</author>
<author pid="99/6879-1">Jia Chen 0001</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="188/1149-1">Yifan Xiong 0001</author>
<author pid="h/AlexanderGHauptmann">Alexander G. Hauptmann</author>
<title>Describing Videos using Multi-modal Fusion.</title>
<pages>1087-1091</pages>
<year>2016</year>
<booktitle>ACM Multimedia</booktitle>
<ee>https://doi.org/10.1145/2964284.2984065</ee>
<crossref>conf/mm/2016</crossref>
<url>db/conf/mm/mm2016.html#JinCCXH16</url>
</inproceedings>
</r>
<r><inproceedings key="conf/pcm/LiJ16" mdate="2025-03-03">
<author orcid="0000-0002-0220-8310" pid="58/5856">Xirong Li 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>Improving Image Captioning by Concept-Based Sentence Reranking.</title>
<pages>231-240</pages>
<year>2016</year>
<booktitle>PCM (2)</booktitle>
<ee>https://doi.org/10.1007/978-3-319-48896-7_23</ee>
<crossref>conf/pcm/2016-2</crossref>
<url>db/conf/pcm/pcm2016-2.html#LiJ16</url>
</inproceedings>
</r>
<r><inproceedings key="conf/trecvid/0001C0L0LPFJSC016" mdate="2020-10-22">
<author pid="62/10704-1">Junwei Liang 0001</author>
<author pid="99/6879-1">Jia Chen 0001</author>
<author pid="154/3943-1">Poyao Huang 0001</author>
<author pid="150/6544">Xuanchong Li</author>
<author pid="22/752-4">Lu Jiang 0004</author>
<author pid="27/3780">Zhenzhong Lan</author>
<author pid="172/0936">Pingbo Pan</author>
<author pid="184/5722">Hehe Fan</author>
<author pid="47/2670">Qin Jin</author>
<author pid="68/365-1">Jiande Sun 0001</author>
<author pid="48/4792">Yang Chen</author>
<author pid="33/4854-1">Yi Yang 0001</author>
<author pid="h/AlexanderGHauptmann">Alexander G. Hauptmann</author>
<title>Informedia @ TRECVID 2016.</title>
<year>2016</year>
<booktitle>TRECVID</booktitle>
<ee type="oa">https://www-nlpir.nist.gov/projects/tvpubs/tv16.papers/inf.pdf</ee>
<crossref>conf/trecvid/2016</crossref>
<url>db/conf/trecvid/trecvid2016.html#0001C0L0LPFJSC016</url>
</inproceedings>
</r>
<r><article publtype="informal" key="journals/corr/LiHXJ16" mdate="2021-09-13">
<author pid="58/5856">Xirong Li 0001</author>
<author pid="167/4776">Yujia Huo</author>
<author pid="96/979">Jieping Xu</author>
<author pid="47/2670">Qin Jin</author>
<title>Detecting Violence in Video using Subclasses.</title>
<year>2016</year>
<volume>abs/1604.08088</volume>
<journal>CoRR</journal>
<ee type="oa">http://arxiv.org/abs/1604.08088</ee>
<url>db/journals/corr/corr1604.html#LiHXJ16</url>
</article>
</r>
<r><article publtype="informal" key="journals/corr/LiJ16a" mdate="2021-09-13">
<author pid="58/5856">Xirong Li 0001</author>
<author pid="47/2670">Qin Jin</author>
<title>Improving Image Captioning by Concept-based Sentence Reranking.</title>
<year>2016</year>
<volume>abs/1605.00855</volume>
<journal>CoRR</journal>
<ee type="oa">http://arxiv.org/abs/1605.00855</ee>
<url>db/journals/corr/corr1605.html#LiJ16a</url>
</article>
</r>
<r><article key="journals/ijon/ChenLJBSY15" mdate="2022-05-30">
<author pid="99/6879-1">Jia Chen 0001</author>
<author pid="82/0-22">Min Li 0022</author>
<author pid="47/2670">Qin Jin</author>
<author pid="31/529">Shenghua Bao</author>
<author pid="87/1363">Zhong Su</author>
<author pid="43/5685-1">Yong Yu 0001</author>
<title>Lead curve detection in drawings with complex cross-points.</title>
<pages>35-46</pages>
<year>2015</year>
<volume>168</volume>
<journal>Neurocomputing</journal>
<ee>https://doi.org/10.1016/j.neucom.2015.06.018</ee>
<url>db/journals/ijon/ijon168.html#ChenLJBSY15</url>
</article>
</r>
<r><article key="journals/jsjkx/JinCL0X15" mdate="2021-09-13">
<author pid="47/2670">Qin Jin</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="58/5856">Xirong Li 0001</author>
<author pid="36/4658-1">Gang Yang 0001</author>
<author pid="96/979">Jieping Xu</author>
<title>&#22522;&#20110;&#22768;&#23398;&#29305;&#24449;&#30340;&#35821;&#35328;&#24773;&#24863;&#35782;&#21035; (Speech Emotion Recognition Based on Acoustic Features).</title>
<pages>24-28</pages>
<year>2015</year>
<volume>42</volume>
<journal>&#35745;&#31639;&#26426;&#31185;&#23398;</journal>
<number>9</number>
<ee type="oa">https://doi.org/10.11896/j.issn.1002-137X.2015.09.005</ee>
<ee type="oa">http://www.jsjkx.com/EN/10.11896/j.issn.1002-137X.2015.09.005</ee>
<url>db/journals/jsjkx/jsjkx42.html#JinCL0X15</url>
</article>
</r>
<r><article key="journals/pvldb/ChenJ15" mdate="2020-04-25">
<author pid="c/ShiminChen">Shimin Chen</author>
<author pid="47/2670">Qin Jin</author>
<title>Persistent B+-Trees in Non-Volatile Main Memory.</title>
<pages>786-797</pages>
<year>2015</year>
<volume>8</volume>
<journal>Proc. VLDB Endow.</journal>
<number>7</number>
<ee type="oa">http://www.vldb.org/pvldb/vol8/p786-chen.pdf</ee>
<ee>https://doi.org/10.14778/2752939.2752947</ee>
<url>db/journals/pvldb/pvldb8.html#ChenJ15</url>
</article>
</r>
<r><article key="journals/tmm/ChenJBSCY15" mdate="2020-10-08">
<author pid="99/6879-1">Jia Chen 0001</author>
<author pid="47/2670">Qin Jin</author>
<author pid="31/529">Shenghua Bao</author>
<author pid="87/1363">Zhong Su</author>
<author pid="c/ShiminChen">Shimin Chen</author>
<author pid="43/5685-1">Yong Yu 0001</author>
<title>Exploitation and Exploration Balanced Hierarchical Summary for Landmark Images.</title>
<pages>1773-1786</pages>
<year>2015</year>
<volume>17</volume>
<journal>IEEE Trans. Multim.</journal>
<number>10</number>
<ee>https://doi.org/10.1109/TMM.2015.2460111</ee>
<url>db/journals/tmm/tmm17.html#ChenJBSCY15</url>
</article>
</r>
<r><inproceedings key="conf/acii/WuJ15" mdate="2023-03-23">
<author pid="19/6993">Huimin Wu</author>
<author pid="47/2670">Qin Jin</author>
<title>Improving emotion classification on Chinese microblog texts with auxiliary cross-domain data.</title>
<pages>821-826</pages>
<year>2015</year>
<booktitle>ACII</booktitle>
<ee>https://doi.org/10.1109/ACII.2015.7344668</ee>
<ee>https://doi.ieeecomputersociety.org/10.1109/ACII.2015.7344668</ee>
<crossref>conf/acii/2015</crossref>
<url>db/conf/acii/acii2015.html#WuJ15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/clef/LiJLLHHLXLX15" mdate="2023-03-10">
<author pid="58/5856">Xirong Li 0001</author>
<author pid="47/2670">Qin Jin</author>
<author pid="138/4239">Shuai Liao</author>
<author pid="62/10704-1">Junwei Liang 0001</author>
<author pid="150/5274">Xixi He</author>
<author pid="167/4776">Yujia Huo</author>
<author pid="166/2924">Weiyu Lan</author>
<author pid="43/5134">Bin Xiao</author>
<author pid="40/7053">Yanxiong Lu</author>
<author pid="96/979">Jieping Xu</author>
<title>RUC-Tencent at ImageCLEF 2015: Concept Detection, Localization and Sentence Generation.</title>
<year>2015</year>
<booktitle>CLEF (Working Notes)</booktitle>
<ee type="oa">https://ceur-ws.org/Vol-1391/38-CR.pdf</ee>
<crossref>conf/clef/2015w</crossref>
<url>db/conf/clef/clef2015w.html#LiJLLHHLXLX15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/LiangJHYXL15" mdate="2023-10-21">
<author orcid="0000-0003-2219-5569" pid="62/10704-1">Junwei Liang 0001</author>
<author pid="47/2670">Qin Jin</author>
<author pid="150/5274">Xixi He</author>
<author pid="36/4658-1">Gang Yang 0001</author>
<author orcid="0000-0001-6228-0317" pid="96/979">Jieping Xu</author>
<author orcid="0000-0002-0220-8310" pid="58/5856">Xirong Li 0001</author>
<title>Detecting semantic concepts in consumer videos using audio.</title>
<pages>2279-2283</pages>
<year>2015</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2015.7178377</ee>
<ee>https://www.wikidata.org/entity/Q121878611</ee>
<crossref>conf/icassp/2015</crossref>
<url>db/conf/icassp/icassp2015.html#LiangJHYXL15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/JinLCW15" mdate="2017-05-19">
<author pid="47/2670">Qin Jin</author>
<author pid="166/6518">Chengxin Li</author>
<author pid="153/0734">Shizhe Chen</author>
<author pid="19/6993">Huimin Wu</author>
<title>Speech emotion recognition with acoustic and lexical features.</title>
<pages>4749-4753</pages>
<year>2015</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2015.7178872</ee>
<crossref>conf/icassp/2015</crossref>
<url>db/conf/icassp/icassp2015.html#JinLCW15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mediaeval/JinLCHLYX15" mdate="2023-03-10">
<author pid="47/2670">Qin Jin</author>
<author pid="58/5856">Xirong Li 0001</author>
<author pid="169/0393">Haibing Cao</author>
<author pid="167/4776">Yujia Huo</author>
<author pid="138/4239">Shuai Liao</author>
<author pid="36/4658-1">Gang Yang 0001</author>
<author pid="96/979">Jieping Xu</author>
<title>RUCMM at MediaEval 2015 Affective Impact of Movies Task: Fusion of Audio and Visual Cues.</title>
<year>2015</year>
<booktitle>MediaEval</booktitle>
<crossref>conf/mediaeval/2015</crossref>
<url>db/conf/mediaeval/mediaeval2015.html#JinLCHLYX15</url>
<ee type="oa">https://ceur-ws.org/Vol-1436/Paper26.pdf</ee>
</inproceedings>
</r>
<r><inproceedings key="conf/mir/JinLHYXL15" mdate="2024-10-06">
<author pid="47/2670">Qin Jin</author>
<author orcid="0000-0003-2219-5569" pid="62/10704-1">Junwei Liang 0001</author>
<author pid="150/5274">Xixi He</author>
<author pid="36/4658-1">Gang Yang 0001</author>
<author orcid="0000-0001-6228-0317" pid="96/979">Jieping Xu</author>
<author orcid="0000-0002-0220-8310" pid="58/5856">Xirong Li 0001</author>
<title>Semantic Concept Annotation For User Generated Videos Using Soundtracks.</title>
<pages>599-602</pages>
<year>2015</year>
<booktitle>ICMR</booktitle>
<ee>https://doi.org/10.1145/2671188.2749388</ee>
<ee>https://www.wikidata.org/entity/Q121878685</ee>
<crossref>conf/mir/2015</crossref>
<url>db/conf/mir/icmr2015.html#JinLHYXL15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/ChenJ15" mdate="2018-11-06">
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<title>Multi-modal Dimensional Emotion Recognition using Recurrent Neural Networks.</title>
<pages>49-56</pages>
<year>2015</year>
<booktitle>AVEC@ACM Multimedia</booktitle>
<ee>https://doi.org/10.1145/2808196.2811638</ee>
<crossref>conf/mm/2015avec</crossref>
<url>db/conf/mm/avec2015.html#ChenJ15</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/ChenJYH15" mdate="2025-01-19">
<author pid="99/6879-1">Jia Chen 0001</author>
<author pid="47/2670">Qin Jin</author>
<author pid="43/5685-1">Yong Yu 0001</author>
<author pid="h/AlexanderGHauptmann">Alexander G. Hauptmann</author>
<title>Image Profiling for History Events on the Fly.</title>
<pages>291-300</pages>
<year>2015</year>
<booktitle>ACM Multimedia</booktitle>
<ee>https://doi.org/10.1145/2733373.2806242</ee>
<ee>https://www.wikidata.org/entity/Q130894096</ee>
<crossref>conf/mm/2015</crossref>
<url>db/conf/mm/mm2015.html#ChenJYH15</url>
</inproceedings>
</r>
<r><article key="journals/fgcs/YenJHZ14" mdate="2026-08-03">
<author pid="20/6973">Neil Y. Yen</author>
<author pid="47/2670">Qin Jin</author>
<author orcid="0000-0002-2440-2771" pid="35/3725">Ching-Hsien Hsu</author>
<author pid="51/1660">Qiangfu Zhao</author>
<title>Special Issue on &#34;Hybrid intelligence for growing internet and its applications&#34;.</title>
<pages>401-403</pages>
<year>2014</year>
<volume>37</volume>
<journal>Future Gener. Comput. Syst.</journal>
<ee>https://doi.org/10.1016/j.future.2014.04.001</ee>
<url>db/journals/fgcs/fgcs37.html#YenJHZ14</url>
</article>
</r>
<r><inproceedings key="conf/apsipa/ZhengJLWB14" mdate="2021-10-14">
<author pid="53/2843">Thomas Fang Zheng</author>
<author pid="47/2670">Qin Jin</author>
<author orcid="0000-0003-4274-7930" pid="153/0735">Lantian Li</author>
<author pid="125/8189-73">Jun Wang 0073</author>
<author pid="153/0709">Fanhu Bie</author>
<title>An overview of robustness related issues in speaker recognition.</title>
<pages>1-10</pages>
<year>2014</year>
<booktitle>APSIPA</booktitle>
<ee>https://doi.org/10.1109/APSIPA.2014.7041826</ee>
<crossref>conf/apsipa/2014</crossref>
<url>db/conf/apsipa/apsipa2014.html#ZhengJLWB14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/clef/LiHYJX14" mdate="2023-03-10">
<author pid="58/5856">Xirong Li 0001</author>
<author pid="150/5274">Xixi He</author>
<author pid="36/4658-1">Gang Yang 0001</author>
<author pid="47/2670">Qin Jin</author>
<author pid="96/979">Jieping Xu</author>
<title>Renmin University of China at ImageCLEF 2014 Scalable Concept Image Annotation.</title>
<pages>380-385</pages>
<year>2014</year>
<booktitle>CLEF (Working Notes)</booktitle>
<ee type="oa">https://ceur-ws.org/Vol-1180/CLEF2014wn-Image-LiEt2014.pdf</ee>
<crossref>conf/clef/2014w</crossref>
<url>db/conf/clef/clef2014w.html#LiHYJX14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icann/YangLXJ14" mdate="2021-09-13">
<author pid="36/4658-1">Gang Yang 0001</author>
<author orcid="0000-0002-0220-8310" pid="58/5856">Xirong Li 0001</author>
<author pid="96/979">Jieping Xu</author>
<author pid="47/2670">Qin Jin</author>
<title>Structure Perturbation Optimization for Hopfield-Type Neural Networks.</title>
<pages>307-314</pages>
<year>2014</year>
<booktitle>ICANN</booktitle>
<ee>https://doi.org/10.1007/978-3-319-11179-7_39</ee>
<crossref>conf/icann/2014</crossref>
<url>db/conf/icann/icann2014.html#YangLXJ14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/iscslp/ChenJLYX14" mdate="2021-09-13">
<author pid="153/0734">Shizhe Chen</author>
<author pid="47/2670">Qin Jin</author>
<author orcid="0000-0002-0220-8310" pid="58/5856">Xirong Li 0001</author>
<author pid="36/4658-1">Gang Yang 0001</author>
<author orcid="0000-0001-6228-0317" pid="96/979">Jieping Xu</author>
<title>Speech emotion classification using acoustic features.</title>
<pages>579-583</pages>
<year>2014</year>
<booktitle>ISCSLP</booktitle>
<ee>https://doi.org/10.1109/ISCSLP.2014.6936664</ee>
<crossref>conf/iscslp/2014</crossref>
<url>db/conf/iscslp/iscslp2014.html#ChenJLYX14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/nlpcc/LiWJ14" mdate="2017-05-21">
<author pid="166/6518">Chengxin Li</author>
<author pid="19/6993">Huimin Wu</author>
<author pid="47/2670">Qin Jin</author>
<title>Emotion Classification of Chinese Microblog Text via Fusion of BoW and eVector Feature Representations.</title>
<pages>217-228</pages>
<year>2014</year>
<booktitle>NLPCC</booktitle>
<ee>https://doi.org/10.1007/978-3-662-45924-9_20</ee>
<crossref>conf/nlpcc/2014</crossref>
<url>db/conf/nlpcc/nlpcc2014.html#LiWJ14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/pcm/HeLYXJ14" mdate="2025-03-03">
<author pid="150/5274">Xixi He</author>
<author orcid="0000-0002-0220-8310" pid="58/5856">Xirong Li 0001</author>
<author pid="36/4658-1">Gang Yang 0001</author>
<author pid="96/979">Jieping Xu</author>
<author pid="47/2670">Qin Jin</author>
<title>Adaptive Tag Selection for Image Annotation.</title>
<pages>11-21</pages>
<year>2014</year>
<booktitle>PCM</booktitle>
<ee>https://doi.org/10.1007/978-3-319-13168-9_2</ee>
<crossref>conf/pcm/2014</crossref>
<url>db/conf/pcm/pcm2014.html#HeLYXJ14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/pcm/LiangJHYXL14" mdate="2025-10-14">
<author orcid="0000-0003-2219-5569" pid="62/10704-1">Junwei Liang 0001</author>
<author pid="47/2670">Qin Jin</author>
<author pid="150/5274">Xixi He</author>
<author pid="36/4658-1">Gang Yang 0001</author>
<author pid="96/979">Jieping Xu</author>
<author orcid="0000-0002-0220-8310" pid="58/5856">Xirong Li 0001</author>
<title>Semantic Concept Annotation of Consumer Videos at Frame-Level Using Audio.</title>
<pages>113-122</pages>
<year>2014</year>
<booktitle>PCM</booktitle>
<ee>https://doi.org/10.1007/978-3-319-13168-9_12</ee>
<ee>https://www.wikidata.org/entity/Q121878688</ee>
<crossref>conf/pcm/2014</crossref>
<url>db/conf/pcm/pcm2014.html#LiangJHYXL14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/sigir/ChenJZBZSY14" mdate="2025-08-05">
<author pid="99/6879-1">Jia Chen 0001</author>
<author pid="47/2670">Qin Jin</author>
<author orcid="0000-0001-5068-025X" pid="73/1947">Shiwan Zhao</author>
<author pid="31/529">Shenghua Bao</author>
<author pid="89/5992-7">Li Zhang 0007</author>
<author pid="87/1363">Zhong Su</author>
<author pid="43/5685-1">Yong Yu 0001</author>
<title>Does product recommendation meet its waterloo in unexplored categories?: no, price comes to help.</title>
<pages>667-676</pages>
<year>2014</year>
<booktitle>SIGIR</booktitle>
<ee>https://doi.org/10.1145/2600428.2609608</ee>
<crossref>conf/sigir/2014</crossref>
<url>db/conf/sigir/sigir2014.html#ChenJZBZSY14</url>
</inproceedings>
</r>
<r><inproceedings key="conf/smc/YangLXJS14" mdate="2021-09-13">
<author pid="36/4658-1">Gang Yang 0001</author>
<author orcid="0000-0002-0220-8310" pid="58/5856">Xirong Li 0001</author>
<author orcid="0000-0001-6228-0317" pid="96/979">Jieping Xu</author>
<author pid="47/2670">Qin Jin</author>
<author pid="31/3925">Hui Sun</author>
<title>A guided Hopfield evolutionary algorithm with local search for maximum clique problem.</title>
<pages>979-982</pages>
<year>2014</year>
<booktitle>SMC</booktitle>
<ee>https://doi.org/10.1109/SMC.2014.6974039</ee>
<crossref>conf/smc/2014</crossref>
<url>db/conf/smc/smc2014.html#YangLXJS14</url>
</inproceedings>
</r>
<r><article publtype="informal" key="journals/corr/HeLYXJ14" mdate="2021-09-13">
<author pid="150/5274">Xixi He</author>
<author pid="58/5856">Xirong Li 0001</author>
<author pid="36/4658-1">Gang Yang 0001</author>
<author pid="96/979">Jieping Xu</author>
<author pid="47/2670">Qin Jin</author>
<title>Adaptive Tag Selection for Image Annotation.</title>
<year>2014</year>
<volume>abs/1409.4995</volume>
<journal>CoRR</journal>
<ee type="oa">http://arxiv.org/abs/1409.4995</ee>
<url>db/journals/corr/corr1409.html#HeLYXJ14</url>
</article>
</r>
<r><inproceedings key="conf/clef/LiLLYJXD13" mdate="2023-03-10">
<author pid="58/5856">Xirong Li 0001</author>
<author pid="138/4239">Shuai Liao</author>
<author pid="22/7702">Binbin Liu</author>
<author pid="36/4658-1">Gang Yang 0001</author>
<author pid="47/2670">Qin Jin</author>
<author pid="96/979">Jieping Xu</author>
<author pid="47/3542-1">Xiaoyong Du 0001</author>
<title>Renmin University of China at ImageCLEF 2013 Scalable Concept Image Annotation.</title>
<year>2013</year>
<booktitle>CLEF (Working Notes)</booktitle>
<ee type="oa">https://ceur-ws.org/Vol-1179/CLEF2013wn-ImageCLEF-LiEt2013.pdf</ee>
<crossref>conf/clef/2013w</crossref>
<url>db/conf/clef/clef2013w.html#LiLLYJXD13</url>
</inproceedings>
</r>
<r><inproceedings key="conf/mm/ChenJZBSY13" mdate="2020-10-08">
<author pid="99/6879-1">Jia Chen 0001</author>
<author pid="47/2670">Qin Jin</author>
<author pid="29/5921">Weipeng Zhang</author>
<author pid="31/529">Shenghua Bao</author>
<author pid="87/1363">Zhong Su</author>
<author pid="43/5685-1">Yong Yu 0001</author>
<title>Tell me what happened here in history.</title>
<pages>467-468</pages>
<year>2013</year>
<booktitle>ACM Multimedia</booktitle>
<ee>https://doi.org/10.1145/2502081.2502272</ee>
<crossref>conf/mm/2013</crossref>
<url>db/conf/mm/mm2013.html#ChenJZBSY13</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/JinSRBDM12" mdate="2023-06-23">
<author pid="47/2670">Qin Jin</author>
<author pid="117/4105">Peter Franz Schulam</author>
<author pid="117/4119">Shourabh Rawat</author>
<author pid="10/4034">Susanne Burger</author>
<author pid="36/3063">Duo Ding</author>
<author pid="26/1652">Florian Metze</author>
<title>Event-based Video Retrieval Using Audio.</title>
<year>2012</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2012-556</ee>
<crossref>conf/interspeech/2012</crossref>
<url>db/conf/interspeech/interspeech2012.html#JinSRBDM12</url>
<pages>2085-2088</pages>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/LaskowskiJ11" mdate="2023-06-23">
<author pid="98/2383">Kornel Laskowski</author>
<author pid="47/2670">Qin Jin</author>
<title>Harmonic Structure Transform for Speaker Recognition.</title>
<pages>365-368</pages>
<year>2011</year>
<booktitle>INTERSPEECH</booktitle>
<crossref>conf/interspeech/2011</crossref>
<url>db/conf/interspeech/interspeech2011.html#LaskowskiJ11</url>
<ee type="oa">https://doi.org/10.21437/Interspeech.2011-131</ee>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/NallasamyGMJSS11" mdate="2023-06-23">
<author pid="76/7551">Udhyakumar Nallasamy</author>
<author pid="42/10648">Michael Garbus</author>
<author pid="26/1652">Florian Metze</author>
<author pid="47/2670">Qin Jin</author>
<author pid="75/5600">Thomas Schaaf</author>
<author pid="s/TanjaSchultz">Tanja Schultz</author>
<title>Analysis of Dialectal Influence in Pan-Arabic ASR.</title>
<pages>1721-1724</pages>
<year>2011</year>
<booktitle>INTERSPEECH</booktitle>
<crossref>conf/interspeech/2011</crossref>
<url>db/conf/interspeech/interspeech2011.html#NallasamyGMJSS11</url>
<ee type="oa">https://doi.org/10.21437/Interspeech.2011-191</ee>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/YangJS11" mdate="2023-06-23">
<author pid="15/3199">Qian Yang</author>
<author pid="47/2670">Qin Jin</author>
<author pid="s/TanjaSchultz">Tanja Schultz</author>
<title>Investigation of Cross-Show Speaker Diarization.</title>
<pages>2925-2928</pages>
<year>2011</year>
<booktitle>INTERSPEECH</booktitle>
<crossref>conf/interspeech/2011</crossref>
<url>db/conf/interspeech/interspeech2011.html#YangJS11</url>
<ee type="oa">https://doi.org/10.21437/Interspeech.2011-732</ee>
</inproceedings>
</r>
<r><inproceedings key="conf/trecvid/BaoZYL0OJTLLGBM11" mdate="2024-08-26">
<author pid="88/1530">Lei Bao</author>
<author pid="30/7663">Longfei Zhang</author>
<author pid="23/7442">Shoou-I Yu</author>
<author pid="27/3780">Zhen-zhong Lan</author>
<author pid="22/752-4">Lu Jiang 0004</author>
<author pid="16/7404">Arnold Overwijk</author>
<author pid="47/2670">Qin Jin</author>
<author pid="181/8511">Shohei Takahashi</author>
<author pid="25/8160">Brian Langner</author>
<author pid="08/8174-1">Yuanpeng Li 0001</author>
<author pid="42/10648">Michael Garbus</author>
<author pid="10/4034">Susanne Burger</author>
<author pid="26/1652">Florian Metze</author>
<author pid="h/AlexanderGHauptmann">Alexander G. Hauptmann</author>
<title>Informedia@TRECVID 2011: Surveillance Event Detection.</title>
<year>2011</year>
<booktitle>TRECVID</booktitle>
<ee type="oa">https://www-nlpir.nist.gov/projects/tvpubs/tv11.papers/cmu.pdf</ee>
<crossref>conf/trecvid/2011</crossref>
<url>db/conf/trecvid/trecvid2011.html#BaoZYL0OJTLLGBM11</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/JinLYLS10" mdate="2021-07-31">
<author pid="47/2670">Qin Jin</author>
<author pid="26/8160">Runxin Li</author>
<author pid="15/3199">Qian Yang</author>
<author pid="98/2383">Kornel Laskowski</author>
<author orcid="0000-0002-9809-7028" pid="s/TanjaSchultz">Tanja Schultz</author>
<title>Speaker identification with distant microphone speech.</title>
<pages>4518-4521</pages>
<year>2010</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2010.5495590</ee>
<ee>https://www.wikidata.org/entity/Q107470626</ee>
<crossref>conf/icassp/2010</crossref>
<url>db/conf/icassp/icassp2010.html#JinLYLS10</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/MetzeHJNS10" mdate="2023-06-23">
<author pid="26/1652">Florian Metze</author>
<author pid="86/8054">Roger Hsiao</author>
<author pid="47/2670">Qin Jin</author>
<author pid="76/7551">Udhyakumar Nallasamy</author>
<author pid="s/TanjaSchultz">Tanja Schultz</author>
<title>The 2010 CMU GALE speech-to-text system.</title>
<pages>1501-1504</pages>
<year>2010</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2010-439</ee>
<crossref>conf/interspeech/2010</crossref>
<url>db/conf/interspeech/interspeech2010.html#MetzeHJNS10</url>
</inproceedings>
</r>
<r><inproceedings key="conf/odyssey/LaskowskiJ10" mdate="2024-07-30">
<author pid="98/2383">Kornel Laskowski</author>
<author pid="47/2670">Qin Jin</author>
<title>Modeling Prosody for Speaker Recognition: Why Estimating Pitch May Be a Red Herring.</title>
<pages>4</pages>
<year>2010</year>
<booktitle>Odyssey</booktitle>
<ee type="oa">https://www.isca-archive.org/odyssey_2010/laskowski10_odyssey.html</ee>
<crossref>conf/odyssey/2010</crossref>
<url>db/conf/odyssey/odyssey2010.html#LaskowskiJ10</url>
</inproceedings>
</r>
<r><inproceedings key="conf/asru/JinTSB09" mdate="2019-12-27">
<author pid="47/2670">Qin Jin</author>
<author pid="99/7822">Arthur R. Toth</author>
<author orcid="0000-0002-9809-7028" pid="s/TanjaSchultz">Tanja Schultz</author>
<author pid="b/AlanWBlack">Alan W. Black</author>
<title>Speaker de-identification via voice transformation.</title>
<pages>529-533</pages>
<year>2009</year>
<booktitle>ASRU</booktitle>
<ee>https://doi.org/10.1109/ASRU.2009.5373356</ee>
<crossref>conf/asru/2009</crossref>
<url>db/conf/asru/asru2009.html#JinTSB09</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/JinTSB09" mdate="2023-03-23">
<author pid="47/2670">Qin Jin</author>
<author pid="99/7822">Arthur R. Toth</author>
<author orcid="0000-0002-9809-7028" pid="s/TanjaSchultz">Tanja Schultz</author>
<author pid="b/AlanWBlack">Alan W. Black</author>
<title>Voice convergin: Speaker de-identification by voice transformation.</title>
<pages>3909-3912</pages>
<year>2009</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2009.4960482</ee>
<ee>https://doi.ieeecomputersociety.org/10.1109/ICASSP.2009.4960482</ee>
<crossref>conf/icassp/2009</crossref>
<url>db/conf/icassp/icassp2009.html#JinTSB09</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/LiMLSZSYTKHPGLDNTEASSJ09" mdate="2025-02-12">
<author orcid="0000-0001-9158-9401" pid="36/4118">Haizhou Li 0001</author>
<author pid="70/6176-1">Bin Ma 0001</author>
<author orcid="0000-0001-9133-3000" pid="35/4621">Kong-Aik Lee</author>
<author pid="22/1460">Hanwu Sun</author>
<author pid="70/5222">Donglai Zhu</author>
<author pid="78/6873">Khe Chai Sim</author>
<author pid="00/3382">Changhuai You</author>
<author pid="09/1880">Rong Tong</author>
<author pid="35/3843">Ismo K&#228;rkk&#228;inen</author>
<author pid="42/3949">Chien-Lin Huang</author>
<author pid="98/6470">Vladimir Pervouchine</author>
<author pid="20/4480">Wu Guo</author>
<author pid="54/8054-1">Yijie Li 0001</author>
<author pid="48/6462-1">Li-Rong Dai 0001</author>
<author pid="31/3678">Mohaddeseh Nosratighods</author>
<author pid="37/8053">Tharmarajah Thiruvaran</author>
<author orcid="0000-0001-6624-5551" pid="92/6424">Julien Epps</author>
<author orcid="0000-0003-4673-6534" pid="24/3343">Eliathamby Ambikairajah</author>
<author orcid="0000-0001-6257-7399" pid="c/ChngEngSiong">Chng Eng Siong</author>
<author orcid="0000-0002-9809-7028" pid="s/TanjaSchultz">Tanja Schultz</author>
<author pid="47/2670">Qin Jin</author>
<title>The I4U system in NIST 2008 speaker recognition evaluation.</title>
<pages>4201-4204</pages>
<year>2009</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2009.4960555</ee>
<ee>https://doi.ieeecomputersociety.org/10.1109/ICASSP.2009.4960555</ee>
<crossref>conf/icassp/2009</crossref>
<url>db/conf/icassp/icassp2009.html#LiMLSZSYTKHPGLDNTEASSJ09</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/LaskowskiJ09" mdate="2023-03-23">
<author pid="98/2383">Kornel Laskowski</author>
<author pid="47/2670">Qin Jin</author>
<title>Modeling instantaneous intonation for speaker identification using the fundamental frequency variation spectrum.</title>
<pages>4541-4544</pages>
<year>2009</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2009.4960640</ee>
<ee>https://doi.ieeecomputersociety.org/10.1109/ICASSP.2009.4960640</ee>
<crossref>conf/icassp/2009</crossref>
<url>db/conf/icassp/icassp2009.html#LaskowskiJ09</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/FuhsJS09" mdate="2023-03-23">
<author pid="25/516">Mark C. Fuhs</author>
<author pid="47/2670">Qin Jin</author>
<author orcid="0000-0002-9809-7028" pid="s/TanjaSchultz">Tanja Schultz</author>
<title>Detecting bandlimited audio in broadcast television shows.</title>
<pages>4589-4592</pages>
<year>2009</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2009.4960652</ee>
<ee>https://doi.ieeecomputersociety.org/10.1109/ICASSP.2009.4960652</ee>
<ee>https://www.wikidata.org/entity/Q107470632</ee>
<crossref>conf/icassp/2009</crossref>
<url>db/conf/icassp/icassp2009.html#FuhsJS09</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/LiSJ09" mdate="2023-06-23">
<author pid="26/8160">Runxin Li</author>
<author pid="s/TanjaSchultz">Tanja Schultz</author>
<author pid="47/2670">Qin Jin</author>
<title>Improving speaker segmentation via speaker identification and text segmentation.</title>
<pages>904-907</pages>
<year>2009</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2009-272</ee>
<crossref>conf/interspeech/2009</crossref>
<url>db/conf/interspeech/interspeech2009.html#LiSJ09</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/WolfelYJS09" mdate="2023-06-23">
<author pid="84/6910">Matthias W&#246;lfel</author>
<author pid="15/3199">Qian Yang</author>
<author pid="47/2670">Qin Jin</author>
<author pid="s/TanjaSchultz">Tanja Schultz</author>
<title>Speaker identification using warped MVDR cepstral features.</title>
<pages>912-915</pages>
<year>2009</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2009-274</ee>
<crossref>conf/interspeech/2009</crossref>
<url>db/conf/interspeech/interspeech2009.html#WolfelYJS09</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/JinTBS08" mdate="2019-12-27">
<author pid="47/2670">Qin Jin</author>
<author pid="99/7822">Arthur R. Toth</author>
<author pid="b/AlanWBlack">Alan W. Black</author>
<author orcid="0000-0002-9809-7028" pid="s/TanjaSchultz">Tanja Schultz</author>
<title>Is voice transformation a threat to speaker identification?</title>
<pages>4845-4848</pages>
<year>2008</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2008.4518742</ee>
<crossref>conf/icassp/2008</crossref>
<url>db/conf/icassp/icassp2008.html#JinTBS08</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/HsiaoFTJS08" mdate="2023-06-23">
<author pid="86/8054">Roger Hsiao</author>
<author pid="25/516">Mark C. Fuhs</author>
<author pid="72/158">Yik-Cheung Tam</author>
<author pid="47/2670">Qin Jin</author>
<author pid="s/TanjaSchultz">Tanja Schultz</author>
<title>The CMU-interACT 2008 Mandarin transcription system.</title>
<pages>1445-1448</pages>
<year>2008</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2008-417</ee>
<crossref>conf/interspeech/2008</crossref>
<url>db/conf/interspeech/interspeech2008.html#HsiaoFTJS08</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/JinS08" mdate="2023-06-23">
<author pid="47/2670">Qin Jin</author>
<author pid="s/TanjaSchultz">Tanja Schultz</author>
<title>Robust far-field speaker identification under mismatched conditions.</title>
<pages>1893-1896</pages>
<year>2008</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2008-501</ee>
<crossref>conf/interspeech/2008</crossref>
<url>db/conf/interspeech/interspeech2008.html#JinS08</url>
</inproceedings>
</r>
<r><article key="journals/taslp/JinSW07" mdate="2020-05-17">
<author pid="47/2670">Qin Jin</author>
<author orcid="0000-0002-9809-7028" pid="s/TanjaSchultz">Tanja Schultz</author>
<author pid="08/2456">Alex Waibel</author>
<title>Far-Field Speaker Recognition.</title>
<pages>2023-2032</pages>
<year>2007</year>
<volume>15</volume>
<journal>IEEE Trans. Speech Audio Process.</journal>
<number>7</number>
<ee>https://doi.org/10.1109/TASL.2007.902876</ee>
<url>db/journals/taslp/taslp15.html#JinSW07</url>
</article>
</r>
<r><inproceedings key="conf/clear/EkenelJFS07" mdate="2021-10-14">
<author orcid="0000-0003-3697-8548" pid="26/1396">Hazim Kemal Ekenel</author>
<author pid="47/2670">Qin Jin</author>
<author pid="23/6655">Mika Fischer</author>
<author pid="31/4699">Rainer Stiefelhagen</author>
<title>ISL Person Identification Systems in the CLEAR 2007 Evaluations.</title>
<pages>256-265</pages>
<year>2007</year>
<booktitle>CLEAR</booktitle>
<ee>https://doi.org/10.1007/978-3-540-68585-2_24</ee>
<crossref>conf/clear/2007</crossref>
<url>db/conf/clear/clear2007.html#EkenelJFS07</url>
</inproceedings>
</r>
<r><inproceedings key="conf/cvpr/EkenelFJS07" mdate="2023-03-24">
<author orcid="0000-0003-3697-8548" pid="26/1396">Hazim Kemal Ekenel</author>
<author pid="23/6655">Mika Fischer</author>
<author pid="47/2670">Qin Jin</author>
<author pid="31/4699">Rainer Stiefelhagen</author>
<title>Multi-modal Person Identification in a Smart Environment.</title>
<year>2007</year>
<crossref>conf/cvpr/2007</crossref>
<booktitle>CVPR</booktitle>
<ee>https://doi.org/10.1109/CVPR.2007.383388</ee>
<ee>https://doi.ieeecomputersociety.org/10.1109/CVPR.2007.383388</ee>
<url>db/conf/cvpr/cvpr2007.html#EkenelFJS07</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icmcs/JinJS07" mdate="2023-03-24">
<author pid="47/2670">Qin Jin</author>
<author pid="17/1054">Szu-Chen Stan Jou</author>
<author pid="s/TanjaSchultz">Tanja Schultz</author>
<title>Whispering Speaker Identification.</title>
<pages>1027-1030</pages>
<year>2007</year>
<booktitle>ICME</booktitle>
<ee>https://doi.org/10.1109/ICME.2007.4284828</ee>
<ee>https://doi.ieeecomputersociety.org/10.1109/ICME.2007.4284828</ee>
<crossref>conf/icmcs/2007</crossref>
<url>db/conf/icmcs/icme2007.html#JinJS07</url>
</inproceedings>
</r>
<r><inproceedings key="conf/clear/EkenelJ06" mdate="2021-10-14">
<author orcid="0000-0003-3697-8548" pid="26/1396">Hazim Kemal Ekenel</author>
<author pid="47/2670">Qin Jin</author>
<title>ISL Person Identification Systems in the CLEAR Evaluations.</title>
<pages>249-257</pages>
<year>2006</year>
<crossref>conf/clear/2006</crossref>
<booktitle>CLEAR</booktitle>
<ee>https://doi.org/10.1007/978-3-540-69568-4_22</ee>
<url>db/conf/clear/clear2006.html#EkenelJ06</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/JinPS06" mdate="2020-06-22">
<author pid="47/2670">Qin Jin</author>
<author pid="91/3754">Yue Pan</author>
<author pid="s/TanjaSchultz">Tanja Schultz</author>
<title>Far-Field Speaker Recognition.</title>
<pages>937-940</pages>
<year>2006</year>
<booktitle>ICASSP (1)</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2006.1660176</ee>
<crossref>conf/icassp/2006</crossref>
<url>db/conf/icassp/icassp2006.html#JinPS06</url>
</inproceedings>
</r>
<r><inproceedings key="conf/trecvid/HauptmannBCCGJL05" mdate="2024-09-19">
<author pid="h/AlexanderGHauptmann">Alexander G. Hauptmann</author>
<author pid="14/5555">Robert V. Baron</author>
<author pid="c/MGChristel">Michael G. Christel</author>
<author pid="263/9828">R. Concescu</author>
<author pid="79/1247">Jiang Gao</author>
<author pid="47/2670">Qin Jin</author>
<author pid="295/0536-3">Wei-Hao Lin 0003</author>
<author pid="263/9885">J.-Y. Pan</author>
<author pid="92/5829">Scott M. Stevens</author>
<author pid="51/1293">Rong Yan</author>
<author pid="y/JunYang3">Jun Yang 0003</author>
<author pid="52/5040">Y. Zhang</author>
<title>CMU Informedia's TRECVID 2005 Skirmishes.</title>
<year>2005</year>
<booktitle>TRECVID</booktitle>
<ee type="oa">https://www-nlpir.nist.gov/projects/tvpubs/tv5.papers/cmu.pdf</ee>
<crossref>conf/trecvid/2005</crossref>
<url>db/conf/trecvid/trecvid2005.html#HauptmannBCCGJL05</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/SoltauYMFJJ04" mdate="2022-08-06">
<author pid="07/2072">Hagen Soltau</author>
<author pid="02/2407-8">Hua Yu 0008</author>
<author pid="26/1652">Florian Metze</author>
<author pid="88/6411">Christian F&#252;gen</author>
<author pid="47/2670">Qin Jin</author>
<author pid="17/1054">Szu-Chen Stan Jou</author>
<title>The 2003 ISL rich transcription system for conversational telephony speech.</title>
<pages>773-776</pages>
<year>2004</year>
<booktitle>ICASSP (1)</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2004.1326100</ee>
<crossref>conf/icassp/2004</crossref>
<url>db/conf/icassp/icassp2004.html#SoltauYMFJJ04</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/JinS04" mdate="2023-06-22">
<author pid="47/2670">Qin Jin</author>
<author pid="s/TanjaSchultz">Tanja Schultz</author>
<title>Speaker segmentation and clustering in meetings.</title>
<year>2004</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2004-249</ee>
<crossref>conf/interspeech/2004</crossref>
<url>db/conf/interspeech/interspeech2004.html#JinS04</url>
<pages>597-600</pages>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/LaskowskiJS04" mdate="2023-06-22">
<author pid="98/2383">Kornel Laskowski</author>
<author pid="47/2670">Qin Jin</author>
<author pid="s/TanjaSchultz">Tanja Schultz</author>
<title>Crosscorrelation-based multispeaker speech activity detection.</title>
<year>2004</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2004-350</ee>
<crossref>conf/interspeech/2004</crossref>
<url>db/conf/interspeech/interspeech2004.html#LaskowskiJS04</url>
<pages>973-976</pages>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/SchultzJLPMF04" mdate="2023-06-22">
<author pid="s/TanjaSchultz">Tanja Schultz</author>
<author pid="47/2670">Qin Jin</author>
<author pid="98/2383">Kornel Laskowski</author>
<author pid="91/3754">Yue Pan</author>
<author pid="26/1652">Florian Metze</author>
<author pid="88/6411">Christian F&#252;gen</author>
<title>Issues in meeting transcription - the ISL meeting transcription system.</title>
<year>2004</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/Interspeech.2004-186</ee>
<crossref>conf/interspeech/2004</crossref>
<url>db/conf/interspeech/interspeech2004.html#SchultzJLPMF04</url>
<pages>1709-1712</pages>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/ReynoldsACNPAJKAMGJX03" mdate="2022-12-07">
<author pid="21/2250">Douglas A. Reynolds</author>
<author pid="62/9243">Walter D. Andrews</author>
<author pid="39/2844">Joseph P. Campbell</author>
<author pid="00/680-1">Jir&#237; Navr&#225;til 0001</author>
<author pid="15/1225">Barbara Peskin</author>
<author orcid="0000-0001-8105-7698" pid="59/974">Andr&#233; Adami</author>
<author pid="47/2670">Qin Jin</author>
<author pid="142/6673">David Klus&#225;cek</author>
<author pid="145/0483">Joy S. Abramson</author>
<author pid="80/5017">Radu Mihaescu</author>
<author pid="07/172">John J. Godfrey</author>
<author pid="32/2452">Douglas A. Jones</author>
<author pid="82/5456">Bing Xiang</author>
<title>The SuperSID project: exploiting high-level information for high-accuracy speaker recognition.</title>
<pages>784-787</pages>
<year>2003</year>
<booktitle>ICASSP (4)</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2003.1202760</ee>
<crossref>conf/icassp/2003</crossref>
<url>db/conf/icassp/icassp2003.html#ReynoldsACNPAJKAMGJX03</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/NavratilJAC03" mdate="2020-06-22">
<author pid="00/680-1">Jir&#237; Navr&#225;til 0001</author>
<author pid="47/2670">Qin Jin</author>
<author pid="62/9243">Walter D. Andrews</author>
<author pid="39/2844">Joseph P. Campbell</author>
<title>Phonetic speaker recognition using maximum-likelihood binary-decision tree models.</title>
<pages>796-799</pages>
<year>2003</year>
<booktitle>ICASSP (4)</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2003.1202763</ee>
<crossref>conf/icassp/2003</crossref>
<url>db/conf/icassp/icassp2003.html#NavratilJAC03</url>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/JinNRCAA03" mdate="2020-06-22">
<author pid="47/2670">Qin Jin</author>
<author pid="00/680-1">Jir&#237; Navr&#225;til 0001</author>
<author pid="21/2250">Douglas A. Reynolds</author>
<author pid="39/2844">Joseph P. Campbell</author>
<author pid="62/9243">Walter D. Andrews</author>
<author pid="145/0483">Joy S. Abramson</author>
<title>Combining cross-stream and time dimensions in phonetic speaker recognition.</title>
<pages>800-803</pages>
<year>2003</year>
<booktitle>ICASSP (4)</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2003.1202764</ee>
<crossref>conf/icassp/2003</crossref>
<url>db/conf/icassp/icassp2003.html#JinNRCAA03</url>
</inproceedings>
</r>
<r><inproceedings key="conf/acl/SchultzJLTW02" mdate="2021-08-06">
<author pid="s/TanjaSchultz">Tanja Schultz</author>
<author pid="47/2670">Qin Jin</author>
<author pid="98/2383">Kornel Laskowski</author>
<author pid="60/5193">Alicia Tribble</author>
<author pid="08/2456">Alex Waibel</author>
<title>Improvements in Non-Verbal Cue Identification Using Multilingual Phone Strings.</title>
<year>2002</year>
<booktitle>Speech-to-Speech Translation@ACL</booktitle>
<ee type="oa">https://aclanthology.org/W02-0714/</ee>
<ee>https://doi.org/10.3115/1118656.1118670</ee>
<crossref>conf/acl/2002stst</crossref>
<url>db/conf/acl/stst2002.html#SchultzJLTW02</url>
<pages>101-78</pages>
</inproceedings>
</r>
<r><inproceedings key="conf/icassp/JinSW02" mdate="2017-05-19">
<author pid="47/2670">Qin Jin</author>
<author pid="s/TanjaSchultz">Tanja Schultz</author>
<author pid="08/2456">Alex Waibel</author>
<title>Speaker identification using multilingual phone strings.</title>
<pages>145-148</pages>
<year>2002</year>
<booktitle>ICASSP</booktitle>
<ee>https://doi.org/10.1109/ICASSP.2002.5743675</ee>
<crossref>conf/icassp/2002</crossref>
<url>db/conf/icassp/icassp2002.html#JinSW02</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/JinSW02" mdate="2023-06-22">
<author pid="47/2670">Qin Jin</author>
<author pid="s/TanjaSchultz">Tanja Schultz</author>
<author pid="08/2456">Alex Waibel</author>
<title>Phonetic speaker identification.</title>
<year>2002</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/ICSLP.2002-409</ee>
<crossref>conf/interspeech/2002</crossref>
<url>db/conf/interspeech/interspeech2002.html#JinSW02</url>
<pages>1345-1348</pages>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/JinW00" mdate="2023-06-22">
<author pid="47/2670">Qin Jin</author>
<author pid="08/2456">Alex Waibel</author>
<title>Application of LDA to speaker recognition.</title>
<pages>250-253</pages>
<year>2000</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/ICSLP.2000-256</ee>
<crossref>conf/interspeech/2000</crossref>
<url>db/conf/interspeech/interspeech2000.html#JinW00</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/JinW00a" mdate="2023-06-22">
<author pid="47/2670">Qin Jin</author>
<author pid="08/2456">Alex Waibel</author>
<title>A na ve de-lambing method for speaker identification.</title>
<pages>466-469</pages>
<year>2000</year>
<booktitle>INTERSPEECH</booktitle>
<ee type="oa">https://doi.org/10.21437/ICSLP.2000-308</ee>
<crossref>conf/interspeech/2000</crossref>
<url>db/conf/interspeech/interspeech2000.html#JinW00a</url>
</inproceedings>
</r>
<r><inproceedings key="conf/interspeech/JinSH98" mdate="2023-06-22">
<author pid="47/2670">Qin Jin</author>
<author pid="14/1217">Luo Si</author>
<author pid="50/4223">Qixiu Hu</author>
<title>A high-performance text-independent speaker identification system based on BCDM.</title>
<year>1998</year>
<booktitle>ICSLP</booktitle>
<ee type="oa">https://doi.org/10.21437/ICSLP.1998-213</ee>
<crossref>conf/interspeech/1998</crossref>
<url>db/conf/interspeech/icslp1998.html#JinSH98</url>
</inproceedings>
</r>
<coauthors n="484" nc="2">
<co c="0"><na f="a/Abramson:Joy_S=" pid="145/0483">Joy S. Abramson</na></co>
<co c="0"><na f="a/Adami:Andr=eacute=" pid="59/974">Andr&#233; Adami</na></co>
<co c="0"><na f="a/Adi:Yossi" pid="171/0957">Yossi Adi</na></co>
<co c="0"><na f="a/Alahari:Karteek" pid="a/KarteekAlahari">Karteek Alahari</na></co>
<co c="0"><na f="a/Alameda=Pineda:Xavier" pid="22/10486">Xavier Alameda-Pineda</na></co>
<co c="0"><na f="a/Albanie:Samuel" pid="188/5765">Samuel Albanie</na></co>
<co c="0"><na f="a/Alharthi:Dareen" pid="358/5007">Dareen Alharthi</na></co>
<co c="0"><na f="a/Ambikairajah:Eliathamby" pid="24/3343">Eliathamby Ambikairajah</na></co>
<co c="0"><na f="a/Andrews:Walter_D=" pid="62/9243">Walter D. Andrews</na></co>
<co c="0"><na f="b/Baali:Massa" pid="253/3004">Massa Baali</na></co>
<co c="0"><na f="b/Bai:Xinyi" pid="221/6115">Xinyi Bai</na></co>
<co c="0"><na f="b/Bao:Lei" pid="88/1530">Lei Bao</na></co>
<co c="0"><na f="b/Bao:Shenghua" pid="31/529">Shenghua Bao</na></co>
<co c="1"><na f="b/Bao:Tengfei" pid="38/8548">Tengfei Bao</na></co>
<co c="0"><na f="b/Baron:Robert_V=" pid="14/5555">Robert V. Baron</na></co>
<co c="0"><na f="b/Barrios:Wayner" pid="208/4288">Wayner Barrios</na></co>
<co c="0"><na f="b/Bie:Fanhu" pid="153/0709">Fanhu Bie</na></co>
<co c="0"><na f="b/Bimbo:Alberto_Del" pid="b/AlbertoDelBimbo">Alberto Del Bimbo</na></co>
<co c="0"><na f="b/Black:Alan_W=" pid="b/AlanWBlack">Alan W. Black</na></co>
<co c="0"><na f="b/Burger:Susanne" pid="10/4034">Susanne Burger</na></co>
<co c="0"><na f="c/Cai:Chenglin" pid="238/3506">Chenglin Cai</na></co>
<co c="0"><na f="c/Campbell:Joseph_P=" pid="39/2844">Joseph P. Campbell</na></co>
<co c="0"><na f="c/Cao:Bin" pid="17/1169">Bin Cao</na></co>
<co c="0"><na f="c/Cao:Haibing" pid="169/0393">Haibing Cao</na></co>
<co c="0"><na f="c/Cao:Shengming" pid="316/8117">Shengming Cao</na></co>
<co c="0"><na f="c/Cao:Zhonghao" pid="363/6984">Zhonghao Cao</na></co>
<co c="0"><na f="c/Chang:Xiaojun" pid="116/8412">Xiaojun Chang</na></co>
<co c="0"><na f="c/Chang:Xuankai" pid="194/1149">Xuankai Chang</na></co>
<co c="0"><na f="c/Chen_0005:Hongyu" pid="28/3046-5">Hongyu Chen 0005</na></co>
<co c="0"><na f="c/Chen_0001:Jia" pid="99/6879-1">Jia Chen 0001</na></co>
<co c="0"><na f="c/Chen:Jieting" pid="274/6597">Jieting Chen</na></co>
<co c="0"><na f="c/Chen:Kehan" pid="139/2261">Kehan Chen</na></co>
<co c="0"><na f="c/Chen_0008:Liangyu" pid="82/353-8">Liangyu Chen 0008</na></co>
<co c="0"><na f="c/Chen:Peng" pid="27/7017">Peng Chen</na></co>
<co c="0"><na f="c/Chen:Shangkui" pid="435/3086">Shangkui Chen</na></co>
<co c="0"><na f="c/Chen:Shimin" pid="c/ShiminChen">Shimin Chen</na></co>
<co c="0"><na f="c/Chen:Shizhe" pid="153/0734">Shizhe Chen</na></co>
<co c="0"><na f="c/Chen:Weijing" pid="197/5539">Weijing Chen</na></co>
<co c="0"><na f="c/Chen:Wenping" pid="04/1620">Wenping Chen</na></co>
<co c="0"><na f="c/Chen:William" pid="82/2443">William Chen</na></co>
<co c="0"><na f="c/Chen:Xiangyu" pid="84/7543">Xiangyu Chen</na></co>
<co c="0"><na f="c/Chen_0001:Xie" pid="86/11429-1">Xie Chen 0001</na></co>
<co c="0"><na f="c/Chen:Yang" pid="48/4792">Yang Chen</na></co>
<co c="0"><na f="c/Chen:Zhansheng" pid="176/4824">Zhansheng Chen</na></co>
<co c="0"><na f="c/Cheng:Yuan" pid="15/5135">Yuan Cheng</na></co>
<co c="0"><na f="c/Christel:Michael_G=" pid="c/MGChristel">Michael G. Christel</na></co>
<co c="0"><na f="c/Concescu:R=" pid="263/9828">R. Concescu</na></co>
<co c="0"><na f="c/Coto:Ernesto" pid="02/4957">Ernesto Coto</na></co>
<co c="0"><na f="c/Cui:Hanwen" pid="260/7814">Hanwen Cui</na></co>
<co c="0"><na f="c/Cui:Kaixu" pid="251/8611">Kaixu Cui</na></co>
<co c="0"><na f="c/Cui:Wanqing" pid="260/2159">Wanqing Cui</na></co>
<co c="0"><na f="d/Dai_0001:Li=Rong" pid="48/6462-1">Li-Rong Dai 0001</na></co>
<co c="0"><na f="d/Dai:Qifeng" pid="157/2769">Qifeng Dai</na></co>
<co c="0"><na f="d/Deng:Ruifan" pid="388/3709">Ruifan Deng</na></co>
<co c="0"><na f="d/Dian:Yujie" pid="187/9165">Yujie Dian</na></co>
<co c="0"><na f="d/Ding:Duo" pid="36/3063">Duo Ding</na></co>
<co c="0"><na f="d/Ding:Junkai" pid="346/0048">Junkai Ding</na></co>
<co c="0"><na f="d/Ding_0002:Ning" pid="04/4910-2">Ning Ding 0002</na></co>
<co c="0"><na f="d/Dou:Zhicheng" pid="18/5740">Zhicheng Dou</na></co>
<co c="0"><na f="d/Du_0001:Xiaoyong" pid="47/3542-1">Xiaoyong Du 0001</na></co>
<co c="0"><na f="d/Du_0011:Yang" pid="51/3199-11">Yang Du 0011</na></co>
<co c="0"><na f="e/Ekenel:Hazim_Kemal" pid="26/1396">Hazim Kemal Ekenel</na></co>
<co c="0"><na f="e/Epps:Julien" pid="92/6424">Julien Epps</na></co>
<co c="0"><na f="f/Fan:Hehe" pid="184/5722">Hehe Fan</na></co>
<co c="0"><na f="f/Fang:Xiangnan" pid="439/1035">Xiangnan Fang</na></co>
<co c="0"><na f="f/Feng:Wenhao" pid="50/9959">Wenhao Feng</na></co>
<co c="0"><na f="f/Feng:Yicheng" pid="340/4016">Yicheng Feng</na></co>
<co c="0"><na f="f/Fischer:Mika" pid="23/6655">Mika Fischer</na></co>
<co c="0"><na f="f/Fu:Jianlong" pid="83/8692">Jianlong Fu</na></co>
<co c="0"><na f="f/Fu:Pei" pid="202/1707">Pei Fu</na></co>
<co c="0"><na f="f/Fu:Yuhang" pid="322/6277">Yuhang Fu</na></co>
<co c="0"><na f="f/F=uuml=gen:Christian" pid="88/6411">Christian F&#252;gen</na></co>
<co c="0"><na f="f/Fuhs:Mark_C=" pid="25/516">Mark C. Fuhs</na></co>
<co c="0"><na f="g/Gabeur:Valentin" pid="246/5141">Valentin Gabeur</na></co>
<co c="0"><na f="g/Gao:Chenqiang" pid="82/9201">Chenqiang Gao</na></co>
<co c="0"><na f="g/Gao:Dongji" pid="226/5309">Dongji Gao</na></co>
<co c="0"><na f="g/Gao:Jiang" pid="79/1247">Jiang Gao</na></co>
<co c="0"><na f="g/Gao:Qixiang" pid="378/3056">Qixiang Gao</na></co>
<co c="0"><na f="g/Gao_0004:Yizhao" pid="132/7629-4">Yizhao Gao 0004</na></co>
<co c="0"><na f="g/Garbus:Michael" pid="42/10648">Michael Garbus</na></co>
<co c="0"><na f="g/Ge:Tiezheng" pid="135/4944">Tiezheng Ge</na></co>
<co c="0"><na f="g/Ghanem:Bernard" pid="37/2516">Bernard Ghanem</na></co>
<co c="0"><na f="g/Godfrey:John_J=" pid="07/172">John J. Godfrey</na></co>
<co c="0"><na f="g/Gong_0001:Zheng" pid="85/5448-1">Zheng Gong 0001</na></co>
<co c="0"><na f="g/Gu:Yuxian" pid="248/2313">Yuxian Gu</na></co>
<co c="0"><na f="g/Guo:Baining" pid="44/3946">Baining Guo</na></co>
<co c="0"><na f="g/Guo:Shuai" pid="07/272">Shuai Guo</na></co>
<co c="0"><na f="g/Guo:Wu" pid="20/4480">Wu Guo</na></co>
<co c="0"><na f="h/Han:Jionghao" pid="378/4662">Jionghao Han</na></co>
<co c="1"><na f="h/Han_0005:Ning" pid="68/690-5">Ning Han 0005</na></co>
<co c="0"><na f="h/Han:Wentao" pid="10/8243">Wentao Han</na></co>
<co c="0"><na f="h/Han_0007:Xu" pid="19/3011-7">Xu Han 0007</na></co>
<co c="0"><na f="h/Han:Yanling" pid="47/468">Yanling Han</na></co>
<co c="0"><na f="h/Hao:Xiaoshuai" pid="271/8403">Xiaoshuai Hao</na></co>
<co c="0" n="2"><na f="h/Hauptmann_0001:Alex" pid="h/AlexanderGHauptmann">Alex Hauptmann 0001</na><na>Alexander G. Hauptmann</na></co>
<co c="0"><na f="h/Hayashi:Tomoki" pid="82/8616">Tomoki Hayashi</na></co>
<co c="0"><na f="h/He:Huiguo" pid="270/6402">Huiguo He</na></co>
<co c="0"><na f="h/He_0001:Liang" pid="42/963-1">Liang He 0001</na></co>
<co c="0"><na f="h/He:Qingrong" pid="376/0782">Qingrong He</na></co>
<co c="0"><na f="h/He:Xixi" pid="150/5274">Xixi He</na></co>
<co c="0"><na f="h/He:Zewen" pid="237/8845">Zewen He</na></co>
<co c="0"><na f="h/Hoi:Steven" pid="324/3832">Steven Hoi</na></co>
<co c="0"><na f="h/Hong:Xin" pid="54/1309">Xin Hong</na></co>
<co c="0"><na f="h/Hong:Zhonghua" pid="142/6255">Zhonghua Hong</na></co>
<co c="0" n="2"><na f="h/Hou:Danyang" pid="228/0778">Danyang Hou</na><na>Dan Yang Hou</na></co>
<co c="0"><na f="h/Hou:Xinglin" pid="319/3945">Xinglin Hou</na></co>
<co c="0"><na f="h/Hsiao:Roger" pid="86/8054">Roger Hsiao</na></co>
<co c="0"><na f="h/Hsu:Ching=Hsien" pid="35/3725">Ching-Hsien Hsu</na></co>
<co c="0"><na f="h/Hu:Anwen" pid="249/1182">Anwen Hu</na></co>
<co c="0"><na f="h/Hu_0001:Di" pid="49/8496-1">Di Hu 0001</na></co>
<co c="0"><na f="h/Hu_0001:Huanran" pid="385/1380-1">Huanran Hu 0001</na></co>
<co c="0"><na f="h/Hu_0003:Jingwen" pid="159/1747-3">Jingwen Hu 0003</na></co>
<co c="0"><na f="h/Hu:Qixiu" pid="50/4223">Qixiu Hu</na></co>
<co c="0"><na f="h/Hu:Shuo" pid="35/7808">Shuo Hu</na></co>
<co c="0"><na f="h/Hu:Ting=Yao" pid="76/10031">Ting-Yao Hu</na></co>
<co c="0"><na f="h/Huang:Chien=Lin" pid="42/3949">Chien-Lin Huang</na></co>
<co c="0"><na f="h/Huang_0007:Dong" pid="94/3756-7">Dong Huang 0007</na></co>
<co c="0"><na f="h/Huang_0002:Fei" pid="h/FeiHuang-2">Fei Huang 0002</na></co>
<co c="0"><na f="h/Huang:Haoyang" pid="248/7736">Haoyang Huang</na></co>
<co c="0"><na f="h/Huang:Minlie" pid="47/6668">Minlie Huang</na></co>
<co c="0" n="2"><na f="h/Huang_0001:Po=Yao" pid="154/3943-1">Po-Yao Huang 0001</na><na>Poyao Huang 0001</na></co>
<co c="0"><na f="h/Huang:Shen" pid="90/839">Shen Huang</na></co>
<co c="0"><na f="h/Huang:Shi=Sheng" pid="58/8294">Shi-Sheng Huang</na></co>
<co c="0"><na f="h/Huang:Zhaopei" pid="305/3439">Zhaopei Huang</na></co>
<co c="0"><na f="h/Huo:Nan" pid="272/8774">Nan Huo</na></co>
<co c="0"><na f="h/Huo:Yujia" pid="167/4776">Yujia Huo</na></co>
<co c="0"><na f="h/Huo:Yuqi" pid="219/6931">Yuqi Huo</na></co>
<co c="0"><na f="i/Idrees:Haroon" pid="19/8535">Haroon Idrees</na></co>
<co c="0"><na f="j/Jia:Jiaya" pid="31/5649">Jiaya Jia</na></co>
<co c="0"><na f="j/Jia:Yanfeng" pid="149/1280">Yanfeng Jia</na></co>
<co c="0"><na f="j/Jiang_0004:Lu" pid="22/752-4">Lu Jiang 0004</na></co>
<co c="0"><na f="j/Jiang:Wenqiang" pid="276/9880">Wenqiang Jiang</na></co>
<co c="0"><na f="j/Jiang:Yudong" pid="150/4256">Yudong Jiang</na></co>
<co c="0"><na f="j/Jiang_0001:Yuning" pid="99/7950-1">Yuning Jiang 0001</na></co>
<co c="0"><na f="j/Jin:Chuhao" pid="287/4999">Chuhao Jin</na></co>
<co c="0"><na f="j/Jing:Tianjiao" pid="402/6306">Tianjiao Jing</na></co>
<co c="0"><na f="j/Jones:Douglas_A=" pid="32/2452">Douglas A. Jones</na></co>
<co c="0"><na f="j/Jou:Szu=Chen_Stan" pid="17/1054">Szu-Chen Stan Jou</na></co>
<co c="0"><na f="j/Ju:Jianzhong" pid="289/4197">Jianzhong Ju</na></co>
<co c="0" n="2"><na f="j/Jung:Jee=Weon" pid="212/6435">Jee-Weon Jung</na><na>Jee-weon Jung</na></co>
<co c="0"><na f="k/Kang:Zhiyu" pid="56/6876">Zhiyu Kang</na></co>
<co c="0"><na f="k/K=auml=rkk=auml=inen:Ismo" pid="35/3843">Ismo K&#228;rkk&#228;inen</na></co>
<co c="0"><na f="k/Ke:Wei" pid="52/7566">Wei Ke</na></co>
<co c="0" n="2"><na f="k/Kitani:Kris_Makoto" pid="42/163">Kris Makoto Kitani</na><na>Kris Kitani</na></co>
<co c="0"><na f="k/Klus=aacute=cek:David" pid="142/6673">David Klus&#225;cek</na></co>
<co c="0"><na f="k/Kong:Quyu" pid="209/9882">Quyu Kong</na></co>
<co c="0"><na f="l/Lai:Shaopeng" pid="285/4166">Shaopeng Lai</na></co>
<co c="0"><na f="l/Lan:Weiyu" pid="166/2924">Weiyu Lan</na></co>
<co c="0"><na f="l/Lan:Yanyan" pid="00/6040">Yanyan Lan</na></co>
<co c="0" n="3"><na f="l/Lan:Zhen=Zhong" pid="27/3780">Zhen-Zhong Lan</na><na>Zhen-zhong Lan</na><na>Zhenzhong Lan</na></co>
<co c="0"><na f="l/Langner:Brian" pid="25/8160">Brian Langner</na></co>
<co c="0"><na f="l/Laptev:Ivan" pid="41/1854">Ivan Laptev</na></co>
<co c="0"><na f="l/Laskowski:Kornel" pid="98/2383">Kornel Laskowski</na></co>
<co c="0"><na f="l/Lee:Kong=Aik" pid="35/4621">Kong-Aik Lee</na></co>
<co c="0"><na f="l/Li:Ao" pid="54/2788">Ao Li</na></co>
<co c="0"><na f="l/Li_0001:Chen" pid="l/ChenLi1">Chen Li 0001</na></co>
<co c="0"><na f="l/Li:Chengxin" pid="166/6518">Chengxin Li</na></co>
<co c="0"><na f="l/Li_0003:Chenliang" pid="52/9457-3">Chenliang Li 0003</na></co>
<co c="0"><na f="l/Li_0001:Haizhou" pid="36/4118">Haizhou Li 0001</na></co>
<co c="0"><na f="l/Li:Huazhe" pid="320/0415">Huazhe Li</na></co>
<co c="0"><na f="l/Li:Jiaze" pid="157/2088">Jiaze Li</na></co>
<co c="0"><na f="l/Li_0001:Junyi" pid="28/6612-1">Junyi Li 0001</na></co>
<co c="0"><na f="l/Li:Lantian" pid="153/0735">Lantian Li</na></co>
<co c="0"><na f="l/Li_0022:Min" pid="82/0-22">Min Li 0022</na></co>
<co c="0"><na f="l/Li:Ruichen" pid="231/3910">Ruichen Li</na></co>
<co c="0"><na f="l/Li:Runxin" pid="26/8160">Runxin Li</na></co>
<co c="1"><na f="l/Li:Wenfeng" pid="20/179">Wenfeng Li</na></co>
<co c="0"><na f="l/Li:Xinrui" pid="151/4579">Xinrui Li</na></co>
<co c="0"><na f="l/Li_0001:Xirong" pid="58/5856">Xirong Li 0001</na></co>
<co c="0"><na f="l/Li:Xuanchong" pid="150/6544">Xuanchong Li</na></co>
<co c="0"><na f="l/Li:Xubin" pid="228/7857">Xubin Li</na></co>
<co c="0"><na f="l/Li_0001:Ya" pid="51/4056-1">Ya Li 0001</na></co>
<co c="0"><na f="l/Li_0001:Yijie" pid="54/8054-1">Yijie Li 0001</na></co>
<co c="0"><na f="l/Li:Yingyan" pid="287/4789">Yingyan Li</na></co>
<co c="0"><na f="l/Li_0001:Yuanpeng" pid="08/8174-1">Yuanpeng Li 0001</na></co>
<co c="0"><na f="l/Li:Zhuoqun" pid="11/2724">Zhuoqun Li</na></co>
<co c="0"><na f="l/Lian:Zhengxuan" pid="419/4007">Zhengxuan Lian</na></co>
<co c="0"><na f="l/Liang:Chao" pid="10/3072">Chao Liang</na></co>
<co c="0"><na f="l/Liang:Jingjun" pid="243/6806">Jingjun Liang</na></co>
<co c="0"><na f="l/Liang_0001:Junwei" pid="62/10704-1">Junwei Liang 0001</na></co>
<co c="0"><na f="l/Liao:Shuai" pid="138/4239">Shuai Liao</na></co>
<co c="0"><na f="l/Lin:Hongpeng" pid="337/9434">Hongpeng Lin</na></co>
<co c="0"><na f="l/Lin:Junqi" pid="268/5425">Junqi Lin</na></co>
<co c="0"><na f="l/Lin:Kejun" pid="251/8268">Kejun Lin</na></co>
<co c="0"><na f="l/Lin:Linli" pid="331/1495">Linli Lin</na></co>
<co c="0"><na f="l/Lin:Pingping" pid="39/7629">Pingping Lin</na></co>
<co c="0" n="2"><na f="l/Lin_0003:Weihao" pid="295/0536-3">Weihao Lin 0003</na><na>Wei-Hao Lin 0003</na></co>
<co c="0"><na f="l/Lin:Xiaozhu" pid="07/1292">Xiaozhu Lin</na></co>
<co c="0" n="2"><na f="l/Lin_0001:Xin" pid="50/3323-1">Xin Lin 0001</na><na>Xin Alex Lin</na></co>
<co c="0"><na f="l/Lin:Yueqian" pid="349/5036">Yueqian Lin</na></co>
<co c="0"><na f="l/Lin:Zhuoran" pid="223/9450">Zhuoran Lin</na></co>
<co c="0" n="2"><na f="l/Ling:Zhen=Hua" pid="70/5210">Zhen-Hua Ling</na><na>Zhenhua Ling</na></co>
<co c="0"><na f="l/Liu:Alexander_H=" pid="227/2380">Alexander H. Liu</na></co>
<co c="0"><na f="l/Liu_0001:Bei" pid="39/3711-1">Bei Liu 0001</na></co>
<co c="0"><na f="l/Liu:Binbin" pid="22/7702">Binbin Liu</na></co>
<co c="0"><na f="l/Liu:Chang" pid="52/5716">Chang Liu</na></co>
<co c="0"><na f="l/Liu:Chen" pid="10/2639">Chen Liu</na></co>
<co c="0"><na f="l/Liu:Chuanhe" pid="227/5169">Chuanhe Liu</na></co>
<co c="0"><na f="l/Liu:Guangzhen" pid="186/6864">Guangzhen Liu</na></co>
<co c="0"><na f="l/Liu:Haibo" pid="83/1694">Haibo Liu</na></co>
<co c="0"><na f="l/Liu_0004:Hongwei" pid="43/5900-4">Hongwei Liu 0004</na></co>
<co c="0"><na f="l/Liu:Hui" pid="93/4010">Hui Liu</na></co>
<co c="0"><na f="l/Liu_0011:Jiang" pid="23/108-11">Jiang Liu 0011</na></co>
<co c="0"><na f="l/Liu:Jiazheng" pid="191/0307">Jiazheng Liu</na></co>
<co c="0"><na f="l/Liu:Jing" pid="72/2590">Jing Liu</na></co>
<co c="0"><na f="l/Liu:Lan" pid="43/6387">Lan Liu</na></co>
<co c="0"><na f="l/Liu:Luoqi" pid="29/8842">Luoqi Liu</na></co>
<co c="0"><na f="l/Liu_0002:Peiyu" pid="85/670-2">Peiyu Liu 0002</na></co>
<co c="0"><na f="l/Liu_0003:Shichao" pid="134/5661-3">Shichao Liu 0003</na></co>
<co c="0"><na f="l/Liu_0001:Si" pid="60/7642">Si Liu 0001</na></co>
<co c="0"><na f="l/Liu:Weiliang" pid="77/6507">Weiliang Liu</na></co>
<co c="0"><na f="l/Liu:Wenhe" pid="184/5777">Wenhe Liu</na></co>
<co c="0"><na f="l/Liu_0036:Xiao" pid="82/1364-36">Xiao Liu 0036</na></co>
<co c="0"><na f="l/Liu:Xiaolong" pid="48/1674">Xiaolong Liu</na></co>
<co c="0"><na f="l/Liu:Xinbi" pid="402/1495">Xinbi Liu</na></co>
<co c="0"><na f="l/Liu:Xuan" pid="13/5407">Xuan Liu</na></co>
<co c="0"><na f="l/Liu_0005:Yang" pid="51/3710-5">Yang Liu 0005</na></co>
<co c="0"><na f="l/Liu_0105:Yang" pid="51/3710-105">Yang Liu 0105</na></co>
<co c="0"><na f="l/Liu:Yuanhang" pid="174/0438">Yuanhang Liu</na></co>
<co c="0"><na f="l/Liu_0003:Yuchen" pid="69/10440-3">Yuchen Liu 0003</na></co>
<co c="0"><na f="l/Liu_0003:Yuqi" pid="35/9071-3">Yuqi Liu 0003</na></co>
<co c="0"><na f="l/Liu_0001:Zhiyuan" pid="53/3245-1">Zhiyuan Liu 0001</na></co>
<co c="0"><na f="l/Lou:Fan" pid="282/9088">Fan Lou</na></co>
<co c="0"><na f="l/Lu:Haoyu" pid="240/2720">Haoyu Lu</na></co>
<co c="0"><na f="l/Lu:Li" pid="416/3281">Li Lu</na></co>
<co c="0"><na f="l/Lu:Qi" pid="41/4012">Qi Lu</na></co>
<co c="0"><na f="l/Lu:Yanxiong" pid="40/7053">Yanxiong Lu</na></co>
<co c="0"><na f="l/Lu_0001:Zhiwu" pid="53/5234">Zhiwu Lu 0001</na></co>
<co c="0"><na f="l/Lu_0002:Zongqing" pid="99/965-2">Zongqing Lu 0002</na></co>
<co c="0"><na f="l/Luan_0001:Jian" pid="61/3233-1">Jian Luan 0001</na></co>
<co c="0"><na f="l/Luo_0011:Hao" pid="14/3727-11">Hao Luo 0011</na></co>
<co c="0"><na f="l/Luo:Jinghan" pid="435/6686">Jinghan Luo</na></co>
<co c="0"><na f="l/Luo:Wei" pid="05/6715">Wei Luo</na></co>
<co c="0"><na f="l/Luo:Zhenbo" pid="152/8206">Zhenbo Luo</na></co>
<co c="0"><na f="m/Ma_0001:Bin" pid="70/6176-1">Bin Ma 0001</na></co>
<co c="0"><na f="m/Ma:Chengcheng" pid="189/3741">Chengcheng Ma</na></co>
<co c="0"><na f="m/Ma:Jianzhe" pid="253/1100">Jianzhe Ma</na></co>
<co c="0"><na f="m/Ma:Yiyang" pid="324/2590">Yiyang Ma</na></co>
<co c="0"><na f="m/Ma_0001:Zejun" pid="18/10648-1">Zejun Ma 0001</na></co>
<co c="0"><na f="m/Magalh=atilde=es:Jo=atilde=o" pid="08/1790">Jo&#227;o Magalh&#227;es</na></co>
<co c="0"><na f="m/Masuyama:Yoshiki" pid="226/5828">Yoshiki Masuyama</na></co>
<co c="0"><na f="m/Mei:Jiahao" pid="354/8774">Jiahao Mei</na></co>
<co c="0"><na f="m/Mei:Yuting" pid="334/2644">Yuting Mei</na></co>
<co c="0"><na f="m/Meng:Liyu" pid="276/9497">Liyu Meng</na></co>
<co c="0"><na f="m/Meng:Wanting" pid="142/6162">Wanting Meng</na></co>
<co c="0"><na f="m/Metze:Florian" pid="26/1652">Florian Metze</na></co>
<co c="0"><na f="m/Miech:Antoine" pid="202/1721">Antoine Miech</na></co>
<co c="0"><na f="m/Mihaescu:Radu" pid="80/5017">Radu Mihaescu</na></co>
<co c="0"><na f="m/Mu:Guankun" pid="187/9132">Guankun Mu</na></co>
<co c="0"><na f="n/Nagrani:Arsha" pid="202/1922">Arsha Nagrani</na></co>
<co c="0"><na f="n/Nallasamy:Udhyakumar" pid="76/7551">Udhyakumar Nallasamy</na></co>
<co c="0"><na f="n/Navr=aacute=til_0001:Jir=iacute=" pid="00/680-1">Jir&#237; Navr&#225;til 0001</na></co>
<co c="0"><na f="n/Nosratighods:Mohaddeseh" pid="31/3678">Mohaddeseh Nosratighods</na></co>
<co c="0"><na f="o/Oria:Vincent" pid="o/VincentOria">Vincent Oria</na></co>
<co c="0"><na f="o/Ou:Zhijian" pid="44/5256">Zhijian Ou</na></co>
<co c="0"><na f="o/Overwijk:Arnold" pid="16/7404">Arnold Overwijk</na></co>
<co c="0"><na f="p/Pan:J==Y=" pid="263/9885">J.-Y. Pan</na></co>
<co c="1"><na f="p/Pan:Keyu" pid="294/1953">Keyu Pan</na></co>
<co c="0"><na f="p/Pan:Pingbo" pid="172/0936">Pingbo Pan</na></co>
<co c="0"><na f="p/Pan:Yue" pid="91/3754">Yue Pan</na></co>
<co c="0"><na f="p/Pervouchine:Vladimir" pid="98/6470">Vladimir Pervouchine</na></co>
<co c="0"><na f="p/Peskin:Barbara" pid="15/1225">Barbara Peskin</na></co>
<co c="0"><na f="q/Qi:Xiaoyu" pid="166/6091">Xiaoyu Qi</na></co>
<co c="0"><na f="q/Qian_0001:Qi" pid="05/2084-1">Qi Qian 0001</na></co>
<co c="0"><na f="q/Qian:Tao" pid="05/5638">Tao Qian</na></co>
<co c="0"><na f="q/Qian:Yanmin" pid="07/8638">Yanmin Qian</na></co>
<co c="0"><na f="q/Qin_0001:Yong" pid="20/4298-1">Yong Qin 0001</na></co>
<co c="0"><na f="q/Qiu:Jiarong" pid="222/3007">Jiarong Qiu</na></co>
<co c="0"><na f="q/Qiu:Jiezhong" pid="152/1733">Jiezhong Qiu</na></co>
<co c="0"><na f="q/Qiu:Xipeng" pid="69/1395">Xipeng Qiu</na></co>
<co c="0"><na f="q/Qu:Tianyuan" pid="366/0733">Tianyuan Qu</na></co>
<co c="0"><na f="r/Raj:Bhiksha" pid="60/3996">Bhiksha Raj</na></co>
<co c="0"><na f="r/Rawat:Shourabh" pid="117/4119">Shourabh Rawat</na></co>
<co c="0"><na f="r/Ren:Yuran" pid="401/8284">Yuran Ren</na></co>
<co c="0"><na f="r/Ren:Zihui" pid="20/7764">Zihui Ren</na></co>
<co c="0"><na f="r/Reynolds:Douglas_A=" pid="21/2250">Douglas A. Reynolds</na></co>
<co c="0"><na f="r/Ruan:Ludan" pid="262/6356">Ludan Ruan</na></co>
<co c="0"><na f="r/Rui:Yong" pid="r/YongRui">Yong Rui</na></co>
<co c="0"><na f="s/Salakhutdinov:Ruslan" pid="62/5884">Ruslan Salakhutdinov</na></co>
<co c="0"><na f="s/Satoh_0001:Shin=ichi" pid="50/290">Shin'ichi Satoh 0001</na></co>
<co c="0"><na f="s/Schaaf:Thomas" pid="75/5600">Thomas Schaaf</na></co>
<co c="0"><na f="s/Schmid:Cordelia" pid="s/CordeliaSchmid">Cordelia Schmid</na></co>
<co c="0" n="2"><na f="s/Schulam_0001:Peter" pid="117/4105">Peter Schulam 0001</na><na>Peter Franz Schulam</na></co>
<co c="0"><na f="s/Schultz:Tanja" pid="s/TanjaSchultz">Tanja Schultz</na></co>
<co c="0"><na f="s/Sebe:Nicu" pid="20/3519">Nicu Sebe</na></co>
<co c="0"><na f="s/Sheikh:Yaser" pid="71/3516">Yaser Sheikh</na></co>
<co c="0"><na f="s/Shi:Jiatong" pid="229/3529">Jiatong Shi</na></co>
<co c="0"><na f="s/Si:Luo" pid="14/1217">Luo Si</na></co>
<co c="0"><na f="s/Sim:Khe_Chai" pid="78/6873">Khe Chai Sim</na></co>
<co c="0"><na f="s/Siong:Chng_Eng" pid="c/ChngEngSiong">Chng Eng Siong</na></co>
<co c="0"><na f="s/Soltau:Hagen" pid="07/2072">Hagen Soltau</na></co>
<co c="0"><na f="s/Song:Kaiqiang" pid="222/2920">Kaiqiang Song</na></co>
<co c="0"><na f="s/Song:Ruihua" pid="s/RuihuaSong">Ruihua Song</na></co>
<co c="0"><na f="s/Song_0003:Yuqing" pid="222/3101">Yuqing Song 0003</na></co>
<co c="0"><na f="s/Song:Zhinan" pid="378/1643">Zhinan Song</na></co>
<co c="0"><na f="s/Srivastava:Tejes" pid="332/5917">Tejes Srivastava</na></co>
<co c="0"><na f="s/Stevens:Scott_M=" pid="92/5829">Scott M. Stevens</na></co>
<co c="0"><na f="s/Stiefelhagen:Rainer" pid="31/4699">Rainer Stiefelhagen</na></co>
<co c="0"><na f="s/Su:Zhong" pid="87/1363">Zhong Su</na></co>
<co c="0"><na f="s/Sukthankar:Rahul" pid="57/3775">Rahul Sukthankar</na></co>
<co c="0"><na f="s/Sun_0002:Chen" pid="01/6072-2">Chen Sun 0002</na></co>
<co c="0"><na f="s/Sun:Hanwu" pid="22/1460">Hanwu Sun</na></co>
<co c="0"><na f="s/Sun:Hui" pid="31/3925">Hui Sun</na></co>
<co c="0"><na f="s/Sun_0001:Jiande" pid="68/365-1">Jiande Sun 0001</na></co>
<co c="0"><na f="s/Sun:Lei" pid="02/2264">Lei Sun</na></co>
<co c="0"><na f="s/Sun:Seson" pid="440/1367">Seson Sun</na></co>
<co c="0"><na f="s/Sun_0001:Xu" pid="37/1971-1">Xu Sun 0001</na></co>
<co c="0"><na f="s/Sun:Yuchong" pid="206/8045">Yuchong Sun</na></co>
<co c="0"><na f="t/Takahashi:Shohei" pid="181/8511">Shohei Takahashi</na></co>
<co c="0"><na f="t/Tam:Yik=Cheung" pid="72/158">Yik-Cheung Tam</na></co>
<co c="0"><na f="t/Tang_0001:Jie" pid="t/JieTang">Jie Tang 0001</na></co>
<co c="0"><na f="t/Tang:Yuxun" pid="358/9271">Yuxun Tang</na></co>
<co c="0"><na f="t/Tang:Zongheng" pid="278/2959">Zongheng Tang</na></co>
<co c="0"><na f="t/Tao_0001:Jianhua" pid="46/2916-1">Jianhua Tao 0001</na></co>
<co c="0"><na f="t/Thiruvaran:Tharmarajah" pid="37/8053">Tharmarajah Thiruvaran</na></co>
<co c="0"><na f="t/Tian:Jinchuan" pid="249/2901">Jinchuan Tian</na></co>
<co c="0"><na f="t/Tian:Junfeng" pid="93/1076">Junfeng Tian</na></co>
<co c="0"><na f="t/Tong:Panrong" pid="223/6918">Panrong Tong</na></co>
<co c="0"><na f="t/Tong:Rong" pid="09/1880">Rong Tong</na></co>
<co c="0"><na f="t/Toni:Laura" pid="81/7871">Laura Toni</na></co>
<co c="0"><na f="t/Toth:Arthur_R=" pid="99/7822">Arthur R. Toth</na></co>
<co c="0"><na f="t/Tribble:Alicia" pid="60/5193">Alicia Tribble</na></co>
<co c="0"><na f="v/Vaibhav:" pid="177/2366">Vaibhav</na></co>
<co c="0"><na f="w/Waibel:Alex" pid="08/2456">Alex Waibel</na></co>
<co c="0"><na f="w/Wang:Biao" pid="15/2887">Biao Wang</na></co>
<co c="0"><na f="w/Wang:Chen" pid="82/4206">Chen Wang</na></co>
<co c="0"><na f="w/Wang:Chuhan" pid="210/1379">Chuhan Wang</na></co>
<co c="0" n="2"><na f="w/Wang:Chun=Ting" pid="239/3587">Chun-Ting Wang</na><na>Chunting Wang</na></co>
<co c="0"><na f="w/Wang:Huihui" pid="240/2025">Huihui Wang</na></co>
<co c="0"><na f="w/Wang_0073:Jun" pid="125/8189-73">Jun Wang 0073</na></co>
<co c="0"><na f="w/Wang:Meng" pid="93/6765">Meng Wang</na></co>
<co c="0"><na f="w/Wang:Mingshuo" pid="141/0952">Mingshuo Wang</na></co>
<co c="0"><na f="w/Wang_0015:Peng" pid="95/4442-15">Peng Wang 0015</na></co>
<co c="0"><na f="w/Wang:Shuai" pid="42/1503">Shuai Wang</na></co>
<co c="0"><na f="w/Wang:Weiying" pid="227/7810">Weiying Wang</na></co>
<co c="0"><na f="w/Wang_0001:Wenxuan" pid="203/1536-1">Wenxuan Wang 0001</na></co>
<co c="0"><na f="w/Wang:Xinchao" pid="23/8015">Xinchao Wang</na></co>
<co c="0"><na f="w/Wang:Xueyan" pid="37/5333">Xueyan Wang</na></co>
<co c="0"><na f="w/Wang:Ye" pid="44/6292">Ye Wang</na></co>
<co c="0"><na f="w/Wang:Yingqiao" pid="430/3481">Yingqiao Wang</na></co>
<co c="0"><na f="w/Wang:Yongcheng" pid="06/4175">Yongcheng Wang</na></co>
<co c="0"><na f="w/Wang_0039:Yue" pid="33/4822-39">Yue Wang 0039</na></co>
<co c="0"><na f="w/Wang:Ziheng" pid="79/10743">Ziheng Wang</na></co>
<co c="0"><na f="w/Wang:Zijia" pid="121/0900">Zijia Wang</na></co>
<co c="0"><na f="w/Watanabe_0001:Shinji" pid="39/3245-1">Shinji Watanabe 0001</na></co>
<co c="0"><na f="w/Wei:Furu" pid="72/5870">Furu Wei</na></co>
<co c="0"><na f="w/Wei:Qianshan" pid="378/1061">Qianshan Wei</na></co>
<co c="0"><na f="w/Wen:Ji=Rong" pid="w/JRWen">Ji-Rong Wen</na></co>
<co c="0"><na f="w/Wen:Jingyuan" pid="287/4833">Jingyuan Wen</na></co>
<co c="0"><na f="w/W=ouml=lfel:Matthias" pid="84/6910">Matthias W&#246;lfel</na></co>
<co c="0"><na f="w/Wu:Guozheng" pid="262/6203">Guozheng Wu</na></co>
<co c="0"><na f="w/Wu:Haibin" pid="95/10219">Haibin Wu</na></co>
<co c="0"><na f="w/Wu:Huimin" pid="19/6993">Huimin Wu</na></co>
<co c="0"><na f="w/Wu:Mengyue" pid="82/2416">Mengyue Wu</na></co>
<co c="0"><na f="w/Wu:Peter" pid="44/3072">Peter Wu</na></co>
<co c="0"><na f="w/Wu_0001:Qi" pid="96/3446-1">Qi Wu 0001</na></co>
<co c="0"><na f="w/Wu:Shaohui" pid="136/9323">Shaohui Wu</na></co>
<co c="0"><na f="w/Wu:Xiaopeng" pid="153/4201">Xiaopeng Wu</na></co>
<co c="0"><na f="w/Wu_0008:Yihan" pid="72/10107-8">Yihan Wu 0008</na></co>
<co c="0"><na f="w/Wu_0002:Yike" pid="246/5764-2">Yike Wu 0002</na></co>
<co c="0"><na f="w/Wu_0001:Yuning" pid="301/4852-1">Yuning Wu 0001</na></co>
<co c="0"><na f="w/Wu_0001:Zhiyong" pid="24/968-1">Zhiyong Wu 0001</na></co>
<co c="0"><na f="x/Xi:Zongzheng" pid="287/4970">Zongzheng Xi</na></co>
<co c="0"><na f="x/Xia:Lujie" pid="352/3374">Lujie Xia</na></co>
<co c="0"><na f="x/Xia:Song" pid="01/9629">Song Xia</na></co>
<co c="0"><na f="x/Xia:Wenke" pid="337/9800">Wenke Xia</na></co>
<co c="0"><na f="x/Xiang:Bing" pid="82/5456">Bing Xiang</na></co>
<co c="0"><na f="x/Xiao:Bin" pid="43/5134">Bin Xiao</na></co>
<co c="0"><na f="x/Xiao_0001:Zihan" pid="319/7252-1">Zihan Xiao 0001</na></co>
<co c="0"><na f="x/Xie:Jun" pid="33/3881">Jun Xie</na></co>
<co c="0"><na f="x/Xie_0001:Lei" pid="70/1741-1">Lei Xie 0001</na></co>
<co c="0"><na f="x/Xiong:Pengfei" pid="48/8617">Pengfei Xiong</na></co>
<co c="0"><na f="x/Xiong_0001:Yifan" pid="188/1149-1">Yifan Xiong 0001</na></co>
<co c="0"><na f="x/Xu:Baogui" pid="287/5028">Baogui Xu</na></co>
<co c="0"><na f="x/Xu:Boshen" pid="293/8958">Boshen Xu</na></co>
<co c="0"><na f="x/Xu:Chaoyi" pid="358/4076">Chaoyi Xu</na></co>
<co c="0"><na f="x/Xu:Fangzheng" pid="257/8301">Fangzheng Xu</na></co>
<co c="0" n="2"><na f="x/Xu:Frank_F=" pid="190/4519">Frank F. Xu</na><na>Frank Xu 0001</na></co>
<co c="0"><na f="x/Xu:Guohai" pid="205/7621">Guohai Xu</na></co>
<co c="0"><na f="x/Xu:Haiweng" pid="426/9329">Haiweng Xu</na></co>
<co c="0"><na f="x/Xu_0001:Haiyang" pid="80/1339-1">Haiyang Xu 0001</na></co>
<co c="0"><na f="x/Xu:Jieping" pid="96/979">Jieping Xu</na></co>
<co c="0"><na f="x/Xu:Luhui" pid="229/5603">Luhui Xu</na></co>
<co c="0"><na f="x/Xu:Nuo" pid="74/6847">Nuo Xu</na></co>
<co c="0"><na f="x/Xu_0003:Yichen" pid="177/6703-3">Yichen Xu 0003</na></co>
<co c="0"><na f="x/Xu:Yixin" pid="04/5753">Yixin Xu</na></co>
<co c="0"><na f="x/Xu:Yuhao" pid="205/8666">Yuhao Xu</na></co>
<co c="0"><na f="y/Yan_0008:Ming" pid="51/5332-8">Ming Yan 0008</na></co>
<co c="0"><na f="y/Yan:Rong" pid="51/1293">Rong Yan</na></co>
<co c="0"><na f="y/Yan:Yiming" pid="121/6500">Yiming Yan</na></co>
<co c="0"><na f="y/Yang:Dingyi" pid="266/2264">Dingyi Yang</na></co>
<co c="0"><na f="y/Yang:Donglu" pid="413/7448">Donglu Yang</na></co>
<co c="0"><na f="y/Yang_0001:Gang" pid="36/4658-1">Gang Yang 0001</na></co>
<co c="0"><na f="y/Yang:Guoxing" pid="271/9521">Guoxing Yang</na></co>
<co c="0"><na f="y/Yang_0005:Huan" pid="86/4843-5">Huan Yang 0005</na></co>
<co c="0"><na f="y/Yang_0003:Jun" pid="y/JunYang3">Jun Yang 0003</na></co>
<co c="0"><na f="y/Yang:Qian" pid="15/3199">Qian Yang</na></co>
<co c="0"><na f="y/Yang:Shan" pid="72/8479">Shan Yang</na></co>
<co c="0"><na f="y/Yang:Shuhu" pid="253/4903">Shuhu Yang</na></co>
<co c="0"><na f="y/Yang:Yantai" pid="331/1645">Yantai Yang</na></co>
<co c="0"><na f="y/Yang_0001:Yi" pid="33/4854-1">Yi Yang 0001</na></co>
<co c="0"><na f="y/Yang:Yueqian" pid="287/5041">Yueqian Yang</na></co>
<co c="0"><na f="y/Yao:Linli" pid="262/6579">Linli Yao</na></co>
<co c="0"><na f="y/Yao_0013:Yuan" pid="25/4120-13">Yuan Yao 0013</na></co>
<co c="0"><na f="y/Ye:Jiabo" pid="304/1336">Jiabo Ye</na></co>
<co c="0"><na f="y/Ye:Qinghao" pid="254/3247">Qinghao Ye</na></co>
<co c="0"><na f="y/Yen:Neil_Y=" pid="20/6973">Neil Y. Yen</na></co>
<co c="0"><na f="y/Yin_0006:Xiang" pid="18/1022-6">Xiang Yin 0006</na></co>
<co c="0"><na f="y/Yip:Jia_Qi" pid="329/5838">Jia Qi Yip</na></co>
<co c="0"><na f="y/Yoo:Donghyun" pid="207/7436">Donghyun Yoo</na></co>
<co c="0" n="2"><na f="y/You:Chang_Huai" pid="00/3382">Chang Huai You</na><na>Changhuai You</na></co>
<co c="0"><na f="y/Yu_0001:Bei" pid="28/4556-1">Bei Yu 0001</na></co>
<co c="0"><na f="y/Yu:Chuan" pid="50/790">Chuan Yu</na></co>
<co c="0"><na f="y/Yu_0008:Hua" pid="02/2407-8">Hua Yu 0008</na></co>
<co c="0"><na f="y/Yu:Mingyang" pid="135/9009">Mingyang Yu</na></co>
<co c="0"><na f="y/Yu:Shoou=I" pid="23/7442">Shoou-I Yu</na></co>
<co c="0"><na f="y/Yu:Yifeng" pid="09/9833">Yifeng Yu</na></co>
<co c="0"><na f="y/Yu_0001:Yong" pid="43/5685-1">Yong Yu 0001</na></co>
<co c="0"><na f="y/Yuan:Haoqi" pid="254/2084">Haoqi Yuan</na></co>
<co c="0"><na f="y/Yuan:Hongliang" pid="178/5601">Hongliang Yuan</na></co>
<co c="0"><na f="y/Yuan:Jinhui" pid="58/3397">Jinhui Yuan</na></co>
<co c="0"><na f="y/Yuan:Nicholas_Jing" pid="131/4855">Nicholas Jing Yuan</na></co>
<co c="0"><na f="y/Yue:Zihao" pid="339/2864">Zihao Yue</na></co>
<co c="0"><na f="z/Zeng:Weishuai" pid="385/3798">Weishuai Zeng</na></co>
<co c="1"><na f="z/Zeng:Yawen" pid="240/5421">Yawen Zeng</na></co>
<co c="0"><na f="z/Zeng:Zhaoyang" pid="207/1887">Zhaoyang Zeng</na></co>
<co c="0"><na f="z/Zhan:Chunru" pid="335/2023">Chunru Zhan</na></co>
<co c="0"><na f="z/Zhang:Ao" pid="187/6243">Ao Zhang</na></co>
<co c="0"><na f="z/Zhang_0071:Bo" pid="36/2259-71">Bo Zhang 0071</na></co>
<co c="0"><na f="z/Zhang:Chao" pid="94/3019">Chao Zhang</na></co>
<co c="0"><na f="z/Zhang:Chunlei" pid="65/6731">Chunlei Zhang</na></co>
<co c="0"><na f="z/Zhang:Dong" pid="68/3245">Dong Zhang</na></co>
<co c="0"><na f="z/Zhang:Dongdong" pid="02/621">Dongdong Zhang</na></co>
<co c="0"><na f="z/Zhang:Fengyuan" pid="192/1145">Fengyuan Zhang</na></co>
<co c="0"><na f="z/Zhang:Haoyu" pid="168/0332">Haoyu Zhang</na></co>
<co c="0"><na f="z/Zhang:Heng" pid="55/826">Heng Zhang</na></co>
<co c="0"><na f="z/Zhang_0011:Ji" pid="86/1953-11">Ji Zhang 0011</na></co>
<co c="0"><na f="z/Zhang:Jianan" pid="69/10223">Jianan Zhang</na></co>
<co c="0"><na f="z/Zhang:Jiawei" pid="10/239">Jiawei Zhang</na></co>
<co c="0"><na f="z/Zhang:Jing" pid="05/3499">Jing Zhang</na></co>
<co c="0"><na f="z/Zhang:Keyi" pid="164/6584">Keyi Zhang</na></co>
<co c="0"><na f="z/Zhang_0007:Li" pid="89/5992-7">Li Zhang 0007</na></co>
<co c="0"><na f="z/Zhang:Liang" pid="50/6759">Liang Zhang</na></co>
<co c="0"><na f="z/Zhang:Longfei" pid="30/7663">Longfei Zhang</na></co>
<co c="0"><na f="z/Zhang:Manli" pid="223/5376">Manli Zhang</na></co>
<co c="0"><na f="z/Zhang:Mengying" pid="133/6346">Mengying Zhang</na></co>
<co c="0"><na f="z/Zhang:Qi" pid="52/323">Qi Zhang</na></co>
<co c="0"><na f="z/Zhang:Qichen" pid="154/7819">Qichen Zhang</na></co>
<co c="0"><na f="z/Zhang:Shilei" pid="10/3523">Shilei Zhang</na></co>
<co c="0"><na f="z/Zhang:Tao" pid="15/4777">Tao Zhang</na></co>
<co c="0"><na f="z/Zhang:Tenggan" pid="305/3353">Tenggan Zhang</na></co>
<co c="0"><na f="z/Zhang_0002:Wanpeng" pid="73/10693-2">Wanpeng Zhang 0002</na></co>
<co c="0"><na f="z/Zhang:Weipeng" pid="29/5921">Weipeng Zhang</na></co>
<co c="0"><na f="z/Zhang:Xinjie" pid="66/10435">Xinjie Zhang</na></co>
<co c="0"><na f="z/Zhang:Xu" pid="98/5660">Xu Zhang</na></co>
<co c="0"><na f="z/Zhang:Y=" pid="52/5040">Y. Zhang</na></co>
<co c="0"><na f="z/Zhang:Yepeng" pid="355/8240">Yepeng Zhang</na></co>
<co c="0"><na f="z/Zhang:Yuanmeng" pid="304/8398">Yuanmeng Zhang</na></co>
<co c="0"><na f="z/Zhang:Yuekai" pid="191/5931">Yuekai Zhang</na></co>
<co c="0"><na f="z/Zhang_0012:Yun" pid="02/6428-12">Yun Zhang 0012</na></co>
<co c="0"><na f="z/Zhang:Zhengyan" pid="23/10446">Zhengyan Zhang</na></co>
<co c="0"><na f="z/Zhang:Zhipeng" pid="49/4941">Zhipeng Zhang</na></co>
<co c="0"><na f="z/Zhao:Jiayi" pid="117/4068">Jiayi Zhao</na></co>
<co c="0"><na f="z/Zhao:Jinming" pid="121/8902">Jinming Zhao</na></co>
<co c="0"><na f="z/Zhao:Qiangfu" pid="51/1660">Qiangfu Zhao</na></co>
<co c="0"><na f="z/Zhao:Shiwan" pid="73/1947">Shiwan Zhao</na></co>
<co c="0"><na f="z/Zhao:Wayne_Xin" pid="52/8700">Wayne Xin Zhao</na></co>
<co c="0"><na f="z/Zhao:Yida" pid="222/2669">Yida Zhao</na></co>
<co c="0"><na f="z/Zhao:Yiwen" pid="122/3760">Yiwen Zhao</na></co>
<co c="0"><na f="z/Zhao:Zihan" pid="216/4838">Zihan Zhao</na></co>
<co c="0"><na f="z/Zheng_0007:Bo" pid="33/1610-7">Bo Zheng 0007</na></co>
<co c="0"><na f="z/Zheng:Sipeng" pid="251/3691">Sipeng Zheng</na></co>
<co c="0"><na f="z/Zheng:Thomas_Fang" pid="53/2843">Thomas Fang Zheng</na></co>
<co c="0"><na f="z/Zheng:Weihao" pid="193/7989">Weihao Zheng</na></co>
<co c="0"><na f="z/Zheng:Zhicheng" pid="50/9014">Zhicheng Zheng</na></co>
<co c="0"><na f="z/Zhou:Hanzhang" pid="295/8180">Hanzhang Zhou</na></co>
<co c="0"><na f="z/Zhou:Hongyu" pid="76/7825">Hongyu Zhou</na></co>
<co c="0"><na f="z/Zhou:Jin" pid="98/2691">Jin Zhou</na></co>
<co c="0"><na f="z/Zhou_0001:Jingren" pid="84/2644-1">Jingren Zhou 0001</na></co>
<co c="0"><na f="z/Zhou:Li" pid="54/40">Li Zhou</na></co>
<co c="0"><na f="z/Zhou:XuanQi" pid="399/8327">XuanQi Zhou</na></co>
<co c="0"><na f="z/Zhu:Donglai" pid="70/5222">Donglai Zhu</na></co>
<co c="0"><na f="z/Zhu_0001:Jun" pid="50/2644-1">Jun Zhu 0001</na></co>
<co c="0"><na f="z/Zisserman:Andrew" pid="z/AndrewZisserman">Andrew Zisserman</na></co>
</coauthors>
</dblpperson>

