Article(id=1212430802814669510, tenantId=1146029695717560320, journalId=1146031591421210625, issueId=1212430797412409505, articleNumber=null, orderNo=21, doi=10.3981/j.issn.1000-7857.2025.05.00059, pmid=null, cstr=null, oa=null, hot=null, price=null, onlineType=0, articleFormat=0, articleType=null, articleTypeStr=research-article, receivedDate=1746979200000, receivedDateStr=2025-05-12, revisedDate=1762099200000, revisedDateStr=2025-11-03, acceptedDate=null, acceptedDateStr=null, onlineDate=1766995629266, onlineDateStr=2025-12-29, pubDate=1764259200000, pubDateStr=2025-11-28, doiRegisterDate=null, doiRegisterDateStr=null, onlineIssueDate=1766764800000, onlineIssueDateStr=2025-12-27, onlineJustAcceptDate=null, onlineJustAcceptDateStr=null, onlineFirstDate=null, onlineFirstDateStr=null, sourceXml=null, magXml=null, createTime=1766995629266, creator=13701087609, updateTime=1774080345048, updator=sys-migrate, issue=Issue{id=1212430797412409505, tenantId=1146029695717560320, journalId=1146031591421210625, year='2025', volume='43', issue='22', pageStart='1', pageEnd='124', issueExtLink='null', onlineDate='null', pubDate='1764259200000', pubDateStr='2025-11-28', beforeIssueId=null, nextIssueId=null, price=null, status=1, issueComplete=1, articleOrder=1, issueType=-1, specialIssue=null, createTime=1766995627976, creator='13701087609', updateTime=1774330566881, updator='13041195026', preIssue=null, nextIssue=null, articleTotal=null, ext={EN=IssueExt(id=1243195761085756072, tenantId=1146029695717560320, journalId=1146031591421210625, issueId=1212430797412409505, language=EN, specialIssueTitle=, coverIllustrator=null, specialIssueEditor=, specialIssueAbout=), CN=IssueExt(id=1243195761085756073, tenantId=1146029695717560320, journalId=1146031591421210625, issueId=1212430797412409505, language=CN, specialIssueTitle=, coverIllustrator=null, specialIssueEditor=, specialIssueAbout=)}, issueFiles=null, downloadFileDto=null}, startPage=98, endPage=107, ext={EN=ArticleExt(id=1212430803221517022, articleId=1212430802814669510, tenantId=1146029695717560320, journalId=1146031591421210625, language=EN, title=Research on depth estimation and portrait segmentation based on diffusion models, columnId=1150494644690366681, journalTitle=Science & Technology Review, columnName=Papers, runingTitle=null, highlight=null, articleAbstract=
While diffusion models have demonstrated remarkable capabilities in generative tasks, their application to visual perception tasks such as depth estimation and portrait segmentation remains underexplored. This paper proposes Diffusion Perception, a unified framework based on diffusion models for high−quality depth estimation and portrait segmentation. By reformulating traditional perception tasks as conditional generation problems, the framework leverages the denoising characteristics of latent diffusion models (LDMs) to optimize prediction results in latent space. The innovative design incorporates three core processing stages: multimodal feature encoding, noise input prediction, and text−controlled feature extraction and reconstruction, enabling the transition of diffusion models from generative paradigms to visual perception task paradigms. Experimental results demonstrate that on our custom depth estimation dataset, the proposed method achieves evaluation metrics of 93.98% Relative Accuracy (RR), 99.61% Plane Estimation Accuracy (Plane), and 93.61% Scene Consistency (Consistence), outperforming existing state−of−the−art depth estimation methods. Furthermore, in portrait segmentation tasks, the method achieves Intersection over Union (IoU) and mean IoU (mIoU) scores of 96.98% and 91.98% respectively, surpassing existing segmentation algorithms. This study provides novel insights into applying diffusion models in visual perception, where their generative paradigm naturally handles prediction uncertainty and is well−suited for robust perception in dynamic environments.
, authors=null, authorsList=Zongbo DONG, Yifan WANG, Lijun WANG, Huchuan LU, authorCompany=null, correspAuthors=Lijun WANG, authorNote=null, correspAuthorsNote=null, copyrightStatement=
All rights reserved. Unauthorized reproduction is prohibited., copyrightOwner=null, extLink=null, articleAbsUrl=null, sourceXml=null, magXml=null, pdfUrl=null, pdf=null, pdfFileSize=null, pdfExtLink=null, richHtmlUrl=null, mobilePdfUrl=null, reviewReport=null, pdfFirstPage=null, abstractGraph=null, abstractGraphContent=null, abstractVideo=null, citation=null, cebUrl=null, magXmlContent=null, mapNumber=null, fund=null), CN=ArticleExt(id=1212430804798575446, articleId=1212430802814669510, tenantId=1146029695717560320, journalId=1146031591421210625, language=CN, title=基于扩散模型的深度估计与人像分割研究, columnId=1146540929516700224, journalTitle=科技导报, columnName=研究论文, runingTitle=null, highlight=null, articleAbstract=
扩散模型在生成式任务中展现出强大的能力,但其在视觉感知任务(如深度估计与人像分割)中的应用仍有待深入探索。提出一种基于扩散模型的统一框架Diffusion Perception,实现高质量深度估计与人像分割。通过将传统感知任务重新定义为条件生成问题,该框架利用潜在扩散模型(LDM)的去噪特性,在潜在空间中优化预测结果。创新性设计3种核心处理阶段:多模态特征编码阶段、噪声输入预测阶段和文本控制特征提取与重建阶段,使扩散模型从生成范式迁移到视觉感知任务范式上。实验表明,这种方法在自建深度估计数据集上对应的评估指标:相对精度(RR)、平面估计精度(plane)和场景一致性(consistence)分别达到了93.98%、99.61%、93.61%,均优于现有先进的深度估计方法。此外,在人像分割任务中,对应的交并比(IOU)与平均交并比(mIOU)分别达到了96.98%、91.98%,均优于现有分割算法。为扩散模型在视觉感知领域的应用提供了新思路,其生成式范式能够自然处理预测不确定性,适用于动态环境下的鲁棒感知任务。
, authors=
, authorsList=董宗博, 王一帆, 王立君, 卢湖川, authorCompany=null, correspAuthors=王立君, authorNote=null, correspAuthorsNote=
, copyrightStatement=
版权所有,未经授权,不得转载。, copyrightOwner=《科技导报》编辑部, extLink=null, articleAbsUrl=null, sourceXml=cOkPhDXIVaXgAnfc0efNWQ==, magXml=cOkPhDXIVaXgAnfc0efNWQ==, pdfUrl=null, pdf=h2cxyhCjJTZWiP3q9iXk0w==, pdfFileSize=1309150, pdfExtLink=null, richHtmlUrl=null, mobilePdfUrl=null, reviewReport=null, pdfFirstPage=null, abstractGraph=DzJhUiFZzstlE1AUt3sGew==, abstractGraphContent=null, abstractVideo=null, citation=null, cebUrl=null, magXmlContent=cmPvXFiaW7bmpInWUDG1gA==, mapNumber=null, fund=null)}, authors=[Author(id=1242146274129027814, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, orderNo=0, firstName=null, middleName=null, lastName=null, nameCn=null, orcid=null, stid=null, country=null, authorPic=null, dead=0, email=1246363088@mail.dlut.edu.cn, emailSecond=null, emailThird=null, correspondingAuthor=0, authorType=1, ext={EN=AuthorExt(id=1242146274196136680, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, authorId=1242146274129027814, language=EN, stringName=Zongbo DONG, firstName=Zongbo, middleName=null, lastName=DONG, prefix=null, suffix=null, authorComment=null, nameInitials=null, affiliation=null, department=null, xref=null, address=Dalian University of Technology School of Future Technology, Dalian 116024, China, bio=null, bioImg=null, bioContent=null, aboutCorrespAuthor=null), CN=AuthorExt(id=1242146274263245545, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, authorId=1242146274129027814, language=CN, stringName=董宗博, firstName=null, middleName=null, lastName=null, prefix=null, suffix=null, authorComment=null, nameInitials=null, affiliation=null, department=null, xref=null, address=大连理工大学未来技术学院,大连 116024, bio={"content":"
董宗博,硕士研究生,研究方向为计算机视觉与深度学习,电子信箱:1246363088@mail.dlut.edu.cn
"}, bioImg=null, bioContent=
董宗博,硕士研究生,研究方向为计算机视觉与深度学习,电子信箱:1246363088@mail.dlut.edu.cn
, aboutCorrespAuthor=null)}, companyList=[AuthorCompany(id=1242146274028364512, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, xref=null, ext=[AuthorCompanyExt(id=1242146274045141729, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, companyId=1242146274028364512, language=EN, country=null, province=null, city=null, postcode=null, companyName=null, departmentName=null, remark=Dalian University of Technology School of Future Technology, Dalian 116024, China), AuthorCompanyExt(id=1242146274066113251, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, companyId=1242146274028364512, language=CN, country=null, province=null, city=null, postcode=null, companyName=null, departmentName=null, remark=大连理工大学未来技术学院,大连 116024)])]), Author(id=1242146274321965803, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, orderNo=1, firstName=null, middleName=null, lastName=null, nameCn=null, orcid=null, stid=null, country=null, authorPic=null, dead=0, email=null, emailSecond=null, emailThird=null, correspondingAuthor=0, authorType=1, ext={EN=AuthorExt(id=1242146274418434798, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, authorId=1242146274321965803, language=EN, stringName=Yifan WANG, firstName=Yifan, middleName=null, lastName=WANG, prefix=null, suffix=null, authorComment=null, nameInitials=null, affiliation=null, department=null, xref=null, address=Dalian University of Technology School of Future Technology, Dalian 116024, China, bio=null, bioImg=null, bioContent=null, aboutCorrespAuthor=null), CN=AuthorExt(id=1242146274472960751, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, authorId=1242146274321965803, language=CN, stringName=王一帆, firstName=null, middleName=null, lastName=null, prefix=null, suffix=null, authorComment=null, nameInitials=null, affiliation=null, department=null, xref=null, address=大连理工大学未来技术学院,大连 116024, bio=null, bioImg=null, bioContent=null, aboutCorrespAuthor=null)}, companyList=[AuthorCompany(id=1242146274028364512, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, xref=null, ext=[AuthorCompanyExt(id=1242146274045141729, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, companyId=1242146274028364512, language=EN, country=null, province=null, city=null, postcode=null, companyName=null, departmentName=null, remark=Dalian University of Technology School of Future Technology, Dalian 116024, China), AuthorCompanyExt(id=1242146274066113251, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, companyId=1242146274028364512, language=CN, country=null, province=null, city=null, postcode=null, companyName=null, departmentName=null, remark=大连理工大学未来技术学院,大连 116024)])]), Author(id=1242146274561041137, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, orderNo=2, firstName=null, middleName=null, lastName=null, nameCn=null, orcid=null, stid=null, country=null, authorPic=null, dead=0, email=ljwang@dlut.edu.cn, emailSecond=null, emailThird=null, correspondingAuthor=1, authorType=1, ext={EN=AuthorExt(id=1242146274628150003, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, authorId=1242146274561041137, language=EN, stringName=Lijun WANG, firstName=Lijun, middleName=null, lastName=WANG, prefix=null, suffix=null, authorComment=null, nameInitials=null, affiliation=null, department=null, xref=
*, address=Dalian University of Technology School of Future Technology, Dalian 116024, China, bio=null, bioImg=null, bioContent=null, aboutCorrespAuthor=null), CN=AuthorExt(id=1242146274691064564, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, authorId=1242146274561041137, language=CN, stringName=王立君, firstName=null, middleName=null, lastName=null, prefix=null, suffix=null, authorComment=null, nameInitials=null, affiliation=null, department=null, xref=
*, address=大连理工大学未来技术学院,大连 116024, bio=null, bioImg=null, bioContent=null, aboutCorrespAuthor=null)}, companyList=[AuthorCompany(id=1242146274028364512, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, xref=null, ext=[AuthorCompanyExt(id=1242146274045141729, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, companyId=1242146274028364512, language=EN, country=null, province=null, city=null, postcode=null, companyName=null, departmentName=null, remark=Dalian University of Technology School of Future Technology, Dalian 116024, China), AuthorCompanyExt(id=1242146274066113251, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, companyId=1242146274028364512, language=CN, country=null, province=null, city=null, postcode=null, companyName=null, departmentName=null, remark=大连理工大学未来技术学院,大连 116024)])]), Author(id=1242146274749784823, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, orderNo=3, firstName=null, middleName=null, lastName=null, nameCn=null, orcid=null, stid=null, country=null, authorPic=null, dead=0, email=null, emailSecond=null, emailThird=null, correspondingAuthor=0, authorType=1, ext={EN=AuthorExt(id=1242146274821087994, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, authorId=1242146274749784823, language=EN, stringName=Huchuan LU, firstName=Huchuan, middleName=null, lastName=LU, prefix=null, suffix=null, authorComment=null, nameInitials=null, affiliation=null, department=null, xref=null, address=Dalian University of Technology School of Future Technology, Dalian 116024, China, bio=null, bioImg=null, bioContent=null, aboutCorrespAuthor=null), CN=AuthorExt(id=1242146274896585467, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, authorId=1242146274749784823, language=CN, stringName=卢湖川, firstName=null, middleName=null, lastName=null, prefix=null, suffix=null, authorComment=null, nameInitials=null, affiliation=null, department=null, xref=null, address=大连理工大学未来技术学院,大连 116024, bio=null, bioImg=null, bioContent=null, aboutCorrespAuthor=null)}, companyList=[AuthorCompany(id=1242146274028364512, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, xref=null, ext=[AuthorCompanyExt(id=1242146274045141729, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, companyId=1242146274028364512, language=EN, country=null, province=null, city=null, postcode=null, companyName=null, departmentName=null, remark=Dalian University of Technology School of Future Technology, Dalian 116024, China), AuthorCompanyExt(id=1242146274066113251, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, companyId=1242146274028364512, language=CN, country=null, province=null, city=null, postcode=null, companyName=null, departmentName=null, remark=大连理工大学未来技术学院,大连 116024)])])], keywords=[Keyword(id=1242146275051774716, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, language=EN, orderNo=1, keyword=diffusion models), Keyword(id=1242146275123077885, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, language=EN, orderNo=2, keyword=depth estimation), Keyword(id=1242146275190186750, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, language=EN, orderNo=3, keyword=portrait segmentation), Keyword(id=1242146275244712703, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, language=EN, orderNo=4, keyword=fully convolutional networks), Keyword(id=1242146275320210176, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, language=EN, orderNo=5, keyword=deep learning), Keyword(id=1242146275395707649, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, language=CN, orderNo=1, keyword=扩散模型), Keyword(id=1242146275483788034, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, language=CN, orderNo=2, keyword=深度估计), Keyword(id=1242146275550896899, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, language=CN, orderNo=3, keyword=人像分割), Keyword(id=1242146275634782980, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, language=CN, orderNo=4, keyword=全卷积神经网络), Keyword(id=1242146275748029189, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, language=CN, orderNo=5, keyword=深度学习)], refs=[Reference(id=1242146278914728733, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[1], rfOrder=0, authorNames=null, journalName=null, refType=null, unstructuredReference=Vaswani A, Shazeer N, Parmar N, et al. Attention is all you need[C]//Advances in Neural Information Processing Systems (NeurIPS). Long Beach, CA, USA: Neural Information Processing Systems Foundation, Inc, 2017: 5998−6008., articleTitle=null, refAbstract=null), Reference(id=1242146278986031902, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[2], rfOrder=1, authorNames=null, journalName=null, refType=null, unstructuredReference=Dosovitskiy A, Beyer L, Kolesnikov A, et al. An image is worth 1 6x16 words: Transformers for image recognition at scale[J]. arXiv: 2020: 2010.11929., articleTitle=null, refAbstract=null), Reference(id=1242146279053140767, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[3], rfOrder=2, authorNames=null, journalName=null, refType=null, unstructuredReference=Krizhevsky A, Sutskever I, Hinton G E. Imagenet classification with deep convolutional neural networks[C]//Advances in Neural Information Processing Systems (NeurIPS). Lake Tahoe, Nevada, USA: Neural Information Processing Systems Foundation, Inc, 2012: 1097−1105., articleTitle=null, refAbstract=null), Reference(id=1242146279128638240, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[4], rfOrder=3, authorNames=null, journalName=null, refType=null, unstructuredReference=Simonyan K, Zisserman A. Very deep convolutional networks for large−scale image recognition[J]. arXiv: 2014: 1409.1556., articleTitle=null, refAbstract=null), Reference(id=1242146279191552801, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[5], rfOrder=4, authorNames=null, journalName=null, refType=null, unstructuredReference=He K M, Zhang X Y, Ren S Q, et al. Deep residual learning for image recognition[C]//Proceedings of IEEE Conference on Computer Vision and Pattern Recognition (CVPR). Piscataway, NJ: IEEE, 2016: 770−778., articleTitle=null, refAbstract=null), Reference(id=1242146279271244578, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[6], rfOrder=5, authorNames=null, journalName=null, refType=null, unstructuredReference=Huang G, Liu Z, Van Der Maaten L, et al. Densely connected convolutional networks[C]//Proceedings of IEEE Conference on Computer Vision and Pattern Recognition (CVPR). Piscataway, NJ: IEEE, 2017: 2261−2269., articleTitle=null, refAbstract=null), Reference(id=1242146279338353445, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[7], rfOrder=6, authorNames=null, journalName=null, refType=null, unstructuredReference=Tan M, Le Q. Efficientnet: Rethinking model scaling for convolutional neural networks[C]//International Conference on Machine Learning. California: PMLR, 2019: 6105−6114., articleTitle=null, refAbstract=null), Reference(id=1242146279397073702, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[8], rfOrder=7, authorNames=null, journalName=null, refType=null, unstructuredReference=Liu Z, Lin Y T, Cao Y, et al. Swin transformer: Hierarchical vision transformer using shifted windows[C]//Proceedings of IEEE/CVF International Conference on Computer Vision (ICCV). Piscataway, NJ: IEEE, 2021: 9992−10002., articleTitle=null, refAbstract=null), Reference(id=1242146279493542695, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=2017, volume=39, issue=6, pageStart=1137, pageEnd=1149, url=null, language=null, rfNumber=[9], rfOrder=8, authorNames=Ren S Q, He K M, Girshick R, journalName=IEEE Transactions on Pattern Analysis and Machine Intelligence, refType=null, unstructuredReference=
Ren S Q,
He K M,
Girshick R,
et al. Faster R−CNN: Towards real−time object detection with region proposal networks[J].
IEEE Transactions on Pattern Analysis and Machine Intelligence,
2017,
39(6): 1137-1149., articleTitle=Faster R−CNN: Towards real−time object detection with region proposal networks, refAbstract=null), Reference(id=1242146279581623081, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[10], rfOrder=9, authorNames=null, journalName=null, refType=null, unstructuredReference=He K M, Gkioxari G, Dollár P, et al. Mask R−CNN[C]//Proceedings of IEEE International Conference on Computer Vision (ICCV). Piscataway, NJ: IEEE, 2017: 2980−2988., articleTitle=null, refAbstract=null), Reference(id=1242146279657120554, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[11], rfOrder=10, authorNames=null, journalName=null, refType=null, unstructuredReference=Cheng B W, Collins M D, Zhu Y K, et al. Panoptic−DeepLab: A simple, strong, and fast baseline for bottom−up panoptic segmentation[C]//Proceedings of IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR). Piscataway, NJ: IEEE, 2020: 12475−12485., articleTitle=null, refAbstract=null), Reference(id=1242146279715840812, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[12], rfOrder=11, authorNames=null, journalName=null, refType=null, unstructuredReference=He K M, Chen X L, Xie S N, et al. Masked autoencoders are scalable vision learners[C]//Proceedings of IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR). Piscataway, NJ: IEEE, 2022: 15979−15988., articleTitle=null, refAbstract=null), Reference(id=1242146279795532589, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[13], rfOrder=12, authorNames=null, journalName=null, refType=null, unstructuredReference=Caron M, Touvron H, Misra I, et al. Emerging properties in self−supervised vision transformers[C]//Proceedings of IEEE/CVF International Conference on Computer Vision (ICCV). Piscataway, NJ: IEEE, 2021: 9630−9640., articleTitle=null, refAbstract=null), Reference(id=1242146279871030062, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=2020, volume=2, issue=null, pageStart=665, pageEnd=673, url=null, language=null, rfNumber=[14], rfOrder=13, authorNames=Geirhos R, Jacobsen J H, Michaelis C, journalName=Nature Machine Intelligence, refType=null, unstructuredReference=
Geirhos R,
Jacobsen J H,
Michaelis C,
et al. Shortcut learning in deep neural networks[J].
Nature Machine Intelligence,
2020,
2: 665-673., articleTitle=Shortcut learning in deep neural networks, refAbstract=null), Reference(id=1242146279967499055, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[15], rfOrder=14, authorNames=null, journalName=null, refType=null, unstructuredReference=Donahue J, Krähenbühl P, Darrell T. Adversarial feature learning[J]. arXiv: 2016: 1605.09782., articleTitle=null, refAbstract=null), Reference(id=1242146280038802224, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=2020, volume=33, issue=null, pageStart=6840, pageEnd=6851, url=null, language=null, rfNumber=[16], rfOrder=15, authorNames=Ho J, Jain A, Abbeel P, journalName=Advances in Neural Information Processing Systems (NeurIPS), refType=null, unstructuredReference=
Ho J,
Jain A,
Abbeel P. Denoising diffusion probabilistic models[J].
Advances in Neural Information Processing Systems (NeurIPS),
2020,
33: 6840-6851., articleTitle=Denoising diffusion probabilistic models, refAbstract=null), Reference(id=1242146280168825650, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[17], rfOrder=16, authorNames=null, journalName=null, refType=null, unstructuredReference=Xie S N, Girshick R, Dollár P, et al. Aggregated residual transformations for deep neural networks[C]//Proceedings of IEEE Conference on Computer Vision and Pattern Recognition (CVPR). Piscataway, NJ: IEEE, 2017: 5987−5995., articleTitle=null, refAbstract=null), Reference(id=1242146280290460469, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[18], rfOrder=17, authorNames=null, journalName=null, refType=null, unstructuredReference=Wang W H, Xie E Z, Li X, et al. Pyramid vision transformer: A versatile backbone for dense prediction without convolutions[C]//Proceedings of IEEE/CVF International Conference on Computer Vision (ICCV). Piscataway, NJ: IEEE, 2021: 548−558., articleTitle=null, refAbstract=null), Reference(id=1242146280361763638, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[19], rfOrder=18, authorNames=null, journalName=null, refType=null, unstructuredReference=Carion N, Massa F, Synnaeve G, et al. End−to−end object detection with transformers[M]//Computer Vision – ECCV 2020. Cham: Springer International Publishing, 2020: 213−229., articleTitle=null, refAbstract=null), Reference(id=1242146280466621239, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[20], rfOrder=19, authorNames=null, journalName=null, refType=null, unstructuredReference=Rombach R, Blattmann A, Lorenz D, et al. High−resolution image synthesis with latent diffusion models[C]//Proceedings of IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR). Piscataway, NJ: IEEE, 2022: 10674−10685., articleTitle=null, refAbstract=null), Reference(id=1242146280542118712, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[21], rfOrder=20, authorNames=null, journalName=null, refType=null, unstructuredReference=Ronneberger O, Fischer P, Brox T. U−Net: Convolutional networks for biomedical image segmentation[M]//Medical Image Computing and Computer−Assisted Intervention – MICCAI 2015. Cham: Springer International Publishing, 2015: 234−241., articleTitle=null, refAbstract=null), Reference(id=1242146280609227577, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[22], rfOrder=21, authorNames=null, journalName=null, refType=null, unstructuredReference=Ke B X, Obukhov A, Huang S Y, et al. Repurposing diffusion−based image generators for monocular depth estimation[C]//Proceedings of IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR). Piscataway, NJ: IEEE, 2024: 9492−9502., articleTitle=null, refAbstract=null), Reference(id=1242146280688919354, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[23], rfOrder=22, authorNames=null, journalName=null, refType=null, unstructuredReference=He J, Li H, Yin W, et al. Lotus: Diffusion−based visual foundation model for high−quality dense prediction[J]. arXiv: 2024: 2409.18124., articleTitle=null, refAbstract=null), Reference(id=1242146280772805435, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[24], rfOrder=23, authorNames=null, journalName=null, refType=null, unstructuredReference=Zhao W L, Rao Y M, Liu Z Y, et al. Unleashing text−to−image diffusion models for visual perception[C]//Proceedings of IEEE/CVF International Conference on Computer Vision (ICCV). Piscataway, NJ: IEEE, 2023: 5706−5716., articleTitle=null, refAbstract=null), Reference(id=1242146280835719997, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[25], rfOrder=24, authorNames=null, journalName=null, refType=null, unstructuredReference=Wu W, Zhao Y, Chen H, et al. Datasetdm: Synthesizing data with perception annotations using diffusion models[C]//Advances in Neural Information Processing Systems (NeurIPS). New Orleans, Louisiana, USA: Neural Information Processing Systems Foundation, Inc, 2023: 54683−54695., articleTitle=null, refAbstract=null), Reference(id=1242146280953160510, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[26], rfOrder=25, authorNames=null, journalName=null, refType=null, unstructuredReference=Kondapaneni N, Marks M, Knott M, et al. Text−image alignment for diffusion−based perception[C]//Proceedings of IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR). Piscataway, NJ: IEEE, 2024: 13883−13893., articleTitle=null, refAbstract=null), Reference(id=1242146281045435201, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[27], rfOrder=26, authorNames=null, journalName=null, refType=null, unstructuredReference=Esser P, Rombach R, Ommer B. Taming transformers for high−resolution image synthesis[C]//Proceedings of IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR). IEEE, 2021: 12873−12883., articleTitle=null, refAbstract=null), Reference(id=1242146281146098499, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[28], rfOrder=27, authorNames=null, journalName=null, refType=null, unstructuredReference=Radford A, Kim J W, Hallacy C, et al. Learning transferable visual models from natural language supervision[C]//International Conference on Machine Learning. Virtual Event: PMLR, 2021: 8748−8763., articleTitle=null, refAbstract=null), Reference(id=1242146281234178886, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[29], rfOrder=28, authorNames=null, journalName=null, refType=null, unstructuredReference=Song J, Meng C, Ermon S. Denoising diffusion implicit models[J]. arXiv: 2020: 2010.02502., articleTitle=null, refAbstract=null), Reference(id=1242146281305482055, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[30], rfOrder=29, authorNames=null, journalName=null, refType=null, unstructuredReference=Ramesh A, Dhariwal P, Nichol A, et al. Hierarchical text−conditional image generation with clip latents[J]. arXiv: 2022: 2204.06125., articleTitle=null, refAbstract=null), Reference(id=1242146282794459978, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[31], rfOrder=30, authorNames=null, journalName=null, refType=null, unstructuredReference=Roberts M, Ramapuram J, Ranjan A, et al. Hypersim: A photorealistic synthetic dataset for holistic indoor scene understanding[C]//Proceedings of IEEE/CVF International Conference on Computer Vision (ICCV). Piscataway, NJ: IEEE, 2021: 10892−10902., articleTitle=null, refAbstract=null), Reference(id=1242146282861568843, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[32], rfOrder=31, authorNames=null, journalName=null, refType=null, unstructuredReference=Cabon Y, Murray N, Humenberger M. Virtual kitti 2[J]. arXiv: 2020: 2001.10773., articleTitle=null, refAbstract=null), Reference(id=1242146282953843532, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[33], rfOrder=32, authorNames=null, journalName=null, refType=null, unstructuredReference=Feng J S, Huang Z L, Kang B Y, et al. Depth anything V2[C]//Proceedings of Advances in Neural Information Processing Systems 37. Vancouver: Neural Information Processing Systems Foundation, Inc. (NeurIPS), 2024: 21875−21911., articleTitle=null, refAbstract=null), Reference(id=1242146283016758093, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[34], rfOrder=33, authorNames=null, journalName=null, refType=null, unstructuredReference=Bochkovskii A, Delaunoy A Ã Ģ, Germain H, et al. Depth pro: Sharp monocular metric depth in less than a second[J]arXiv: 2024: 2410.02073., articleTitle=null, refAbstract=null), Reference(id=1242146283096449870, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[35], rfOrder=34, authorNames=null, journalName=null, refType=null, unstructuredReference=Zhang X, Ke B, Riemenschneider H, et al. Betterdepth: Plug−and−play diffusion refiner for zero−shot monocular depth estimation[J]. arXiv: 2024: 2407.17952., articleTitle=null, refAbstract=null), Reference(id=1242146283159364431, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[36], rfOrder=35, authorNames=null, journalName=null, refType=null, unstructuredReference=Kirillov A, Mintun E, Ravi N, et al. Segment anything[C]//Proceedings of IEEE/CVF International Conference on Computer Vision (ICCV). Piscataway, NJ: IEEE, 2023: 3992−4003., articleTitle=null, refAbstract=null), Reference(id=1242146283234861907, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, doi=null, pmid=null, pmcid=null, year=null, volume=null, issue=null, pageStart=null, pageEnd=null, url=null, language=null, rfNumber=[37], rfOrder=36, authorNames=null, journalName=null, refType=null, unstructuredReference=Oquab M, Darcet T, Moutakanni T, et al. Dinov2: Learning robust visual features without supervision[J]. arXiv: 2023: 2304.07193., articleTitle=null, refAbstract=null)], funds=[Fund(id=1242146278616933144, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, awardId=62276045, language=CN, fundingSource=国家自然科学基金(62276045), fundOrder=null, country=null), Fund(id=1242146278675653402, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, awardId=62422610, language=CN, fundingSource=国家自然科学基金(62422610), fundOrder=null, country=null), Fund(id=1242146278746956572, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, awardId=U23A20386, language=CN, fundingSource=国家自然科学基金(U23A20386), fundOrder=null, country=null)], companyList=[AuthorCompany(id=1242146274028364512, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, xref=null, ext=[AuthorCompanyExt(id=1242146274045141729, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, companyId=1242146274028364512, language=EN, country=null, province=null, city=null, postcode=null, companyName=null, departmentName=null, remark=Dalian University of Technology School of Future Technology, Dalian 116024, China), AuthorCompanyExt(id=1242146274066113251, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, companyId=1242146274028364512, language=CN, country=null, province=null, city=null, postcode=null, companyName=null, departmentName=null, remark=大连理工大学未来技术学院,大连 116024)])], figs=[ArticleFig(id=1242146276016464646, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, language=EN, label=null, caption=null, figureFileSmall=ZngYnpuKSgWajMedjTJtdA==, figureFileBig=EUWEUnP3Cyob2xQuQQCxDQ==, tableContent=null), ArticleFig(id=1242146276079379207, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, language=CN, label=图1, caption=
基于Transformer的编码器与解码器架构, figureFileSmall=ZngYnpuKSgWajMedjTJtdA==, figureFileBig=EUWEUnP3Cyob2xQuQQCxDQ==, tableContent=null), ArticleFig(id=1242146276205208329, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, language=EN, label=null, caption=null, figureFileSmall=ML9+6DyvBaElyFjz0wjqTA==, figureFileBig=FA6638RhSAAy3sgvx3lMJA==, tableContent=null), ArticleFig(id=1242146276280705802, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, language=CN, label=图2, caption=
Stable Diffusion模型架构, figureFileSmall=ML9+6DyvBaElyFjz0wjqTA==, figureFileBig=FA6638RhSAAy3sgvx3lMJA==, tableContent=null), ArticleFig(id=1242146276343620363, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, language=EN, label=null, caption=null, figureFileSmall=UHyP58lNeNY6EDVHeXKaTQ==, figureFileBig=iMYIcSdJNh9cmPDD4mWVaQ==, tableContent=null), ArticleFig(id=1242146276452672268, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, language=CN, label=图3, caption=
Diffusion Perception算法训练流程图, figureFileSmall=UHyP58lNeNY6EDVHeXKaTQ==, figureFileBig=iMYIcSdJNh9cmPDD4mWVaQ==, tableContent=null), ArticleFig(id=1242146276523975437, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, language=EN, label=null, caption=null, figureFileSmall=uaRRbeJNO87gV/Lg907vHQ==, figureFileBig=DUoVEbiS5EE3QhlyQtL+TQ==, tableContent=null), ArticleFig(id=1242146276582695694, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, language=CN, label=图4, caption=
Diffusion Perception算法推理流程, figureFileSmall=uaRRbeJNO87gV/Lg907vHQ==, figureFileBig=DUoVEbiS5EE3QhlyQtL+TQ==, tableContent=null), ArticleFig(id=1242146276691747599, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, language=EN, label=null, caption=null, figureFileSmall=null, figureFileBig=null, tableContent=
| 数据集类型 | 数据集名称 | 总数据量 | 训练集 | 测试集 | 场景覆盖 |
| 深度估计 | Hypersim | 54000 | 43200 | 10800 | 室内场景 |
| Virtual KITTI | 20000 | 16000 | 4000 | 室外场景 |
| 人像分割 | 构建数据集 | 65000 | 52000 | 13000 | 室内外多光照 |
), ArticleFig(id=1242146276754662160, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, language=CN, label=表1, caption=
深度估计与人像分割数据集
, figureFileSmall=null, figureFileBig=null, tableContent=
| 数据集类型 | 数据集名称 | 总数据量 | 训练集 | 测试集 | 场景覆盖 |
| 深度估计 | Hypersim | 54000 | 43200 | 10800 | 室内场景 |
| Virtual KITTI | 20000 | 16000 | 4000 | 室外场景 |
| 人像分割 | 构建数据集 | 65000 | 52000 | 13000 | 室内外多光照 |
), ArticleFig(id=1242146276817576721, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, language=EN, label=null, caption=null, figureFileSmall=null, figureFileBig=null, tableContent=
| 模型 | 测试项目/% |
| RR | Plane | Consistence |
| Marigold | 88.89 | 99.06 | 88.72 |
| Lotus | 83.91 | 97.30 | 87.29 |
| Depth Anything v2 | 81.73 | 98.74 | 85.82 |
| DepthPro | 83.47 | 98.63 | 86.29 |
| BetterDepth | 85.42 | 99.38 | 86.67 |
| Diffusion Perception | 93.98 | 99.61 | 93.61 |
), ArticleFig(id=1242146276884685586, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, language=CN, label=表2, caption=
不同方法深度估计结果
, figureFileSmall=null, figureFileBig=null, tableContent=
| 模型 | 测试项目/% |
| RR | Plane | Consistence |
| Marigold | 88.89 | 99.06 | 88.72 |
| Lotus | 83.91 | 97.30 | 87.29 |
| Depth Anything v2 | 81.73 | 98.74 | 85.82 |
| DepthPro | 83.47 | 98.63 | 86.29 |
| BetterDepth | 85.42 | 99.38 | 86.67 |
| Diffusion Perception | 93.98 | 99.61 | 93.61 |
), ArticleFig(id=1242146278369469206, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, language=EN, label=null, caption=null, figureFileSmall=null, figureFileBig=null, tableContent=
| 模型 | 测试项目/% |
| IoU | mIoU |
| SAM | 95.85 | 90.60 |
| Dino v2 | 94.78 | 84.70 |
| ViT | 93.33 | 80.76 |
| Diffusion Perception | 96.98 | 91.98 |
), ArticleFig(id=1242146278465938199, tenantId=1146029695717560320, journalId=1146031591421210625, articleId=1212430802814669510, language=CN, label=表3, caption=
深度估计与人像分割数据集
, figureFileSmall=null, figureFileBig=null, tableContent=
| 模型 | 测试项目/% |
| IoU | mIoU |
| SAM | 95.85 | 90.60 |
| Dino v2 | 94.78 | 84.70 |
| ViT | 93.33 | 80.76 |
| Diffusion Perception | 96.98 | 91.98 |
)], attaches=null, journal=Journal(id=1125356956822126595, delFlag=0, nameCn=科技导报, nameEn=Science & Technology Review, nameHistory1=null, nameHistory2=null, issn=1000-7857, eissn=, cn=11-1421/N, coden=null, periodic=3, language=CN, oaType=0, ccby=null, superviseOffice=null, ownerOffice=null, pubOffice=null, editorOffice=null, officeType=null, aims=null, clcCode=null, officeProv=null, officeCity=null, officeAddr=null, officeZip=null, officeEmail=null, officePhone=null, editDirector=null, officeDirector=null, officeDirectorPhone=null, officeStaffNum=null, officeEmpNum=null, coverPicUrl=wfghvu3bhh/dKxuZ+ucVHA==, journalPrice=null, startedYear=null, abbrevIsoEn=Sci Technol Rev, journalRemark=null, publicationField=null, createdTime=null, updatedTime=1784015846012, createdBy=null, updatedBy=13041195026, firstLetterCn=K, firstLetterEn=K, subjectCode=Natural Sciences, subjectName=自然科学, subjectCodeEn=Natural Sciences, subjectNameEn=null, picCn=wfghvu3bhh/dKxuZ+ucVHA==, picEn=yjSfclmpNm7ihn9NbTZ69g==, jcr=null, cjcr=null, exts=[JournalExt(id=1283818766098219763, language=CN, name=科技导报, nameHistory1=null, nameHistory2=null, managedBy=中国科学技术协会, sponsoredBy=中国科学技术协会, publishedBy=科技导报社, editorOffice=, officeProv=null, officeCity=null, officeAddr=, officeZip=, editDirector=, officeDirector=null, officePhone=null, coverPicUrl=null, journalRemark=, submitArticleUrl=null, websiteUrl=http://www.kjdb.org/CN/home, createdTime=1784015846037, updatedTime=1784015846037, createdBy=13041195026, updatedBy=13041195026, submissionGuidelinesUrl=http://www.kjdb.org/CN/column/column7.shtml, submissionAuthorUrl=https://kjdbauthor.cast.org.cn/webm, submissionEditorUrl=https://kjdbeditor.cast.org.cn/webm/, submissionReviewUrl=https://kjdbauthor.cast.org.cn/webm, submissionCeEditorUrl=https://kjdbeditor.cast.org.cn/webm/, submissionAeEditorUrl=https://kjdbeditor.cast.org.cn/webm/, option={"copyright":""}), JournalExt(id=1283818766144357108, language=EN, name=Science & Technology Review, nameHistory1=null, nameHistory2=null, managedBy=, sponsoredBy=, publishedBy=, editorOffice=, officeProv=null, officeCity=null, officeAddr=, officeZip=, editDirector=, officeDirector=null, officePhone=null, coverPicUrl=null, journalRemark=, submitArticleUrl=null, websiteUrl=http://www.kjdb.org/EN/home, createdTime=1784015846048, updatedTime=1784015846048, createdBy=13041195026, updatedBy=13041195026, submissionGuidelinesUrl=http://www.kjdb.org/EN/column/column7.shtml, submissionAuthorUrl=https://kjdbauthor.manuscriptcloud.com/login, submissionEditorUrl=https://kjdbeditor.manuscriptcloud.com/login, submissionReviewUrl=https://kjdbauthor.manuscriptcloud.com/login, submissionCeEditorUrl=https://kjdbeditor.manuscriptcloud.com/login, submissionAeEditorUrl=https://kjdbeditor.manuscriptcloud.com/login, option={"copyright":""})], databaseList=null, tenantJournalId=1146031591421210625, websiteList=[Website(id=1146104741081231361, webName=null, webTitle=null, webDomain=null, webCopyrigh=null, webIpcNo=null, seoTitle=null, seoKeywords=null, seoDescription=null, tenantJournalId=null, journalId=1146031591421210625, journalNameCn=null, journalNameEn=null, grayFlag=null, tenantId=1146029695717560320, platformId=null, journalGroupId=null, journalGroupNameCn=null, journalGroupNameEn=null, type=1, domain=https://castjournals.cast.org.cn/joweb/kjdb/CN, language=CN, createTime=1751182263881, createBy=18614031015, updateTime=1751778001962, updateBy=18614031015, name=科技导报, tplId=1146099689490845704, title=科技导报, delFlag=0, indexPage=/home, props=[WebsiteProps(id=1148021146403992296, tenantId=1146029695717560320, journalId=null, journalGroupId=null, siteId=1146104741081231361, code=articleTextType, value=kx, createTime=1751639170504, updateTime=1751639170504, creator=18614031015, updator=18614031015), WebsiteProps(id=1148021146378826469, tenantId=1146029695717560320, journalId=null, journalGroupId=null, siteId=1146104741081231361, code=banner, value=null, createTime=1751639170498, updateTime=1751639170498, creator=18614031015, updator=18614031015), WebsiteProps(id=1148021146366243556, tenantId=1146029695717560320, journalId=null, journalGroupId=null, siteId=1146104741081231361, code=logo, value=https://castjournals.cast.org.cn/joweb/kjdb/CN/file/pic?fileId=9GHSf7eGlIPH0Tv/OOdstA==, createTime=1751639170495, updateTime=1751639170495, creator=18614031015, updator=18614031015), WebsiteProps(id=1148021146395603687, tenantId=1146029695717560320, journalId=null, journalGroupId=null, siteId=1146104741081231361, code=picServerUrl, value=https://castjournals.cast.org.cn/joweb/kjdb/CN/file/pic, createTime=1751639170502, updateTime=1751639170502, creator=18614031015, updator=18614031015), WebsiteProps(id=1148021146387215078, tenantId=1146029695717560320, journalId=null, journalGroupId=null, siteId=1146104741081231361, code=staticResourcePath, value=https://castjournals.cast.org.cn/joweb/cast_kjdb_cn_619/, createTime=1751639170500, updateTime=1751639170500, creator=18614031015, updator=18614031015)]), Website(id=1146105254833139715, webName=null, webTitle=null, webDomain=null, webCopyrigh=null, webIpcNo=null, seoTitle=null, seoKeywords=null, seoDescription=null, tenantJournalId=null, journalId=1146031591421210625, journalNameCn=null, journalNameEn=null, grayFlag=null, tenantId=1146029695717560320, platformId=null, journalGroupId=null, journalGroupNameCn=null, journalGroupNameEn=null, type=1, domain=https://castjournals.cast.org.cn/joweb/kjdb/EN, language=EN, createTime=1751182386363, createBy=18614031015, updateTime=1753500121937, updateBy=18614031015, name=科技导报, tplId=1146101810881728533, title=Science & Technology Review, delFlag=0, indexPage=/home, props=[WebsiteProps(id=1155838567709528217, tenantId=1146029695717560320, journalId=null, journalGroupId=null, siteId=1146105254833139715, code=articleTextType, value=kx, createTime=1753502988984, updateTime=1753502988984, creator=18614031015, updator=18614031015), WebsiteProps(id=1155838567692750998, tenantId=1146029695717560320, journalId=null, journalGroupId=null, siteId=1146105254833139715, code=banner, value=null, createTime=1753502988980, updateTime=1753502988980, creator=18614031015, updator=18614031015), WebsiteProps(id=1155838567688556693, tenantId=1146029695717560320, journalId=null, journalGroupId=null, siteId=1146105254833139715, code=logo, value=https://castjournals.cast.org.cn/joweb/kjdb/EN/file/pic?fileId=9GHSf7eGlIPH0Tv/OOdstA==, createTime=1753502988979, updateTime=1753502988979, creator=18614031015, updator=18614031015), WebsiteProps(id=1155838567705333912, tenantId=1146029695717560320, journalId=null, journalGroupId=null, siteId=1146105254833139715, code=picServerUrl, value=https://castjournals.cast.org.cn/joweb/kjdb/EN/file/pic, createTime=1753502988983, updateTime=1753502988983, creator=18614031015, updator=18614031015), WebsiteProps(id=1155838567701139607, tenantId=1146029695717560320, journalId=null, journalGroupId=null, siteId=1146105254833139715, code=staticResourcePath, value=https://castjournals.cast.org.cn/joweb/cast_kjdb_en_623/, createTime=1753502988982, updateTime=1753502988982, creator=18614031015, updator=18614031015)])], journalTitle=科技导报, weixinUrl=null, journalUrl=null, iacademicId=null, status=1, seqNo=null, journalTitleEn=Science & Technology Review, journalPhotoCn=wfghvu3bhh/dKxuZ+ucVHA==, journalPhotoEn=yjSfclmpNm7ihn9NbTZ69g==, journalFirstLetter=K, journalRecommend=null, journalNew=null, journalCollection=1, jcrJf=null, cjcrJf=0.91, jcrJfStr=null, cjcrJfStr=null, submissionFirstDecision=null, sciSubjectClassification=null, casSubjectClassification=null, citeScore=null, totalCitationFrequency=null, icpCode=null, psCode=null, advertisingLicenseCode=null, copyrightInformation=null, country=null, option=, provinceCode=null, provinceName=null, collectFlag=false, interPubPlatform=, interPubPlatformUrl=null), detailUrlCn=https://castjournals.cast.org.cn/joweb/kjdb/CN/10.3981/j.issn.1000-7857.2025.05.00059, detailUrlEn=https://castjournals.cast.org.cn/joweb/kjdb/EN/10.3981/j.issn.1000-7857.2025.05.00059, pdfUrlCn=https://castjournals.cast.org.cn/joweb/kjdb/CN/PDF/10.3981/j.issn.1000-7857.2025.05.00059, pdfUrlEn=https://castjournals.cast.org.cn/joweb/kjdb/EN/PDF/10.3981/j.issn.1000-7857.2025.05.00059, aliStartDate=null, aliEndDate=null, collectionFlag=false, citedCount=null, citedUrl=null, previewStatus=0, delFlag=0, hasFullText=1, orderTime=1764259200000, fullTextJson=null, articleText=null, reference=null)