<?xml version="1.0" encoding="UTF-8"?>
<article xmlns:xlink="http://www.w3.org/1999/xlink" xml:lang="en" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance">
  <front>
    <journal-meta>
      <journal-id journal-id-type="ojs">JDF</journal-id>
      <journal-title-group>
        <journal-title xml:lang="en">Journal of Digital Frontier</journal-title>
      </journal-title-group>
      <publisher>
        <publisher-name>Digital Intelligence Press</publisher-name>
        <publisher-loc>
          <country>HK</country>
          <uri>https://dipscie.com/index.php/JDF/index</uri>
        </publisher-loc>
      </publisher>
      <issn pub-type="epub">3135-7695</issn>
      <self-uri xlink:href="https://dipscie.com/index.php/JDF"/>
    </journal-meta>
    <article-meta>
      <article-id pub-id-type="publisher-id">22</article-id>
      <article-categories>
        <subj-group xml:lang="en" subj-group-type="heading">
          <subject>Articles</subject>
        </subj-group>
      </article-categories>
      <title-group>
        <article-title xml:lang="en">&lt;bold&gt;Computational Modeling of Soundscape and Visual Imagery in Beijing Old City Communities Based on Multi-modal Data&lt;/bold&gt;</article-title>
      </title-group>
      <contrib-group content-type="author">
        <contrib>
          <name-alternatives>
            <name name-style="western" specific-use="primary">
              <surname>Shen</surname>
              <given-names>Ziyan</given-names>
            </name>
          </name-alternatives>
          <email>shenzy@buct.edu.cn</email>
          <xref ref-type="aff" rid="aff-1"/>
        </contrib>
        <contrib>
          <name-alternatives>
            <name name-style="western" specific-use="primary">
              <surname>Qi</surname>
              <given-names>Xuanyi</given-names>
            </name>
          </name-alternatives>
          <email>qixysylvia1995@163.com</email>
          <xref ref-type="aff" rid="aff-2"/>
        </contrib>
        <contrib corresp="yes">
          <name-alternatives>
            <name name-style="western" specific-use="primary">
              <surname>Liu</surname>
              <given-names>Mengyu</given-names>
            </name>
          </name-alternatives>
          <email>daimou971231@naver.com</email>
          <xref ref-type="aff" rid="aff-3"/>
        </contrib>
      </contrib-group>
      <aff id="aff-1">
        <institution content-type="orgname">College of mechanical and electrical engineering, Beijing University of Chemical Technology, Chaoyang District, Beijing, China</institution>
      </aff>
      <aff id="aff-2">
        <institution content-type="orgname">Department of Art and Design, Beijing City University, Shunyi District, Beijing, China</institution>
      </aff>
      <aff id="aff-3">
        <institution content-type="orgname">Digital Media Art Specialty, Department of Art and Design, Beijing City University, Shunyi District, Beijing, China</institution>
      </aff>
      <pub-date date-type="pub" publication-format="epub">
        <day>06</day>
        <month>08</month>
        <year>2026</year>
      </pub-date>
      <fpage>61</fpage>
      <lpage>82</lpage>
      <pub-history>
        <event event-type="received">
          <event-desc>Received: <date date-type="received" iso-8601-date="2026-08-06T12:55:59+00:00"><day>6</day><month>8</month><year>2026</year></date></event-desc>
        </event>
      </pub-history>
      <permissions>
        <copyright-statement>Copyright (c) 2026 Ziyan Shen, Xuanyi Qi, Mengyu Liu (Author)</copyright-statement>
        <copyright-year>2026</copyright-year>
        <copyright-holder>Ziyan Shen, Xuanyi Qi, Mengyu Liu (Author)</copyright-holder>
        <license xlink:href="https://creativecommons.org/licenses/by/4.0">
          <license-p>&lt;a rel="license" href="https://creativecommons.org/licenses/by/4.0/"&gt;&lt;img alt="Creative Commons License" src="//i.creativecommons.org/l/by/4.0/88x31.png" /&gt;&lt;/a&gt;&lt;p&gt;This work is licensed under a &lt;a rel="license" href="https://creativecommons.org/licenses/by/4.0/"&gt;Creative Commons Attribution 4.0 International License&lt;/a&gt;.&lt;/p&gt;</license-p>
        </license>
      </permissions>
      <self-uri xlink:href="https://dipscie.com/index.php/JDF/article/view/10.67541_jdf2603"/>
      <kwd-group xml:lang="en">
        <kwd>Soundscape; Visual imagery; Multimodal learning; Cross-modal correlation; Beijing old city community</kwd>
      </kwd-group>
      <counts>
        <page-count count="22"/>
      </counts>
      <custom-meta-group/>
    </article-meta>
  </front>
  <body/>
  <back>
    <ref-list>
      <ref id="R1">
        <mixed-citation>[1] Wang, S., Zhang, J., Wang, F., &amp; Dong, Y. (2023). How to achieve a balance between functional improvement and heritage conservation? A case study on the renewal of old Beijing city. Sustainable Cities and Society, 98, 104790. https://doi.org/10.1016/j.scs.2023.104790</mixed-citation>
      </ref>
      <ref id="R2">
        <mixed-citation>[2] Zhang, R., Martí Casanovas, M., Bosch González, M., &amp; Sun, S. (2024). Revitalizing heritage: The role of urban morphology in creating public value in China’s historic districts. Land, 13(11), 1919. https://doi.org/10.3390/land13111919</mixed-citation>
      </ref>
      <ref id="R3">
        <mixed-citation>[3] Amir, S., Sadoway, D., &amp; Dommaraju, P. (2023). Taming the noise: Soundscape and livability in a technocratic city-state. East Asian Science, Technology and Society: An International Journal, 17(1), 88-104. https://doi.org/10.1080/18752160.2021.1936749</mixed-citation>
      </ref>
      <ref id="R4">
        <mixed-citation>[4] Gil-Sayas, S., Di Pierro, G., Tansini, A., Serra, S., Currò, D., Broatch, A., &amp; Fontaras, G. (2024). Energy consumption of mobile air-conditioning systems in electrified vehicles under different ambient temperatures. International Journal of Engine Research, 25(2), 293-304. https://dx.doi.org/10.1177/14680874231171303</mixed-citation>
      </ref>
      <ref id="R5">
        <mixed-citation>[5] Bahgat, G., Al-Makhlasawy, R. M., Khairy, M., Nour, M., &amp; Abdelfattah, A. (2026). Energy saving for air conditioning devices based on the amalgamation of intelligent models and occupancy detection. Journal of Electrical Systems and Information Technology, 13(1), 78. https://doi.org/10.1186/s43067-026-00379-1</mixed-citation>
      </ref>
      <ref id="R6">
        <mixed-citation>[6] Zhang, D., Ni, J., &amp; Shi, X. (2024). Study of low-temperature energy consumption optimization of battery electric vehicle air conditioning systems considering blower efficiency. Processes, 12(7), 1495. https://doi.org/10.3390/pr12071495</mixed-citation>
      </ref>
      <ref id="R7">
        <mixed-citation>[7] Muhammad, K., Hussain, T., Ullah, H., Del Ser, J., Rezaei, M., Kumar, N., ... &amp; De Albuquerque, V. H. C. (2022). Vision-based semantic segmentation in scene understanding for autonomous driving: Recent achievements, challenges, and outlooks. IEEE Transactions on Intelligent Transportation Systems, 23(12), 22694-22715. https://doi.org/10.1109/TITS.2022.3207665</mixed-citation>
      </ref>
      <ref id="R8">
        <mixed-citation>[8] Emek Soylu, B., Guzel, M. S., Bostanci, G. E., Ekinci, F., Asuroglu, T., &amp; Acici, K. (2023). Deep-learning-based approaches for semantic segmentation of natural scene images: A review. Electronics, 12(12), 2730. https://doi.org/10.3390/electronics12122730</mixed-citation>
      </ref>
      <ref id="R9">
        <mixed-citation>[9] Liang, Y., Mitchell, A., Kang, J., &amp; Aletta, F. (2026). A Review of Soundscape Datasets: Challenges and Prospects for Multimodal Research. IEEE Transactions on Affective Computing. https://doi.org/10.1109/TAFFC.2026.3659084</mixed-citation>
      </ref>
      <ref id="R10">
        <mixed-citation>[10] Chen, M., Lin, Z., Song, X., Luo, Y., Duan, X., Li, S., &amp; Xie, C. (2026). Theoretical Framework, Technical Evolution, and Future Prospects of Cross-Modal Mapping and Controllable Image Generation Under Multi-Source Heterogeneous Collaboration. Sensors (Basel, Switzerland), 26(10), 2972. https://doi.org/10.3390/s26102972</mixed-citation>
      </ref>
      <ref id="R11">
        <mixed-citation>[11] Kothinti, S. R., &amp; Elhilali, M. (2023). Are acoustics enough? Semantic effects on auditory salience in natural scenes. Frontiers in Psychology, 14, 1276237. https://doi.org/10.3389/fpsyg.2023.1276237</mixed-citation>
      </ref>
      <ref id="R12">
        <mixed-citation>[12] Sun, J., Deng, L., Afouras, T., Owens, A., &amp; Davis, A. (2023). Eventfulness for interactive video alignment. ACM Transactions on Graphics (TOG), 42(4), 1-10. https://doi.org/10.1145/3592118</mixed-citation>
      </ref>
      <ref id="R13">
        <mixed-citation>[13] Zhuang, Y., Kang, Y., Fei, T., Bian, M., &amp; Du, Y. (2024). From hearing to seeing: Linking auditory and visual place perceptions with soundscape-to-image generative artificial intelligence. Computers, Environment and Urban Systems, 110, 102122. https://doi.org/10.1016/j.compenvurbsys.2024.102122</mixed-citation>
      </ref>
      <ref id="R14">
        <mixed-citation>[14] Chen, P., Huang, X., Fei, T., &amp; Wang, S. (2026). Cross‐Modal Urban Sensing: Evaluating Sound–Vision Alignment Across Street‐Level and Aerial Imagery. Transactions in GIS, 30(2), e70246. https://doi.org/10.1111/tgis.70246</mixed-citation>
      </ref>
      <ref id="R15">
        <mixed-citation>[15] Zhang, Y., Wu, M., &amp; Cai, X. (2025). A dynamic cross-modal learning framework for joint text-to-audio grounding and acoustic scene classification in smart city environments. Digital Signal Processing, 167, 105444. https://doi.org/10.1016/j.dsp.2025.105444</mixed-citation>
      </ref>
      <ref id="R16">
        <mixed-citation>[16] Wang, T., Li, F., Zhu, L., Li, J., Zhang, Z., &amp; Shen, H. T. (2025). Cross-modal retrieval: a systematic review of methods and future directions. Proceedings of the IEEE, 112(11), 1716-1754. https://doi.org/10.48550/arXiv.2308.14263</mixed-citation>
      </ref>
      <ref id="R17">
        <mixed-citation>[17] Yang, Y., Wang, D., Hu, Y., &amp; Liang, L. (2025, June). Research on Time Synchronization Technology of Non-safety DCS System Based on IEEE1588v2 Timing Protocol. In International Conference on Nuclear Engineering (pp. 163-176). Singapore: Springer Nature Singapore. https://doi.org/10.1007/978-981-95-2921-6_13</mixed-citation>
      </ref>
      <ref id="R18">
        <mixed-citation>[18] Patoli, A. A., &amp; Fortino, G. (2025). FPGA-based system implementation of IEEE 1588 precision time protocol: A review. IEEE Sensors Journal, 25(11), 18624-18642. https://doi.org/10.1109/JSEN.2025.3557277</mixed-citation>
      </ref>
      <ref id="R19">
        <mixed-citation>[19] Iturbe-Martin, Z., Martín-Garín, A., &amp; Casado-Rezola, A. (2026). Soundscape-Informed Urban Planning and Architecture in Historic Centers: A Multi-Layer Method for Soundscape Characterization Applied to Bilbao Old Town. Applied Sciences, 16(8), 3630. https://doi.org/10.3390/app16083630</mixed-citation>
      </ref>
      <ref id="R20">
        <mixed-citation>[20] Versümer, S., Steffens, J., &amp; Weinzierl, S. (2023). Day-to-day loudness assessments of indoor soundscapes: Exploring the impact of loudness indicators, person, and situation. The Journal of the Acoustical Society of America, 153(5), 2956-2956. https://doi.org/10.1121/10.0019413</mixed-citation>
      </ref>
      <ref id="R21">
        <mixed-citation>[21] Chen, Y., Lin, Y., Xu, R., &amp; Vela, P. A. (2023, October). Wdiscood: Out-of-distribution detection via whitened linear discriminant analysis. In 2023 IEEE/CVF International Conference on Computer Vision (ICCV) (pp. 5275-5284). IEEE. https://doi.org/10.48550/arXiv.2303.07543</mixed-citation>
      </ref>
      <ref id="R22">
        <mixed-citation>[22] Gao, S., Yang, K., Shi, H., Wang, K., &amp; Bai, J. (2022). Review on panoramic imaging and its applications in scene understanding. IEEE Transactions on Instrumentation and Measurement, 71, 1-34. https://doi.org/10.1109/TIM.2022.3216675</mixed-citation>
      </ref>
      <ref id="R23">
        <mixed-citation>[23] Taha, K. (2026). Generative AI for multimodal content: a survey with empirical and experimental evaluations. Artificial Intelligence Review, 59(6), 140. https://doi.org/10.1007/s10462-026-11525-6</mixed-citation>
      </ref>
      <ref id="R24">
        <mixed-citation>[24] Yan, T., Zhao, S., Hu, M., Wang, M., Zhang, X., Luo, Z., &amp; Wang, M. (2024). HCL: A hierarchical contrastive learning framework for zero-shot relation extraction. IEEE Transactions on Neural Networks and Learning Systems, 36(3), 5694-5705. https://doi.org/10.1109/TNNLS.2024.3379527</mixed-citation>
      </ref>
      <ref id="R25">
        <mixed-citation>[25] Zhong, B., Wang, P., &amp; Wang, X. (2024, August). Ts-HCL: hierarchical layer-wise contrastive learning for unsupervised domain adaptation on time-series. In Asia-Pacific Web (APWeb) and Web-Age Information Management (WAIM) Joint International Conference on Web and Big Data (pp. 31-45). Singapore: Springer Nature Singapore. https://doi.org/10.1007/978-981-97-7238-4_3</mixed-citation>
      </ref>
      <ref id="R26">
        <mixed-citation>[26] Xie, H., He, Y., Wu, X., &amp; Lu, Y. (2022). Interplay between auditory and visual environments in historic districts: A big data approach based on social media. Environment and Planning B: Urban Analytics and City Science, 49(4), 1245-1265. https://doi.org/10.1177/23998083211059838</mixed-citation>
      </ref>
      <ref id="R27">
        <mixed-citation>[27] Yao, X. W., Zheng, J. M., Li, Q., Yang, K. H., Yao, Z. H., &amp; Shang, Y. H. (2025, September). MSTDFN: Multi-modal Spatio-Temporal Dynamic Fusion Network for Traffic Flow Prediction. In China Conference on Wireless Sensor Networks (pp. 189-202). Singapore: Springer Nature Singapore. https://doi.org/10.1007/978-981-95-9615-7_12</mixed-citation>
      </ref>
      <ref id="R28">
        <mixed-citation>[28] Zhao, T., Chen, G., Suraphee, S., Phoophiwfa, T., &amp; Busababodhin, P. (2025). A hybrid TCN-XGBoost model for agricultural product market price forecasting. PLoS One, 20(5), e0322496. https://doi.org/10.1371/journal.pone.0322496</mixed-citation>
      </ref>
      <ref id="R29">
        <mixed-citation>[29] Zhao, T., Chen, G., Pang, C., &amp; Busababodhin, P. (2025). Application and performance optimization of SLHS-TCN-XGBoost model in power demand forecasting. Comput. Model. Eng. Sci, 143(3), 2883-2917. https://doi.org/10.32604/cmes.2025.066442</mixed-citation>
      </ref>
      <ref id="R30">
        <mixed-citation>[30] Zhao, T., Chen, G., Pang, C., Seenoi, P., Papukdee, N., &amp; Busababodhin, P. (2025). Time-lapse earthquake difference prediction based on physics-informed long short-term memory coupled with interpretability boosting. Journal of Seismic Exploration, 34(3), 25. https://doi.org/10.36922/JSE025310049</mixed-citation>
      </ref>
      <ref id="R31">
        <mixed-citation>[31] Zhao, T., Chen, G., Pang, C., Li, L., &amp; Busababodhin, P. (2026). Forecasting global agricultural trade imbalances using a hybrid deep learning and gradient boosting framework. Discover Computing, 29(1), 443. https://doi.org/10.1007/s10791-026-10367-8</mixed-citation>
      </ref>
    </ref-list>
  </back>
</article>
