Wei Rao

21 papers C 2Journal 10Unranked 9
YearRankTypeTitle / Venue / Authors
2026 J jnl
J. Circuits Syst. Comput.
Wei Rao, Hongbo Wang, Yibo Cui, Haoqian Chen, Hao Zhou
2026 J jnl
PeerJ Comput. Sci.
Zeyang Tang, Wei Rao, Lei Hu, Qibiao Hu, Yang Chen, Yiwen Li
2026 J jnl
IEEE Access
Zeyang Tang, Xin Mei, Ling Ruan, Wanting Deng, Yibo Cui, Wei Rao, Yiwen Li
2025 J jnl
J. Netw. Comput. Appl.
Hongjian Li, Wei Rao, Baojian Hu, Yu Tian, Jie Shen
2024 J jnl
Qual. Reliab. Eng. Int.
Guoqiang Liu, Zezhou Sun, Wei Rao, Jie Dong
2023 J jnl
IEEE Trans. Instrum. Meas.
Yeqi Hu, Wei Rao, Lin Qi, Junyu Dong, Jinzhen Cai, Hao Fan
2022 J jnl
J. Sensors
Chang Liu, Wei Rao, Jin Wang, Zeyang Tang, Jie Wang, Li Tian, Liang Zhou, Jiangpei Xu
2022 J jnl
IEEE Trans. Geosci. Remote. Sens.
Zhaojin Li, Bo Wu, Wai Chung Liu, Long Chen, Hongliang Li, Jie Dong, Wei Rao, Dong Wang, Qingyu Meng, Jihong Dong
2021 C conf
ASRU
Jiangyu Han, Wei Rao, Yanhua Long, Jiaen Liang
2021 conf
ICAIIS
Zeyang Tang, Wenshuo Wang, Wei Rao, Dongsheng Li, Heng Liu, Youjun Deng, Qian Jiang
2021 conf
NEMS
Dawei Wang, Chennan Lu, Xiaohong Wang, Wei Rao
2020 conf
ISPA/BDCloud/SocialCom/SustainCom
Ruixuan Zhao, Shenghang Liu, Wei Rao, Hao Xu, Chongning Na, Ruoyu Li
2019 C conf
ICDS
Li-Peng Zhu, Wei Rao, Junfeng Qiao, Sen Pan
2019 J jnl
Inf.
Qiuzheng Li, Zuopeng Justin Zhang, Wei Rao, Wenwen Xu, Lijia Jiang
2019 conf
DPTA
Wei Rao, Jian Chen
2017 conf
LSMS/ICSEE (3)
Wei Rao, Jing Jiang, Ming Yang, Wei Peng, Aihua Zhou
2016 conf
DSP
Wei Rao
2009 conf
SSME
Wei Rao, Fei Xia, Yuanyuan Liu, Wen-qun Tan, Hui-jun Xu, Kang-ming Yuan, Jian-Bing Liu
2008 conf
FSKD (2)
Xing-Wang Zhang, Wei Rao
2007 J jnl
Appl. Math. Comput.
Yecai Guo, Wei Rao, Yinge Han
2005 conf
FSKD (1)
Yecai Guo, Wei Rao, Yi Guo, Wei Ma
redb/extractors/pe_extractors/pe_resources.py
← Index redb/extractors/pe_extractors/pe_resources.py python
from hashlib import sha256
import inspect
from datetime import datetime, timezone
from typing import Any

import magic
from magika import Magika
import pefile
from pefile import UnicodeStringWrapperPostProcessor

from redb.extractors.enum import Tag
from redb.extractors.pe_extractor import PEExtractor
from redb.models.dataclasses import PEResource


class PEResourceExtractor(PEExtractor):

    def __init__(
        self,
        filepath,
        log,
        exporters=None,
        index_prefix=None,
        elastic_index=None,
        known_benign=False,
        known_malicious=False,
        pe=None,
    ):
        super().__init__(
            filepath,
            log,
            exporters,
            index_prefix,
            elastic_index,
            known_benign,
            known_malicious,
            pe,
        )
        self.elastic_index = self.index_prefix + "-pe_resources"
        self.log.debug(inspect.currentframe().f_code.co_name)

    def tag(self):
        return Tag.PE_RESOURCE.value

    def _extract_resources(self):
        """
        Returns:
        resources: a list of dictionaries, one per each resources type found.
                    each dictionary the key represents the name of the content,
                    which is the value itself.
                    Empty list if no resources present.
        """
        self.log.debug(inspect.currentframe().f_code.co_name)
        resources_list = []
        try:
            if hasattr(self.pe, "DIRECTORY_ENTRY_RESOURCE"):
                for resource_type in self.pe.DIRECTORY_ENTRY_RESOURCE.entries:
                    # if resource_type.name is not None:
                    #     name = resource_type.name
                    # else:
                    #     name = pefile.RESOURCE_TYPE.get(resource_type.struct.Id)
                    # if not name:
                    #     name = resource_type.struct.Id
                    name = (
                        resource_type.name
                        if resource_type.name is not None
                        else pefile.RESOURCE_TYPE.get(resource_type.struct.Id)
                    )
                    if isinstance(name, UnicodeStringWrapperPostProcessor):
                        name = name.decode()
                    try:
                        if hasattr(resource_type, "directory"):
                            for resource_id in resource_type.directory.entries:
                                if hasattr(resource_id, "directory"):
                                    for resource_lang in resource_id.directory.entries:
                                        rsrc_data = self.pe.get_data(
                                            resource_lang.data.struct.OffsetToData,
                                            resource_lang.data.struct.Size,
                                        )
                                        file_type = magic.from_buffer(rsrc_data)
                                        magik = Magika().identify_bytes(rsrc_data).output.label

                                        rsrc_entropy = (
                                            "%.2f"
                                            % pefile.SectionStructure.entropy_H(
                                                self.pe, rsrc_data
                                            )
                                        )
                                        rsrc_sha256 = sha256(rsrc_data).hexdigest()
                                        lang = pefile.LANG.get(
                                            resource_lang.data.lang, "*unknown*"
                                        )
                                        sublang = pefile.get_sublang_name_for_lang(
                                            resource_lang.data.lang,
                                            resource_lang.data.sublang,
                                        )
                                        pe_resource = PEResource(
                                            _id=rsrc_sha256,
                                            resource_type=name,
                                            resource_entropy=rsrc_entropy,
                                            resource_sha256=rsrc_sha256,
                                            resource_filetype=file_type,
                                            resource_magika=magik,
                                            resource_language=lang,
                                            resource_rva=resource_lang.data.struct.OffsetToData,
                                            resource_size=resource_lang.data.struct.Size,
                                            resource_sub_lang=sublang,
                                        )
                                        resources_list.append(pe_resource)
                    except Exception as e:
                        self.log.warning(
                            f"Continue after Error in {self.hash.sha256}: {resource_type.name} "
                            f"Exception: {e}",
                            stack_info=True,
                        )
                        # resources_list.append({f"{e} - {resource_type.name}"})
                        continue
        except Exception as e:
            self.log.exception(
                f"Extract exports error {self.hash.sha256} Exception: {e}"
            )
        self.log.debug(f"Resource list {resources_list}")
        return resources_list

    def extract(self):
        try:
            self.log.debug(inspect.currentframe().f_code.co_name)
            resources = self._extract_resources()
            # self.export_to_elastic(resources)  # Let the exporters handle this
            return resources
        except Exception as e:
            self.log.error(f"Extract resources error {self.hash.sha256} Exception: {e}")
            return None

    def prepare_export_data(self, exporter_type: str) -> Any:
        if exporter_type == "ElasticsearchExporter":
            return self.extract()
        elif exporter_type == "ClickHouseExporter":
            resources = self.extract()
            if resources is None:
                return None
            
            data = []
            current_time = datetime.now(timezone.utc)
            
            for resource in resources:
                data.append([
                    self.sha256,                    # sha256
                    self.md5,                       # md5
                    self.sha1,                      # sha1
                    resource.resource_type,         # resource_type
                    resource.resource_entropy,      # resource_entropy
                    resource.resource_sha256,       # resource_sha256
                    resource.resource_filetype,     # resource_filetype
                    resource.resource_magika,       # resource_magika
                    resource.resource_language,     # resource_language
                    resource.resource_sub_lang,     # resource_sub_lang
                    resource.resource_size,         # resource_size
                    resource.resource_rva,          # resource_rva
                    current_time                    # analysis_date
                ])
            
            column_names = [
                'sha256', 'md5', 'sha1', 'resource_type', 'resource_entropy',
                'resource_sha256', 'resource_filetype', 'resource_magika',
                'resource_language', 'resource_sub_lang', 'resource_size',
                'resource_rva', 'analysis_date'
            ]
            
            if not data:
                return None

            column_type_names = [
                'FixedString(64)', 'FixedString(32)', 'FixedString(40)',
                'LowCardinality(Nullable(String))', 'Float64',
                'FixedString(64)', 'LowCardinality(Nullable(String))', 'LowCardinality(Nullable(String))',
                'LowCardinality(Nullable(String))', 'LowCardinality(Nullable(String))', 'UInt64',
                'UInt64', 'DateTime64(3, \'UTC\')'
            ]

            return (data, column_names, column_type_names)

    def get_clickhouse_table(self) -> str:
        return "redb_pe_resources"