Xiaolei Li

26 papers A 2Journal 13Unranked 11
YearRankTypeTitle / Venue / Authors
2025 J jnl
Biomed. Signal Process. Control.
Xiaolei Li, Duanwei Ma, Hao Zhang, Xiao Jia, Chuanpeng Li, Ran Song, Wei Zhang
2025 J jnl
IEEE Trans. Robotics
Jiyu Cheng, Junhui Fan, Xiaolei Li, Paul L. Rosin, Yibin Li, Wei Zhang
2025 J jnl
IEEE Trans. Instrum. Meas.
Tiyu Fang, Mingxin Zhang, Ran Song, Xiaolei Li, Zhiyuan Wei, Wei Zhang
2025 J jnl
IEEE Trans Autom. Sci. Eng.
Ning Yang, Fei Lu, Xiaolei Li, Guohui Tian, Zhongyang Li, Tengfan Fu
2025 J jnl
IEEE Trans. Image Process.
Mingxin Zhang, Fuxiang Feng, Xing Fang, Lin Zhang, Youmei Zhang, Xiaolei Li, Wei Zhang
2024 A conf
IROS
Dayou Li, Chenkun Zhao, Shuo Yang, Ran Song, Xiaolei Li, Wei Zhang
2024 J jnl
CoRR
Dayou Li, Chenkun Zhao, Shuo Yang, Ran Song, Xiaolei Li, Wei Zhang
2024 J jnl
IEEE Trans. Multim.
Lin Zhang, Yifan Wang, Ran Song, Mingxin Zhang, Xiaolei Li, Wei Zhang
2024 conf
VSIP
Kaiwen Li, Tiyu Fang, Jianguo Fan, Jinqiu Fan, Jiacheng Wang, Zhiyuan Wei, Xiaolei Li
2023 conf
VSIP
Zhenling Li, Futao Liu, Yuezhong Wan, Weida Cao, Fukui Wang, Xiaolei Li
2023 J jnl
IEEE Intell. Transp. Syst. Mag.
Qing Song, Xiaolei Li, Chao Gao, Zhen Shen, Gang Xiong
2023 conf
VSIP
Baosheng Li, Jishui Han, Yuan Cheng, Chong Tan, Peng Qi, Jianping Zhang, Xiaolei Li
2023 conf
VSIP
Baosheng Li, Jishui Han, Zhenliang Qi, Liqiang Gao, Ruijie Duan, Tongtong Wang, Ran Song, Xiaolei Li
2023 J jnl
IEEE Trans. Multim.
Lin Zhang, Mingxin Zhang, Ran Song, Ziying Zhao, Xiaolei Li
2022 conf
VSIP
Futao Liu, Zhenling Li, Yuezhong Wan, Weida Cao, Duanwei Ma, Xiaolei Li, Peng Yan
2022 conf
VSIP
Baosheng Li, Jishui Han, Zhenliang Qi, Liqiang Gao, Ruijie Duan, Tongtong Wang, Peng Yan, Ran Song, Xiaolei Li
2022 J jnl
Pattern Recognit.
Guotao Wu, Ran Song, Mingxin Zhang, Xiaolei Li, Paul L. Rosin
2022 conf
VSIP
Baosheng Li, Jishui Han, Yuan Cheng, Chong Tan, Peng Qi, Jianping Zhang, Xiaolei Li
2021 conf
VSIP
Baosheng Li, Peng Qi, Jian Wang, Chong Tan, Runze Qi, Xiaolei Li
2021 J jnl
Pattern Recognit.
Zhijie Wang, Ran Song, Peng Duan, Xiaolei Li
2021 J jnl
IEEE Trans. Image Process.
Tianjiao Li, Wei Zhang, Ran Song, Zhiheng Li, Jun Liu, Xiaolei Li, Shijian Lu
2021 conf
VSIP
Baosheng Li, Chong Tan, Jian Wang, Runze Qi, Peng Qi, Xiaolei Li
2021 J jnl
Pattern Recognit. Lett.
Lei Qu, Wan Wan, Kaixuan Guo, Yu Liu, Jun Tang, Xiaolei Li, Jun Wu
2020 conf
ICRAI
Xiaolei Li, Yuxin Wu, Quanchao Jia
2019 conf
ROBIO
Mingxin Zhang, Zhijie Wang, Tiezhu Sun, Xiaolei Li
2006 A conf
IROS
Fei Lu, Mumin Song, Guohui Tian, Xiaolei Li
redb/extractors/pe_extractors/pe_resources.py
← Index redb/extractors/pe_extractors/pe_resources.py python
from hashlib import sha256
import inspect
from datetime import datetime, timezone
from typing import Any

import magic
from magika import Magika
import pefile
from pefile import UnicodeStringWrapperPostProcessor

from redb.extractors.enum import Tag
from redb.extractors.pe_extractor import PEExtractor
from redb.models.dataclasses import PEResource


class PEResourceExtractor(PEExtractor):

    def __init__(
        self,
        filepath,
        log,
        exporters=None,
        index_prefix=None,
        elastic_index=None,
        known_benign=False,
        known_malicious=False,
        pe=None,
    ):
        super().__init__(
            filepath,
            log,
            exporters,
            index_prefix,
            elastic_index,
            known_benign,
            known_malicious,
            pe,
        )
        self.elastic_index = self.index_prefix + "-pe_resources"
        self.log.debug(inspect.currentframe().f_code.co_name)

    def tag(self):
        return Tag.PE_RESOURCE.value

    def _extract_resources(self):
        """
        Returns:
        resources: a list of dictionaries, one per each resources type found.
                    each dictionary the key represents the name of the content,
                    which is the value itself.
                    Empty list if no resources present.
        """
        self.log.debug(inspect.currentframe().f_code.co_name)
        resources_list = []
        try:
            if hasattr(self.pe, "DIRECTORY_ENTRY_RESOURCE"):
                for resource_type in self.pe.DIRECTORY_ENTRY_RESOURCE.entries:
                    # if resource_type.name is not None:
                    #     name = resource_type.name
                    # else:
                    #     name = pefile.RESOURCE_TYPE.get(resource_type.struct.Id)
                    # if not name:
                    #     name = resource_type.struct.Id
                    name = (
                        resource_type.name
                        if resource_type.name is not None
                        else pefile.RESOURCE_TYPE.get(resource_type.struct.Id)
                    )
                    if isinstance(name, UnicodeStringWrapperPostProcessor):
                        name = name.decode()
                    try:
                        if hasattr(resource_type, "directory"):
                            for resource_id in resource_type.directory.entries:
                                if hasattr(resource_id, "directory"):
                                    for resource_lang in resource_id.directory.entries:
                                        rsrc_data = self.pe.get_data(
                                            resource_lang.data.struct.OffsetToData,
                                            resource_lang.data.struct.Size,
                                        )
                                        file_type = magic.from_buffer(rsrc_data)
                                        magik = Magika().identify_bytes(rsrc_data).output.label

                                        rsrc_entropy = (
                                            "%.2f"
                                            % pefile.SectionStructure.entropy_H(
                                                self.pe, rsrc_data
                                            )
                                        )
                                        rsrc_sha256 = sha256(rsrc_data).hexdigest()
                                        lang = pefile.LANG.get(
                                            resource_lang.data.lang, "*unknown*"
                                        )
                                        sublang = pefile.get_sublang_name_for_lang(
                                            resource_lang.data.lang,
                                            resource_lang.data.sublang,
                                        )
                                        pe_resource = PEResource(
                                            _id=rsrc_sha256,
                                            resource_type=name,
                                            resource_entropy=rsrc_entropy,
                                            resource_sha256=rsrc_sha256,
                                            resource_filetype=file_type,
                                            resource_magika=magik,
                                            resource_language=lang,
                                            resource_rva=resource_lang.data.struct.OffsetToData,
                                            resource_size=resource_lang.data.struct.Size,
                                            resource_sub_lang=sublang,
                                        )
                                        resources_list.append(pe_resource)
                    except Exception as e:
                        self.log.warning(
                            f"Continue after Error in {self.hash.sha256}: {resource_type.name} "
                            f"Exception: {e}",
                            stack_info=True,
                        )
                        # resources_list.append({f"{e} - {resource_type.name}"})
                        continue
        except Exception as e:
            self.log.exception(
                f"Extract exports error {self.hash.sha256} Exception: {e}"
            )
        self.log.debug(f"Resource list {resources_list}")
        return resources_list

    def extract(self):
        try:
            self.log.debug(inspect.currentframe().f_code.co_name)
            resources = self._extract_resources()
            # self.export_to_elastic(resources)  # Let the exporters handle this
            return resources
        except Exception as e:
            self.log.error(f"Extract resources error {self.hash.sha256} Exception: {e}")
            return None

    def prepare_export_data(self, exporter_type: str) -> Any:
        if exporter_type == "ElasticsearchExporter":
            return self.extract()
        elif exporter_type == "ClickHouseExporter":
            resources = self.extract()
            if resources is None:
                return None
            
            data = []
            current_time = datetime.now(timezone.utc)
            
            for resource in resources:
                data.append([
                    self.sha256,                    # sha256
                    self.md5,                       # md5
                    self.sha1,                      # sha1
                    resource.resource_type,         # resource_type
                    resource.resource_entropy,      # resource_entropy
                    resource.resource_sha256,       # resource_sha256
                    resource.resource_filetype,     # resource_filetype
                    resource.resource_magika,       # resource_magika
                    resource.resource_language,     # resource_language
                    resource.resource_sub_lang,     # resource_sub_lang
                    resource.resource_size,         # resource_size
                    resource.resource_rva,          # resource_rva
                    current_time                    # analysis_date
                ])
            
            column_names = [
                'sha256', 'md5', 'sha1', 'resource_type', 'resource_entropy',
                'resource_sha256', 'resource_filetype', 'resource_magika',
                'resource_language', 'resource_sub_lang', 'resource_size',
                'resource_rva', 'analysis_date'
            ]
            
            if not data:
                return None

            column_type_names = [
                'FixedString(64)', 'FixedString(32)', 'FixedString(40)',
                'LowCardinality(Nullable(String))', 'Float64',
                'FixedString(64)', 'LowCardinality(Nullable(String))', 'LowCardinality(Nullable(String))',
                'LowCardinality(Nullable(String))', 'LowCardinality(Nullable(String))', 'UInt64',
                'UInt64', 'DateTime64(3, \'UTC\')'
            ]

            return (data, column_names, column_type_names)

    def get_clickhouse_table(self) -> str:
        return "redb_pe_resources"