Wei Gong

39 papers Journal 39
YearRankTypeTitle / Venue / Authors
2025 J jnl
J. Sci. Comput.
Wei Gong, Zhiyu Tan
2025 J jnl
CoRR
Jing Li, Xindi Hu, Helin Gong, Wei Gong, Shengfeng Zhu
2025 J jnl
CoRR
Qianqian Wu, Rongfang Gong, Wei Gong, Ziyi Zhang, Shengfeng Zhu
2025 J jnl
Comput. Optim. Appl.
Wei Gong, Le Liu
2024 J jnl
CoRR
Wei Gong, Ziyi Zhang
2024 J jnl
CoRR
Wei Gong, Dongdong Liang
2024 J jnl
SIAM J. Sci. Comput.
Xuejian Li, Wei Gong, Xiaoming He, Tao Lin
2023 J jnl
CoRR
Dongdong Liang, Wei Gong, Xiaoping Xie
2023 J jnl
CoRR
Wei Gong, Zhiyu Tan
2023 J jnl
CoRR
Wei Gong, Le Liu
2023 J jnl
Adv. Comput. Math.
Wei Gong, Dongdong Liang, Xiaoping Xie
2022 J jnl
CoRR
Gang Chen, Wei Gong, Mariano Mateos, John R. Singler, Yangwen Zhang
2022 J jnl
SIAM J. Numer. Anal.
Wei Gong, Mariano Mateos, John R. Singler, Yangwen Zhang
2022 J jnl
CoRR
Wei Gong, Felix Kwok, Zhiyu Tan
2022 J jnl
J. Comput. Appl. Math.
Yue Shen, Wei Gong, Ningning Yan
2022 J jnl
J. Sci. Comput.
Kaiye Zhou, Wei Gong
2022 J jnl
SIAM J. Sci. Comput.
Wei Gong, Jiajie Li, Shengfeng Zhu
2022 J jnl
J. Sci. Comput.
Wei Gong, Buyang Li, Huanhuan Yang
2022 J jnl
SIAM J. Appl. Math.
Lili Chang, Wei Gong, Zhen Jin, Gui-Quan Sun
2021 J jnl
SIAM J. Numer. Anal.
Wei Gong, Shengfeng Zhu
2018 J jnl
SIAM J. Numer. Anal.
Wei Gong, Weiwei Hu, Mariano Mateos, John R. Singler, Xiao Zhang, Yangwen Zhang
2018 J jnl
Numer. Linear Algebra Appl.
Wei Gong, Zhiyu Tan, Shuo Zhang
2017 J jnl
J. Sci. Comput.
Wei Gong, Hehu Xie, Ningning Yan
2017 J jnl
Numerische Mathematik
Wei Gong, Ningning Yan
2016 J jnl
SIAM J. Numer. Anal.
Wei Gong, Ningning Yan
2016 J jnl
J. Sci. Comput.
Wei Gong, Michael Hinze, Zhaojie Zhou
2016 J jnl
Comput. Math. Appl.
Zhaojie Zhou, Wei Gong
2015 J jnl
SIAM J. Sci. Comput.
Wei Gong, Hehu Xie, Ningning Yan
2015 J jnl
Math. Comput. Simul.
Dongyang Shi, Qili Tang, Wei Gong
2014 J jnl
SIAM J. Control. Optim.
Wei Gong, Michael Hinze, Zhaojie Zhou
2014 J jnl
SIAM J. Control. Optim.
Wei Gong, Gengsheng Wang, Ningning Yan
2013 J jnl
Math. Comput.
Wei Gong
2013 J jnl
Comput. Optim. Appl.
Wei Gong, Michael Hinze
2012 J jnl
J. Num. Math.
Wei Gong, Michael Hinze, Zhaojie Zhou
2011 J jnl
J. Sci. Comput.
Wei Gong, Ningning Yan
2011 J jnl
Math. Comput. Model.
Dongyang Shi, Wei Gong, Jincheng Ren
2011 J jnl
SIAM J. Control. Optim.
Wei Gong, Ningning Yan
2011 J jnl
J. Comput. Appl. Math.
Wei Gong, Ningning Yan
2009 J jnl
J. Syst. Sci. Complex.
Dongyang Shi, Wei Gong
redb/extractors/pe_extractors/pe_resources.py
← Index redb/extractors/pe_extractors/pe_resources.py python
from hashlib import sha256
import inspect
from datetime import datetime, timezone
from typing import Any

import magic
from magika import Magika
import pefile
from pefile import UnicodeStringWrapperPostProcessor

from redb.extractors.enum import Tag
from redb.extractors.pe_extractor import PEExtractor
from redb.models.dataclasses import PEResource


class PEResourceExtractor(PEExtractor):

    def __init__(
        self,
        filepath,
        log,
        exporters=None,
        index_prefix=None,
        elastic_index=None,
        known_benign=False,
        known_malicious=False,
        pe=None,
    ):
        super().__init__(
            filepath,
            log,
            exporters,
            index_prefix,
            elastic_index,
            known_benign,
            known_malicious,
            pe,
        )
        self.elastic_index = self.index_prefix + "-pe_resources"
        self.log.debug(inspect.currentframe().f_code.co_name)

    def tag(self):
        return Tag.PE_RESOURCE.value

    def _extract_resources(self):
        """
        Returns:
        resources: a list of dictionaries, one per each resources type found.
                    each dictionary the key represents the name of the content,
                    which is the value itself.
                    Empty list if no resources present.
        """
        self.log.debug(inspect.currentframe().f_code.co_name)
        resources_list = []
        try:
            if hasattr(self.pe, "DIRECTORY_ENTRY_RESOURCE"):
                for resource_type in self.pe.DIRECTORY_ENTRY_RESOURCE.entries:
                    # if resource_type.name is not None:
                    #     name = resource_type.name
                    # else:
                    #     name = pefile.RESOURCE_TYPE.get(resource_type.struct.Id)
                    # if not name:
                    #     name = resource_type.struct.Id
                    name = (
                        resource_type.name
                        if resource_type.name is not None
                        else pefile.RESOURCE_TYPE.get(resource_type.struct.Id)
                    )
                    if isinstance(name, UnicodeStringWrapperPostProcessor):
                        name = name.decode()
                    try:
                        if hasattr(resource_type, "directory"):
                            for resource_id in resource_type.directory.entries:
                                if hasattr(resource_id, "directory"):
                                    for resource_lang in resource_id.directory.entries:
                                        rsrc_data = self.pe.get_data(
                                            resource_lang.data.struct.OffsetToData,
                                            resource_lang.data.struct.Size,
                                        )
                                        file_type = magic.from_buffer(rsrc_data)
                                        magik = Magika().identify_bytes(rsrc_data).output.label

                                        rsrc_entropy = (
                                            "%.2f"
                                            % pefile.SectionStructure.entropy_H(
                                                self.pe, rsrc_data
                                            )
                                        )
                                        rsrc_sha256 = sha256(rsrc_data).hexdigest()
                                        lang = pefile.LANG.get(
                                            resource_lang.data.lang, "*unknown*"
                                        )
                                        sublang = pefile.get_sublang_name_for_lang(
                                            resource_lang.data.lang,
                                            resource_lang.data.sublang,
                                        )
                                        pe_resource = PEResource(
                                            _id=rsrc_sha256,
                                            resource_type=name,
                                            resource_entropy=rsrc_entropy,
                                            resource_sha256=rsrc_sha256,
                                            resource_filetype=file_type,
                                            resource_magika=magik,
                                            resource_language=lang,
                                            resource_rva=resource_lang.data.struct.OffsetToData,
                                            resource_size=resource_lang.data.struct.Size,
                                            resource_sub_lang=sublang,
                                        )
                                        resources_list.append(pe_resource)
                    except Exception as e:
                        self.log.warning(
                            f"Continue after Error in {self.hash.sha256}: {resource_type.name} "
                            f"Exception: {e}",
                            stack_info=True,
                        )
                        # resources_list.append({f"{e} - {resource_type.name}"})
                        continue
        except Exception as e:
            self.log.exception(
                f"Extract exports error {self.hash.sha256} Exception: {e}"
            )
        self.log.debug(f"Resource list {resources_list}")
        return resources_list

    def extract(self):
        try:
            self.log.debug(inspect.currentframe().f_code.co_name)
            resources = self._extract_resources()
            # self.export_to_elastic(resources)  # Let the exporters handle this
            return resources
        except Exception as e:
            self.log.error(f"Extract resources error {self.hash.sha256} Exception: {e}")
            return None

    def prepare_export_data(self, exporter_type: str) -> Any:
        if exporter_type == "ElasticsearchExporter":
            return self.extract()
        elif exporter_type == "ClickHouseExporter":
            resources = self.extract()
            if resources is None:
                return None
            
            data = []
            current_time = datetime.now(timezone.utc)
            
            for resource in resources:
                data.append([
                    self.sha256,                    # sha256
                    self.md5,                       # md5
                    self.sha1,                      # sha1
                    resource.resource_type,         # resource_type
                    resource.resource_entropy,      # resource_entropy
                    resource.resource_sha256,       # resource_sha256
                    resource.resource_filetype,     # resource_filetype
                    resource.resource_magika,       # resource_magika
                    resource.resource_language,     # resource_language
                    resource.resource_sub_lang,     # resource_sub_lang
                    resource.resource_size,         # resource_size
                    resource.resource_rva,          # resource_rva
                    current_time                    # analysis_date
                ])
            
            column_names = [
                'sha256', 'md5', 'sha1', 'resource_type', 'resource_entropy',
                'resource_sha256', 'resource_filetype', 'resource_magika',
                'resource_language', 'resource_sub_lang', 'resource_size',
                'resource_rva', 'analysis_date'
            ]
            
            if not data:
                return None

            column_type_names = [
                'FixedString(64)', 'FixedString(32)', 'FixedString(40)',
                'LowCardinality(Nullable(String))', 'Float64',
                'FixedString(64)', 'LowCardinality(Nullable(String))', 'LowCardinality(Nullable(String))',
                'LowCardinality(Nullable(String))', 'LowCardinality(Nullable(String))', 'UInt64',
                'UInt64', 'DateTime64(3, \'UTC\')'
            ]

            return (data, column_names, column_type_names)

    def get_clickhouse_table(self) -> str:
        return "redb_pe_resources"