N. Hemachandra

28 papers A* 1B 1Misc 1Journal 18Unranked 7
YearRankTypeTitle / Venue / Authors
2022 conf
VALUETOOLS
Anirban Mitra, Manu K. Gupta, N. Hemachandra
2022 B conf
WiOpt
Yashvardhan Didwania, Jayakrishnan Nair, N. Hemachandra
2022 J jnl
CoRR
Yashvardhan Didwania, Jayakrishnan Nair, N. Hemachandra
2021 conf
IEEM
Devanand R, Tushar Shekhar, Ashutosh Mahajan, N. Hemachandra
2020 A* conf
AAAI
Aditya Petety, Sandhya Tripathi, N. Hemachandra
2020 J jnl
Queueing Syst. Theory Appl.
N. Hemachandra, Kishor Patil, Sandhya Tripathi
2020 J jnl
CoRR
Sandhya Tripathi, N. Hemachandra
2020 conf
IEEE BigData
Sandhya Tripathi, N. Hemachandra, Prashant Trivedi
2020 J jnl
CoRR
Sandhya Tripathi, N. Hemachandra, Prashant Trivedi
2019 J jnl
CoRR
Aditya Petety, Sandhya Tripathi, N. Hemachandra
2018 conf
COMAD/CODS
Sandhya Tripathi, N. Hemachandra
2016 J jnl
CoRR
Vikas Vikram Singh, N. Hemachandra
2015 J jnl
IGTR
Vikas Vikram Singh, N. Hemachandra
2014 J jnl
Oper. Res. Lett.
Vikas Vikram Singh, N. Hemachandra
2013 J jnl
IGTR
Vikas Vikram Singh, N. Hemachandra, K. S. Mallikarjuna Rao
2013 J jnl
IGTR
Krishna Chaitanya Vanam, N. Hemachandra
2012 J jnl
CoRR
Vikas Vikram Singh, N. Hemachandra
2012 J jnl
Comput. Networks
Koteswara Rao Vemu, Shalabh Bhatnagar, N. Hemachandra
2011 J jnl
IEEE Trans Autom. Sci. Eng.
Shalabh Bhatnagar, Vivek Kumar Mishra, N. Hemachandra
2011 J jnl
ACM Trans. Model. Comput. Simul.
Shalabh Bhatnagar, N. Hemachandra, Vivek Kumar Mishra
2010 J jnl
Eur. J. Oper. Res.
Sudhir K. Sinha, N. Rangaraj, N. Hemachandra
2007 Misc conf
ICDCIT
Koteswara Rao Vemu, Shalabh Bhatnagar, N. Hemachandra
2007 conf
CDC
Vivek Kumar Mishra, Shalabh Bhatnagar, N. Hemachandra
2007 conf
CDC
Koteswara Rao Vemu, Shalabh Bhatnagar, N. Hemachandra
2007 conf
CEC/EEE
Raghav Kumar Gautam, N. Hemachandra, Y. Narahari, V. Hastagiri Prakash
2003 J jnl
Comput. Oper. Res.
N. Hemachandra, Sarat Kumar Eedupuganti
2000 J jnl
Queueing Syst. Theory Appl.
N. Hemachandra, Y. Narahari
1997 J jnl
Comput. Oper. Res.
Y. Narahari, N. Hemachandra, M. S. Gaur
redb/extractors/pe_extractors/pe_resources.py
← Index redb/extractors/pe_extractors/pe_resources.py python
from hashlib import sha256
import inspect
from datetime import datetime, timezone
from typing import Any

import magic
from magika import Magika
import pefile
from pefile import UnicodeStringWrapperPostProcessor

from redb.extractors.enum import Tag
from redb.extractors.pe_extractor import PEExtractor
from redb.models.dataclasses import PEResource


class PEResourceExtractor(PEExtractor):

    def __init__(
        self,
        filepath,
        log,
        exporters=None,
        index_prefix=None,
        elastic_index=None,
        known_benign=False,
        known_malicious=False,
        pe=None,
    ):
        super().__init__(
            filepath,
            log,
            exporters,
            index_prefix,
            elastic_index,
            known_benign,
            known_malicious,
            pe,
        )
        self.elastic_index = self.index_prefix + "-pe_resources"
        self.log.debug(inspect.currentframe().f_code.co_name)

    def tag(self):
        return Tag.PE_RESOURCE.value

    def _extract_resources(self):
        """
        Returns:
        resources: a list of dictionaries, one per each resources type found.
                    each dictionary the key represents the name of the content,
                    which is the value itself.
                    Empty list if no resources present.
        """
        self.log.debug(inspect.currentframe().f_code.co_name)
        resources_list = []
        try:
            if hasattr(self.pe, "DIRECTORY_ENTRY_RESOURCE"):
                for resource_type in self.pe.DIRECTORY_ENTRY_RESOURCE.entries:
                    # if resource_type.name is not None:
                    #     name = resource_type.name
                    # else:
                    #     name = pefile.RESOURCE_TYPE.get(resource_type.struct.Id)
                    # if not name:
                    #     name = resource_type.struct.Id
                    name = (
                        resource_type.name
                        if resource_type.name is not None
                        else pefile.RESOURCE_TYPE.get(resource_type.struct.Id)
                    )
                    if isinstance(name, UnicodeStringWrapperPostProcessor):
                        name = name.decode()
                    try:
                        if hasattr(resource_type, "directory"):
                            for resource_id in resource_type.directory.entries:
                                if hasattr(resource_id, "directory"):
                                    for resource_lang in resource_id.directory.entries:
                                        rsrc_data = self.pe.get_data(
                                            resource_lang.data.struct.OffsetToData,
                                            resource_lang.data.struct.Size,
                                        )
                                        file_type = magic.from_buffer(rsrc_data)
                                        magik = Magika().identify_bytes(rsrc_data).output.label

                                        rsrc_entropy = (
                                            "%.2f"
                                            % pefile.SectionStructure.entropy_H(
                                                self.pe, rsrc_data
                                            )
                                        )
                                        rsrc_sha256 = sha256(rsrc_data).hexdigest()
                                        lang = pefile.LANG.get(
                                            resource_lang.data.lang, "*unknown*"
                                        )
                                        sublang = pefile.get_sublang_name_for_lang(
                                            resource_lang.data.lang,
                                            resource_lang.data.sublang,
                                        )
                                        pe_resource = PEResource(
                                            _id=rsrc_sha256,
                                            resource_type=name,
                                            resource_entropy=rsrc_entropy,
                                            resource_sha256=rsrc_sha256,
                                            resource_filetype=file_type,
                                            resource_magika=magik,
                                            resource_language=lang,
                                            resource_rva=resource_lang.data.struct.OffsetToData,
                                            resource_size=resource_lang.data.struct.Size,
                                            resource_sub_lang=sublang,
                                        )
                                        resources_list.append(pe_resource)
                    except Exception as e:
                        self.log.warning(
                            f"Continue after Error in {self.hash.sha256}: {resource_type.name} "
                            f"Exception: {e}",
                            stack_info=True,
                        )
                        # resources_list.append({f"{e} - {resource_type.name}"})
                        continue
        except Exception as e:
            self.log.exception(
                f"Extract exports error {self.hash.sha256} Exception: {e}"
            )
        self.log.debug(f"Resource list {resources_list}")
        return resources_list

    def extract(self):
        try:
            self.log.debug(inspect.currentframe().f_code.co_name)
            resources = self._extract_resources()
            # self.export_to_elastic(resources)  # Let the exporters handle this
            return resources
        except Exception as e:
            self.log.error(f"Extract resources error {self.hash.sha256} Exception: {e}")
            return None

    def prepare_export_data(self, exporter_type: str) -> Any:
        if exporter_type == "ElasticsearchExporter":
            return self.extract()
        elif exporter_type == "ClickHouseExporter":
            resources = self.extract()
            if resources is None:
                return None
            
            data = []
            current_time = datetime.now(timezone.utc)
            
            for resource in resources:
                data.append([
                    self.sha256,                    # sha256
                    self.md5,                       # md5
                    self.sha1,                      # sha1
                    resource.resource_type,         # resource_type
                    resource.resource_entropy,      # resource_entropy
                    resource.resource_sha256,       # resource_sha256
                    resource.resource_filetype,     # resource_filetype
                    resource.resource_magika,       # resource_magika
                    resource.resource_language,     # resource_language
                    resource.resource_sub_lang,     # resource_sub_lang
                    resource.resource_size,         # resource_size
                    resource.resource_rva,          # resource_rva
                    current_time                    # analysis_date
                ])
            
            column_names = [
                'sha256', 'md5', 'sha1', 'resource_type', 'resource_entropy',
                'resource_sha256', 'resource_filetype', 'resource_magika',
                'resource_language', 'resource_sub_lang', 'resource_size',
                'resource_rva', 'analysis_date'
            ]
            
            if not data:
                return None

            column_type_names = [
                'FixedString(64)', 'FixedString(32)', 'FixedString(40)',
                'LowCardinality(Nullable(String))', 'Float64',
                'FixedString(64)', 'LowCardinality(Nullable(String))', 'LowCardinality(Nullable(String))',
                'LowCardinality(Nullable(String))', 'LowCardinality(Nullable(String))', 'UInt64',
                'UInt64', 'DateTime64(3, \'UTC\')'
            ]

            return (data, column_names, column_type_names)

    def get_clickhouse_table(self) -> str:
        return "redb_pe_resources"