Xi Wang

15 papers Journal 13Unranked 1
YearRankTypeTitle / Venue / Authors
2025 J jnl
Inf. Sci.
Deguang Wang, Jiahan He, Xi Wang, ZhiWu Li
2024 J jnl
IEEE Syst. J.
Deguang Wang, Xi Wang, Jing Yang, ZhiWu Li
2023 book
Xi Wang, ZhiWu Li
2022 J jnl
J. Frankl. Inst.
Deguang Wang, Xi Wang, Jing Yang, ZhiWu Li
2021 J jnl
IEEE Trans. Autom. Control.
Xi Wang, ZhiWu Li, W. M. Wonham
2020 J jnl
IEEE Trans. Ind. Informatics
Deguang Wang, Xi Wang, ZhiWu Li
2020 J jnl
Inf. Sci.
Deguang Wang, Xi Wang, ZhiWu Li
2019 J jnl
Discret. Event Dyn. Syst.
Xi Wang, ZhiWu Li, Thomas Moor
2019 J jnl
IEEE Trans Autom. Sci. Eng.
Chan Gu, Xi Wang, ZhiWu Li
2018 J jnl
Autom.
Xi Wang, ZhiWu Li, W. M. Wonham
2018 J jnl
Inf. Sci.
Chan Gu, Xi Wang, ZhiWu Li, Naiqi Wu
2017 J jnl
IEEE Trans. Syst. Man Cybern. Syst.
Xi Wang, ZhiWu Li, Walter Murray Wonham
2016 J jnl
IEEE Trans. Ind. Informatics
Xi Wang, ZhiWu Li, W. M. Wonham
2015 J jnl
IEEE Trans Autom. Sci. Eng.
Xi Wang, Imen Khemaissia, Mohamed Khalgui, ZhiWu Li, Olfa Mosbahi, MengChu Zhou
2011 conf
PECCS
Xi Wang, Mohamed Khalgui, ZhiWu Li
redb/extractors/decompiler/bninja/analysis/strings.py
← Index redb/extractors/decompiler/bninja/analysis/strings.py python
from collections import Counter
import math

class StringAnalysis:
    def __init__(self, bv, functions):
        self.bv = bv
        self.functions = functions

    def entropy(self, s: str) -> float:
        """Compute Shannon entropy of a string."""
        if not s:
            return 0.0
        freq = Counter(s)
        length = len(s)
        return -sum((count / length) * math.log2(count / length) for count in freq.values())

    def analyze(self):
        """
        Extract unique strings from the binary.

        Deduplicates by (string, encoding) within the same binary, keeping the
        first occurrence (lowest offset). Cross-binary deduplication and
        aggregation is handled by ClickHouse materialized views.
        """
        strings = {}

        # Sort strings by their starting address
        sorted_entries = sorted(self.bv.strings, key=lambda e: e.start)

        for entry in sorted_entries:
            # Key is the string and its encoding
            key = (entry.value, entry.type.name)

            # Skip if this string (value + encoding) was already added.
            # Because entries are sorted by address, the first one is always kept.
            if key in strings:
                continue

            # Store only the first occurrence with schema-matching field names
            # entry.length is the raw byte length, len(entry.value) is decoded string length
            string_entry = {
                "string": entry.value,
                "string_raw": entry.raw,
                "string_encoding": entry.type.name,
                "string_offset": entry.start,
                "string_length": len(entry.value),
                "string_raw_length": entry.length,
                "string_entropy": self.entropy(entry.value),
            }

            strings[key] = string_entry

        # Return as list for export compatibility
        return list(strings.values())