Warwick Irwin

19 papers B 3C 3Journal 1Unranked 12
YearRankTypeTitle / Venue / Authors
2011 B conf
ENASE
Janina Voigt, Warwick Irwin, Neville Churcher
2011 conf
ENASE (Selected Papers)
Janina Voigt, Warwick Irwin, Neville Churcher
2011 conf
ACSC
James Ashford, Neville Churcher, Warwick Irwin
2011 conf
ACSC
Joshua Oosterman, Warwick Irwin, Neville Churcher
2010 B conf
ENASE
Janina Voigt, Warwick Irwin, Neville Churcher
2010 conf
Australian Software Engineering Conference
Matthew Harward, Warwick Irwin, Neville Churcher
2007 conf
ASWEC
Neville Churcher, Sarah Frater, Cong Phuoc Huynh, Warwick Irwin
2007 J jnl
Int. J. Comput. Support. Collab. Learn.
Nilufar Baghaei, Antonija Mitrovic, Warwick Irwin
2006 conf
ASWEC
Blair Neate, Warwick Irwin, Neville Churcher
2005 C conf
ICCE
Nilufar Baghaei, Antonija Mitrovic, Warwick Irwin
2005 B conf
PROFES
Tony Dale, Neville Churcher, Warwick Irwin
2005 C conf
APSEC
Carl Cook, Warwick Irwin, Neville Churcher
2005 conf
APVIS
Neville Churcher, Warwick Irwin
2005 conf
Australian Software Engineering Conference
Warwick Irwin, Carl Cook, Neville I. Churcher
2004 conf
InVis.au
Neville Churcher, Warwick Irwin, Carl Cook
2004 C conf
APSEC
Carl Cook, Neville Churcher, Warwick Irwin
2003 conf
IEEE METRICS
Warwick Irwin, Neville I. Churcher
2003 conf
InVis.au
Neville Churcher, Warwick Irwin, Ronald D. Kriz
2001 conf
VIP
Warwick Irwin, Neville Churcher
redb/extractor_registry.py
← Index redb/extractor_registry.py python
"""
Extractor registry with lazy loading by file type.

Groups extractors by file type (pe, elf, macho, apk) and only imports
the relevant group when that file type is first encountered. This avoids
loading heavy dependencies (pefile, lief, androguard, etc.) into workers
that don't need them.
"""
import importlib
import logging

logger = logging.getLogger(__name__)

# Registry: group name -> list of (module_path, class_names)
_REGISTRY = {
    "pe": [
        ("redb.extractors.pe_extractors", [
            "PEFeaturesExtractor",
            "PEImportExtractor",
            "PEResourceExtractor",
            "PEOverlayExtractor",
            "PESectionExtractor",
            "PESignatureExtractor",
            "PEExtraFindings",
            "PEInconstistencyTestsExtractor",
            "PEDotNetExtractor",
        ]),
    ],
    "elf": [
        ("redb.extractors.elf_extractors", [
            "ELFFeaturesExtractor",
            "ELFSegmentExtractor",
            "ELFSectionExtractor",
            "ELFDependencyExtractor",
            "ELFSymbolExtractor",
            "ELFImportExtractor",
            "ELFExportExtractor",
            "ELFRelocationExtractor",
            "ELFNotesExtractor",
        ]),
    ],
    "macho": [
        ("redb.extractors.macho_extractors", [
            "MachOFeaturesExtractor",
            "MachOSegmentExtractor",
            "MachOImportExtractor",
            "MachOExportExtractor",
            "MachODylibExtractor",
            "MachOSignatureExtractor",
        ]),
    ],
    "apk": [
        ("redb.extractors.apk_extractors", [
            "APKFeaturesExtractor",
            "APKManifestExtractor",
            "APKPermissionsExtractor",
            "APKSignatureExtractor",
            "APKDexExtractor",
            "APKResourceExtractor",
            "APKNativeLibExtractor",
            "APKInconsistencyTestsExtractor",
        ]),
        ("redb.extractors.decompiler.DecompileAPK", [
            "DecompileAPK",
        ]),
    ],
    "js": [
        ("redb.extractors.js_extractors", [
            "JSFeaturesExtractor",
            "JSSuspiciousAPIsExtractor",
            "JSStringsExtractor",
            "JSDeobfuscationExtractor",
            "JSContentExtractor",
        ]),
    ],
}

# Map filetype labels (from Magika) to registry group names
FILETYPE_TO_GROUP = {
    "pebin": "pe",
    "elf": "elf",
    "macho": "macho",
    "apk": "apk",
    "javascript": "js",
}

# Cache: group name -> {class_name: class}
_group_cache = {}


def _load_group(group_name):
    """Import all extractors for a group. Cached after first call."""
    if group_name in _group_cache:
        return _group_cache[group_name]

    logger.debug(f"Loading extractor group: {group_name}")

    # Suppress androguard logging before importing APK extractors
    if group_name == "apk":
        from loguru import logger as loguru_logger
        loguru_logger.disable("androguard")

    classes = {}
    for module_path, class_names in _REGISTRY[group_name]:
        mod = importlib.import_module(module_path)
        for name in class_names:
            classes[name] = getattr(mod, name)

    _group_cache[group_name] = classes
    return classes


def get_filetype_modules(filetype):
    """Get the list of format-specific extractor classes for a filetype.

    Returns only analysis extractors (excludes decompilers like DecompileAPK).
    """
    group = FILETYPE_TO_GROUP.get(filetype)
    if not group:
        return []
    classes = _load_group(group)
    return [cls for name, cls in classes.items() if not name.startswith("Decompile")]


def get_extractor_class(name):
    """Look up a single extractor class by name across all groups.

    Checks cached groups first, then loads groups on demand.
    """
    # Check already-loaded groups first
    for group_classes in _group_cache.values():
        if name in group_classes:
            return group_classes[name]

    # Search all groups (triggers lazy loading)
    for group_name in _REGISTRY:
        classes = _load_group(group_name)
        if name in classes:
            return classes[name]

    return None