☰ GDG /
Docstrings (001–005)
#002 gdtest_google #003 gdtest_sphinx #001 gdtest_minimal #004 gdtest_nodocs #005 gdtest_mixed_docs
Layouts (006–013)
#007 gdtest_python_layout #008 gdtest_lib_layout #006 gdtest_src_layout #009 gdtest_hatch #010 gdtest_setuptools_find #011 gdtest_setup_cfg #012 gdtest_setup_py #013 gdtest_auto_discover
Exports (014–017)
#014 gdtest_no_all #016 gdtest_config_exclude #015 gdtest_all_concat #017 gdtest_auto_exclude
Object Types (018–027)
#018 gdtest_small_class #020 gdtest_dataclasses #019 gdtest_big_class #021 gdtest_enums #022 gdtest_typed_containers #023 gdtest_protocols #024 gdtest_descriptors #026 gdtest_nested_class #025 gdtest_dunders #027 gdtest_constants
Directives (028–032)
#028 gdtest_seealso #029 gdtest_nodoc #030 gdtest_user_guide_auto #031 gdtest_user_guide_sections #032 gdtest_user_guide_subdirs
User Guide (033–038)
#033 gdtest_user_guide_explicit #034 gdtest_user_guide_custom_dir #035 gdtest_user_guide_hyphen #036 gdtest_readme_rst #037 gdtest_index_qmd #038 gdtest_index_md
Landing Pages (039–043)
#039 gdtest_no_readme #040 gdtest_index_wins #041 gdtest_index_frontmatter #042 gdtest_full_extras #043 gdtest_github_contrib
Extras & Config (044–050)
#044 gdtest_cli_click #048 gdtest_name_mismatch #046 gdtest_explicit_ref #045 gdtest_cli_nested #047 gdtest_kitchen_sink #049 gdtest_src_big_class #050 gdtest_google_big_class
Cross-Dimension (051–065)
#051 gdtest_user_guide_cli #053 gdtest_src_no_all #052 gdtest_explicit_big_class #055 gdtest_google_seealso #054 gdtest_extras_guide #056 gdtest_setup_cfg_src #057 gdtest_exclude_cli #058 gdtest_src_explicit_ref #059 gdtest_async_funcs #060 gdtest_generators #061 gdtest_overloads #062 gdtest_abstract_props #063 gdtest_multi_inherit #064 gdtest_slots_class #065 gdtest_frozen_dc
API Patterns (066–077)
#066 gdtest_generics #067 gdtest_context_mgr #068 gdtest_decorators #069 gdtest_exceptions #070 gdtest_reexports #071 gdtest_many_exports #072 gdtest_deep_nesting #073 gdtest_long_docs #074 gdtest_many_guides #076 gdtest_flit #077 gdtest_pdm #075 gdtest_many_big_classes
Scale & Stress (078–082)
#078 gdtest_namespace #079 gdtest_monorepo #080 gdtest_multi_module #081 gdtest_src_legacy #082 gdtest_empty_module
Build Systems (083–088)
#083 gdtest_all_private #084 gdtest_duplicate_all #085 gdtest_badge_readme #086 gdtest_math_docs #087 gdtest_mixed_guide_ext #088 gdtest_unicode_docs
Edge Cases (089–095)
#089 gdtest_config_all_on #090 gdtest_config_display #091 gdtest_config_minimal #092 gdtest_config_parser #093 gdtest_config_extra_keys #094 gdtest_github_icon #095 gdtest_source_branch
Config Matrix (096–100)
#096 gdtest_source_path #097 gdtest_source_title #098 gdtest_source_disabled #099 gdtest_sidebar_disabled #100 gdtest_sidebar_min_items
Config Options (101–125)
#102 gdtest_cli_name #103 gdtest_dynamic_false #101 gdtest_sidebar_float #104 gdtest_parser_google #105 gdtest_parser_sphinx #106 gdtest_display_name #107 gdtest_funding #108 gdtest_authors_multi #109 gdtest_no_darkmode #110 gdtest_exclude_list #111 gdtest_jupyter_kernel #112 gdtest_config_sections #113 gdtest_config_ug_string #114 gdtest_config_ug_list #115 gdtest_config_changelog #116 gdtest_config_reference #117 gdtest_config_combo_a
#121 gdtest_config_combo_e FAIL
#118 gdtest_config_combo_b #119 gdtest_config_combo_c #120 gdtest_config_combo_d #122 gdtest_config_combo_f #123 gdtest_attribution_on #124 gdtest_attribution_off #125 gdtest_rst_versionadded
Docstring Richness (126–150)
#126 gdtest_rst_deprecated #127 gdtest_rst_note #128 gdtest_rst_warning #130 gdtest_rst_caution #129 gdtest_rst_tip #131 gdtest_rst_danger #132 gdtest_rst_important #134 gdtest_directives #133 gdtest_rst_mixed_dirs #135 gdtest_sphinx_func_role #136 gdtest_sphinx_class_role #137 gdtest_sphinx_exc_role #138 gdtest_sphinx_meth_role #139 gdtest_sphinx_mixed_roles #140 gdtest_numpy_rich #141 gdtest_google_rich #142 gdtest_sphinx_rich #143 gdtest_docstring_examples #144 gdtest_examples_rst_repro #145 gdtest_docstring_notes #146 gdtest_docstring_warnings #147 gdtest_docstring_references #148 gdtest_docstring_seealso #150 gdtest_docstring_tables #149 gdtest_docstring_math
UG Variations (151–165)
#151 gdtest_docstring_combo #152 gdtest_ug_auto #153 gdtest_ug_numbered #154 gdtest_ug_sections_fm #155 gdtest_ug_subdirs #156 gdtest_ug_custom_dir #157 gdtest_ug_deep_nest #158 gdtest_ug_mixed_ext #160 gdtest_ug_explicit_order #159 gdtest_ug_many_pages #162 gdtest_ug_no_frontmatter #161 gdtest_ug_single_page #164 gdtest_ug_with_images #163 gdtest_ug_with_code #165 gdtest_ug_hyphen_dir
Custom Sections (166–175)
#166 gdtest_ug_combo #167 gdtest_sec_examples #168 gdtest_sec_tutorials #169 gdtest_sec_recipes #170 gdtest_sec_blog #171 gdtest_sec_faq #172 gdtest_sec_multi #173 gdtest_sec_navbar_after #174 gdtest_sec_with_ug #175 gdtest_sec_with_ref
Reference Config (176–185)
#176 gdtest_sec_deep #177 gdtest_sec_index_opt #178 gdtest_sec_index_hero #179 gdtest_sec_sidebar_single #180 gdtest_custom_passthrough_navbar #181 gdtest_custom_raw_navbar_after #182 gdtest_custom_mixed_modes #183 gdtest_custom_nested_combo #184 gdtest_custom_basename_output #185 gdtest_custom_nested_output
Site Theming (186–195)
#186 gdtest_custom_missing_dir_combo #187 gdtest_ref_explicit #188 gdtest_ref_members_false #189 gdtest_ref_mixed #190 gdtest_ref_reorder #192 gdtest_ref_single_section #191 gdtest_ref_sectioned #193 gdtest_ref_module_expand #194 gdtest_ref_big_class #195 gdtest_ref_multi_big
Stress Tests (196–200)
#196 gdtest_ref_title #197 gdtest_theme_cosmo #198 gdtest_theme_lumen #199 gdtest_theme_cerulean #200 gdtest_toc_disabled #201 gdtest_toc_depth #202 gdtest_toc_title #203 gdtest_site_combo #204 gdtest_display_badges #205 gdtest_display_authors #206 gdtest_display_funding #207 gdtest_stress_all_config #208 gdtest_stress_all_docstr #209 gdtest_stress_all_ug #210 gdtest_stress_all_sections #212 gdtest_src_google_seealso #211 gdtest_stress_everything #213 gdtest_hatch_nodoc #214 gdtest_pdm_big_class #215 gdtest_flit_enums #216 gdtest_namespace_ug #217 gdtest_ug_subdir_numbered #218 gdtest_homepage_ug #220 gdtest_logo #221 gdtest_hero_basic #222 gdtest_hero_readme_badges #219 gdtest_long_names #223 gdtest_hero_disabled #224 gdtest_hero_custom #225 gdtest_hero_wordmark #226 gdtest_hero_no_logo #227 gdtest_hero_explicit_badges #228 gdtest_hero_index_qmd #229 gdtest_hero_auto_logo #230 gdtest_md_disabled #231 gdtest_md_no_widget #232 gdtest_announce_simple #233 gdtest_announce_dict #234 gdtest_announce_disabled #235 gdtest_gradient_sky #236 gdtest_gradient_peach #237 gdtest_gradient_prism #238 gdtest_gradient_lilac #239 gdtest_gradient_slate #240 gdtest_gradient_honey #241 gdtest_gradient_dusk #242 gdtest_gradient_mint #243 gdtest_gradient_navbar #244 gdtest_gradient_both #245 gdtest_gradient_mixed #246 gdtest_gradient_no_dismiss #247 gdtest_header_text #248 gdtest_header_list #249 gdtest_header_file #250 gdtest_navbar_color #251 gdtest_navbar_color_light #252 gdtest_navbar_color_dark #253 gdtest_navbar_color_same #254 gdtest_navbar_color_split #256 gdtest_stress_everything_q #255 gdtest_kitchen_sink_q #257 gdtest_seealso_desc #258 gdtest_numpy_seealso_desc #259 gdtest_interlinks_prose #260 gdtest_autolink #261 gdtest_skill_default #262 gdtest_skill_curated #264 gdtest_skill_disabled #263 gdtest_skill_config #265 gdtest_skill_rich #266 gdtest_skill_combo #267 gdtest_skill_complex #268 gdtest_i18n_french #269 gdtest_i18n_japanese #270 gdtest_i18n_arabic #271 gdtest_code_cells #273 gdtest_page_tags #272 gdtest_nav_icons #274 gdtest_page_status #277 gdtest_homepage_ug_subdirs #275 gdtest_tag_location #276 gdtest_icon_shortcode #278 gdtest_gt_tables #281 gdtest_homepage_wide #282 gdtest_interlinks_userguide #279 gdtest_scale_to_fit #280 gdtest_scale_min_scale #283 gdtest_code_span_headings #286 gdtest_namespace_src #284 gdtest_sec_blog_user_index #285 gdtest_sec_dir_titles #287 gdtest_auto_include #288 gdtest_no_auto_exclude #290 gdtest_tbl_shortcode #292 gdtest_hr_shortcode #293 gdtest_accent_color #294 gdtest_keys_shortcode #291 gdtest_tbl_explorer #289 gdtest_tbl_preview #296 gdtest_inline_always #295 gdtest_inline_methods #297 gdtest_inline_never #298 gdtest_inline_threshold #299 gdtest_ref_inherited_explicit #300 gdtest_ref_include_inherited #304 gdtest_lightbox #302 gdtest_details_shortcode #303 gdtest_termshow #301 gdtest_mock_code #305 gdtest_hero_no_name #306 gdtest_sec_nested_tags #307 gdtest_sec_xref_subdirs #308 gdtest_bibliography #311 gdtest_go_cli #309 gdtest_bibliography_csl #310 gdtest_custom_css #312 gdtest_ug_dark_assets #313 gdtest_code_include #314 gdtest_ug_mixed_subdir_order #315 gdtest_type_aliases
314/315 built ⏱ 1m 35s 🧪 9/24

gdtest-long-names

Test sidebar wrapping with long object names

Installation

pip install gdtest-long-names

Get Started

Source files
📁 gdtest_long_names/
📄 __init__.py
"""Package with deliberately long object names."""

__version__ = "0.1.0"

from gdtest_long_names.store import (
    BaseDocumentStore,
    DuckDBDocumentStore,
    PostgreSQLDocumentStore,
)
from gdtest_long_names.embedding import (
    EmbeddingProvider,
    OpenAIEmbeddingProvider,
    CohereEmbeddingProvider,
)
from gdtest_long_names.chunker import (
    BaseChunkerStrategy,
    MarkdownChunkerStrategy,
)
from gdtest_long_names.types import (
    RetrievedDocumentChunk,
    DocumentMetadataConfig,
    EmbeddingVectorResult,
)
from gdtest_long_names.plaintext import (
    documentstorewithvectorsearchcapabilities,
    EMBEDDINGPROVIDERWITHBATCHPROCESSINGSUPPORT,
    Chunkerstrategywithoverlapdetection,
)

__all__ = [
    "BaseDocumentStore",
    "DuckDBDocumentStore",
    "PostgreSQLDocumentStore",
    "EmbeddingProvider",
    "OpenAIEmbeddingProvider",
    "CohereEmbeddingProvider",
    "BaseChunkerStrategy",
    "MarkdownChunkerStrategy",
    "RetrievedDocumentChunk",
    "DocumentMetadataConfig",
    "EmbeddingVectorResult",
    "documentstorewithvectorsearchcapabilities",
    "EMBEDDINGPROVIDERWITHBATCHPROCESSINGSUPPORT",
    "Chunkerstrategywithoverlapdetection",
]
📄 chunker.py
"""Chunker strategy implementations."""


class BaseChunkerStrategy:
    """
    Abstract base class for document chunking strategies.

    Parameters
    ----------
    max_chunk_size
        Maximum size of each chunk in characters.
    overlap_size
        Number of overlapping characters between chunks.
    """

    def __init__(self, max_chunk_size: int = 1000, overlap_size: int = 200):
        self.max_chunk_size = max_chunk_size
        self.overlap_size = overlap_size

    def chunk_document_content(self, content: str) -> list:
        """Split document content into chunks."""
        return []

    def calculate_optimal_boundaries(self, content: str) -> list:
        """Find optimal chunk boundary positions."""
        return []


class MarkdownChunkerStrategy(BaseChunkerStrategy):
    """
    Markdown-aware chunking strategy that respects heading boundaries.

    Parameters
    ----------
    max_chunk_size
        Maximum size of each chunk in characters.
    overlap_size
        Number of overlapping characters between chunks.
    preserve_code_blocks
        Whether to keep code blocks intact.
    """

    def __init__(self, max_chunk_size: int = 1000, overlap_size: int = 200, preserve_code_blocks: bool = True):
        super().__init__(max_chunk_size, overlap_size)
        self.preserve_code_blocks = preserve_code_blocks

    def split_by_heading_hierarchy(self, content: str) -> list:
        """Split content by markdown heading hierarchy."""
        return []

    def merge_undersized_fragments(self, chunks: list) -> list:
        """Merge chunks that are too small to stand alone."""
        return []
📄 embedding.py
"""Embedding provider implementations."""


class EmbeddingProvider:
    """
    Base class for embedding providers.

    Parameters
    ----------
    model_name
        Name of the embedding model.
    """

    def __init__(self, model_name: str):
        self.model_name = model_name

    def generate_embeddings(self, texts: list) -> list:
        """Generate embeddings for a list of texts."""
        return []


class OpenAIEmbeddingProvider(EmbeddingProvider):
    """
    OpenAI embedding provider using text-embedding models.

    Parameters
    ----------
    model_name
        Name of the OpenAI model.
    api_key
        OpenAI API key.
    """

    def __init__(self, model_name: str = "text-embedding-3-small", api_key: str = ""):
        super().__init__(model_name)
        self.api_key = api_key

    def generate_embeddings_batch(self, texts: list, batch_size: int = 100) -> list:
        """Generate embeddings in batches to handle rate limits."""
        return []

    def calculate_token_usage(self, texts: list) -> int:
        """Calculate total token usage for a list of texts."""
        return 0


class CohereEmbeddingProvider(EmbeddingProvider):
    """
    Cohere embedding provider with input type support.

    Parameters
    ----------
    model_name
        Name of the Cohere model.
    input_type
        Type of input for embedding.
    """

    def __init__(self, model_name: str = "embed-english-v3.0", input_type: str = "search_document"):
        super().__init__(model_name)
        self.input_type = input_type

    def generate_with_input_type(self, texts: list, input_type: str) -> list:
        """Generate embeddings with specific input type."""
        return []

    def get_supported_languages(self) -> list:
        """Return list of supported languages."""
        return []
📄 plaintext.py
"""Classes with long plain-text names (no special characters)."""


class documentstorewithvectorsearchcapabilities:
    """
    A store for documents supporting vector search.

    This class name is entirely lowercase with no separators,
    underscores, dots, or camelCase transitions.

    Parameters
    ----------
    connectionstring
        Database connection string.
    vectordimension
        Dimensionality of stored vectors.
    """

    def __init__(self, connectionstring: str, vectordimension: int = 1536):
        self.connectionstring = connectionstring
        self.vectordimension = vectordimension

    def insertdocumentswithembeddings(self, docs: list) -> int:
        """Insert documents along with their embedding vectors."""
        return 0

    def searchbyvectorsimilarity(self, query: str, topk: int = 10) -> list:
        """Search for documents by vector similarity."""
        return []

    def rebuildvectorsearchindex(self) -> None:
        """Rebuild the internal vector search index."""
        pass

    def deletedocumentsbyidentifier(self, docid: str) -> bool:
        """Delete a document by its unique identifier."""
        return False

    def countdocumentsincollection(self) -> int:
        """Return the total number of documents stored."""
        return 0

    def exportcollectiontojsonlines(self, filepath: str) -> int:
        """Export all documents to a JSON Lines file."""
        return 0


class EMBEDDINGPROVIDERWITHBATCHPROCESSINGSUPPORT:
    """
    All-uppercase embedding provider class.

    This class name is entirely uppercase with no separators,
    underscores, dots, or camelCase transitions.

    Parameters
    ----------
    MODELIDENTIFIER
        Identifier for the embedding model.
    BATCHLIMIT
        Maximum batch size for processing.
    """

    def __init__(self, MODELIDENTIFIER: str, BATCHLIMIT: int = 100):
        self.MODELIDENTIFIER = MODELIDENTIFIER
        self.BATCHLIMIT = BATCHLIMIT

    def GENERATEEMBEDDINGSFROMTEXTINPUT(self, texts: list) -> list:
        """Generate embeddings from a list of text inputs."""
        return []

    def CALCULATETOKENCOUNTFORTEXTS(self, texts: list) -> int:
        """Calculate total token count for the given texts."""
        return 0

    def RETRIEVEMODELCONFIGURATION(self) -> dict:
        """Retrieve the current model configuration."""
        return {}

    def VALIDATEINPUTTEXTLENGTHS(self, texts: list) -> bool:
        """Validate that all input texts are within length limits."""
        return True

    def EXPORTEMBEDDINGSTOFILE(self, filepath: str) -> int:
        """Export computed embeddings to a file."""
        return 0

    def RESETINTERNALBATCHCOUNTER(self) -> None:
        """Reset the internal batch processing counter."""
        pass


class Chunkerstrategywithoverlapdetection:
    """
    Initial-cap chunker strategy class.

    This class name starts with an uppercase letter and the rest
    is entirely lowercase, with no other separators.

    Parameters
    ----------
    maxchunksize
        Maximum size of each chunk in characters.
    overlapsize
        Number of overlapping characters between chunks.
    """

    def __init__(self, maxchunksize: int = 1000, overlapsize: int = 200):
        self.maxchunksize = maxchunksize
        self.overlapsize = overlapsize

    def splitcontentintochunks(self, content: str) -> list:
        """Split document content into overlapping chunks."""
        return []

    def detectoverlapboundaries(self, content: str) -> list:
        """Detect optimal overlap boundary positions."""
        return []

    def mergeundersizedfragments(self, chunks: list) -> list:
        """Merge fragments that are too small to stand alone."""
        return []

    def calculateoverlappercentage(self, chunks: list) -> float:
        """Calculate the average overlap percentage between chunks."""
        return 0.0

    def exportchunkswithoverlap(self, filepath: str) -> int:
        """Export chunks with overlap markers to a file."""
        return 0

    def resetinternalchunkcache(self) -> None:
        """Reset the internal chunk processing cache."""
        pass
📄 store.py
"""Document store implementations."""


class BaseDocumentStore:
    """
    Abstract base class for document stores.

    Parameters
    ----------
    connection_string
        Database connection string.
    """

    def __init__(self, connection_string: str):
        self.connection_string = connection_string

    def connect_to_database(self) -> None:
        """Establish connection to the underlying database."""
        pass

    def create_collection(self, name: str) -> None:
        """Create a new document collection."""
        pass


class DuckDBDocumentStore(BaseDocumentStore):
    """
    DuckDB-backed document store with vector search.

    Parameters
    ----------
    connection_string
        Database connection string.
    index_type
        Type of vector index to use.
    """

    def __init__(self, connection_string: str, index_type: str = "hnsw"):
        super().__init__(connection_string)
        self.index_type = index_type

    def upsert_documents(self, docs: list) -> int:
        """Insert or update documents in the store."""
        return 0

    def ingest_from_directory(self, path: str) -> int:
        """Ingest all documents from a directory."""
        return 0

    def retrieve_by_similarity(self, query: str, top_k: int = 10) -> list:
        """Retrieve documents by vector similarity search."""
        return []

    def retrieve_by_bm25_score(self, query: str, top_k: int = 10) -> list:
        """Retrieve documents using BM25 text scoring."""
        return []

    def retrieve_hybrid_combination(self, query: str, top_k: int = 10) -> list:
        """Retrieve using hybrid vector + BM25 combination."""
        return []

    def build_vector_index(self) -> None:
        """Build or rebuild the vector similarity index."""
        pass

    def get_collection_size(self) -> int:
        """Return the number of documents in the store."""
        return 0


class PostgreSQLDocumentStore(BaseDocumentStore):
    """
    PostgreSQL-backed document store with pgvector.

    Parameters
    ----------
    connection_string
        Database connection string.
    embedding_dimension
        Dimensionality of embedding vectors.
    """

    def __init__(self, connection_string: str, embedding_dimension: int = 1536):
        super().__init__(connection_string)
        self.embedding_dimension = embedding_dimension

    def upsert_with_embeddings(self, docs: list, embeddings: list) -> int:
        """Insert or update documents with precomputed embeddings."""
        return 0

    def retrieve_nearest_neighbors(self, embedding: list, top_k: int = 10) -> list:
        """Retrieve documents using nearest neighbor search."""
        return []

    def create_ivfflat_index(self, num_lists: int = 100) -> None:
        """Create an IVFFlat index for approximate search."""
        pass

    def vacuum_analyze_table(self) -> None:
        """Run VACUUM ANALYZE on the document table."""
        pass
📄 types.py
"""Type definitions and data containers."""

from dataclasses import dataclass


@dataclass
class RetrievedDocumentChunk:
    """
    A document chunk returned from a retrieval query.

    Parameters
    ----------
    content
        The text content of the chunk.
    similarity_score
        Cosine similarity score (0 to 1).
    document_id
        Identifier of the source document.
    """

    content: str
    similarity_score: float
    document_id: str


@dataclass
class DocumentMetadataConfig:
    """
    Configuration for document metadata extraction.

    Parameters
    ----------
    extract_title
        Whether to extract document titles.
    extract_author
        Whether to extract author information.
    custom_metadata_fields
        Additional metadata fields to extract.
    """

    extract_title: bool = True
    extract_author: bool = True
    custom_metadata_fields: list = None

    def __post_init__(self):
        if self.custom_metadata_fields is None:
            self.custom_metadata_fields = []


@dataclass
class EmbeddingVectorResult:
    """
    Result container for embedding vector operations.

    Parameters
    ----------
    vectors
        List of embedding vectors.
    model_name
        Name of the model used.
    token_count
        Total tokens processed.
    """

    vectors: list
    model_name: str
    token_count: int
📄 great-docs.yml
reference:
  sections:
    - title: Document Stores
      desc: Backend storage systems for documents and embeddings.
      contents:
        - BaseDocumentStore
        - DuckDBDocumentStore
        - PostgreSQLDocumentStore
    - title: Embedding Providers
      desc: Services for generating vector embeddings.
      contents:
        - EmbeddingProvider
        - OpenAIEmbeddingProvider
        - CohereEmbeddingProvider
    - title: Chunker Strategies
      desc: Strategies for splitting documents into chunks.
      contents:
        - BaseChunkerStrategy
        - MarkdownChunkerStrategy
    - title: Data Types
      desc: Type definitions and result containers.
      contents:
        - RetrievedDocumentChunk
        - DocumentMetadataConfig
        - EmbeddingVectorResult
    - title: Plain Text Names
      desc: Classes with long names containing no special characters.
      contents:
        - documentstorewithvectorsearchcapabilities
        - EMBEDDINGPROVIDERWITHBATCHPROCESSINGSUPPORT
        - Chunkerstrategywithoverlapdetection
sidebar_filter:
  enabled: true
  min_items: 1