Files
e2db474707 release(text-splitters): 1.1.3 (#41001)
Bump `langchain-text-splitters` from 1.1.2 to 1.1.3. Changes since
[1.1.2](https://github.com/langchain-ai/langchain/releases/tag/langchain-text-splitters%3D%3D1.1.2)
are itemized below; dependency bumps and lockfile-only updates are
excluded.

### Fixes

- Restore lazy imports for heavy optional dependencies, with
import-isolation and missing-dependency regression coverage
([#35469](https://github.com/langchain-ai/langchain/pull/35469)).
- Raise a descriptive `TypeError` for unsupported
`RecursiveJsonSplitter` inputs rather than silently returning an empty
result; top-level lists require `convert_lists=True`, while `None` still
returns `[]`
([#39238](https://github.com/langchain-ai/langchain/pull/39238)).
- Remove invalid or duplicate Kotlin, Rust, and Haskell separators
([#37039](https://github.com/langchain-ai/langchain/pull/37039)).
- Remove incorrect C# `implements` and Elixir `while` separators
([#37037](https://github.com/langchain-ai/langchain/pull/37037)).
- Avoid `None` metadata keys when
`ExperimentalMarkdownSyntaxTextSplitter` has no header mapping
([#34545](https://github.com/langchain-ai/langchain/pull/34545)).
- Clarify the existing `HTMLHeaderTextSplitter.split_text_from_url`
deprecation warning: fetch HTML separately and use `split_text`
([#37164](https://github.com/langchain-ai/langchain/pull/37164)).

### Tooling, packaging, and documentation

- Replace `mypy` with `ty` and refactor tokenizer/HTML helpers. Also fix
`SentenceTransformersTokenTextSplitter` for models without a maximum
token limit: honor explicit `tokens_per_chunk`, or raise a clear
`ValueError` when omitted. Correct the GPT-4.1-mini encoding expectation
to `o200k_base`
([#38658](https://github.com/langchain-ai/langchain/pull/38658)).
- Update the token-splitter integration-test model from GPT-3.5 Turbo to
GPT-4.1-mini
([#38042](https://github.com/langchain-ai/langchain/pull/38042)).
- Tighten tokenizer/spaCy annotations and replace deprecated
`load_module()` in the import-check script with module-spec loading
([#40085](https://github.com/langchain-ai/langchain/pull/40085),
non-dependency changes only).
- Document existing support for `None` Markdown header names with a
targeted type-checker suppression; no runtime change
([#40566](https://github.com/langchain-ai/langchain/pull/40566),
non-dependency change only).

## References
- Slack thread:
https://langchain.slack.com/archives/C0C5950ARJT/p1790955367956069

Made by [Open SWE](https://github.com/langchain-ai/open-swe) · [view
thread](https://openswe.langchain.dev/agents/aadecb8c-db4f-537a-a174-fe630f165f81)
· openai:gpt-6-astra (medium)

Co-authored-by: Mason Daugherty <mdrxy@users.noreply.github.com>
Co-authored-by: open-swe[bot] <open-swe@users.noreply.github.com>
2026-10-02 11:44:24 -04:00

144 lines
4.5 KiB
TOML

[build-system]
requires = ["hatchling"]
build-backend = "hatchling.build"
[project]
name = "langchain-text-splitters"
description = "LangChain text splitting utilities"
license = { text = "MIT" }
readme = "README.md"
classifiers = [
"Development Status :: 5 - Production/Stable",
"Intended Audience :: Developers",
"License :: OSI Approved :: MIT License",
"Programming Language :: Python :: 3",
"Programming Language :: Python :: 3.10",
"Programming Language :: Python :: 3.11",
"Programming Language :: Python :: 3.12",
"Programming Language :: Python :: 3.13",
"Programming Language :: Python :: 3.14",
"Topic :: Scientific/Engineering :: Artificial Intelligence",
"Topic :: Software Development :: Libraries :: Python Modules",
"Topic :: Text Processing",
]
version = "1.1.3"
requires-python = ">=3.10.0,<4.0.0"
dependencies = [
"langchain-core>=1.4.7,<2.0.0",
]
[project.urls]
Homepage = "https://docs.langchain.com/"
Documentation = "https://docs.langchain.com/"
Repository = "https://github.com/langchain-ai/langchain"
Issues = "https://github.com/langchain-ai/langchain/issues"
Changelog = "https://github.com/langchain-ai/langchain/releases?q=%22langchain-text-splitters%22"
Twitter = "https://x.com/langchain_oss"
Slack = "https://www.langchain.com/join-community"
Reddit = "https://www.reddit.com/r/LangChain/"
[dependency-groups]
lint = [
"ruff>=0.15.0,<0.17.0",
]
typing = [
"ty>=0.0.56,<0.1.0",
"lxml-stubs>=0.5.1,<1.0.0",
"types-requests>=2.31.0.20240218,<3.0.0.0",
"spacy>=3.8.13,<4.0.0",
"nltk>=3.9.1,<4.0.0",
"transformers>=4.51.3,<6.0.0",
"sentence-transformers>=5.3.0,<7.0.0",
"tiktoken>=0.8.0,<1.0.0",
]
dev = [
"jupyter<2.0.0,>=1.0.0",
]
test = [
"pytest>=9.0.3,<10.0.0",
"freezegun>=1.2.2,<2.0.0",
"pytest-mock>=3.10.0,<4.0.0",
"pytest-watcher>=0.3.4,<1.0.0",
"pytest-asyncio>=1.3.0,<2.0.0",
"pytest-socket>=0.7.0,<1.0.0",
"pytest-xdist<4.0.0,>=3.6.1",
]
test_integration = [
"spacy>=3.8.13,!=3.8.14,<4.0.0",
"nltk>=3.9.1,<4.0.0",
"transformers>=4.51.3,<6.0.0",
"sentence-transformers>=5.3.0,<7.0.0",
"tiktoken>=0.8.0,<1.0.0",
"en-core-web-sm",
]
[tool.uv]
constraint-dependencies = ["pygments>=2.20.0"] # CVE-2026-4539
[tool.uv.sources]
langchain-core = { path = "../core", editable = true }
en-core-web-sm = { url = "https://github.com/explosion/spacy-models/releases/download/en_core_web_sm-3.8.0/en_core_web_sm-3.8.0-py3-none-any.whl" }
[tool.ty.rules]
all = "error"
[tool.ty.analysis]
allowed-unresolved-imports = ["konlpy.**"]
[tool.ruff.format]
docstring-code-format = true
[tool.ruff.lint]
select = [ "ALL",]
ignore = [
"C90", # McCabe complexity
"COM812", # Messes with the formatter
"CPY", # No copyright
"FIX002", # Line contains TODO
"PERF203", # Rarely useful
"PLR09", # Too many something (arg, statements, etc)
"TD002", # Missing author in TODO
"TD003", # Missing issue link in TODO
]
unfixable = [
"B028", # People should intentionally tune the stacklevel
]
flake8-annotations.allow-star-arg-any = true
flake8-annotations.mypy-init-return = true
flake8-type-checking.runtime-evaluated-base-classes = ["pydantic.BaseModel","langchain_core.load.serializable.Serializable","langchain_core.runnables.base.RunnableSerializable"]
pep8-naming.classmethod-decorators = [ "classmethod", "langchain_core.utils.pydantic.pre_init", "pydantic.field_validator", "pydantic.v1.root_validator",]
[tool.ruff.lint.pydocstyle]
convention = "google"
ignore-var-parameters = true # ignore missing documentation for *args and **kwargs parameters
[tool.ruff.lint.flake8-tidy-imports]
ban-relative-imports = "all"
[tool.ruff.lint.per-file-ignores]
"scripts/**" = [
"D1", # Docstrings not mandatory in scripts
"INP001", # Not a package
"S311" # Standard pseudo-random generators are not suitable for cryptographic purposes
]
"tests/**" = [
"D1", # Docstrings not mandatory in tests
"PLR2004", # Magic value comparisons
"S101", # Tests need assertions
"S311", # Standard pseudo-random generators are not suitable for cryptographic purposes
"SLF001" # Private member access in tests
]
[tool.coverage.run]
omit = ["tests/*"]
[tool.pytest.ini_options]
addopts = "--strict-markers --strict-config --durations=5"
markers = [
"requires: mark tests as requiring a specific library",
"compile: mark placeholder test used to compile integration tests without running them",
]
asyncio_mode = "auto"