PyPI - ultralytics-opencv-headless - Versions diffs - 8.3.242__py3-none-any.whl - Mend

ultralytics-opencv-headless 8.3.242__py3-none-any.whl

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.

Files changed (298) hide show

tests/__init__.py +23 -0
tests/conftest.py +59 -0
tests/test_cli.py +131 -0
tests/test_cuda.py +216 -0
tests/test_engine.py +157 -0
tests/test_exports.py +309 -0
tests/test_integrations.py +151 -0
tests/test_python.py +777 -0
tests/test_solutions.py +371 -0
ultralytics/__init__.py +48 -0
ultralytics/assets/bus.jpg +0 -0
ultralytics/assets/zidane.jpg +0 -0
ultralytics/cfg/__init__.py +1026 -0
ultralytics/cfg/datasets/Argoverse.yaml +78 -0
ultralytics/cfg/datasets/DOTAv1.5.yaml +37 -0
ultralytics/cfg/datasets/DOTAv1.yaml +36 -0
ultralytics/cfg/datasets/GlobalWheat2020.yaml +68 -0
ultralytics/cfg/datasets/HomeObjects-3K.yaml +32 -0
ultralytics/cfg/datasets/ImageNet.yaml +2025 -0
ultralytics/cfg/datasets/Objects365.yaml +447 -0
ultralytics/cfg/datasets/SKU-110K.yaml +58 -0
ultralytics/cfg/datasets/VOC.yaml +102 -0
ultralytics/cfg/datasets/VisDrone.yaml +87 -0
ultralytics/cfg/datasets/african-wildlife.yaml +25 -0
ultralytics/cfg/datasets/brain-tumor.yaml +22 -0
ultralytics/cfg/datasets/carparts-seg.yaml +44 -0
ultralytics/cfg/datasets/coco-pose.yaml +64 -0
ultralytics/cfg/datasets/coco.yaml +118 -0
ultralytics/cfg/datasets/coco128-seg.yaml +101 -0
ultralytics/cfg/datasets/coco128.yaml +101 -0
ultralytics/cfg/datasets/coco8-grayscale.yaml +103 -0
ultralytics/cfg/datasets/coco8-multispectral.yaml +104 -0
ultralytics/cfg/datasets/coco8-pose.yaml +47 -0
ultralytics/cfg/datasets/coco8-seg.yaml +101 -0
ultralytics/cfg/datasets/coco8.yaml +101 -0
ultralytics/cfg/datasets/construction-ppe.yaml +32 -0
ultralytics/cfg/datasets/crack-seg.yaml +22 -0
ultralytics/cfg/datasets/dog-pose.yaml +52 -0
ultralytics/cfg/datasets/dota8-multispectral.yaml +38 -0
ultralytics/cfg/datasets/dota8.yaml +35 -0
ultralytics/cfg/datasets/hand-keypoints.yaml +50 -0
ultralytics/cfg/datasets/kitti.yaml +27 -0
ultralytics/cfg/datasets/lvis.yaml +1240 -0
ultralytics/cfg/datasets/medical-pills.yaml +21 -0
ultralytics/cfg/datasets/open-images-v7.yaml +663 -0
ultralytics/cfg/datasets/package-seg.yaml +22 -0
ultralytics/cfg/datasets/signature.yaml +21 -0
ultralytics/cfg/datasets/tiger-pose.yaml +41 -0
ultralytics/cfg/datasets/xView.yaml +155 -0
ultralytics/cfg/default.yaml +130 -0
ultralytics/cfg/models/11/yolo11-cls-resnet18.yaml +17 -0
ultralytics/cfg/models/11/yolo11-cls.yaml +33 -0
ultralytics/cfg/models/11/yolo11-obb.yaml +50 -0
ultralytics/cfg/models/11/yolo11-pose.yaml +51 -0
ultralytics/cfg/models/11/yolo11-seg.yaml +50 -0
ultralytics/cfg/models/11/yolo11.yaml +50 -0
ultralytics/cfg/models/11/yoloe-11-seg.yaml +48 -0
ultralytics/cfg/models/11/yoloe-11.yaml +48 -0
ultralytics/cfg/models/12/yolo12-cls.yaml +32 -0
ultralytics/cfg/models/12/yolo12-obb.yaml +48 -0
ultralytics/cfg/models/12/yolo12-pose.yaml +49 -0
ultralytics/cfg/models/12/yolo12-seg.yaml +48 -0
ultralytics/cfg/models/12/yolo12.yaml +48 -0
ultralytics/cfg/models/rt-detr/rtdetr-l.yaml +53 -0
ultralytics/cfg/models/rt-detr/rtdetr-resnet101.yaml +45 -0
ultralytics/cfg/models/rt-detr/rtdetr-resnet50.yaml +45 -0
ultralytics/cfg/models/rt-detr/rtdetr-x.yaml +57 -0
ultralytics/cfg/models/v10/yolov10b.yaml +45 -0
ultralytics/cfg/models/v10/yolov10l.yaml +45 -0
ultralytics/cfg/models/v10/yolov10m.yaml +45 -0
ultralytics/cfg/models/v10/yolov10n.yaml +45 -0
ultralytics/cfg/models/v10/yolov10s.yaml +45 -0
ultralytics/cfg/models/v10/yolov10x.yaml +45 -0
ultralytics/cfg/models/v3/yolov3-spp.yaml +49 -0
ultralytics/cfg/models/v3/yolov3-tiny.yaml +40 -0
ultralytics/cfg/models/v3/yolov3.yaml +49 -0
ultralytics/cfg/models/v5/yolov5-p6.yaml +62 -0
ultralytics/cfg/models/v5/yolov5.yaml +51 -0
ultralytics/cfg/models/v6/yolov6.yaml +56 -0
ultralytics/cfg/models/v8/yoloe-v8-seg.yaml +48 -0
ultralytics/cfg/models/v8/yoloe-v8.yaml +48 -0
ultralytics/cfg/models/v8/yolov8-cls-resnet101.yaml +28 -0
ultralytics/cfg/models/v8/yolov8-cls-resnet50.yaml +28 -0
ultralytics/cfg/models/v8/yolov8-cls.yaml +32 -0
ultralytics/cfg/models/v8/yolov8-ghost-p2.yaml +58 -0
ultralytics/cfg/models/v8/yolov8-ghost-p6.yaml +60 -0
ultralytics/cfg/models/v8/yolov8-ghost.yaml +50 -0
ultralytics/cfg/models/v8/yolov8-obb.yaml +49 -0
ultralytics/cfg/models/v8/yolov8-p2.yaml +57 -0
ultralytics/cfg/models/v8/yolov8-p6.yaml +59 -0
ultralytics/cfg/models/v8/yolov8-pose-p6.yaml +60 -0
ultralytics/cfg/models/v8/yolov8-pose.yaml +50 -0
ultralytics/cfg/models/v8/yolov8-rtdetr.yaml +49 -0
ultralytics/cfg/models/v8/yolov8-seg-p6.yaml +59 -0
ultralytics/cfg/models/v8/yolov8-seg.yaml +49 -0
ultralytics/cfg/models/v8/yolov8-world.yaml +51 -0
ultralytics/cfg/models/v8/yolov8-worldv2.yaml +49 -0
ultralytics/cfg/models/v8/yolov8.yaml +49 -0
ultralytics/cfg/models/v9/yolov9c-seg.yaml +41 -0
ultralytics/cfg/models/v9/yolov9c.yaml +41 -0
ultralytics/cfg/models/v9/yolov9e-seg.yaml +64 -0
ultralytics/cfg/models/v9/yolov9e.yaml +64 -0
ultralytics/cfg/models/v9/yolov9m.yaml +41 -0
ultralytics/cfg/models/v9/yolov9s.yaml +41 -0
ultralytics/cfg/models/v9/yolov9t.yaml +41 -0
ultralytics/cfg/trackers/botsort.yaml +21 -0
ultralytics/cfg/trackers/bytetrack.yaml +12 -0
ultralytics/data/__init__.py +26 -0
ultralytics/data/annotator.py +66 -0
ultralytics/data/augment.py +2801 -0
ultralytics/data/base.py +435 -0
ultralytics/data/build.py +437 -0
ultralytics/data/converter.py +855 -0
ultralytics/data/dataset.py +834 -0
ultralytics/data/loaders.py +704 -0
ultralytics/data/scripts/download_weights.sh +18 -0
ultralytics/data/scripts/get_coco.sh +61 -0
ultralytics/data/scripts/get_coco128.sh +18 -0
ultralytics/data/scripts/get_imagenet.sh +52 -0
ultralytics/data/split.py +138 -0
ultralytics/data/split_dota.py +344 -0
ultralytics/data/utils.py +798 -0
ultralytics/engine/__init__.py +1 -0
ultralytics/engine/exporter.py +1574 -0
ultralytics/engine/model.py +1124 -0
ultralytics/engine/predictor.py +508 -0
ultralytics/engine/results.py +1522 -0
ultralytics/engine/trainer.py +974 -0
ultralytics/engine/tuner.py +448 -0
ultralytics/engine/validator.py +384 -0
ultralytics/hub/__init__.py +166 -0
ultralytics/hub/auth.py +151 -0
ultralytics/hub/google/__init__.py +174 -0
ultralytics/hub/session.py +422 -0
ultralytics/hub/utils.py +162 -0
ultralytics/models/__init__.py +9 -0
ultralytics/models/fastsam/__init__.py +7 -0
ultralytics/models/fastsam/model.py +79 -0
ultralytics/models/fastsam/predict.py +169 -0
ultralytics/models/fastsam/utils.py +23 -0
ultralytics/models/fastsam/val.py +38 -0
ultralytics/models/nas/__init__.py +7 -0
ultralytics/models/nas/model.py +98 -0
ultralytics/models/nas/predict.py +56 -0
ultralytics/models/nas/val.py +38 -0
ultralytics/models/rtdetr/__init__.py +7 -0
ultralytics/models/rtdetr/model.py +63 -0
ultralytics/models/rtdetr/predict.py +88 -0
ultralytics/models/rtdetr/train.py +89 -0
ultralytics/models/rtdetr/val.py +216 -0
ultralytics/models/sam/__init__.py +25 -0
ultralytics/models/sam/amg.py +275 -0
ultralytics/models/sam/build.py +365 -0
ultralytics/models/sam/build_sam3.py +377 -0
ultralytics/models/sam/model.py +169 -0
ultralytics/models/sam/modules/__init__.py +1 -0
ultralytics/models/sam/modules/blocks.py +1067 -0
ultralytics/models/sam/modules/decoders.py +495 -0
ultralytics/models/sam/modules/encoders.py +794 -0
ultralytics/models/sam/modules/memory_attention.py +298 -0
ultralytics/models/sam/modules/sam.py +1160 -0
ultralytics/models/sam/modules/tiny_encoder.py +979 -0
ultralytics/models/sam/modules/transformer.py +344 -0
ultralytics/models/sam/modules/utils.py +512 -0
ultralytics/models/sam/predict.py +3940 -0
ultralytics/models/sam/sam3/__init__.py +3 -0
ultralytics/models/sam/sam3/decoder.py +546 -0
ultralytics/models/sam/sam3/encoder.py +529 -0
ultralytics/models/sam/sam3/geometry_encoders.py +415 -0
ultralytics/models/sam/sam3/maskformer_segmentation.py +286 -0
ultralytics/models/sam/sam3/model_misc.py +199 -0
ultralytics/models/sam/sam3/necks.py +129 -0
ultralytics/models/sam/sam3/sam3_image.py +339 -0
ultralytics/models/sam/sam3/text_encoder_ve.py +307 -0
ultralytics/models/sam/sam3/vitdet.py +547 -0
ultralytics/models/sam/sam3/vl_combiner.py +160 -0
ultralytics/models/utils/__init__.py +1 -0
ultralytics/models/utils/loss.py +466 -0
ultralytics/models/utils/ops.py +315 -0
ultralytics/models/yolo/__init__.py +7 -0
ultralytics/models/yolo/classify/__init__.py +7 -0
ultralytics/models/yolo/classify/predict.py +90 -0
ultralytics/models/yolo/classify/train.py +202 -0
ultralytics/models/yolo/classify/val.py +216 -0
ultralytics/models/yolo/detect/__init__.py +7 -0
ultralytics/models/yolo/detect/predict.py +122 -0
ultralytics/models/yolo/detect/train.py +227 -0
ultralytics/models/yolo/detect/val.py +507 -0
ultralytics/models/yolo/model.py +430 -0
ultralytics/models/yolo/obb/__init__.py +7 -0
ultralytics/models/yolo/obb/predict.py +56 -0
ultralytics/models/yolo/obb/train.py +79 -0
ultralytics/models/yolo/obb/val.py +302 -0
ultralytics/models/yolo/pose/__init__.py +7 -0
ultralytics/models/yolo/pose/predict.py +65 -0
ultralytics/models/yolo/pose/train.py +110 -0
ultralytics/models/yolo/pose/val.py +248 -0
ultralytics/models/yolo/segment/__init__.py +7 -0
ultralytics/models/yolo/segment/predict.py +109 -0
ultralytics/models/yolo/segment/train.py +69 -0
ultralytics/models/yolo/segment/val.py +307 -0
ultralytics/models/yolo/world/__init__.py +5 -0
ultralytics/models/yolo/world/train.py +173 -0
ultralytics/models/yolo/world/train_world.py +178 -0
ultralytics/models/yolo/yoloe/__init__.py +22 -0
ultralytics/models/yolo/yoloe/predict.py +162 -0
ultralytics/models/yolo/yoloe/train.py +287 -0
ultralytics/models/yolo/yoloe/train_seg.py +122 -0
ultralytics/models/yolo/yoloe/val.py +206 -0
ultralytics/nn/__init__.py +27 -0
ultralytics/nn/autobackend.py +958 -0
ultralytics/nn/modules/__init__.py +182 -0
ultralytics/nn/modules/activation.py +54 -0
ultralytics/nn/modules/block.py +1947 -0
ultralytics/nn/modules/conv.py +669 -0
ultralytics/nn/modules/head.py +1183 -0
ultralytics/nn/modules/transformer.py +793 -0
ultralytics/nn/modules/utils.py +159 -0
ultralytics/nn/tasks.py +1768 -0
ultralytics/nn/text_model.py +356 -0
ultralytics/py.typed +1 -0
ultralytics/solutions/__init__.py +41 -0
ultralytics/solutions/ai_gym.py +108 -0
ultralytics/solutions/analytics.py +264 -0
ultralytics/solutions/config.py +107 -0
ultralytics/solutions/distance_calculation.py +123 -0
ultralytics/solutions/heatmap.py +125 -0
ultralytics/solutions/instance_segmentation.py +86 -0
ultralytics/solutions/object_blurrer.py +89 -0
ultralytics/solutions/object_counter.py +190 -0
ultralytics/solutions/object_cropper.py +87 -0
ultralytics/solutions/parking_management.py +280 -0
ultralytics/solutions/queue_management.py +93 -0
ultralytics/solutions/region_counter.py +133 -0
ultralytics/solutions/security_alarm.py +151 -0
ultralytics/solutions/similarity_search.py +219 -0
ultralytics/solutions/solutions.py +828 -0
ultralytics/solutions/speed_estimation.py +114 -0
ultralytics/solutions/streamlit_inference.py +260 -0
ultralytics/solutions/templates/similarity-search.html +156 -0
ultralytics/solutions/trackzone.py +88 -0
ultralytics/solutions/vision_eye.py +67 -0
ultralytics/trackers/__init__.py +7 -0
ultralytics/trackers/basetrack.py +115 -0
ultralytics/trackers/bot_sort.py +257 -0
ultralytics/trackers/byte_tracker.py +469 -0
ultralytics/trackers/track.py +116 -0
ultralytics/trackers/utils/__init__.py +1 -0
ultralytics/trackers/utils/gmc.py +339 -0
ultralytics/trackers/utils/kalman_filter.py +482 -0
ultralytics/trackers/utils/matching.py +154 -0
ultralytics/utils/__init__.py +1450 -0
ultralytics/utils/autobatch.py +118 -0
ultralytics/utils/autodevice.py +205 -0
ultralytics/utils/benchmarks.py +728 -0
ultralytics/utils/callbacks/__init__.py +5 -0
ultralytics/utils/callbacks/base.py +233 -0
ultralytics/utils/callbacks/clearml.py +146 -0
ultralytics/utils/callbacks/comet.py +625 -0
ultralytics/utils/callbacks/dvc.py +197 -0
ultralytics/utils/callbacks/hub.py +110 -0
ultralytics/utils/callbacks/mlflow.py +134 -0
ultralytics/utils/callbacks/neptune.py +126 -0
ultralytics/utils/callbacks/platform.py +73 -0
ultralytics/utils/callbacks/raytune.py +42 -0
ultralytics/utils/callbacks/tensorboard.py +123 -0
ultralytics/utils/callbacks/wb.py +188 -0
ultralytics/utils/checks.py +998 -0
ultralytics/utils/cpu.py +85 -0
ultralytics/utils/dist.py +123 -0
ultralytics/utils/downloads.py +529 -0
ultralytics/utils/errors.py +35 -0
ultralytics/utils/events.py +113 -0
ultralytics/utils/export/__init__.py +7 -0
ultralytics/utils/export/engine.py +237 -0
ultralytics/utils/export/imx.py +315 -0
ultralytics/utils/export/tensorflow.py +231 -0
ultralytics/utils/files.py +219 -0
ultralytics/utils/git.py +137 -0
ultralytics/utils/instance.py +484 -0
ultralytics/utils/logger.py +444 -0
ultralytics/utils/loss.py +849 -0
ultralytics/utils/metrics.py +1560 -0
ultralytics/utils/nms.py +337 -0
ultralytics/utils/ops.py +664 -0
ultralytics/utils/patches.py +201 -0
ultralytics/utils/plotting.py +1045 -0
ultralytics/utils/tal.py +403 -0
ultralytics/utils/torch_utils.py +984 -0
ultralytics/utils/tqdm.py +440 -0
ultralytics/utils/triton.py +112 -0
ultralytics/utils/tuner.py +160 -0
ultralytics_opencv_headless-8.3.242.dist-info/METADATA +374 -0
ultralytics_opencv_headless-8.3.242.dist-info/RECORD +298 -0
ultralytics_opencv_headless-8.3.242.dist-info/WHEEL +5 -0
ultralytics_opencv_headless-8.3.242.dist-info/entry_points.txt +3 -0
ultralytics_opencv_headless-8.3.242.dist-info/licenses/LICENSE +661 -0
ultralytics_opencv_headless-8.3.242.dist-info/top_level.txt +1 -0

ultralytics/data/scripts/download_weights.sh ADDED Viewed

@@ -0,0 +1,18 @@
+#!/bin/bash
+# Ultralytics 🚀 AGPL-3.0 License - https://ultralytics.com/license
+# Download latest models from https://github.com/ultralytics/assets/releases
+# Example usage: bash ultralytics/data/scripts/download_weights.sh
+# parent
+# └── weights
+#     ├── yolov8n.pt  ← downloads here
+#     ├── yolov8s.pt
+#     └── ...
+python << EOF
+from ultralytics.utils.downloads import attempt_download_asset
+assets = [f"yolov8{size}{suffix}.pt" for size in "nsmlx" for suffix in ("", "-cls", "-seg", "-pose")]
+for x in assets:
+    attempt_download_asset(f"weights/{x}")
+EOF

ultralytics/data/scripts/get_coco.sh ADDED Viewed

@@ -0,0 +1,61 @@
+#!/bin/bash
+# Ultralytics 🚀 AGPL-3.0 License - https://ultralytics.com/license
+# Download COCO 2017 dataset https://cocodataset.org
+# Example usage: bash data/scripts/get_coco.sh
+# parent
+# ├── ultralytics
+# └── datasets
+#     └── coco  ← downloads here
+# Arguments (optional) Usage: bash data/scripts/get_coco.sh --train --val --test --segments
+if [ "$#" -gt 0 ]; then
+  for opt in "$@"; do
+    case "${opt}" in
+      --train) train=true ;;
+      --val) val=true ;;
+      --test) test=true ;;
+      --segments) segments=true ;;
+      --sama) sama=true ;;
+    esac
+  done
+else
+  train=true
+  val=true
+  test=false
+  segments=false
+  sama=false
+fi
+# Download/unzip labels
+d='../datasets' # unzip directory
+url=https://github.com/ultralytics/assets/releases/download/v0.0.0/
+if [ "$segments" == "true" ]; then
+  f='coco2017labels-segments.zip' # 169 MB
+elif [ "$sama" == "true" ]; then
+  f='coco2017labels-segments-sama.zip' # 199 MB https://www.sama.com/sama-coco-dataset/
+else
+  f='coco2017labels.zip' # 46 MB
+fi
+echo 'Downloading' $url$f ' ...'
+curl -L $url$f -o $f -# && unzip -q $f -d $d && rm $f &
+# Download/unzip images
+d='../datasets/coco/images' # unzip directory
+url=http://images.cocodataset.org/zips/
+if [ "$train" == "true" ]; then
+  f='train2017.zip' # 19G, 118k images
+  echo 'Downloading' $url$f '...'
+  curl -L $url$f -o $f -# && unzip -q $f -d $d && rm $f &
+fi
+if [ "$val" == "true" ]; then
+  f='val2017.zip' # 1G, 5k images
+  echo 'Downloading' $url$f '...'
+  curl -L $url$f -o $f -# && unzip -q $f -d $d && rm $f &
+fi
+if [ "$test" == "true" ]; then
+  f='test2017.zip' # 7G, 41k images (optional)
+  echo 'Downloading' $url$f '...'
+  curl -L $url$f -o $f -# && unzip -q $f -d $d && rm $f &
+fi
+wait # finish background tasks

ultralytics/data/scripts/get_coco128.sh ADDED Viewed

@@ -0,0 +1,18 @@
+#!/bin/bash
+# Ultralytics 🚀 AGPL-3.0 License - https://ultralytics.com/license
+# Download COCO128 dataset https://www.kaggle.com/ultralytics/coco128 (first 128 images from COCO train2017)
+# Example usage: bash data/scripts/get_coco128.sh
+# parent
+# ├── ultralytics
+# └── datasets
+#     └── coco128  ← downloads here
+# Download/unzip images and labels
+d='../datasets' # unzip directory
+url=https://github.com/ultralytics/assets/releases/download/v0.0.0/
+f='coco128.zip' # or 'coco128-segments.zip', 68 MB
+echo 'Downloading' $url$f ' ...'
+curl -L $url$f -o $f -# && unzip -q $f -d $d && rm $f &
+wait # finish background tasks

ultralytics/data/scripts/get_imagenet.sh ADDED Viewed

@@ -0,0 +1,52 @@
+#!/bin/bash
+# Ultralytics 🚀 AGPL-3.0 License - https://ultralytics.com/license
+# Download ILSVRC2012 ImageNet dataset https://image-net.org
+# Example usage: bash data/scripts/get_imagenet.sh
+# parent
+# ├── ultralytics
+# └── datasets
+#     └── imagenet  ← downloads here
+# Arguments (optional) Usage: bash data/scripts/get_imagenet.sh --train --val
+if [ "$#" -gt 0 ]; then
+  for opt in "$@"; do
+    case "${opt}" in
+      --train) train=true ;;
+      --val) val=true ;;
+    esac
+  done
+else
+  train=true
+  val=true
+fi
+# Make dir
+d='../datasets/imagenet' # unzip directory
+mkdir -p $d && cd $d
+# Download/unzip train
+if [ "$train" == "true" ]; then
+  wget https://image-net.org/data/ILSVRC/2012/ILSVRC2012_img_train.tar # download 138G, 1281167 images
+  mkdir train && mv ILSVRC2012_img_train.tar train/ && cd train
+  tar -xf ILSVRC2012_img_train.tar && rm -f ILSVRC2012_img_train.tar
+  find . -name "*.tar" | while read NAME; do
+    mkdir -p "${NAME%.tar}"
+    tar -xf "${NAME}" -C "${NAME%.tar}"
+    rm -f "${NAME}"
+  done
+  cd ..
+fi
+# Download/unzip val
+if [ "$val" == "true" ]; then
+  wget https://image-net.org/data/ILSVRC/2012/ILSVRC2012_img_val.tar # download 6.3G, 50000 images
+  mkdir val && mv ILSVRC2012_img_val.tar val/ && cd val && tar -xf ILSVRC2012_img_val.tar
+  wget -qO- https://raw.githubusercontent.com/soumith/imagenetloader.torch/master/valprep.sh | bash # move into subdirs
+fi
+# Delete corrupted image (optional: PNG under JPEG name that may cause dataloaders to fail)
+# rm train/n04266014/n04266014_10835.JPEG
+# TFRecords (optional)
+# wget https://raw.githubusercontent.com/tensorflow/models/master/research/slim/datasets/imagenet_lsvrc_2015_synsets.txt

ultralytics/data/split.py ADDED Viewed

@@ -0,0 +1,138 @@
+# Ultralytics 🚀 AGPL-3.0 License - https://ultralytics.com/license
+from __future__ import annotations
+import random
+import shutil
+from pathlib import Path
+from ultralytics.data.utils import IMG_FORMATS, img2label_paths
+from ultralytics.utils import DATASETS_DIR, LOGGER, TQDM
+def split_classify_dataset(source_dir: str | Path, train_ratio: float = 0.8) -> Path:
+    """Split classification dataset into train and val directories in a new directory.
+    Creates a new directory '{source_dir}_split' with train/val subdirectories, preserving the original class structure
+    with an 80/20 split by default.
+    Directory structure:
+        Before:
+            caltech/
+            ├── class1/
+            │   ├── img1.jpg
+            │   ├── img2.jpg
+            │   └── ...
+            ├── class2/
+            │   ├── img1.jpg
+            │   └── ...
+            └── ...
+        After:
+            caltech_split/
+            ├── train/
+            │   ├── class1/
+            │   │   ├── img1.jpg
+            │   │   └── ...
+            │   ├── class2/
+            │   │   ├── img1.jpg
+            │   │   └── ...
+            │   └── ...
+            └── val/
+                ├── class1/
+                │   ├── img2.jpg
+                │   └── ...
+                ├── class2/
+                │   └── ...
+                └── ...
+    Args:
+        source_dir (str | Path): Path to classification dataset root directory.
+        train_ratio (float): Ratio for train split, between 0 and 1.
+    Returns:
+        (Path): Path to the created split directory.
+    Examples:
+        Split dataset with default 80/20 ratio
+        >>> split_classify_dataset("path/to/caltech")
+        Split with custom ratio
+        >>> split_classify_dataset("path/to/caltech", 0.75)
+    """
+    source_path = Path(source_dir)
+    split_path = Path(f"{source_path}_split")
+    train_path, val_path = split_path / "train", split_path / "val"
+    # Create directory structure
+    split_path.mkdir(exist_ok=True)
+    train_path.mkdir(exist_ok=True)
+    val_path.mkdir(exist_ok=True)
+    # Process class directories
+    class_dirs = [d for d in source_path.iterdir() if d.is_dir()]
+    total_images = sum(len(list(d.glob("*.*"))) for d in class_dirs)
+    stats = f"{len(class_dirs)} classes, {total_images} images"
+    LOGGER.info(f"Splitting {source_path} ({stats}) into {train_ratio:.0%} train, {1 - train_ratio:.0%} val...")
+    for class_dir in class_dirs:
+        # Create class directories
+        (train_path / class_dir.name).mkdir(exist_ok=True)
+        (val_path / class_dir.name).mkdir(exist_ok=True)
+        # Split and copy files
+        image_files = list(class_dir.glob("*.*"))
+        random.shuffle(image_files)
+        split_idx = int(len(image_files) * train_ratio)
+        for img in image_files[:split_idx]:
+            shutil.copy2(img, train_path / class_dir.name / img.name)
+        for img in image_files[split_idx:]:
+            shutil.copy2(img, val_path / class_dir.name / img.name)
+    LOGGER.info(f"Split complete in {split_path} ✅")
+    return split_path
+def autosplit(
+    path: Path = DATASETS_DIR / "coco8/images",
+    weights: tuple[float, float, float] = (0.9, 0.1, 0.0),
+    annotated_only: bool = False,
+) -> None:
+    """Automatically split a dataset into train/val/test splits and save the resulting splits into autosplit_*.txt
+    files.
+    Args:
+        path (Path): Path to images directory.
+        weights (tuple): Train, validation, and test split fractions.
+        annotated_only (bool): If True, only images with an associated txt file are used.
+    Examples:
+        Split images with default weights
+        >>> from ultralytics.data.split import autosplit
+        >>> autosplit()
+        Split with custom weights and annotated images only
+        >>> autosplit(path="path/to/images", weights=(0.8, 0.15, 0.05), annotated_only=True)
+    """
+    path = Path(path)  # images dir
+    files = sorted(x for x in path.rglob("*.*") if x.suffix[1:].lower() in IMG_FORMATS)  # image files only
+    n = len(files)  # number of files
+    random.seed(0)  # for reproducibility
+    indices = random.choices([0, 1, 2], weights=weights, k=n)  # assign each image to a split
+    txt = ["autosplit_train.txt", "autosplit_val.txt", "autosplit_test.txt"]  # 3 txt files
+    for x in txt:
+        if (path.parent / x).exists():
+            (path.parent / x).unlink()  # remove existing
+    LOGGER.info(f"Autosplitting images from {path}" + ", using *.txt labeled images only" * annotated_only)
+    for i, img in TQDM(zip(indices, files), total=n):
+        if not annotated_only or Path(img2label_paths([str(img)])[0]).exists():  # check label
+            with open(path.parent / txt[i], "a", encoding="utf-8") as f:
+                f.write(f"./{img.relative_to(path.parent).as_posix()}" + "\n")  # add image to txt file
+if __name__ == "__main__":
+    split_classify_dataset("caltech101")

ultralytics/data/split_dota.py ADDED Viewed

@@ -0,0 +1,344 @@
+# Ultralytics 🚀 AGPL-3.0 License - https://ultralytics.com/license
+from __future__ import annotations
+import itertools
+from glob import glob
+from math import ceil
+from pathlib import Path
+from typing import Any
+import cv2
+import numpy as np
+from PIL import Image
+from ultralytics.data.utils import exif_size, img2label_paths
+from ultralytics.utils import TQDM
+from ultralytics.utils.checks import check_requirements
+def bbox_iof(polygon1: np.ndarray, bbox2: np.ndarray, eps: float = 1e-6) -> np.ndarray:
+    """Calculate Intersection over Foreground (IoF) between polygons and bounding boxes.
+    Args:
+        polygon1 (np.ndarray): Polygon coordinates with shape (N, 8).
+        bbox2 (np.ndarray): Bounding boxes with shape (N, 4).
+        eps (float, optional): Small value to prevent division by zero.
+    Returns:
+        (np.ndarray): IoF scores with shape (N, 1) or (N, M) if bbox2 is (M, 4).
+    Notes:
+        Polygon format: [x1, y1, x2, y2, x3, y3, x4, y4].
+        Bounding box format: [x_min, y_min, x_max, y_max].
+    """
+    check_requirements("shapely>=2.0.0")
+    from shapely.geometry import Polygon
+    polygon1 = polygon1.reshape(-1, 4, 2)
+    lt_point = np.min(polygon1, axis=-2)  # left-top
+    rb_point = np.max(polygon1, axis=-2)  # right-bottom
+    bbox1 = np.concatenate([lt_point, rb_point], axis=-1)
+    lt = np.maximum(bbox1[:, None, :2], bbox2[..., :2])
+    rb = np.minimum(bbox1[:, None, 2:], bbox2[..., 2:])
+    wh = np.clip(rb - lt, 0, np.inf)
+    h_overlaps = wh[..., 0] * wh[..., 1]
+    left, top, right, bottom = (bbox2[..., i] for i in range(4))
+    polygon2 = np.stack([left, top, right, top, right, bottom, left, bottom], axis=-1).reshape(-1, 4, 2)
+    sg_polys1 = [Polygon(p) for p in polygon1]
+    sg_polys2 = [Polygon(p) for p in polygon2]
+    overlaps = np.zeros(h_overlaps.shape)
+    for p in zip(*np.nonzero(h_overlaps)):
+        overlaps[p] = sg_polys1[p[0]].intersection(sg_polys2[p[-1]]).area
+    unions = np.array([p.area for p in sg_polys1], dtype=np.float32)
+    unions = unions[..., None]
+    unions = np.clip(unions, eps, np.inf)
+    outputs = overlaps / unions
+    if outputs.ndim == 1:
+        outputs = outputs[..., None]
+    return outputs
+def load_yolo_dota(data_root: str, split: str = "train") -> list[dict[str, Any]]:
+    """Load DOTA dataset annotations and image information.
+    Args:
+        data_root (str): Data root directory.
+        split (str, optional): The split data set, could be 'train' or 'val'.
+    Returns:
+        (list[dict[str, Any]]): List of annotation dictionaries containing image information.
+    Notes:
+        The directory structure assumed for the DOTA dataset:
+            - data_root
+                - images
+                    - train
+                    - val
+                - labels
+                    - train
+                    - val
+    """
+    assert split in {"train", "val"}, f"Split must be 'train' or 'val', not {split}."
+    im_dir = Path(data_root) / "images" / split
+    assert im_dir.exists(), f"Can't find {im_dir}, please check your data root."
+    im_files = glob(str(Path(data_root) / "images" / split / "*"))
+    lb_files = img2label_paths(im_files)
+    annos = []
+    for im_file, lb_file in zip(im_files, lb_files):
+        w, h = exif_size(Image.open(im_file))
+        with open(lb_file, encoding="utf-8") as f:
+            lb = [x.split() for x in f.read().strip().splitlines() if len(x)]
+            lb = np.array(lb, dtype=np.float32)
+        annos.append(dict(ori_size=(h, w), label=lb, filepath=im_file))
+    return annos
+def get_windows(
+    im_size: tuple[int, int],
+    crop_sizes: tuple[int, ...] = (1024,),
+    gaps: tuple[int, ...] = (200,),
+    im_rate_thr: float = 0.6,
+    eps: float = 0.01,
+) -> np.ndarray:
+    """Get the coordinates of sliding windows for image cropping.
+    Args:
+        im_size (tuple[int, int]): Original image size, (H, W).
+        crop_sizes (tuple[int, ...], optional): Crop size of windows.
+        gaps (tuple[int, ...], optional): Gap between crops.
+        im_rate_thr (float, optional): Threshold of windows areas divided by image areas.
+        eps (float, optional): Epsilon value for math operations.
+    Returns:
+        (np.ndarray): Array of window coordinates of shape (N, 4) where each row is [x_start, y_start, x_stop, y_stop].
+    """
+    h, w = im_size
+    windows = []
+    for crop_size, gap in zip(crop_sizes, gaps):
+        assert crop_size > gap, f"invalid crop_size gap pair [{crop_size} {gap}]"
+        step = crop_size - gap
+        xn = 1 if w <= crop_size else ceil((w - crop_size) / step + 1)
+        xs = [step * i for i in range(xn)]
+        if len(xs) > 1 and xs[-1] + crop_size > w:
+            xs[-1] = w - crop_size
+        yn = 1 if h <= crop_size else ceil((h - crop_size) / step + 1)
+        ys = [step * i for i in range(yn)]
+        if len(ys) > 1 and ys[-1] + crop_size > h:
+            ys[-1] = h - crop_size
+        start = np.array(list(itertools.product(xs, ys)), dtype=np.int64)
+        stop = start + crop_size
+        windows.append(np.concatenate([start, stop], axis=1))
+    windows = np.concatenate(windows, axis=0)
+    im_in_wins = windows.copy()
+    im_in_wins[:, 0::2] = np.clip(im_in_wins[:, 0::2], 0, w)
+    im_in_wins[:, 1::2] = np.clip(im_in_wins[:, 1::2], 0, h)
+    im_areas = (im_in_wins[:, 2] - im_in_wins[:, 0]) * (im_in_wins[:, 3] - im_in_wins[:, 1])
+    win_areas = (windows[:, 2] - windows[:, 0]) * (windows[:, 3] - windows[:, 1])
+    im_rates = im_areas / win_areas
+    if not (im_rates > im_rate_thr).any():
+        max_rate = im_rates.max()
+        im_rates[abs(im_rates - max_rate) < eps] = 1
+    return windows[im_rates > im_rate_thr]
+def get_window_obj(anno: dict[str, Any], windows: np.ndarray, iof_thr: float = 0.7) -> list[np.ndarray]:
+    """Get objects for each window based on IoF threshold."""
+    h, w = anno["ori_size"]
+    label = anno["label"]
+    if len(label):
+        label[:, 1::2] *= w
+        label[:, 2::2] *= h
+        iofs = bbox_iof(label[:, 1:], windows)
+        # Unnormalized and misaligned coordinates
+        return [(label[iofs[:, i] >= iof_thr]) for i in range(len(windows))]  # window_anns
+    else:
+        return [np.zeros((0, 9), dtype=np.float32) for _ in range(len(windows))]  # window_anns
+def crop_and_save(
+    anno: dict[str, Any],
+    windows: np.ndarray,
+    window_objs: list[np.ndarray],
+    im_dir: str,
+    lb_dir: str,
+    allow_background_images: bool = True,
+) -> None:
+    """Crop images and save new labels for each window.
+    Args:
+        anno (dict[str, Any]): Annotation dict, including 'filepath', 'label', 'ori_size' as its keys.
+        windows (np.ndarray): Array of windows coordinates with shape (N, 4).
+        window_objs (list[np.ndarray]): A list of labels inside each window.
+        im_dir (str): The output directory path of images.
+        lb_dir (str): The output directory path of labels.
+        allow_background_images (bool, optional): Whether to include background images without labels.
+    Notes:
+        The directory structure assumed for the DOTA dataset:
+            - data_root
+                - images
+                    - train
+                    - val
+                - labels
+                    - train
+                    - val
+    """
+    im = cv2.imread(anno["filepath"])
+    name = Path(anno["filepath"]).stem
+    for i, window in enumerate(windows):
+        x_start, y_start, x_stop, y_stop = window.tolist()
+        new_name = f"{name}__{x_stop - x_start}__{x_start}___{y_start}"
+        patch_im = im[y_start:y_stop, x_start:x_stop]
+        ph, pw = patch_im.shape[:2]
+        label = window_objs[i]
+        if len(label) or allow_background_images:
+            cv2.imwrite(str(Path(im_dir) / f"{new_name}.jpg"), patch_im)
+        if len(label):
+            label[:, 1::2] -= x_start
+            label[:, 2::2] -= y_start
+            label[:, 1::2] /= pw
+            label[:, 2::2] /= ph
+            with open(Path(lb_dir) / f"{new_name}.txt", "w", encoding="utf-8") as f:
+                for lb in label:
+                    formatted_coords = [f"{coord:.6g}" for coord in lb[1:]]
+                    f.write(f"{int(lb[0])} {' '.join(formatted_coords)}\n")
+def split_images_and_labels(
+    data_root: str,
+    save_dir: str,
+    split: str = "train",
+    crop_sizes: tuple[int, ...] = (1024,),
+    gaps: tuple[int, ...] = (200,),
+) -> None:
+    """Split both images and labels for a given dataset split.
+    Args:
+        data_root (str): Root directory of the dataset.
+        save_dir (str): Directory to save the split dataset.
+        split (str, optional): The split data set, could be 'train' or 'val'.
+        crop_sizes (tuple[int, ...], optional): Tuple of crop sizes.
+        gaps (tuple[int, ...], optional): Tuple of gaps between crops.
+    Notes:
+        The directory structure assumed for the DOTA dataset:
+            - data_root
+                - images
+                    - split
+                - labels
+                    - split
+        and the output directory structure is:
+            - save_dir
+                - images
+                    - split
+                - labels
+                    - split
+    """
+    im_dir = Path(save_dir) / "images" / split
+    im_dir.mkdir(parents=True, exist_ok=True)
+    lb_dir = Path(save_dir) / "labels" / split
+    lb_dir.mkdir(parents=True, exist_ok=True)
+    annos = load_yolo_dota(data_root, split=split)
+    for anno in TQDM(annos, total=len(annos), desc=split):
+        windows = get_windows(anno["ori_size"], crop_sizes, gaps)
+        window_objs = get_window_obj(anno, windows)
+        crop_and_save(anno, windows, window_objs, str(im_dir), str(lb_dir))
+def split_trainval(
+    data_root: str, save_dir: str, crop_size: int = 1024, gap: int = 200, rates: tuple[float, ...] = (1.0,)
+) -> None:
+    """Split train and val sets of DOTA dataset with multiple scaling rates.
+    Args:
+        data_root (str): Root directory of the dataset.
+        save_dir (str): Directory to save the split dataset.
+        crop_size (int, optional): Base crop size.
+        gap (int, optional): Base gap between crops.
+        rates (tuple[float, ...], optional): Scaling rates for crop_size and gap.
+    Notes:
+        The directory structure assumed for the DOTA dataset:
+            - data_root
+                - images
+                    - train
+                    - val
+                - labels
+                    - train
+                    - val
+        and the output directory structure is:
+            - save_dir
+                - images
+                    - train
+                    - val
+                - labels
+                    - train
+                    - val
+    """
+    crop_sizes, gaps = [], []
+    for r in rates:
+        crop_sizes.append(int(crop_size / r))
+        gaps.append(int(gap / r))
+    for split in {"train", "val"}:
+        split_images_and_labels(data_root, save_dir, split, crop_sizes, gaps)
+def split_test(
+    data_root: str, save_dir: str, crop_size: int = 1024, gap: int = 200, rates: tuple[float, ...] = (1.0,)
+) -> None:
+    """Split test set of DOTA dataset, labels are not included within this set.
+    Args:
+        data_root (str): Root directory of the dataset.
+        save_dir (str): Directory to save the split dataset.
+        crop_size (int, optional): Base crop size.
+        gap (int, optional): Base gap between crops.
+        rates (tuple[float, ...], optional): Scaling rates for crop_size and gap.
+    Notes:
+        The directory structure assumed for the DOTA dataset:
+            - data_root
+                - images
+                    - test
+        and the output directory structure is:
+            - save_dir
+                - images
+                    - test
+    """
+    crop_sizes, gaps = [], []
+    for r in rates:
+        crop_sizes.append(int(crop_size / r))
+        gaps.append(int(gap / r))
+    save_dir = Path(save_dir) / "images" / "test"
+    save_dir.mkdir(parents=True, exist_ok=True)
+    im_dir = Path(data_root) / "images" / "test"
+    assert im_dir.exists(), f"Can't find {im_dir}, please check your data root."
+    im_files = glob(str(im_dir / "*"))
+    for im_file in TQDM(im_files, total=len(im_files), desc="test"):
+        w, h = exif_size(Image.open(im_file))
+        windows = get_windows((h, w), crop_sizes=crop_sizes, gaps=gaps)
+        im = cv2.imread(im_file)
+        name = Path(im_file).stem
+        for window in windows:
+            x_start, y_start, x_stop, y_stop = window.tolist()
+            new_name = f"{name}__{x_stop - x_start}__{x_start}___{y_start}"
+            patch_im = im[y_start:y_stop, x_start:x_stop]
+            cv2.imwrite(str(save_dir / f"{new_name}.jpg"), patch_im)
+if __name__ == "__main__":
+    split_trainval(data_root="DOTAv2", save_dir="DOTAv2-split")
+    split_test(data_root="DOTAv2", save_dir="DOTAv2-split")