{"metadata":{"kernelspec":{"language":"python","display_name":"Python 3","name":"python3"},"language_info":{"name":"python","version":"3.11.13","mimetype":"text/x-python","codemirror_mode":{"name":"ipython","version":3},"pygments_lexer":"ipython3","nbconvert_exporter":"python","file_extension":".py"},"kaggle":{"accelerator":"nvidiaTeslaT4","dataSources":[{"sourceId":4104,"databundleVersionId":46661,"sourceType":"competition"},{"sourceId":7866129,"sourceType":"datasetVersion","datasetId":4614938}],"dockerImageVersionId":31090,"isInternetEnabled":true,"language":"python","sourceType":"notebook","isGpuEnabled":true}},"nbformat_minor":4,"nbformat":4,"cells":[{"cell_type":"code","source":"!nvidia-smi","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-06T08:47:02.241761Z","iopub.execute_input":"2025-08-06T08:47:02.242472Z","iopub.status.idle":"2025-08-06T08:47:02.471999Z","shell.execute_reply.started":"2025-08-06T08:47:02.242438Z","shell.execute_reply":"2025-08-06T08:47:02.471119Z"}},"outputs":[],"execution_count":null},{"cell_type":"markdown","source":"# Libraries","metadata":{}},{"cell_type":"code","source":"import pandas as pd\nfrom glob import glob\nimport matplotlib.pyplot as plt\nimport cv2\nimport tensorflow as tf\nimport numpy as np\nimport os\nfrom sklearn.model_selection import train_test_split\nfrom tensorflow.keras import layers, models, callbacks, mixed_precision","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-06T08:47:03.329608Z","iopub.execute_input":"2025-08-06T08:47:03.330354Z","iopub.status.idle":"2025-08-06T08:47:24.921157Z","shell.execute_reply.started":"2025-08-06T08:47:03.330322Z","shell.execute_reply":"2025-08-06T08:47:24.920272Z"}},"outputs":[],"execution_count":null},{"cell_type":"code","source":"import wandb\nfrom kaggle_secrets import UserSecretsClient\nuser_secrets = UserSecretsClient()\nsecret_value_0 = user_secrets.get_secret(\"WANDB_API_KEY\")\n\nwandb.login(key=secret_value_0)","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-06T08:47:24.922368Z","iopub.execute_input":"2025-08-06T08:47:24.923006Z","iopub.status.idle":"2025-08-06T08:47:35.186789Z","shell.execute_reply.started":"2025-08-06T08:47:24.922980Z","shell.execute_reply":"2025-08-06T08:47:35.185886Z"}},"outputs":[],"execution_count":null},{"cell_type":"code","source":"# import tensorflow as tf\nprint(tf.__version__)\n# Check if a GPU is available\ngpus = tf.config.list_physical_devices('GPU')\nif gpus:\n    try:\n        # Set memory growth to prevent TensorFlow from allocating all memory at once\n        for gpu in gpus:\n            tf.config.experimental.set_memory_growth(gpu, True)\n        print(\"GPU is available and will be used.\")\n    except RuntimeError as e:\n        # Memory growth must be set before GPUs have been initialized\n        print(f\"Error: {e}\")\nelse:\n    print(\"GPU is not available, running on CPU.\")\n\ntf.test.is_gpu_available()\n\npolicy = mixed_precision.Policy('mixed_float16')\nmixed_precision.set_global_policy(policy)\nprint(\"Mixed precision enabled:\", policy)","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-06T08:47:35.187745Z","iopub.execute_input":"2025-08-06T08:47:35.188697Z","iopub.status.idle":"2025-08-06T08:47:36.784120Z","shell.execute_reply.started":"2025-08-06T08:47:35.188675Z","shell.execute_reply":"2025-08-06T08:47:36.782957Z"}},"outputs":[],"execution_count":null},{"cell_type":"markdown","source":"# Imports","metadata":{}},{"cell_type":"code","source":"IMG_SIZE = 224\nBATCH_SIZE = 8\nEPOCHS = 3\n\n\ntrain_csv_path = \"/kaggle/input/diabetic-retinopathy-detection/trainLabels.csv.zip\"\ntrain_images_dir = \"/kaggle/input/diabetic-retinopathy-train-unzipped/train/\"\ntest_images_dir = \"/kaggle/input/diabetic-retinopathy-test-unzipped/test/\"\nsubmission_csv_path = \"/kaggle/input/diabetic-retinopathy-detection/sampleSubmission.csv.zip\"","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-06T08:47:45.927014Z","iopub.execute_input":"2025-08-06T08:47:45.927353Z","iopub.status.idle":"2025-08-06T08:47:45.932307Z","shell.execute_reply.started":"2025-08-06T08:47:45.927328Z","shell.execute_reply":"2025-08-06T08:47:45.931418Z"}},"outputs":[],"execution_count":null},{"cell_type":"code","source":"img_number = 10016\n\ndf = pd.read_csv(train_csv_path)\ndf_train = df[:img_number].copy()\ndf_train.head()","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-06T08:47:57.977342Z","iopub.execute_input":"2025-08-06T08:47:57.977758Z","iopub.status.idle":"2025-08-06T08:47:58.067509Z","shell.execute_reply.started":"2025-08-06T08:47:57.977727Z","shell.execute_reply":"2025-08-06T08:47:58.066641Z"}},"outputs":[],"execution_count":null},{"cell_type":"markdown","source":"**Train Test Split**","metadata":{}},{"cell_type":"code","source":"# Convert multiclass level to binary: 0 = No DR, 1 = DR Present\ndf_train[\"Label\"] = df_train[\"level\"].apply(lambda x: False if x == 0 else True)\n\ndf_train[\"filepath\"] = df_train[\"image\"].apply(lambda x: os.path.join(train_images_dir, f\"{x}.jpeg\"))\n\ntrain_df, val_df = train_test_split(df_train, test_size=0.4, stratify=df_train[\"Label\"], random_state=42)\nval_df, test_df = train_test_split(val_df, test_size=0.5, stratify=val_df[\"Label\"], random_state=42)\n","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-06T08:48:03.800830Z","iopub.execute_input":"2025-08-06T08:48:03.801123Z","iopub.status.idle":"2025-08-06T08:48:03.847523Z","shell.execute_reply.started":"2025-08-06T08:48:03.801102Z","shell.execute_reply":"2025-08-06T08:48:03.846272Z"}},"outputs":[],"execution_count":null},{"cell_type":"code","source":"print(train_df.shape, val_df.shape, test_df.shape)\n# train_df.filepath[0]\ntrain_df","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-06T08:48:05.852896Z","iopub.execute_input":"2025-08-06T08:48:05.853236Z","iopub.status.idle":"2025-08-06T08:48:05.868148Z","shell.execute_reply.started":"2025-08-06T08:48:05.853211Z","shell.execute_reply":"2025-08-06T08:48:05.867460Z"}},"outputs":[],"execution_count":null},{"cell_type":"markdown","source":"# Visualization","metadata":{}},{"cell_type":"code","source":"import cv2\nimport matplotlib.pyplot as plt\n\nplt.figure(figsize=(15, 15))\n\n# Show first 25 images from train_df\nfor idx, (i, row) in enumerate(train_df.head(25).iterrows()):\n    plt.subplot(5, 5, idx + 1)\n    plt.xticks([])\n    plt.yticks([])\n    plt.grid(False)\n    \n    # Load and preprocess image\n    img_path = row['filepath']\n    img = cv2.imread(img_path)\n    \n    if img is None:\n        print(f\"Warning: Image not found at {img_path}\")\n        continue\n    \n    img = cv2.resize(img, (224, 224))\n    img = cv2.cvtColor(img, cv2.COLOR_BGR2RGB)\n    \n    # Set label from binary column\n    label = row.get('binary_label', row.get('Label', None))\n    label_text = 'Yes' if label else 'No'\n    \n    plt.imshow(img)\n    plt.xlabel(label_text, fontsize=12, color='green' if label else 'red')\n\nplt.tight_layout()\nplt.suptitle(\"Sample Images with Binary DR Labels\", fontsize=20, y=1.02)\nplt.show()\n","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-06T06:56:32.033573Z","iopub.execute_input":"2025-08-06T06:56:32.033859Z","iopub.status.idle":"2025-08-06T06:56:36.069631Z","shell.execute_reply.started":"2025-08-06T06:56:32.033838Z","shell.execute_reply":"2025-08-06T06:56:36.068810Z"}},"outputs":[],"execution_count":null},{"cell_type":"code","source":"","metadata":{"trusted":true},"outputs":[],"execution_count":null},{"cell_type":"markdown","source":"# Evaluation Functions","metadata":{}},{"cell_type":"markdown","source":"## Results","metadata":{}},{"cell_type":"code","source":"# from sklearn.metrics import (\n#     accuracy_score, precision_score, recall_score, f1_score, roc_auc_score,\n#     average_precision_score, confusion_matrix, cohen_kappa_score,\n#     precision_recall_curve, roc_curve\n# )\n# import tensorflow as tf\n# import pandas as pd\n# import numpy as np\n# import os\n\n# def Results(model, model_name,test_set, save_csv=False):\n#     # Evaluate the model\n#     results = model.evaluate(test_set)\n#     print(\"Evaluation results: \", results)\n\n#     y_true = []\n#     y_pred_probs = []\n\n#     print(\"Loading....\", flush=True)\n\n#     for images, labels in test_set:\n#         y_true.extend(labels.numpy())\n#         y_pred_probs.extend(model.predict(images, verbose=0).flatten())\n\n#     y_true = tf.convert_to_tensor(y_true)\n#     y_pred_probs = tf.convert_to_tensor(y_pred_probs)\n#     y_pred = (y_pred_probs > 0.5).numpy().astype(\"int32\")\n\n#     acc = accuracy_score(y_true, y_pred)\n#     precision = precision_score(y_true, y_pred, zero_division=1)\n#     recall = recall_score(y_true, y_pred)\n#     f1 = f1_score(y_true, y_pred)\n#     roc_auc = roc_auc_score(y_true, y_pred_probs)\n#     prc_auc = average_precision_score(y_true, y_pred_probs)\n#     conf_matrix = confusion_matrix(y_true, y_pred)\n#     kappa = cohen_kappa_score(y_true, y_pred)\n#     tn, fp, fn, tp = conf_matrix.ravel()\n#     npv = tn / (tn + fn)\n\n#     print(f\"Accuracy: {acc}\")\n#     print(f\"Precision: {precision}\")\n#     print(f\"Recall: {recall}\")\n#     print(f\"F1 Score: {f1}\")\n#     print(f\"ROC AUC: {roc_auc}\")\n#     print(f\"PRC AUC: {prc_auc}\")\n#     print(f\"Confusion Matrix: \\n{conf_matrix}\")\n#     print(f\"Kappa Coefficient: {kappa}\")\n#     print(f\"NPV: {npv}\")\n\n#     if save_csv:\n#         # Compute ROC and Precision-Recall curves\n#         fpr, tpr, _ = roc_curve(y_true, y_pred_probs)\n#         prec, rec, _ = precision_recall_curve(y_true, y_pred_probs)\n\n#         # Prepare DataFrame\n#         max_len = max(len(fpr), len(prec))\n#         fpr = np.pad(fpr, (0, max_len - len(fpr)), 'constant', constant_values=np.nan)\n#         tpr = np.pad(tpr, (0, max_len - len(tpr)), 'constant', constant_values=np.nan)\n#         prec = np.pad(prec, (0, max_len - len(prec)), 'constant', constant_values=np.nan)\n#         rec = np.pad(rec, (0, max_len - len(rec)), 'constant', constant_values=np.nan)\n\n#         df = pd.DataFrame({\n#             'FPR': fpr,\n#             'TPR': tpr,\n#             'Precision': prec,\n#             'Recall': rec\n#         })\n\n\n#         file_name = model_name\n#         df.to_csv(file_name, index=False)\n#         print(f\"Saved metrics to: {file_name}\")\n","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-06T06:42:13.061168Z","iopub.execute_input":"2025-08-06T06:42:13.061510Z","iopub.status.idle":"2025-08-06T06:42:13.066232Z","shell.execute_reply.started":"2025-08-06T06:42:13.061480Z","shell.execute_reply":"2025-08-06T06:42:13.065605Z"}},"outputs":[],"execution_count":null},{"cell_type":"code","source":"from sklearn.metrics import (\n    accuracy_score, precision_score, recall_score, f1_score, roc_auc_score,\n    average_precision_score, confusion_matrix, cohen_kappa_score,\n    precision_recall_curve, roc_curve\n)\nimport tensorflow as tf\nimport pandas as pd\nimport numpy as np\nimport os\nimport wandb\nimport matplotlib.pyplot as plt\n\ndef Results(model, test_set, model_name=\"Model\", save_csv=False, wandb_log=False):\n    # Evaluate the model\n    results = model.evaluate(test_set)\n    print(\"Evaluation results: \", results)\n\n    y_true = []\n    y_pred_probs = []\n\n    print(\"Loading....\", flush=True)\n\n    for images, labels in test_set:\n        y_true.extend(labels.numpy())\n        y_pred_probs.extend(model.predict(images, verbose=0).flatten())\n\n    y_true = tf.convert_to_tensor(y_true)\n    y_pred_probs = tf.convert_to_tensor(y_pred_probs)\n    y_pred = (y_pred_probs > 0.5).numpy().astype(\"int32\")\n\n    acc = accuracy_score(y_true, y_pred)\n    precision = precision_score(y_true, y_pred, zero_division=1)\n    recall = recall_score(y_true, y_pred)\n    f1 = f1_score(y_true, y_pred)\n    roc_auc = roc_auc_score(y_true, y_pred_probs)\n    prc_auc = average_precision_score(y_true, y_pred_probs)\n    conf_matrix = confusion_matrix(y_true, y_pred)\n    kappa = cohen_kappa_score(y_true, y_pred)\n    tn, fp, fn, tp = conf_matrix.ravel()\n    npv = tn / (tn + fn)\n\n    # Print metrics\n    print(f\"Accuracy: {acc}\")\n    print(f\"Precision: {precision}\")\n    print(f\"Recall: {recall}\")\n    print(f\"F1 Score: {f1}\")\n    print(f\"ROC AUC: {roc_auc}\")\n    print(f\"PRC AUC: {prc_auc}\")\n    print(f\"Confusion Matrix: \\n{conf_matrix}\")\n    print(f\"Kappa Coefficient: {kappa}\")\n    print(f\"NPV: {npv}\")\n\n    # Log metrics to W&B\n    if wandb_log:\n        wandb.log({\n            \"Accuracy\": acc,\n            \"Precision\": precision,\n            \"Recall\": recall,\n            \"F1 Score\": f1,\n            \"ROC AUC\": roc_auc,\n            \"PRC AUC\": prc_auc,\n            \"Kappa\": kappa,\n            \"NPV\": npv\n        })\n\n    if save_csv or wandb_log:\n        # Compute ROC and Precision-Recall curves\n        fpr, tpr, _ = roc_curve(y_true, y_pred_probs)\n        prec, rec, _ = precision_recall_curve(y_true, y_pred_probs)\n\n        # Pad arrays to same length\n        max_len = max(len(fpr), len(prec))\n        fpr = np.pad(fpr, (0, max_len - len(fpr)), 'constant', constant_values=np.nan)\n        tpr = np.pad(tpr, (0, max_len - len(tpr)), 'constant', constant_values=np.nan)\n        prec = np.pad(prec, (0, max_len - len(prec)), 'constant', constant_values=np.nan)\n        rec = np.pad(rec, (0, max_len - len(rec)), 'constant', constant_values=np.nan)\n\n        df = pd.DataFrame({\n            'FPR': fpr,\n            'TPR': tpr,\n            'Precision': prec,\n            'Recall': rec\n        })\n\n        file_name = model_name+'.metrices.csv'\n        df.to_csv(file_name, index=False)\n        print(f\"Saved metrics to: {file_name}\")\n\n        if wandb_log:\n            wandb.save(file_name)\n\n            # Plot ROC Curve\n            plt.figure()\n            plt.plot(fpr, tpr, label=f'ROC AUC = {roc_auc:.2f}')\n            plt.xlabel(\"False Positive Rate\")\n            plt.ylabel(\"True Positive Rate\")\n            plt.title(\"ROC Curve\")\n            plt.legend()\n            wandb.log({\"ROC Curve\": wandb.Image(plt)})\n            plt.close()\n\n            # Plot PR Curve\n            plt.figure()\n            plt.plot(rec, prec, label=f'PRC AUC = {prc_auc:.2f}')\n            plt.xlabel(\"Recall\")\n            plt.ylabel(\"Precision\")\n            plt.title(\"Precision-Recall Curve\")\n            plt.legend()\n            wandb.log({\"PR Curve\": wandb.Image(plt)})\n            plt.close()\n","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-06T08:48:13.796417Z","iopub.execute_input":"2025-08-06T08:48:13.796745Z","iopub.status.idle":"2025-08-06T08:48:13.811036Z","shell.execute_reply.started":"2025-08-06T08:48:13.796703Z","shell.execute_reply":"2025-08-06T08:48:13.810209Z"}},"outputs":[],"execution_count":null},{"cell_type":"markdown","source":"## MACS & FLOPS","metadata":{}},{"cell_type":"code","source":"import tensorflow as tf\nfrom tensorflow.python.framework.convert_to_constants import convert_variables_to_constants_v2_as_graph\n\ndef get_flops(model, input_shape):\n    # Convert Keras model to ConcreteFunction\n    concrete_func = tf.function(model).get_concrete_function(tf.TensorSpec(input_shape, tf.float32))\n    \n    # Convert ConcreteFunction to a frozen graph\n    frozen_func, graph_def = convert_variables_to_constants_v2_as_graph(concrete_func)\n    \n    # Create a session to run the frozen graph\n    with tf.compat.v1.Session(graph=tf.Graph()) as sess:\n        tf.import_graph_def(graph_def, name=\"\")\n        graph = tf.compat.v1.get_default_graph()\n        \n        # Calculate FLOPs using TensorFlow profiler\n        run_meta = tf.compat.v1.RunMetadata()\n        opts = tf.compat.v1.profiler.ProfileOptionBuilder.float_operation()\n        flops = tf.compat.v1.profiler.profile(graph=graph, run_meta=run_meta, cmd='op', options=opts)\n        \n    return flops.total_float_ops","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-06T08:48:15.311280Z","iopub.execute_input":"2025-08-06T08:48:15.311958Z","iopub.status.idle":"2025-08-06T08:48:15.317584Z","shell.execute_reply.started":"2025-08-06T08:48:15.311935Z","shell.execute_reply":"2025-08-06T08:48:15.316922Z"}},"outputs":[],"execution_count":null},{"cell_type":"code","source":"# # Define input shape\n# input_shape = (1, 224, 224,3 )\n# # Calculate FLOPs\n# flops = get_flops(model, input_shape)\n# macs = flops // 2  # MACs is half the FLOPs for Conv2D\n\n# print(f\"MACs: {macs}\")\n# print(f\"FLOPs: {flops}\")","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-06T08:48:22.899956Z","iopub.execute_input":"2025-08-06T08:48:22.900721Z","iopub.status.idle":"2025-08-06T08:48:22.905210Z","shell.execute_reply.started":"2025-08-06T08:48:22.900687Z","shell.execute_reply":"2025-08-06T08:48:22.904407Z"}},"outputs":[],"execution_count":null},{"cell_type":"markdown","source":"## Parameters","metadata":{}},{"cell_type":"code","source":"import numpy as np\n\ndef model_parametersInfo(model):\n    \"\"\"\n    Prints total, trainable, and non-trainable parameters of a Keras model.\n    \"\"\"\n    total_params = model.count_params()\n    trainable_params = np.sum([np.prod(v.shape) for v in model.trainable_variables])\n    non_trainable_params = np.sum([np.prod(v.shape) for v in model.non_trainable_variables])\n\n    print(f\"Total params: {total_params:,}\")\n    print(f\"Trainable params: {trainable_params:,}\")\n    print(f\"Non-trainable params: {non_trainable_params:,}\")\n","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-06T08:48:24.511322Z","iopub.execute_input":"2025-08-06T08:48:24.511631Z","iopub.status.idle":"2025-08-06T08:48:24.516731Z","shell.execute_reply.started":"2025-08-06T08:48:24.511609Z","shell.execute_reply":"2025-08-06T08:48:24.516042Z"}},"outputs":[],"execution_count":null},{"cell_type":"markdown","source":"## Inference Report","metadata":{}},{"cell_type":"code","source":"import tensorflow as tf\nimport psutil\nimport os\nimport time\n\ndef inference_Report(model, test_set):\n    process = psutil.Process(os.getpid())\n\n    if tf.config.list_physical_devices('GPU'):\n        gpu_enabled = True\n        gpu_device = 'GPU:0'\n    else:\n        gpu_enabled = False\n\n    print(\"🔍 Capturing memory stats and inference time for 100 images...\\n\")\n    # --- Before Inference ---\n    ram_before = process.memory_info().rss\n    gpu_before = tf.config.experimental.get_memory_info(gpu_device)['current'] if gpu_enabled else 0\n\n    total_images = 0\n    start_time = time.time()\n\n    for images, _ in test_set:\n        for i in range(images.shape[0]):\n            if total_images >= 100:\n                break\n            image = tf.expand_dims(images[i], axis=0)\n            _ = model(image, training=False)\n            total_images += 1\n        if total_images >= 100:\n            break\n\n    end_time = time.time()\n\n    # --- After Inference ---\n    ram_after = process.memory_info().rss\n    gpu_after = tf.config.experimental.get_memory_info(gpu_device)['current'] if gpu_enabled else 0\n\n    # --- Results ---\n    print(\"📊 Inference Resource Usage Summary (100 images):\")\n    print(f\"CPU RAM Before: {ram_before / (1024 ** 2):.2f} MB\")\n    print(f\"CPU RAM After : {ram_after / (1024 ** 2):.2f} MB\")\n    print(f\"CPU RAM Used  : {(ram_after - ram_before) / (1024 ** 2):.2f} MB\\n\")\n\n    if gpu_enabled:\n        print(f\"GPU Mem Before: {gpu_before / (1024 ** 2):.2f} MB\")\n        print(f\"GPU Mem After : {gpu_after / (1024 ** 2):.2f} MB\")\n        print(f\"GPU Mem Used  : {(gpu_after - gpu_before) / (1024 ** 2):.2f} MB\\n\")\n    else:\n        print(\"GPU Not Detected.\\n\")\n    time_per_img = (end_time - start_time)/100\n    print(f\"🕒 Inference Time per image: {time_per_img:.4f} seconds\")\n\n","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-06T08:48:27.411618Z","iopub.execute_input":"2025-08-06T08:48:27.411919Z","iopub.status.idle":"2025-08-06T08:48:27.420689Z","shell.execute_reply.started":"2025-08-06T08:48:27.411898Z","shell.execute_reply":"2025-08-06T08:48:27.419704Z"}},"outputs":[],"execution_count":null},{"cell_type":"markdown","source":"# Data Pipeline","metadata":{}},{"cell_type":"code","source":"def retreive_dataset(set_name):\n    images,labels=[],[]\n    def imgResize(img,width,height):\n        imgResize = cv2.resize(img,(width,height))\n        return imgResize\n    \n    for (img, imclass) in zip(set_name['filepath'], set_name['Label']):\n        img = cv2.imread(img)\n\n        # img = cv2.cvtColor(img, cv2.COLOR_BGR2RGB)   #BGR to RGB\n        img = cv2.resize(img, (224, 224))    \n        # img = img.astype(np.float32) / 255.0   #preprocess\n        images.append(img)\n        if(imclass==True):\n            labels.append(1)\n        else:\n            labels.append(0)\n    print(f\"Done\")\n    return np.array(images),np.array(labels)","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-06T08:48:31.907641Z","iopub.execute_input":"2025-08-06T08:48:31.907955Z","iopub.status.idle":"2025-08-06T08:48:31.913896Z","shell.execute_reply.started":"2025-08-06T08:48:31.907933Z","shell.execute_reply":"2025-08-06T08:48:31.913190Z"}},"outputs":[],"execution_count":null},{"cell_type":"code","source":"X_train,y_train=retreive_dataset(train_df)\nX_val,y_val=retreive_dataset(val_df)\nX_test,y_test=retreive_dataset(test_df)","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-06T08:48:32.575166Z","iopub.execute_input":"2025-08-06T08:48:32.576006Z","iopub.status.idle":"2025-08-06T08:59:42.603352Z","shell.execute_reply.started":"2025-08-06T08:48:32.575972Z","shell.execute_reply":"2025-08-06T08:59:42.602447Z"}},"outputs":[],"execution_count":null},{"cell_type":"code","source":"\n\nindex = 0  # Change this to view another image\nimage = X_train[index]\nlabel = y_train[index]\n\nplt.imshow(cv2.cvtColor(image, cv2.COLOR_BGR2RGB))\nplt.title(f\"Label: {'DR Present' if label == 1 else 'No DR'}\")\nplt.axis('off')\nplt.show()","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-06T08:59:42.604561Z","iopub.execute_input":"2025-08-06T08:59:42.604852Z","iopub.status.idle":"2025-08-06T08:59:43.164429Z","shell.execute_reply.started":"2025-08-06T08:59:42.604833Z","shell.execute_reply":"2025-08-06T08:59:43.163532Z"}},"outputs":[],"execution_count":null},{"cell_type":"code","source":"train_set_raw=tf.data.Dataset.from_tensor_slices((X_train,y_train))\nvalid_set_raw=tf.data.Dataset.from_tensor_slices((X_val,y_val))\ntest_set_raw=tf.data.Dataset.from_tensor_slices((X_test,y_test))","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-06T08:59:43.165333Z","iopub.execute_input":"2025-08-06T08:59:43.165548Z","iopub.status.idle":"2025-08-06T08:59:46.872994Z","shell.execute_reply.started":"2025-08-06T08:59:43.165532Z","shell.execute_reply":"2025-08-06T08:59:46.872146Z"}},"outputs":[],"execution_count":null},{"cell_type":"markdown","source":"# ConvNeXt","metadata":{}},{"cell_type":"code","source":"tf.keras.backend.clear_session()  # extra code – resets layer name counter\n\nbatch_size = 32\npreprocess = tf.keras.applications.convnext.preprocess_input\ntrain_set = train_set_raw.map(lambda X, y: (preprocess(tf.cast(X, tf.float32)), y))\ntrain_set = train_set.shuffle(1000, seed=42).batch(batch_size).prefetch(1)\nvalid_set = valid_set_raw.map(lambda X, y: (preprocess(tf.cast(X, tf.float32)), y)).batch(batch_size)\ntest_set = test_set_raw.map(lambda X, y: (preprocess(tf.cast(X, tf.float32)), y)).batch(batch_size)","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-06T08:59:46.875548Z","iopub.execute_input":"2025-08-06T08:59:46.875811Z","iopub.status.idle":"2025-08-06T08:59:49.972528Z","shell.execute_reply.started":"2025-08-06T08:59:46.875791Z","shell.execute_reply":"2025-08-06T08:59:49.971831Z"}},"outputs":[],"execution_count":null},{"cell_type":"code","source":"plt.figure(figsize=(12, 12))\nfor X_batch, y_batch in train_set.take(1):\n    for index in range(9):\n        plt.subplot(3, 3, index + 1)\n        img = (X_batch[index].numpy())\n        img = cv2.cvtColor(img, cv2.COLOR_BGR2RGB)\n        plt.imshow(img /np.max(img)) # rescale to 0–1 for imshow()\n        if(y_batch[index]==1):\n            classt= 'DR Present'\n        else:\n            classt= \"No DR\"\n        plt.title(f\"Class: {classt}\")\n        plt.axis(\"off\")\n\nplt.show()","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-05T10:17:41.663723Z","iopub.execute_input":"2025-08-05T10:17:41.664302Z","iopub.status.idle":"2025-08-05T10:17:43.455806Z","shell.execute_reply.started":"2025-08-05T10:17:41.664250Z","shell.execute_reply":"2025-08-05T10:17:43.455000Z"}},"outputs":[],"execution_count":null},{"cell_type":"code","source":"tf.random.set_seed(42)  # Ensures reproducibility\n\ntf.keras.backend.clear_session()\n# Ensures reproducibility\n\n\nbase_model = tf.keras.applications.ConvNeXtBase(weights=\"imagenet\", include_top=False)\navg = tf.keras.layers.GlobalAveragePooling2D()(base_model.output)\noutput = tf.keras.layers.Dense(1, activation=\"sigmoid\")(avg)\nmodel = tf.keras.Model(inputs=base_model.input, outputs=output)\nfor layer in base_model.layers:\n    layer.trainable = False","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-06T08:59:49.973309Z","iopub.execute_input":"2025-08-06T08:59:49.973518Z","iopub.status.idle":"2025-08-06T08:59:55.808591Z","shell.execute_reply.started":"2025-08-06T08:59:49.973501Z","shell.execute_reply":"2025-08-06T08:59:55.807877Z"}},"outputs":[],"execution_count":null},{"cell_type":"code","source":"optimizer = tf.keras.optimizers.Adam(learning_rate=0.0001)\nmodel.compile(loss=\"binary_crossentropy\", optimizer=optimizer,\n              metrics=[\"accuracy\"])","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-06T08:59:55.809440Z","iopub.execute_input":"2025-08-06T08:59:55.809743Z","iopub.status.idle":"2025-08-06T08:59:55.825735Z","shell.execute_reply.started":"2025-08-06T08:59:55.809718Z","shell.execute_reply":"2025-08-06T08:59:55.824903Z"}},"outputs":[],"execution_count":null},{"cell_type":"code","source":"","metadata":{"trusted":true},"outputs":[],"execution_count":null},{"cell_type":"code","source":"wandb.init(\n    project=\"Diabetic Retinopathy Detection\",   # change this\n    name=\"ConvNeXtBase\",        # optional\n    config={                       # optional: log hyperparameters\n        \"epochs\": 10,\n        \"batch_size\": 32,\n        \"learning_rate\": 0.001\n    }\n)","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-05T10:45:19.486646Z","iopub.execute_input":"2025-08-05T10:45:19.487162Z","iopub.status.idle":"2025-08-05T10:45:26.304800Z","shell.execute_reply.started":"2025-08-05T10:45:19.487139Z","shell.execute_reply":"2025-08-05T10:45:26.304210Z"}},"outputs":[],"execution_count":null},{"cell_type":"code","source":"model_name = \"DR_ConvNeXtBase\"\nmodel_name","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-06T08:59:55.826588Z","iopub.execute_input":"2025-08-06T08:59:55.826854Z","iopub.status.idle":"2025-08-06T08:59:55.832590Z","shell.execute_reply.started":"2025-08-06T08:59:55.826825Z","shell.execute_reply":"2025-08-06T08:59:55.831901Z"}},"outputs":[],"execution_count":null},{"cell_type":"code","source":"from tensorflow.keras.callbacks import ModelCheckpoint, CSVLogger\n# from wandb.keras import WandbMetricsLogger, WandbModelCheckpoint\n\n\n\n# checkpoint_callback = ModelCheckpoint(\n#     filepath='best_model.keras',  # Path to save the model\n#     monitor='val_accuracy',        # yMetric to monitor\n#     save_best_only=True,       # Save only the best model\n#     save_weights_only=False,   # Save the entire model, not just weights\n#     mode='max',                # Mode to minimize the monitored metric\n#     verbose=1                  # Verbosity mode\n# )\ncsv_logger = CSVLogger(model_name+'.csv',append = True)\n\n# callbacks = [checkpoint_callback,csv_logger]\ncallbacks = [csv_logger,wandb.keras.WandbMetricsLogger(log_freq=5)]","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-05T10:45:26.311985Z","iopub.execute_input":"2025-08-05T10:45:26.312137Z","iopub.status.idle":"2025-08-05T10:45:26.329408Z","shell.execute_reply.started":"2025-08-05T10:45:26.312125Z","shell.execute_reply":"2025-08-05T10:45:26.328678Z"}},"outputs":[],"execution_count":null},{"cell_type":"code","source":"","metadata":{"trusted":true},"outputs":[],"execution_count":null},{"cell_type":"code","source":"history = model.fit(train_set, validation_data=valid_set, epochs=10, callbacks= callbacks)","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-05T10:45:31.542576Z","iopub.execute_input":"2025-08-05T10:45:31.542828Z","iopub.status.idle":"2025-08-05T11:03:23.664021Z","shell.execute_reply.started":"2025-08-05T10:45:31.542811Z","shell.execute_reply":"2025-08-05T11:03:23.663382Z"}},"outputs":[],"execution_count":null},{"cell_type":"code","source":"# model.save_weights(model_name+'.weights.h5')\n\nmodel.load_weights('/kaggle/working/DR_ConvNeXtBase.weights.h5')","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-06T08:59:55.833428Z","iopub.execute_input":"2025-08-06T08:59:55.833637Z","iopub.status.idle":"2025-08-06T08:59:55.851609Z","shell.execute_reply.started":"2025-08-06T08:59:55.833620Z","shell.execute_reply":"2025-08-06T08:59:55.850818Z"}},"outputs":[],"execution_count":null},{"cell_type":"code","source":"Results(model,test_set,model_name=model_name,save_csv=True,wandb_log=False)","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-05T11:03:31.869720Z","iopub.execute_input":"2025-08-05T11:03:31.869991Z","iopub.status.idle":"2025-08-05T11:04:30.296922Z","shell.execute_reply.started":"2025-08-05T11:03:31.869970Z","shell.execute_reply":"2025-08-05T11:04:30.295925Z"}},"outputs":[],"execution_count":null},{"cell_type":"code","source":"wandb.save(f\"/kaggle/working/{model_name}.metrices.csv\", policy=\"now\")\nwandb.save(f\"/kaggle/working/{model_name}.weights.h5\", policy=\"now\")\n","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-05T11:04:30.298280Z","iopub.execute_input":"2025-08-05T11:04:30.298574Z","iopub.status.idle":"2025-08-05T11:04:30.307969Z","shell.execute_reply.started":"2025-08-05T11:04:30.298545Z","shell.execute_reply":"2025-08-05T11:04:30.307338Z"}},"outputs":[],"execution_count":null},{"cell_type":"code","source":"inference_Report(model,test_set)\nmodel_parametersInfo(model)","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-05T11:08:34.063466Z","iopub.execute_input":"2025-08-05T11:08:34.063904Z","iopub.status.idle":"2025-08-05T11:14:05.092438Z","shell.execute_reply.started":"2025-08-05T11:08:34.063878Z","shell.execute_reply":"2025-08-05T11:14:05.091612Z"}},"outputs":[],"execution_count":null},{"cell_type":"code","source":"# Define input shape\ninput_shape = (1, 224, 224,3 )\n# Calculate FLOPs\nflops = get_flops(model, input_shape)\nmacs = flops // 2  # MACs is half the FLOPs for Conv2D\n\nprint(f\"MACs: {macs}\")\nprint(f\"FLOPs: {flops}\")","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-05T11:14:05.093719Z","iopub.execute_input":"2025-08-05T11:14:05.093985Z","iopub.status.idle":"2025-08-05T11:14:26.619524Z","shell.execute_reply.started":"2025-08-05T11:14:05.093961Z","shell.execute_reply":"2025-08-05T11:14:26.618840Z"}},"outputs":[],"execution_count":null},{"cell_type":"code","source":"wandb.finish()","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-05T11:14:26.620235Z","iopub.execute_input":"2025-08-05T11:14:26.620457Z","iopub.status.idle":"2025-08-05T11:14:26.970348Z","shell.execute_reply.started":"2025-08-05T11:14:26.620441Z","shell.execute_reply":"2025-08-05T11:14:26.969760Z"}},"outputs":[],"execution_count":null},{"cell_type":"markdown","source":"# ConvNeXt fine tuned","metadata":{}},{"cell_type":"code","source":"n =len(base_model.layers)\nL = int(0.3*n)  # 30% of n\nf\"{L} trainable of {n} layers\"","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-06T08:59:55.852379Z","iopub.execute_input":"2025-08-06T08:59:55.852685Z","iopub.status.idle":"2025-08-06T08:59:55.873268Z","shell.execute_reply.started":"2025-08-06T08:59:55.852642Z","shell.execute_reply":"2025-08-06T08:59:55.872316Z"}},"outputs":[],"execution_count":null},{"cell_type":"code","source":"for layer in base_model.layers[-L:]:\n    layer.trainable = True","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-06T08:59:55.877457Z","iopub.execute_input":"2025-08-06T08:59:55.877699Z","iopub.status.idle":"2025-08-06T08:59:55.898585Z","shell.execute_reply.started":"2025-08-06T08:59:55.877664Z","shell.execute_reply":"2025-08-06T08:59:55.897641Z"}},"outputs":[],"execution_count":null},{"cell_type":"code","source":"optimizer = tf.keras.optimizers.Adam(learning_rate=0.0001)\nmodel.compile(loss=\"binary_crossentropy\", optimizer=optimizer,\n              metrics=[\"accuracy\"])","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-06T08:59:55.899397Z","iopub.execute_input":"2025-08-06T08:59:55.899695Z","iopub.status.idle":"2025-08-06T08:59:55.923422Z","shell.execute_reply.started":"2025-08-06T08:59:55.899667Z","shell.execute_reply":"2025-08-06T08:59:55.922619Z"}},"outputs":[],"execution_count":null},{"cell_type":"code","source":"wandb.init(\n    project=\"Diabetic Retinopathy Detection\", \n    name=\"CovNeXtBase_ft\",   # Change this to \"VGG_Model_Run\", etc. for next model\n    resume=\"allow\"\n)\n","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-06T08:59:55.924392Z","iopub.execute_input":"2025-08-06T08:59:55.925239Z","iopub.status.idle":"2025-08-06T09:00:07.913627Z","shell.execute_reply.started":"2025-08-06T08:59:55.925215Z","shell.execute_reply":"2025-08-06T09:00:07.912812Z"}},"outputs":[],"execution_count":null},{"cell_type":"code","source":"model_name = model_name+'_ft'\nmodel_name","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-06T09:00:07.914586Z","iopub.execute_input":"2025-08-06T09:00:07.914809Z","iopub.status.idle":"2025-08-06T09:00:07.920780Z","shell.execute_reply.started":"2025-08-06T09:00:07.914792Z","shell.execute_reply":"2025-08-06T09:00:07.920169Z"}},"outputs":[],"execution_count":null},{"cell_type":"code","source":"from tensorflow.keras.callbacks import ModelCheckpoint, CSVLogger\n\ncsv_logger = CSVLogger(model_name+'.csv',append = True)\n\ncallbacks = [csv_logger,wandb.keras.WandbMetricsLogger(log_freq=5)]","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-06T09:00:07.921711Z","iopub.execute_input":"2025-08-06T09:00:07.922001Z","iopub.status.idle":"2025-08-06T09:00:09.264465Z","shell.execute_reply.started":"2025-08-06T09:00:07.921977Z","shell.execute_reply":"2025-08-06T09:00:09.263779Z"}},"outputs":[],"execution_count":null},{"cell_type":"code","source":"history = model.fit(train_set, validation_data=valid_set, epochs=10, callbacks= callbacks)\n","metadata":{"trusted":true,"execution":{"iopub.status.busy":"2025-08-06T09:00:09.265475Z","iopub.execute_input":"2025-08-06T09:00:09.266122Z"}},"outputs":[],"execution_count":null},{"cell_type":"code","source":"model.save_weights(model_name+'.weights.h5')\n\n# model.load_weights('/kaggle/working/DR_Xception.weights.h5')","metadata":{"trusted":true},"outputs":[],"execution_count":null},{"cell_type":"code","source":"\nResults(model,test_set,model_name=model_name,save_csv=True,wandb_log=True)","metadata":{"trusted":true},"outputs":[],"execution_count":null},{"cell_type":"code","source":"wandb.save(f\"/kaggle/working/{model_name}.metrices.csv\", policy=\"now\")\nwandb.save(f\"/kaggle/working/{model_name}.weights.h5\", policy=\"now\")\nprint(f\"saved to wandb {model_name}\")","metadata":{"trusted":true},"outputs":[],"execution_count":null},{"cell_type":"code","source":"inference_Report(model,test_set)\nmodel_parametersInfo(model)","metadata":{"trusted":true},"outputs":[],"execution_count":null},{"cell_type":"code","source":"# Define input shape\ninput_shape = (1, 224, 224,3 )\n# Calculate FLOPs\nflops = get_flops(model, input_shape)\nmacs = flops // 2  # MACs is half the FLOPs for Conv2D\n\nprint(f\"MACs: {macs}\")\nprint(f\"FLOPs: {flops}\")","metadata":{"trusted":true},"outputs":[],"execution_count":null},{"cell_type":"code","source":"wandb.finish()","metadata":{"trusted":true},"outputs":[],"execution_count":null},{"cell_type":"code","source":"","metadata":{"trusted":true},"outputs":[],"execution_count":null}]}