{"metadata":{"kernelspec":{"language":"python","display_name":"Python 3","name":"python3"},"language_info":{"pygments_lexer":"ipython3","nbconvert_exporter":"python","version":"3.6.4","file_extension":".py","codemirror_mode":{"name":"ipython","version":3},"name":"python","mimetype":"text/x-python"}},"nbformat_minor":4,"nbformat":4,"cells":[{"cell_type":"code","source":"# This Python 3 environment comes with many helpful analytics libraries installed\n# It is defined by the kaggle/python Docker image: https://github.com/kaggle/docker-python\n# For example, here's several helpful packages to load\n\nimport numpy as np # linear algebra\nimport pandas as pd # data processing, CSV file I/O (e.g. pd.read_csv)\n\n# Input data files are available in the read-only \"../input/\" directory\n# For example, running this (by clicking run or pressing Shift+Enter) will list all files under the input directory\n\nimport os\nfor dirname, _, filenames in os.walk('/kaggle/input'):\n    for filename in filenames:\n        print(os.path.join(dirname, filename))\n\n# You can write up to 20GB to the current directory (/kaggle/working/) that gets preserved as output when you create a version using \"Save & Run All\" \n# You can also write temporary files to /kaggle/temp/, but they won't be saved outside of the current session","metadata":{"_uuid":"8f2839f25d086af736a60e9eeb907d3b93b6e0e5","_cell_guid":"b1076dfc-b9ad-4769-8c92-a6c4dae69d19","execution":{"iopub.status.busy":"2022-11-01T11:21:16.119701Z","iopub.execute_input":"2022-11-01T11:21:16.120126Z","iopub.status.idle":"2022-11-01T11:21:16.608654Z","shell.execute_reply.started":"2022-11-01T11:21:16.120093Z","shell.execute_reply":"2022-11-01T11:21:16.607914Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"# U-Net Original Research Paper\n### Paper Name: U-Net: Convolutional Networks for Biomedical Image Segmentation\n### Paper Link: https://arxiv.org/pdf/1505.04597.pdf","metadata":{}},{"cell_type":"markdown","source":"# Configs","metadata":{}},{"cell_type":"code","source":"class ROOTDIR:\n    train = \"./train_val/images\"\n    train_mask = \"./train_val/masks\"","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:16.610138Z","iopub.execute_input":"2022-11-01T11:21:16.610939Z","iopub.status.idle":"2022-11-01T11:21:16.615286Z","shell.execute_reply.started":"2022-11-01T11:21:16.6109Z","shell.execute_reply":"2022-11-01T11:21:16.614602Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"# General Imports","metadata":{}},{"cell_type":"code","source":"import os\nimport cv2\nimport zipfile\nimport numpy as np\nimport pandas as pd\nfrom PIL import Image\nfrom tqdm.auto import tqdm\nimport matplotlib.pyplot as plt","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:16.616754Z","iopub.execute_input":"2022-11-01T11:21:16.617421Z","iopub.status.idle":"2022-11-01T11:21:16.625613Z","shell.execute_reply.started":"2022-11-01T11:21:16.617345Z","shell.execute_reply":"2022-11-01T11:21:16.624667Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"# Extracting Files","metadata":{}},{"cell_type":"code","source":"dirs = [\"../input/semantic-segmentation-of-underwater-imagery-suim/train_val/images\",\n        \"../input/semantic-segmentation-of-underwater-imagery-suim/train_val/masks\"\n        ]\n\n#for i in tqdm(dirs):\n   # with zipfile.ZipFile(i) as z:\n    #    z.extractall()\n    #","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:16.627977Z","iopub.execute_input":"2022-11-01T11:21:16.628613Z","iopub.status.idle":"2022-11-01T11:21:16.634996Z","shell.execute_reply.started":"2022-11-01T11:21:16.628578Z","shell.execute_reply":"2022-11-01T11:21:16.634241Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"# Working with CSV","metadata":{}},{"cell_type":"code","source":"#df = pd.read_csv(\"../input/people-clothing-segmentation/labels.csv\")\n#df.head()","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:16.636259Z","iopub.execute_input":"2022-11-01T11:21:16.636953Z","iopub.status.idle":"2022-11-01T11:21:16.644447Z","shell.execute_reply.started":"2022-11-01T11:21:16.636919Z","shell.execute_reply":"2022-11-01T11:21:16.643639Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"#df.info()","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:16.645987Z","iopub.execute_input":"2022-11-01T11:21:16.646725Z","iopub.status.idle":"2022-11-01T11:21:16.65501Z","shell.execute_reply.started":"2022-11-01T11:21:16.646685Z","shell.execute_reply":"2022-11-01T11:21:16.654237Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"# Working with Train Images and Masks","metadata":{}},{"cell_type":"code","source":"train_img_lst =os.listdir(\"../input/semantic-segmentation-of-underwater-imagery-suim/train_val/images\")#./train\"\ntrain_mask_lst = os.listdir(\"../input/semantic-segmentation-of-underwater-imagery-suim/train_val/masks\") # \"./train_masks\"","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:16.656298Z","iopub.execute_input":"2022-11-01T11:21:16.656947Z","iopub.status.idle":"2022-11-01T11:21:16.666214Z","shell.execute_reply.started":"2022-11-01T11:21:16.656912Z","shell.execute_reply":"2022-11-01T11:21:16.665396Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"print(train_mask_lst[:5])\nprint(train_img_lst[:5])","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:16.667164Z","iopub.execute_input":"2022-11-01T11:21:16.668547Z","iopub.status.idle":"2022-11-01T11:21:16.675635Z","shell.execute_reply.started":"2022-11-01T11:21:16.66852Z","shell.execute_reply":"2022-11-01T11:21:16.67479Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"print(len(train_mask_lst))\nprint(len(train_img_lst))","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:16.676954Z","iopub.execute_input":"2022-11-01T11:21:16.677337Z","iopub.status.idle":"2022-11-01T11:21:16.685835Z","shell.execute_reply.started":"2022-11-01T11:21:16.677302Z","shell.execute_reply":"2022-11-01T11:21:16.68492Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"#### Sorting to make sure we get right image and right mask","metadata":{}},{"cell_type":"code","source":"sorted_train_mask_lst = sorted(train_mask_lst)","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:16.989287Z","iopub.execute_input":"2022-11-01T11:21:16.989957Z","iopub.status.idle":"2022-11-01T11:21:16.996954Z","shell.execute_reply.started":"2022-11-01T11:21:16.989923Z","shell.execute_reply":"2022-11-01T11:21:16.996168Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"sorted_train_img_lst = sorted(train_img_lst)","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:17.103959Z","iopub.execute_input":"2022-11-01T11:21:17.104331Z","iopub.status.idle":"2022-11-01T11:21:17.110187Z","shell.execute_reply.started":"2022-11-01T11:21:17.104303Z","shell.execute_reply":"2022-11-01T11:21:17.1093Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"print(sorted_train_mask_lst[:16])\nprint(sorted_train_img_lst[:16])","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:17.111758Z","iopub.execute_input":"2022-11-01T11:21:17.112544Z","iopub.status.idle":"2022-11-01T11:21:17.119995Z","shell.execute_reply.started":"2022-11-01T11:21:17.112506Z","shell.execute_reply":"2022-11-01T11:21:17.118857Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"# Visualizing Images with their Mask\n### Making sure images and mask are paired correctly.","metadata":{}},{"cell_type":"code","source":"def show_images(imgs_lst,masks_lst,loops=2):\n    for i in range(loops):\n        img_path = os.path.join(\"../input/semantic-segmentation-of-underwater-imagery-suim/train_val/images\",imgs_lst[i])\n        mask_path = os.path.join(\"../input/semantic-segmentation-of-underwater-imagery-suim/train_val/masks\",masks_lst[i])\n        img = Image.open(img_path)\n        mask = Image.open(mask_path)\n        print(img_path)\n        print(img.size)\n        print(type(img))\n        plt.imshow(img)\n        plt.show()\n        print(mask_path)\n        print(mask.size)\n        plt.imshow(mask)\n        plt.show()\n        print(\"----------------------------------------------------\")\n\nshow_images(sorted_train_img_lst, sorted_train_mask_lst)","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:17.123528Z","iopub.execute_input":"2022-11-01T11:21:17.123788Z","iopub.status.idle":"2022-11-01T11:21:18.110397Z","shell.execute_reply.started":"2022-11-01T11:21:17.123763Z","shell.execute_reply":"2022-11-01T11:21:18.109356Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"# PyTorch Imports","metadata":{}},{"cell_type":"code","source":"import torch\nimport torchvision\nimport torch.nn as nn\nimport albumentations as A\nimport torch.optim as optim\nfrom torchvision import models\nimport torch.nn.functional as F\nfrom torch.optim import lr_scheduler\nimport torchvision.datasets as datasets\nimport torchvision.transforms as transforms\nfrom torchvision.datasets import ImageFolder\nfrom albumentations.pytorch import ToTensorV2 \nfrom torch.utils.data import DataLoader, Dataset","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:18.112398Z","iopub.execute_input":"2022-11-01T11:21:18.112808Z","iopub.status.idle":"2022-11-01T11:21:18.122486Z","shell.execute_reply.started":"2022-11-01T11:21:18.112771Z","shell.execute_reply":"2022-11-01T11:21:18.121755Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"# PyTorch Configuration","metadata":{}},{"cell_type":"code","source":"class CFG:\n    device = torch.device(\"cuda\" if torch.cuda.is_available() else \"cpu\")\n    split_pct = 0.2\n    learning_rate = 3e-4\n    batch_size = 4\n    epochs = 3","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:18.123864Z","iopub.execute_input":"2022-11-01T11:21:18.12454Z","iopub.status.idle":"2022-11-01T11:21:18.137261Z","shell.execute_reply.started":"2022-11-01T11:21:18.124503Z","shell.execute_reply":"2022-11-01T11:21:18.13652Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"seed = 123\nnp.random.seed(seed)\ntorch.manual_seed(seed)","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:18.139322Z","iopub.execute_input":"2022-11-01T11:21:18.139912Z","iopub.status.idle":"2022-11-01T11:21:18.14909Z","shell.execute_reply.started":"2022-11-01T11:21:18.139879Z","shell.execute_reply":"2022-11-01T11:21:18.148416Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"CFG.device","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:18.150288Z","iopub.execute_input":"2022-11-01T11:21:18.150895Z","iopub.status.idle":"2022-11-01T11:21:18.162633Z","shell.execute_reply.started":"2022-11-01T11:21:18.150862Z","shell.execute_reply":"2022-11-01T11:21:18.16179Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"# Working with data","metadata":{}},{"cell_type":"markdown","source":"### Shuffling the data.","metadata":{}},{"cell_type":"code","source":"permuted_train_img_lst = np.random.permutation(np.array(sorted_train_img_lst))\npermuted_train_mask_lst = [x.replace(\".jpg\", \".bmp\") for x in permuted_train_img_lst]\nprint(permuted_train_img_lst[:5])\nprint(permuted_train_mask_lst[:5])","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:18.163955Z","iopub.execute_input":"2022-11-01T11:21:18.164228Z","iopub.status.idle":"2022-11-01T11:21:18.175324Z","shell.execute_reply.started":"2022-11-01T11:21:18.164201Z","shell.execute_reply":"2022-11-01T11:21:18.174552Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"show_images(permuted_train_img_lst,permuted_train_mask_lst)","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:18.176777Z","iopub.execute_input":"2022-11-01T11:21:18.177441Z","iopub.status.idle":"2022-11-01T11:21:19.867875Z","shell.execute_reply.started":"2022-11-01T11:21:18.1774Z","shell.execute_reply":"2022-11-01T11:21:19.867039Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"### Splitting into Training and Validation","metadata":{}},{"cell_type":"code","source":"length = len(permuted_train_img_lst)\nprint(length*0.2) # convert this to int","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:19.870472Z","iopub.execute_input":"2022-11-01T11:21:19.870756Z","iopub.status.idle":"2022-11-01T11:21:19.875334Z","shell.execute_reply.started":"2022-11-01T11:21:19.870728Z","shell.execute_reply":"2022-11-01T11:21:19.874553Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"train_images_list = permuted_train_img_lst[int(CFG.split_pct*len(permuted_train_img_lst)) :]\ntrain_masks_list = permuted_train_mask_lst[int(CFG.split_pct*len(permuted_train_mask_lst)) :]\nprint(len(train_masks_list))\n\nval_images_list = permuted_train_img_lst[: int(CFG.split_pct*len(permuted_train_img_lst))]\nval_masks_list = permuted_train_mask_lst[: int(CFG.split_pct*len(permuted_train_mask_lst))]\nprint(len(val_masks_list))\n\n# 4071+1017=5088 (split includes all items)","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:19.876694Z","iopub.execute_input":"2022-11-01T11:21:19.877255Z","iopub.status.idle":"2022-11-01T11:21:19.888145Z","shell.execute_reply.started":"2022-11-01T11:21:19.877218Z","shell.execute_reply":"2022-11-01T11:21:19.887043Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"### Visualizing Train Dataset","metadata":{}},{"cell_type":"code","source":"show_images(train_images_list,train_masks_list)","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:19.892007Z","iopub.execute_input":"2022-11-01T11:21:19.892268Z","iopub.status.idle":"2022-11-01T11:21:20.989687Z","shell.execute_reply.started":"2022-11-01T11:21:19.892245Z","shell.execute_reply":"2022-11-01T11:21:20.988075Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"### Visualizing Validation Dataset","metadata":{}},{"cell_type":"code","source":"show_images(val_images_list,val_masks_list)","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:20.991462Z","iopub.execute_input":"2022-11-01T11:21:20.991909Z","iopub.status.idle":"2022-11-01T11:21:22.405144Z","shell.execute_reply.started":"2022-11-01T11:21:20.991864Z","shell.execute_reply":"2022-11-01T11:21:22.404423Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"# Dataset Class","metadata":{}},{"cell_type":"code","source":"class CarvanaDataset(Dataset):\n    def __init__(self,img_list,mask_list,transform=None):\n        self.img_list = img_list\n        self.mask_list = mask_list\n        self.transform = transform\n        \n    def __len__(self):\n        return len(self.img_list)\n    \n    def __getitem__(self,index):\n        img_path = os.path.join(\"../input/semantic-segmentation-of-underwater-imagery-suim/train_val/images\",self.img_list[index])\n        mask_path = os.path.join(\"../input/semantic-segmentation-of-underwater-imagery-suim/train_val/masks\",self.mask_list[index])\n        img = Image.open(img_path)\n        mask = Image.open(mask_path)\n        img = np.array(img)\n        mask = np.array(mask)\n        mask[mask==255.0] = 1.0\n        #img_mask_dict = {\"image\": img, \"mask\": mask}\n        \n        if self.transform:\n            augmentation = self.transform(image=img, mask=mask)\n            img = augmentation[\"image\"]\n            mask = augmentation[\"mask\"]\n            mask = torch.unsqueeze(mask,0)\n            #transformations = self.transform(image=img, mask=mask)\n            #img = transformations[\"image\"]\n            #mask = transformations[\"mask\"]\n            \n        return img,mask","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:22.408793Z","iopub.execute_input":"2022-11-01T11:21:22.411176Z","iopub.status.idle":"2022-11-01T11:21:22.423939Z","shell.execute_reply.started":"2022-11-01T11:21:22.411134Z","shell.execute_reply":"2022-11-01T11:21:22.42312Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"train_transform = A.Compose([A.Resize(572,572), \n                             A.Rotate(limit=15,p=0.1),\n                             A.HorizontalFlip(p=0.5),\n                             A.Normalize(mean=(0,0,0),std=(1,1,1),max_pixel_value=255),\n                             ToTensorV2()])\n\nval_transform = A.Compose([A.Resize(572,572),\n                           A.Normalize(mean=(0,0,0),std=(1,1,1),max_pixel_value=255),\n                           ToTensorV2()])","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:22.427023Z","iopub.execute_input":"2022-11-01T11:21:22.427498Z","iopub.status.idle":"2022-11-01T11:21:22.437294Z","shell.execute_reply.started":"2022-11-01T11:21:22.427453Z","shell.execute_reply":"2022-11-01T11:21:22.436526Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"train_dataset = CarvanaDataset(train_images_list, train_masks_list, transform = train_transform)\nval_dataset = CarvanaDataset(val_images_list, val_masks_list, transform = train_transform)","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:22.440024Z","iopub.execute_input":"2022-11-01T11:21:22.443157Z","iopub.status.idle":"2022-11-01T11:21:22.45059Z","shell.execute_reply.started":"2022-11-01T11:21:22.443119Z","shell.execute_reply":"2022-11-01T11:21:22.449829Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"idx = 200\nimg,mask = train_dataset[idx]","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:22.451857Z","iopub.execute_input":"2022-11-01T11:21:22.452231Z","iopub.status.idle":"2022-11-01T11:21:22.515386Z","shell.execute_reply.started":"2022-11-01T11:21:22.452196Z","shell.execute_reply":"2022-11-01T11:21:22.514669Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"mask.shape","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:22.517101Z","iopub.execute_input":"2022-11-01T11:21:22.517704Z","iopub.status.idle":"2022-11-01T11:21:22.523111Z","shell.execute_reply.started":"2022-11-01T11:21:22.517668Z","shell.execute_reply":"2022-11-01T11:21:22.522339Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"img.max()","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:22.52447Z","iopub.execute_input":"2022-11-01T11:21:22.525038Z","iopub.status.idle":"2022-11-01T11:21:22.541317Z","shell.execute_reply.started":"2022-11-01T11:21:22.525003Z","shell.execute_reply":"2022-11-01T11:21:22.540237Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"def show_single_img(img,mask,index=None,train=True):\n    if index:\n        if train:\n            img,mask = train_dataset[index]\n        else:\n            img,mask = val_dataset[index]\n    plt.imshow(img.permute(1,2,0),cmap=\"gray\")  # Convert (3, 572, 572) -> (572, 572, 3)\n    plt.show()\n    plt.imshow(mask.permute(1,2,0), cmap=\"gray\")  # Convert (1, 572, 572) -> (572, 572, 1)\n    print(mask.shape)\n    plt.show()\n    ","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:22.542649Z","iopub.execute_input":"2022-11-01T11:21:22.543109Z","iopub.status.idle":"2022-11-01T11:21:22.549398Z","shell.execute_reply.started":"2022-11-01T11:21:22.543066Z","shell.execute_reply":"2022-11-01T11:21:22.548578Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"print(\"---------------Train---------------\")\nshow_single_img(img,mask,index=15,train=False)\nprint(\"---------------Validation---------------\")\nshow_single_img(img,mask,index=15,train=True)","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:22.550759Z","iopub.execute_input":"2022-11-01T11:21:22.551328Z","iopub.status.idle":"2022-11-01T11:21:22.904043Z","shell.execute_reply.started":"2022-11-01T11:21:22.551294Z","shell.execute_reply":"2022-11-01T11:21:22.902234Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"# Dataloader","metadata":{}},{"cell_type":"code","source":"train_dataloader = DataLoader(train_dataset,batch_size=CFG.batch_size,shuffle=True)\nval_dataloader = DataLoader(val_dataset,batch_size=CFG.batch_size,shuffle=False)","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:25:35.065004Z","iopub.execute_input":"2022-11-01T11:25:35.065423Z","iopub.status.idle":"2022-11-01T11:25:35.070285Z","shell.execute_reply.started":"2022-11-01T11:25:35.065379Z","shell.execute_reply":"2022-11-01T11:25:35.069565Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"a = iter(train_dataloader)\nimg,mask = a.next()\nprint(img.shape,mask.shape)","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:25:35.414715Z","iopub.execute_input":"2022-11-01T11:25:35.415018Z","iopub.status.idle":"2022-11-01T11:25:35.69117Z","shell.execute_reply.started":"2022-11-01T11:25:35.414991Z","shell.execute_reply":"2022-11-01T11:25:35.690316Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"### Utility Functions","metadata":{}},{"cell_type":"code","source":"def double_conv(in_ch, out_ch):\n    conv = nn.Sequential(\n        nn.Conv2d(in_channels=in_ch,out_channels=out_ch,kernel_size=3,stride=1,padding=1),\n        nn.BatchNorm2d(out_ch),                                                            \n        nn.ReLU(inplace=True),\n        nn.Conv2d(in_channels=out_ch,out_channels=out_ch,kernel_size=3,stride=1,padding=1), \n        nn.BatchNorm2d(out_ch),                                                            \n        nn.ReLU(inplace=True)\n    )\n    \n    return conv\n\n#def cropper(og_tensor, target_tensor):\n#    og_shape = og_tensor.shape[2]\n#    target_shape = target_tensor.shape[2]\n#    delta = (og_shape - target_shape) // 2\n#    cropped_og_tensor = og_tensor[:,:,delta:og_shape-delta,delta:og_shape-delta]\n#    return cropped_og_tensor\n \n    \ndef padder(left_tensor, right_tensor): \n    # left_tensor is the tensor on the encoder side of UNET\n    # right_tensor is the tensor on the decoder side  of the UNET\n    \n    if left_tensor.shape != right_tensor.shape:\n        padded = torch.zeros(left_tensor.shape)\n        padded[:, :, :right_tensor.shape[2], :right_tensor.shape[3]] = right_tensor\n        return padded.to(CFG.device)\n    \n    return right_tensor.to(CFG.device)\n    ","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:27:09.600924Z","iopub.execute_input":"2022-11-01T11:27:09.601313Z","iopub.status.idle":"2022-11-01T11:27:09.609179Z","shell.execute_reply.started":"2022-11-01T11:27:09.601274Z","shell.execute_reply":"2022-11-01T11:27:09.607999Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"# UNET MODEL FROM SCRATCH","metadata":{}},{"cell_type":"code","source":"class UNET(nn.Module):\n    def __init__(self,in_chnls, n_classes):\n        super(UNET,self).__init__()\n        \n        self.in_chnls = in_chnls\n        self.n_classes = n_classes\n        \n        self.max_pool = nn.MaxPool2d(kernel_size=2,stride=2)\n        \n        self.down_conv_1 = double_conv(in_ch=self.in_chnls,out_ch=64)\n        self.down_conv_2 = double_conv(in_ch=64,out_ch=128)\n        self.down_conv_3 = double_conv(in_ch=128,out_ch=256)\n        self.down_conv_4 = double_conv(in_ch=256,out_ch=512)\n        self.down_conv_5 = double_conv(in_ch=512,out_ch=1024)\n        #print(self.down_conv_1)\n        \n        self.up_conv_trans_1 = nn.ConvTranspose2d(in_channels=1024,out_channels=512,kernel_size=2,stride=2)\n        self.up_conv_trans_2 = nn.ConvTranspose2d(in_channels=512,out_channels=256,kernel_size=2,stride=2)\n        self.up_conv_trans_3 = nn.ConvTranspose2d(in_channels=256,out_channels=128,kernel_size=2,stride=2)\n        self.up_conv_trans_4 = nn.ConvTranspose2d(in_channels=128,out_channels=64,kernel_size=2,stride=2)\n        \n        self.up_conv_1 = double_conv(in_ch=1024,out_ch=512)\n        self.up_conv_2 = double_conv(in_ch=512,out_ch=256)\n        self.up_conv_3 = double_conv(in_ch=256,out_ch=128)\n        self.up_conv_4 = double_conv(in_ch=128,out_ch=64)\n        \n        self.conv_1x1 = nn.Conv2d(in_channels=64,out_channels=self.n_classes,kernel_size=1,stride=1)\n        \n    def forward(self,x):\n        \n        # encoding\n        x1 = self.down_conv_1(x)\n        #print(\"X1\", x1.shape)\n        p1 = self.max_pool(x1)\n        #print(\"p1\", p1.shape)\n        x2 = self.down_conv_2(p1)\n        #print(\"X2\", x2.shape)\n        p2 = self.max_pool(x2)\n        #print(\"p2\", p2.shape)\n        x3 = self.down_conv_3(p2)\n        #print(\"X2\", x3.shape)\n        p3 = self.max_pool(x3)\n        #print(\"p3\", p3.shape)\n        x4 = self.down_conv_4(p3)\n        #print(\"X4\", x4.shape)\n        p4 = self.max_pool(x4)\n        #print(\"p4\", p4.shape)\n        x5 = self.down_conv_5(p4)\n        #print(\"X5\", x5.shape)\n        \n        # decoding\n        d1 = self.up_conv_trans_1(x5)  # up transpose convolution (\"up sampling\" as called in UNET paper)\n        pad1 = padder(x4,d1) # padding d1 to match x4 shape\n        cat1 = torch.cat([x4,pad1],dim=1) # concatenating padded d1 and x4 on channel dimension(dim 1) [batch(dim 0),channel(dim 1),height(dim 2),width(dim 3)]\n        uc1 = self.up_conv_1(cat1) # 1st up double convolution\n        \n        d2 = self.up_conv_trans_2(uc1)\n        pad2 = padder(x3,d2)\n        cat2 = torch.cat([x3,pad2],dim=1)\n        uc2 = self.up_conv_2(cat2)\n        \n        d3 = self.up_conv_trans_3(uc2)\n        pad3 = padder(x2,d3)\n        cat3 = torch.cat([x2,pad3],dim=1)\n        uc3 = self.up_conv_3(cat3)\n        \n        d4 = self.up_conv_trans_4(uc3)\n        pad4 = padder(x1,d4)\n        cat4 = torch.cat([x1,pad4],dim=1)\n        uc4 = self.up_conv_4(cat4)\n        \n        conv_1x1 = self.conv_1x1(uc4)\n        return conv_1x1\n        #print(conv_1x1.shape)","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:27:14.706631Z","iopub.execute_input":"2022-11-01T11:27:14.707008Z","iopub.status.idle":"2022-11-01T11:27:14.912279Z","shell.execute_reply.started":"2022-11-01T11:27:14.706972Z","shell.execute_reply":"2022-11-01T11:27:14.911274Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"# Training and Validation","metadata":{}},{"cell_type":"markdown","source":"### Train Function\n","metadata":{}},{"cell_type":"code","source":"def train_model(model,dataloader,criterion,optimizer):\n    model.train()\n    train_running_loss = 0.0\n    for j,img_mask in enumerate(tqdm(dataloader)):\n        img = img_mask[0].float().to(CFG.device)\n        #print(\" ----- IMAGE -----\")\n        #print(img)\n        mask = img_mask[1].float().to(CFG.device)\n        #print(\" ----- MASK -----\")\n        #print(mask)\n        \n        y_pred = model(img)\n        #print(\" ----- Y PRED -----\")\n        #print(y_pred)\n        #print(\" ----- Y PRED SHAPE -----\")#\n        #print(y_pred.shape)\n        optimizer.zero_grad()\n        \n        loss = criterion(y_pred,mask)\n        \n        train_running_loss += loss.item() * CFG.batch_size\n        \n        loss.backward()\n        optimizer.step()\n        \n    train_loss = train_running_loss / (j+1)\n    return train_loss","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:22.913729Z","iopub.status.idle":"2022-11-01T11:21:22.914564Z","shell.execute_reply.started":"2022-11-01T11:21:22.914311Z","shell.execute_reply":"2022-11-01T11:21:22.914336Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"### Validation Function","metadata":{}},{"cell_type":"code","source":"def val_model(model,dataloader,criterion,optimizer):\n    model.eval()\n    val_running_loss = 0\n    with torch.no_grad():\n        for j,img_mask in enumerate(tqdm(dataloader)):\n            img = img_mask[0].float().to(CFG.device)\n            mask = img_mask[1].float().to(CFG.device)\n            y_pred = model(img)\n            loss = criterion(y_pred,mask)\n            \n            val_running_loss += loss.item() * CFG.batch_size\n            \n        val_loss = val_running_loss / (j+1)\n    return val_loss","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:27:23.55026Z","iopub.execute_input":"2022-11-01T11:27:23.550824Z","iopub.status.idle":"2022-11-01T11:27:23.559339Z","shell.execute_reply.started":"2022-11-01T11:27:23.550788Z","shell.execute_reply":"2022-11-01T11:27:23.558592Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"model = UNET(in_chnls = 3, n_classes = 1).to(CFG.device)\noptimizer = optim.Adam(model.parameters(), lr = CFG.learning_rate)\ncriterion = nn.BCEWithLogitsLoss()\ntrain_loss_lst = []\nval_loss_lst = []  ","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:27:27.874867Z","iopub.execute_input":"2022-11-01T11:27:27.875223Z","iopub.status.idle":"2022-11-01T11:27:31.244979Z","shell.execute_reply.started":"2022-11-01T11:27:27.875191Z","shell.execute_reply":"2022-11-01T11:27:31.24408Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"### Train and Validation Loop","metadata":{}},{"cell_type":"code","source":"for i in tqdm(range(CFG.epochs)):\n    train_loss = train_model(model=model,dataloader=train_dataloader,criterion=criterion,optimizer=optimizer)\n    val_loss = val_model(model=model,dataloader=val_dataloader,criterion=criterion,optimizer=optimizer)\n    train_loss_lst.append(train_loss)\n    val_loss_lst.append(val_loss)\n    print(f\" Train Loss : {train_loss:.4f}\")\n    print(f\" Validation Loss : {val_loss:.4f}\")","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:27:33.81501Z","iopub.execute_input":"2022-11-01T11:27:33.815393Z","iopub.status.idle":"2022-11-01T11:27:39.683147Z","shell.execute_reply.started":"2022-11-01T11:27:33.815339Z","shell.execute_reply":"2022-11-01T11:27:39.681778Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"### Training and Validation Loss Plot","metadata":{}},{"cell_type":"code","source":"plt.plot(train_loss_lst, color=\"green\", label='train loss')\nplt.plot(val_loss_lst, color=\"red\", label='validation loss')\nplt.xlabel(\"epochs\")\nplt.ylabel(\"loss\")\nplt.legend()\nplt.show()","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:22.921665Z","iopub.status.idle":"2022-11-01T11:21:22.922486Z","shell.execute_reply.started":"2022-11-01T11:21:22.922236Z","shell.execute_reply":"2022-11-01T11:21:22.92226Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"# Saving Model","metadata":{}},{"cell_type":"code","source":"TRAINED_FILE = \"./unet_scratch.pth\"","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:22.923643Z","iopub.status.idle":"2022-11-01T11:21:22.924441Z","shell.execute_reply.started":"2022-11-01T11:21:22.924188Z","shell.execute_reply":"2022-11-01T11:21:22.924213Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"torch.save(model.state_dict(), TRAINED_FILE)","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:22.925722Z","iopub.status.idle":"2022-11-01T11:21:22.926562Z","shell.execute_reply.started":"2022-11-01T11:21:22.926323Z","shell.execute_reply":"2022-11-01T11:21:22.926345Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"from IPython.display import FileLink\nFileLink(TRAINED_FILE)","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:22.927718Z","iopub.status.idle":"2022-11-01T11:21:22.928538Z","shell.execute_reply.started":"2022-11-01T11:21:22.928287Z","shell.execute_reply":"2022-11-01T11:21:22.928311Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"# Testing","metadata":{}},{"cell_type":"code","source":"trained_model = UNET(in_chnls = 3, n_classes = 1)","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:22.92971Z","iopub.status.idle":"2022-11-01T11:21:22.930529Z","shell.execute_reply.started":"2022-11-01T11:21:22.930279Z","shell.execute_reply":"2022-11-01T11:21:22.930303Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"UNET_TRAINED = \"../input/unet-4-epoch-trained/unet_scratch.pth\"","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:22.931689Z","iopub.status.idle":"2022-11-01T11:21:22.932509Z","shell.execute_reply.started":"2022-11-01T11:21:22.93226Z","shell.execute_reply":"2022-11-01T11:21:22.932284Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"trained_model.load_state_dict(torch.load(UNET_TRAINED))","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:22.933686Z","iopub.status.idle":"2022-11-01T11:21:22.934496Z","shell.execute_reply.started":"2022-11-01T11:21:22.934247Z","shell.execute_reply":"2022-11-01T11:21:22.934271Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"trained_model = trained_model.to(\"cuda\")\ntrained_model.eval()\n","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:22.935667Z","iopub.status.idle":"2022-11-01T11:21:22.936473Z","shell.execute_reply.started":"2022-11-01T11:21:22.936223Z","shell.execute_reply":"2022-11-01T11:21:22.936247Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"img_path = \"../input/carvana-image-masking-challenge/29bb3ece3180_11.jpg\"\n\nimg = cv2.imread(img_path)","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:22.937629Z","iopub.status.idle":"2022-11-01T11:21:22.938478Z","shell.execute_reply.started":"2022-11-01T11:21:22.938208Z","shell.execute_reply":"2022-11-01T11:21:22.938232Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"plt.imshow(img)\nplt.show()","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:22.939645Z","iopub.status.idle":"2022-11-01T11:21:22.940461Z","shell.execute_reply.started":"2022-11-01T11:21:22.940213Z","shell.execute_reply":"2022-11-01T11:21:22.940237Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"test_transform = A.Compose([A.Resize(572,572),\n                           A.Normalize(mean=(0,0,0),std=(1,1,1),max_pixel_value=255),\n                           ToTensorV2()])","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:22.941621Z","iopub.status.idle":"2022-11-01T11:21:22.942432Z","shell.execute_reply.started":"2022-11-01T11:21:22.942188Z","shell.execute_reply":"2022-11-01T11:21:22.942212Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"test_image = test_transform(image = img)\n\nprint(test_image)\n\nprint(test_image[\"image\"].dtype)\nprint(test_image[\"image\"].shape)\n\nimg = test_image[\"image\"].unsqueeze(0)\nprint(img.shape)\n\nimg = img.to(\"cuda\")","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:22.943616Z","iopub.status.idle":"2022-11-01T11:21:22.944519Z","shell.execute_reply.started":"2022-11-01T11:21:22.94424Z","shell.execute_reply":"2022-11-01T11:21:22.944264Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"pred = trained_model(img)\npred.shape","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:22.945604Z","iopub.status.idle":"2022-11-01T11:21:22.946417Z","shell.execute_reply.started":"2022-11-01T11:21:22.946171Z","shell.execute_reply":"2022-11-01T11:21:22.946194Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"mask = pred.squeeze(0).cpu().detach().numpy()\nprint(mask.shape)\nmask = mask.transpose(1,2,0)\nprint(mask.shape)","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:22.947564Z","iopub.status.idle":"2022-11-01T11:21:22.948406Z","shell.execute_reply.started":"2022-11-01T11:21:22.948136Z","shell.execute_reply":"2022-11-01T11:21:22.94816Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"display_test_img = test_image[\"image\"].cpu().detach().numpy()\nprint(display_test_img.shape)\ndisplay_test_img = display_test_img.transpose(1,2,0)\ndisplay_test_img.shape","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:22.949566Z","iopub.status.idle":"2022-11-01T11:21:22.950356Z","shell.execute_reply.started":"2022-11-01T11:21:22.950137Z","shell.execute_reply":"2022-11-01T11:21:22.95016Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"mask[mask < 0]=0\nmask[mask > 0]=1","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:22.951527Z","iopub.status.idle":"2022-11-01T11:21:22.952332Z","shell.execute_reply.started":"2022-11-01T11:21:22.952104Z","shell.execute_reply":"2022-11-01T11:21:22.952129Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"print(\"-------Original Image-------\")\nplt.imshow(display_test_img, cmap=\"gray\")\nplt.show()\nprint(\"-------Image Mask-------\")\nplt.imshow(mask,cmap=\"gray\")\nplt.show()","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:21:22.95353Z","iopub.status.idle":"2022-11-01T11:21:22.95433Z","shell.execute_reply.started":"2022-11-01T11:21:22.954098Z","shell.execute_reply":"2022-11-01T11:21:22.954122Z"},"trusted":true},"execution_count":null,"outputs":[]}]}