{"metadata":{"kernelspec":{"language":"python","display_name":"Python 3","name":"python3"},"language_info":{"pygments_lexer":"ipython3","nbconvert_exporter":"python","version":"3.6.4","file_extension":".py","codemirror_mode":{"name":"ipython","version":3},"name":"python","mimetype":"text/x-python"}},"nbformat_minor":4,"nbformat":4,"cells":[{"cell_type":"code","source":"# This Python 3 environment comes with many helpful analytics libraries installed\n# It is defined by the kaggle/python Docker image: https://github.com/kaggle/docker-python\n# For example, here's several helpful packages to load\n\nimport numpy as np # linear algebra\nimport pandas as pd # data processing, CSV file I/O (e.g. pd.read_csv)\n\n# Input data files are available in the read-only \"../input/\" directory\n# For example, running this (by clicking run or pressing Shift+Enter) will list all files under the input directory\n\nimport os\nfor dirname, _, filenames in os.walk('/kaggle/input'):\n    for filename in filenames:\n        print(os.path.join(dirname, filename))\n\n# You can write up to 20GB to the current directory (/kaggle/working/) that gets preserved as output when you create a version using \"Save & Run All\" \n# You can also write temporary files to /kaggle/temp/, but they won't be saved outside of the current session","metadata":{"_uuid":"8f2839f25d086af736a60e9eeb907d3b93b6e0e5","_cell_guid":"b1076dfc-b9ad-4769-8c92-a6c4dae69d19","trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"# U-Net Original Research Paper\n### Paper Name: U-Net: Convolutional Networks for Biomedical Image Segmentation\n### Paper Link: https://arxiv.org/pdf/1505.04597.pdf","metadata":{}},{"cell_type":"markdown","source":"# Configs","metadata":{}},{"cell_type":"code","source":"class ROOTDIR:\n    train = \"./train\"\n    train_mask = \"./train_masks\"","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:14.471203Z","iopub.execute_input":"2022-11-01T11:03:14.473909Z","iopub.status.idle":"2022-11-01T11:03:14.478534Z","shell.execute_reply.started":"2022-11-01T11:03:14.473866Z","shell.execute_reply":"2022-11-01T11:03:14.477663Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"# General Imports","metadata":{}},{"cell_type":"code","source":"import os\nimport cv2\nimport zipfile\nimport numpy as np\nimport pandas as pd\nfrom PIL import Image\nfrom tqdm.auto import tqdm\nimport matplotlib.pyplot as plt","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:14.48014Z","iopub.execute_input":"2022-11-01T11:03:14.480807Z","iopub.status.idle":"2022-11-01T11:03:14.487737Z","shell.execute_reply.started":"2022-11-01T11:03:14.48077Z","shell.execute_reply":"2022-11-01T11:03:14.486902Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"# Extracting Files","metadata":{}},{"cell_type":"code","source":"dirs = [\"../input/semantic-segmentation-of-underwater-imagery-suim/train_val/images\",\n        \"../input/semantic-segmentation-of-underwater-imagery-suim/train_val/masks\"]\n\n    ","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:14.490153Z","iopub.execute_input":"2022-11-01T11:03:14.490867Z","iopub.status.idle":"2022-11-01T11:03:14.496687Z","shell.execute_reply.started":"2022-11-01T11:03:14.49083Z","shell.execute_reply":"2022-11-01T11:03:14.495987Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"# Working with CSV","metadata":{}},{"cell_type":"code","source":"#df = pd.read_csv(\"./metadata.csv\")\n#df.head()","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:14.497875Z","iopub.execute_input":"2022-11-01T11:03:14.498762Z","iopub.status.idle":"2022-11-01T11:03:14.505551Z","shell.execute_reply.started":"2022-11-01T11:03:14.498727Z","shell.execute_reply":"2022-11-01T11:03:14.504784Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"#df.info()","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:14.507038Z","iopub.execute_input":"2022-11-01T11:03:14.507407Z","iopub.status.idle":"2022-11-01T11:03:14.515584Z","shell.execute_reply.started":"2022-11-01T11:03:14.507372Z","shell.execute_reply":"2022-11-01T11:03:14.514861Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"# Working with Train Images and Masks","metadata":{}},{"cell_type":"code","source":"train_img_lst = os.listdir(ROOTDIR.train) # \"./train\"\ntrain_mask_lst = os.listdir(ROOTDIR.train_mask) # \"./train_masks\"","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:14.516492Z","iopub.execute_input":"2022-11-01T11:03:14.518693Z","iopub.status.idle":"2022-11-01T11:03:14.530474Z","shell.execute_reply.started":"2022-11-01T11:03:14.516763Z","shell.execute_reply":"2022-11-01T11:03:14.529668Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"print(train_mask_lst[:5])\nprint(train_img_lst[:5])","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:14.533381Z","iopub.execute_input":"2022-11-01T11:03:14.533671Z","iopub.status.idle":"2022-11-01T11:03:14.540589Z","shell.execute_reply.started":"2022-11-01T11:03:14.533642Z","shell.execute_reply":"2022-11-01T11:03:14.539604Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"print(len(train_mask_lst))\nprint(len(train_img_lst))","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:14.542094Z","iopub.execute_input":"2022-11-01T11:03:14.542442Z","iopub.status.idle":"2022-11-01T11:03:14.548368Z","shell.execute_reply.started":"2022-11-01T11:03:14.542409Z","shell.execute_reply":"2022-11-01T11:03:14.547538Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"#### Sorting to make sure we get right image and right mask","metadata":{}},{"cell_type":"code","source":"sorted_train_mask_lst = sorted(train_mask_lst)","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:14.551784Z","iopub.execute_input":"2022-11-01T11:03:14.552123Z","iopub.status.idle":"2022-11-01T11:03:14.558852Z","shell.execute_reply.started":"2022-11-01T11:03:14.552098Z","shell.execute_reply":"2022-11-01T11:03:14.558029Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"sorted_train_img_lst = sorted(train_img_lst)","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:14.560158Z","iopub.execute_input":"2022-11-01T11:03:14.560573Z","iopub.status.idle":"2022-11-01T11:03:14.569821Z","shell.execute_reply.started":"2022-11-01T11:03:14.560537Z","shell.execute_reply":"2022-11-01T11:03:14.569043Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"print(sorted_train_mask_lst[:16])\nprint(sorted_train_img_lst[:16])","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:14.572582Z","iopub.execute_input":"2022-11-01T11:03:14.572898Z","iopub.status.idle":"2022-11-01T11:03:14.58105Z","shell.execute_reply.started":"2022-11-01T11:03:14.572873Z","shell.execute_reply":"2022-11-01T11:03:14.580094Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"# Visualizing Images with their Mask\n### Making sure images and mask are paired correctly.","metadata":{}},{"cell_type":"code","source":"def show_images(imgs_lst,masks_lst,loops=2):\n    for i in range(loops):\n        img_path = os.path.join(ROOTDIR.train,imgs_lst[i])\n        mask_path = os.path.join(ROOTDIR.train_mask,masks_lst[i])\n        img = Image.open(img_path)\n        mask = Image.open(mask_path)\n        print(img_path)\n        print(img.size)\n        print(type(img))\n        plt.imshow(img)\n        plt.show()\n        print(mask_path)\n        print(mask.size)\n        plt.imshow(mask)\n        plt.show()\n        print(\"----------------------------------------------------\")\n\nshow_images(sorted_train_img_lst, sorted_train_mask_lst)","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:14.582692Z","iopub.execute_input":"2022-11-01T11:03:14.583038Z","iopub.status.idle":"2022-11-01T11:03:16.274369Z","shell.execute_reply.started":"2022-11-01T11:03:14.583003Z","shell.execute_reply":"2022-11-01T11:03:16.27354Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"# PyTorch Imports","metadata":{}},{"cell_type":"code","source":"import torch\nimport torchvision\nimport torch.nn as nn\nimport albumentations as A\nimport torch.optim as optim\nfrom torchvision import models\nimport torch.nn.functional as F\nfrom torch.optim import lr_scheduler\nimport torchvision.datasets as datasets\nimport torchvision.transforms as transforms\nfrom torchvision.datasets import ImageFolder\nfrom albumentations.pytorch import ToTensorV2 \nfrom torch.utils.data import DataLoader, Dataset","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:16.275774Z","iopub.execute_input":"2022-11-01T11:03:16.276314Z","iopub.status.idle":"2022-11-01T11:03:16.282598Z","shell.execute_reply.started":"2022-11-01T11:03:16.276273Z","shell.execute_reply":"2022-11-01T11:03:16.281673Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"# PyTorch Configuration","metadata":{}},{"cell_type":"code","source":"class CFG:\n    device = torch.device(\"cuda\" if torch.cuda.is_available() else \"cpu\")\n    split_pct = 0.2\n    learning_rate = 3e-4\n    batch_size = 4\n    epochs = 3","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:16.283919Z","iopub.execute_input":"2022-11-01T11:03:16.284287Z","iopub.status.idle":"2022-11-01T11:03:16.293631Z","shell.execute_reply.started":"2022-11-01T11:03:16.28425Z","shell.execute_reply":"2022-11-01T11:03:16.292881Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"seed = 123\nnp.random.seed(seed)\ntorch.manual_seed(seed)","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:16.29503Z","iopub.execute_input":"2022-11-01T11:03:16.29541Z","iopub.status.idle":"2022-11-01T11:03:16.305349Z","shell.execute_reply.started":"2022-11-01T11:03:16.295372Z","shell.execute_reply":"2022-11-01T11:03:16.304482Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"CFG.device","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:16.307927Z","iopub.execute_input":"2022-11-01T11:03:16.308679Z","iopub.status.idle":"2022-11-01T11:03:16.320804Z","shell.execute_reply.started":"2022-11-01T11:03:16.308601Z","shell.execute_reply":"2022-11-01T11:03:16.319764Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"# Working with data","metadata":{}},{"cell_type":"markdown","source":"### Shuffling the data.","metadata":{}},{"cell_type":"code","source":"permuted_train_img_lst = np.random.permutation(np.array(sorted_train_img_lst))\npermuted_train_mask_lst = [x.replace(\".jpg\", \"_mask.gif\") for x in permuted_train_img_lst]\nprint(permuted_train_img_lst[:5])\nprint(permuted_train_mask_lst[:5])","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:16.322085Z","iopub.execute_input":"2022-11-01T11:03:16.32255Z","iopub.status.idle":"2022-11-01T11:03:16.332876Z","shell.execute_reply.started":"2022-11-01T11:03:16.322514Z","shell.execute_reply":"2022-11-01T11:03:16.332059Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"show_images(permuted_train_img_lst,permuted_train_mask_lst)","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:16.333968Z","iopub.execute_input":"2022-11-01T11:03:16.334581Z","iopub.status.idle":"2022-11-01T11:03:18.172464Z","shell.execute_reply.started":"2022-11-01T11:03:16.334544Z","shell.execute_reply":"2022-11-01T11:03:18.17162Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"### Splitting into Training and Validation","metadata":{}},{"cell_type":"code","source":"length = len(permuted_train_img_lst)\nprint(length*0.2) # convert this to int","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:18.174085Z","iopub.execute_input":"2022-11-01T11:03:18.174749Z","iopub.status.idle":"2022-11-01T11:03:18.180966Z","shell.execute_reply.started":"2022-11-01T11:03:18.174708Z","shell.execute_reply":"2022-11-01T11:03:18.179995Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"train_images_list = permuted_train_img_lst[int(CFG.split_pct*len(permuted_train_img_lst)) :]\ntrain_masks_list = permuted_train_mask_lst[int(CFG.split_pct*len(permuted_train_mask_lst)) :]\nprint(len(train_masks_list))\n\nval_images_list = permuted_train_img_lst[: int(CFG.split_pct*len(permuted_train_img_lst))]\nval_masks_list = permuted_train_mask_lst[: int(CFG.split_pct*len(permuted_train_mask_lst))]\nprint(len(val_masks_list))\n\n# 4071+1017=5088 (split includes all items)","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:18.182554Z","iopub.execute_input":"2022-11-01T11:03:18.182932Z","iopub.status.idle":"2022-11-01T11:03:18.191146Z","shell.execute_reply.started":"2022-11-01T11:03:18.182896Z","shell.execute_reply":"2022-11-01T11:03:18.190379Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"### Visualizing Train Dataset","metadata":{}},{"cell_type":"code","source":"show_images(train_images_list,train_masks_list)","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:18.192208Z","iopub.execute_input":"2022-11-01T11:03:18.19324Z","iopub.status.idle":"2022-11-01T11:03:19.877275Z","shell.execute_reply.started":"2022-11-01T11:03:18.193203Z","shell.execute_reply":"2022-11-01T11:03:19.876329Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"### Visualizing Validation Dataset","metadata":{}},{"cell_type":"code","source":"show_images(val_images_list,val_masks_list)","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:19.87865Z","iopub.execute_input":"2022-11-01T11:03:19.879101Z","iopub.status.idle":"2022-11-01T11:03:22.090398Z","shell.execute_reply.started":"2022-11-01T11:03:19.879063Z","shell.execute_reply":"2022-11-01T11:03:22.089534Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"# Dataset Class","metadata":{}},{"cell_type":"code","source":"class CarvanaDataset(Dataset):\n    def __init__(self,img_list,mask_list,transform=None):\n        self.img_list = img_list\n        self.mask_list = mask_list\n        self.transform = transform\n        \n    def __len__(self):\n        return len(self.img_list)\n    \n    def __getitem__(self,index):\n        img_path = os.path.join(ROOTDIR.train,self.img_list[index])\n        mask_path = os.path.join(ROOTDIR.train_mask,self.mask_list[index])\n        img = Image.open(img_path)\n        mask = Image.open(mask_path)\n        img = np.array(img)\n        mask = np.array(mask)\n        mask[mask==255.0] = 1.0\n        #img_mask_dict = {\"image\": img, \"mask\": mask}\n        \n        if self.transform:\n            augmentation = self.transform(image=img, mask=mask)\n            img = augmentation[\"image\"]\n            mask = augmentation[\"mask\"]\n            mask = torch.unsqueeze(mask,0)\n            #transformations = self.transform(image=img, mask=mask)\n            #img = transformations[\"image\"]\n            #mask = transformations[\"mask\"]\n            \n        return img,mask\n    ","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:22.091624Z","iopub.execute_input":"2022-11-01T11:03:22.092579Z","iopub.status.idle":"2022-11-01T11:03:22.101699Z","shell.execute_reply.started":"2022-11-01T11:03:22.092538Z","shell.execute_reply":"2022-11-01T11:03:22.100911Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"train_transform = A.Compose([A.Resize(572,572), \n                             A.Rotate(limit=15,p=0.1),\n                             A.HorizontalFlip(p=0.5),\n                             A.Normalize(mean=(0,0,0),std=(1,1,1),max_pixel_value=255),\n                             ToTensorV2()])\n\nval_transform = A.Compose([A.Resize(572,572),\n                           A.Normalize(mean=(0,0,0),std=(1,1,1),max_pixel_value=255),\n                           ToTensorV2()])","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:22.104565Z","iopub.execute_input":"2022-11-01T11:03:22.104888Z","iopub.status.idle":"2022-11-01T11:03:22.113321Z","shell.execute_reply.started":"2022-11-01T11:03:22.10486Z","shell.execute_reply":"2022-11-01T11:03:22.112486Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"train_dataset = CarvanaDataset(train_images_list, train_masks_list, transform = train_transform)\nval_dataset = CarvanaDataset(val_images_list, val_masks_list, transform = train_transform)","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:22.11448Z","iopub.execute_input":"2022-11-01T11:03:22.11507Z","iopub.status.idle":"2022-11-01T11:03:22.12657Z","shell.execute_reply.started":"2022-11-01T11:03:22.115031Z","shell.execute_reply":"2022-11-01T11:03:22.125795Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"idx = 200\nimg,mask = train_dataset[idx]","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:22.128059Z","iopub.execute_input":"2022-11-01T11:03:22.128404Z","iopub.status.idle":"2022-11-01T11:03:22.183336Z","shell.execute_reply.started":"2022-11-01T11:03:22.12837Z","shell.execute_reply":"2022-11-01T11:03:22.182511Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"mask.shape","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:22.188719Z","iopub.execute_input":"2022-11-01T11:03:22.189034Z","iopub.status.idle":"2022-11-01T11:03:22.194377Z","shell.execute_reply.started":"2022-11-01T11:03:22.189007Z","shell.execute_reply":"2022-11-01T11:03:22.193581Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"img.max()","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:22.195464Z","iopub.execute_input":"2022-11-01T11:03:22.196061Z","iopub.status.idle":"2022-11-01T11:03:22.207146Z","shell.execute_reply.started":"2022-11-01T11:03:22.196024Z","shell.execute_reply":"2022-11-01T11:03:22.20606Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"def show_single_img(img,mask,index=None,train=True):\n    if index:\n        if train:\n            img,mask = train_dataset[index]\n        else:\n            img,mask = val_dataset[index]\n    plt.imshow(img.permute(1,2,0),cmap=\"gray\")  # Convert (3, 572, 572) -> (572, 572, 3)\n    plt.show()\n    plt.imshow(mask.permute(1,2,0), cmap=\"gray\")  # Convert (1, 572, 572) -> (572, 572, 1)\n    print(mask.shape)\n    plt.show()\n    ","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:22.208702Z","iopub.execute_input":"2022-11-01T11:03:22.209139Z","iopub.status.idle":"2022-11-01T11:03:22.216126Z","shell.execute_reply.started":"2022-11-01T11:03:22.209104Z","shell.execute_reply":"2022-11-01T11:03:22.215249Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"print(\"---------------Train---------------\")\nshow_single_img(img,mask,index=15,train=False)\nprint(\"---------------Validation---------------\")\nshow_single_img(img,mask,index=15,train=True)","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:22.217315Z","iopub.execute_input":"2022-11-01T11:03:22.21798Z","iopub.status.idle":"2022-11-01T11:03:23.158741Z","shell.execute_reply.started":"2022-11-01T11:03:22.217877Z","shell.execute_reply":"2022-11-01T11:03:23.157846Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"# Dataloader","metadata":{}},{"cell_type":"code","source":"train_dataloader = DataLoader(train_dataset,batch_size=CFG.batch_size,shuffle=True)\nval_dataloader = DataLoader(val_dataset,batch_size=CFG.batch_size,shuffle=False)","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:23.15992Z","iopub.execute_input":"2022-11-01T11:03:23.16033Z","iopub.status.idle":"2022-11-01T11:03:23.165588Z","shell.execute_reply.started":"2022-11-01T11:03:23.160291Z","shell.execute_reply":"2022-11-01T11:03:23.164516Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"a = iter(train_dataloader)\nimg,mask = a.next()\nprint(img.shape,mask.shape)","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:23.166769Z","iopub.execute_input":"2022-11-01T11:03:23.167502Z","iopub.status.idle":"2022-11-01T11:03:23.372841Z","shell.execute_reply.started":"2022-11-01T11:03:23.167464Z","shell.execute_reply":"2022-11-01T11:03:23.371861Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"### Utility Functions","metadata":{}},{"cell_type":"code","source":"def double_conv(in_ch, out_ch):\n    conv = nn.Sequential(\n        nn.Conv2d(in_channels=in_ch,out_channels=out_ch,kernel_size=3,stride=1,padding=1),\n        nn.BatchNorm2d(out_ch),                                                            \n        nn.ReLU(inplace=True),\n        nn.Conv2d(in_channels=out_ch,out_channels=out_ch,kernel_size=3,stride=1,padding=1), \n        nn.BatchNorm2d(out_ch),                                                            \n        nn.ReLU(inplace=True)\n    )\n    \n    return conv\n\n#def cropper(og_tensor, target_tensor):\n#    og_shape = og_tensor.shape[2]\n#    target_shape = target_tensor.shape[2]\n#    delta = (og_shape - target_shape) // 2\n#    cropped_og_tensor = og_tensor[:,:,delta:og_shape-delta,delta:og_shape-delta]\n#    return cropped_og_tensor\n \n    \ndef padder(left_tensor, right_tensor): \n    # left_tensor is the tensor on the encoder side of UNET\n    # right_tensor is the tensor on the decoder side  of the UNET\n    \n    if left_tensor.shape != right_tensor.shape:\n        padded = torch.zeros(left_tensor.shape)\n        padded[:, :, :right_tensor.shape[2], :right_tensor.shape[3]] = right_tensor\n        return padded.to(CFG.device)\n    \n    return right_tensor.to(CFG.device)\n    ","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:23.374107Z","iopub.execute_input":"2022-11-01T11:03:23.375012Z","iopub.status.idle":"2022-11-01T11:03:23.384513Z","shell.execute_reply.started":"2022-11-01T11:03:23.374971Z","shell.execute_reply":"2022-11-01T11:03:23.383653Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"# UNET MODEL FROM SCRATCH","metadata":{}},{"cell_type":"code","source":"class UNET(nn.Module):\n    def __init__(self,in_chnls, n_classes):\n        super(UNET,self).__init__()\n        \n        self.in_chnls = in_chnls\n        self.n_classes = n_classes\n        \n        self.max_pool = nn.MaxPool2d(kernel_size=2,stride=2)\n        \n        self.down_conv_1 = double_conv(in_ch=self.in_chnls,out_ch=64)\n        self.down_conv_2 = double_conv(in_ch=64,out_ch=128)\n        self.down_conv_3 = double_conv(in_ch=128,out_ch=256)\n        self.down_conv_4 = double_conv(in_ch=256,out_ch=512)\n        self.down_conv_5 = double_conv(in_ch=512,out_ch=1024)\n        #print(self.down_conv_1)\n        \n        self.up_conv_trans_1 = nn.ConvTranspose2d(in_channels=1024,out_channels=512,kernel_size=2,stride=2)\n        self.up_conv_trans_2 = nn.ConvTranspose2d(in_channels=512,out_channels=256,kernel_size=2,stride=2)\n        self.up_conv_trans_3 = nn.ConvTranspose2d(in_channels=256,out_channels=128,kernel_size=2,stride=2)\n        self.up_conv_trans_4 = nn.ConvTranspose2d(in_channels=128,out_channels=64,kernel_size=2,stride=2)\n        \n        self.up_conv_1 = double_conv(in_ch=1024,out_ch=512)\n        self.up_conv_2 = double_conv(in_ch=512,out_ch=256)\n        self.up_conv_3 = double_conv(in_ch=256,out_ch=128)\n        self.up_conv_4 = double_conv(in_ch=128,out_ch=64)\n        \n        self.conv_1x1 = nn.Conv2d(in_channels=64,out_channels=self.n_classes,kernel_size=1,stride=1)\n        \n    def forward(self,x):\n        \n        # encoding\n        x1 = self.down_conv_1(x)\n        #print(\"X1\", x1.shape)\n        p1 = self.max_pool(x1)\n        #print(\"p1\", p1.shape)\n        x2 = self.down_conv_2(p1)\n        #print(\"X2\", x2.shape)\n        p2 = self.max_pool(x2)\n        #print(\"p2\", p2.shape)\n        x3 = self.down_conv_3(p2)\n        #print(\"X2\", x3.shape)\n        p3 = self.max_pool(x3)\n        #print(\"p3\", p3.shape)\n        x4 = self.down_conv_4(p3)\n        #print(\"X4\", x4.shape)\n        p4 = self.max_pool(x4)\n        #print(\"p4\", p4.shape)\n        x5 = self.down_conv_5(p4)\n        #print(\"X5\", x5.shape)\n        \n        # decoding\n        d1 = self.up_conv_trans_1(x5)  # up transpose convolution (\"up sampling\" as called in UNET paper)\n        pad1 = padder(x4,d1) # padding d1 to match x4 shape\n        cat1 = torch.cat([x4,pad1],dim=1) # concatenating padded d1 and x4 on channel dimension(dim 1) [batch(dim 0),channel(dim 1),height(dim 2),width(dim 3)]\n        uc1 = self.up_conv_1(cat1) # 1st up double convolution\n        \n        d2 = self.up_conv_trans_2(uc1)\n        pad2 = padder(x3,d2)\n        cat2 = torch.cat([x3,pad2],dim=1)\n        uc2 = self.up_conv_2(cat2)\n        \n        d3 = self.up_conv_trans_3(uc2)\n        pad3 = padder(x2,d3)\n        cat3 = torch.cat([x2,pad3],dim=1)\n        uc3 = self.up_conv_3(cat3)\n        \n        d4 = self.up_conv_trans_4(uc3)\n        pad4 = padder(x1,d4)\n        cat4 = torch.cat([x1,pad4],dim=1)\n        uc4 = self.up_conv_4(cat4)\n        \n        conv_1x1 = self.conv_1x1(uc4)\n        return conv_1x1\n        #print(conv_1x1.shape)","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:23.385776Z","iopub.execute_input":"2022-11-01T11:03:23.386269Z","iopub.status.idle":"2022-11-01T11:03:23.403314Z","shell.execute_reply.started":"2022-11-01T11:03:23.386232Z","shell.execute_reply":"2022-11-01T11:03:23.402442Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"# Training and Validation","metadata":{}},{"cell_type":"markdown","source":"### Train Function\n","metadata":{}},{"cell_type":"code","source":"def train_model(model,dataloader,criterion,optimizer):\n    model.train()\n    train_running_loss = 0.0\n    for j,img_mask in enumerate(tqdm(dataloader)):\n        img = img_mask[0].float().to(CFG.device)\n        #print(\" ----- IMAGE -----\")\n        #print(img)\n        mask = img_mask[1].float().to(CFG.device)\n        #print(\" ----- MASK -----\")\n        #print(mask)\n        \n        y_pred = model(img)\n        #print(\" ----- Y PRED -----\")\n        #print(y_pred)\n        #print(\" ----- Y PRED SHAPE -----\")#\n        #print(y_pred.shape)\n        optimizer.zero_grad()\n        \n        loss = criterion(y_pred,mask)\n        \n        train_running_loss += loss.item() * CFG.batch_size\n        \n        loss.backward()\n        optimizer.step()\n        \n    train_loss = train_running_loss / (j+1)\n    return train_loss","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:23.404542Z","iopub.execute_input":"2022-11-01T11:03:23.404882Z","iopub.status.idle":"2022-11-01T11:03:23.416883Z","shell.execute_reply.started":"2022-11-01T11:03:23.404855Z","shell.execute_reply":"2022-11-01T11:03:23.416163Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"### Validation Function","metadata":{}},{"cell_type":"code","source":"def val_model(model,dataloader,criterion,optimizer):\n    model.eval()\n    val_running_loss = 0\n    with torch.no_grad():\n        for j,img_mask in enumerate(tqdm(dataloader)):\n            img = img_mask[0].float().to(CFG.device)\n            mask = img_mask[1].float().to(CFG.device)\n            y_pred = model(img)\n            loss = criterion(y_pred,mask)\n            \n            val_running_loss += loss.item() * CFG.batch_size\n            \n        val_loss = val_running_loss / (j+1)\n    return val_loss","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:23.41819Z","iopub.execute_input":"2022-11-01T11:03:23.41879Z","iopub.status.idle":"2022-11-01T11:03:23.427744Z","shell.execute_reply.started":"2022-11-01T11:03:23.418751Z","shell.execute_reply":"2022-11-01T11:03:23.427158Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"model = UNET(in_chnls = 3, n_classes = 1).to(CFG.device)\noptimizer = optim.Adam(model.parameters(), lr = CFG.learning_rate)\ncriterion = nn.BCEWithLogitsLoss()\ntrain_loss_lst = []\nval_loss_lst = []  ","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:23.428719Z","iopub.execute_input":"2022-11-01T11:03:23.429533Z","iopub.status.idle":"2022-11-01T11:03:23.695585Z","shell.execute_reply.started":"2022-11-01T11:03:23.429496Z","shell.execute_reply":"2022-11-01T11:03:23.69487Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"### Train and Validation Loop","metadata":{}},{"cell_type":"code","source":"for i in tqdm(range(CFG.epochs)):\n    train_loss = train_model(model=model,dataloader=train_dataloader,criterion=criterion,optimizer=optimizer)\n    val_loss = val_model(model=model,dataloader=val_dataloader,criterion=criterion,optimizer=optimizer)\n    train_loss_lst.append(train_loss)\n    val_loss_lst.append(val_loss)\n    print(f\" Train Loss : {train_loss:.4f}\")\n    print(f\" Validation Loss : {val_loss:.4f}\")","metadata":{"execution":{"iopub.status.busy":"2022-11-01T11:03:23.696629Z","iopub.execute_input":"2022-11-01T11:03:23.696981Z"},"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"### Training and Validation Loss Plot","metadata":{}},{"cell_type":"code","source":"plt.plot(train_loss_lst, color=\"green\", label='train loss')\nplt.plot(val_loss_lst, color=\"red\", label='validation loss')\nplt.xlabel(\"epochs\")\nplt.ylabel(\"loss\")\nplt.legend()\nplt.show()","metadata":{"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"# Saving Model","metadata":{}},{"cell_type":"code","source":"TRAINED_FILE = \"./unet_scratch.pth\"","metadata":{"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"torch.save(model.state_dict(), TRAINED_FILE)","metadata":{"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"from IPython.display import FileLink\nFileLink(TRAINED_FILE)","metadata":{"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"markdown","source":"# Testing","metadata":{}},{"cell_type":"code","source":"trained_model = UNET(in_chnls = 3, n_classes = 1)","metadata":{"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"UNET_TRAINED = \"../input/unet-4-epoch-trained/unet_scratch.pth\"","metadata":{"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"trained_model.load_state_dict(torch.load(UNET_TRAINED))","metadata":{"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"trained_model = trained_model.to(\"cuda\")\ntrained_model.eval()\n","metadata":{"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"img_path = \"../input/carvana-image-masking-challenge/29bb3ece3180_11.jpg\"\n\nimg = cv2.imread(img_path)","metadata":{"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"plt.imshow(img)\nplt.show()","metadata":{"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"test_transform = A.Compose([A.Resize(572,572),\n                           A.Normalize(mean=(0,0,0),std=(1,1,1),max_pixel_value=255),\n                           ToTensorV2()])","metadata":{"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"test_image = test_transform(image = img)\n\nprint(test_image)\n\nprint(test_image[\"image\"].dtype)\nprint(test_image[\"image\"].shape)\n\nimg = test_image[\"image\"].unsqueeze(0)\nprint(img.shape)\n\nimg = img.to(\"cuda\")","metadata":{"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"pred = trained_model(img)\npred.shape","metadata":{"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"mask = pred.squeeze(0).cpu().detach().numpy()\nprint(mask.shape)\nmask = mask.transpose(1,2,0)\nprint(mask.shape)","metadata":{"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"display_test_img = test_image[\"image\"].cpu().detach().numpy()\nprint(display_test_img.shape)\ndisplay_test_img = display_test_img.transpose(1,2,0)\ndisplay_test_img.shape","metadata":{"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"mask[mask < 0]=0\nmask[mask > 0]=1","metadata":{"trusted":true},"execution_count":null,"outputs":[]},{"cell_type":"code","source":"print(\"-------Original Image-------\")\nplt.imshow(display_test_img, cmap=\"gray\")\nplt.show()\nprint(\"-------Image Mask-------\")\nplt.imshow(mask,cmap=\"gray\")\nplt.show()","metadata":{"trusted":true},"execution_count":null,"outputs":[]}]}